Johnny221B commited on
Commit
c86ee86
·
verified ·
1 Parent(s): 10e71a3

Promote validation-selected Qwen2.5-3B-Instruct PersonaMem-32K reader

Browse files
README.md CHANGED
@@ -41,6 +41,18 @@ Both memory generation and answer decoding are greedy; trials use different fixe
41
 
42
  Download `personamem-32k/qwen3-4b/` from the default branch for this updated reader. To reproduce the original paper release, set `revision="d39842eeca8132288c5e2814fdd66733fa92fc3d"` when downloading. The previous weights remain available at [that fixed revision](https://huggingface.co/Johnny221B/memfold/tree/d39842eeca8132288c5e2814fdd66733fa92fc3d). This post-release update does not replace historical paper results.
43
 
 
 
 
 
 
 
 
 
 
 
 
 
44
  ## Components and names
45
 
46
  MemFold uses one workflow:
 
41
 
42
  Download `personamem-32k/qwen3-4b/` from the default branch for this updated reader. To reproduce the original paper release, set `revision="d39842eeca8132288c5e2814fdd66733fa92fc3d"` when downloading. The previous weights remain available at [that fixed revision](https://huggingface.co/Johnny221B/memfold/tree/d39842eeca8132288c5e2814fdd66733fa92fc3d). This post-release update does not replace historical paper results.
43
 
44
+ ## Updated Qwen2.5-3B-Instruct checkpoint
45
+
46
+ The default `personamem-32k/qwen2.5-3b/` reader is now **current-ratio-epoch-4**, selected by five-trial validation mean. Only the final OPD + GRPO phase was retrained; its compressor is unchanged.
47
+
48
+ | Split | Mean accuracy | Best trial |
49
+ |---|---:|---:|
50
+ | Previous reader validation | 60.0% | 70.0% |
51
+ | Updated reader validation | 67.6% | 76.0% |
52
+ | Updated reader test | 61.2% | 66.0% |
53
+
54
+ Memory generation and answers use greedy decoding; trials change answer-option order. Validation best scores are not test scores. The test split was evaluated after checkpoint selection. [Settings and all validation results](personamem-32k/qwen2.5-3b/optimization.json). Previous weights remain at revision `10e71a32171ccde6f3a53bc8ee5b407264321c3d`. This is a post-release update, not a replacement for historical paper results.
55
+
56
  ## Components and names
57
 
58
  MemFold uses one workflow:
VALIDATION.json CHANGED
@@ -190,5 +190,83 @@
190
  "seconds": 50.37944988813251
191
  },
192
  "protocol": "OPTIMIZATION.json"
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
193
  }
194
  }
 
190
  "seconds": 50.37944988813251
191
  },
192
  "protocol": "OPTIMIZATION.json"
193
+ },
194
+ "optimized_defaults": {
195
+ "personamem-32k/qwen2.5-3b": {
196
+ "source_revision": "45b50a33bbcf55013f0c720a6980ea4fcc5e5db3",
197
+ "validation": {
198
+ "trials": [
199
+ {
200
+ "trial": 0,
201
+ "correct": 32,
202
+ "total": 50,
203
+ "accuracy": 0.64
204
+ },
205
+ {
206
+ "trial": 1,
207
+ "correct": 34,
208
+ "total": 50,
209
+ "accuracy": 0.68
210
+ },
211
+ {
212
+ "trial": 2,
213
+ "correct": 34,
214
+ "total": 50,
215
+ "accuracy": 0.68
216
+ },
217
+ {
218
+ "trial": 3,
219
+ "correct": 38,
220
+ "total": 50,
221
+ "accuracy": 0.76
222
+ },
223
+ {
224
+ "trial": 4,
225
+ "correct": 31,
226
+ "total": 50,
227
+ "accuracy": 0.62
228
+ }
229
+ ],
230
+ "mean_accuracy": 0.6759999999999999,
231
+ "best_accuracy": 0.76
232
+ },
233
+ "test": {
234
+ "trials": [
235
+ {
236
+ "trial": 0,
237
+ "correct": 32,
238
+ "total": 50,
239
+ "accuracy": 0.64
240
+ },
241
+ {
242
+ "trial": 1,
243
+ "correct": 28,
244
+ "total": 50,
245
+ "accuracy": 0.56
246
+ },
247
+ {
248
+ "trial": 2,
249
+ "correct": 30,
250
+ "total": 50,
251
+ "accuracy": 0.6
252
+ },
253
+ {
254
+ "trial": 3,
255
+ "correct": 30,
256
+ "total": 50,
257
+ "accuracy": 0.6
258
+ },
259
+ {
260
+ "trial": 4,
261
+ "correct": 33,
262
+ "total": 50,
263
+ "accuracy": 0.66
264
+ }
265
+ ],
266
+ "mean_accuracy": 0.6120000000000001,
267
+ "best_accuracy": 0.66
268
+ },
269
+ "details": "personamem-32k/qwen2.5-3b/optimization.json"
270
+ }
271
  }
272
  }
manifest.json CHANGED
@@ -178,16 +178,16 @@
178
  "path": "personamem-32k/qwen2.5-3b/reader/adapter_model.safetensors",
179
  "role": "reader_lora",
180
  "size": 119801528,
181
- "sha256": "ecf6da9d83cb82890aeffd4079d86880b29c3a62d7e8cf919329d47ae7197f71",
182
- "original_sha256": "ecf6da9d83cb82890aeffd4079d86880b29c3a62d7e8cf919329d47ae7197f71",
183
  "metadata_only_repack": false
184
  },
185
  {
186
  "path": "personamem-32k/qwen2.5-3b/reader/adapter_config.json",
187
  "role": "reader_config",
188
  "size": 975,
189
- "sha256": "521c2e7073bfd1a9f621ec16624cfaf915b8d2b91b86eafdaf297d8b8e5bd36c",
190
- "original_sha256": "35c900dd3a881bb6fd6e6740cebc4a8277be859b547379ae5cb240e0c36e7822",
191
  "metadata_only_repack": true
192
  },
193
  {
@@ -411,7 +411,8 @@
411
  "Only selected MemFold main-result reader weights and inference dependencies are included.",
412
  "PersonaMem bridge tensors are unchanged; training-only metadata is omitted.",
413
  "LoCoMo transfer used a shared Qwen3 writer and bounded-text fusion.",
414
- "2026-09-29: PersonaMem-32K Qwen3-4B reader updated to validation-selected more-grpo epoch 5 (615 steps). Original paper release remains at revision d39842eeca8132288c5e2814fdd66733fa92fc3d. See OPTIMIZATION.json."
 
415
  ],
416
  "default_qwen3_4b_personamem32k": {
417
  "configuration": "more-grpo",
@@ -419,5 +420,14 @@
419
  "steps": 615,
420
  "source_revision": "01e46fd762eeff06efaba52d6170d7daea2f0da2",
421
  "original_release_revision": "d39842eeca8132288c5e2814fdd66733fa92fc3d"
 
 
 
 
 
 
 
 
 
422
  }
423
  }
 
178
  "path": "personamem-32k/qwen2.5-3b/reader/adapter_model.safetensors",
179
  "role": "reader_lora",
180
  "size": 119801528,
181
+ "sha256": "4fdf8e12b6211301883425c5d792843698d9d9b884ce391a9380d76e112d7a4d",
182
+ "original_sha256": "4fdf8e12b6211301883425c5d792843698d9d9b884ce391a9380d76e112d7a4d",
183
  "metadata_only_repack": false
184
  },
185
  {
186
  "path": "personamem-32k/qwen2.5-3b/reader/adapter_config.json",
187
  "role": "reader_config",
188
  "size": 975,
189
+ "sha256": "e3631907424d1fe0f9adc9c104405b661f89d35adc662bd4fee1e59fc86a9380",
190
+ "original_sha256": "164804506e6dfd2b3ec68f9c33b7b1c93126554bcd78b44d24a7966db5ec761a",
191
  "metadata_only_repack": true
192
  },
193
  {
 
411
  "Only selected MemFold main-result reader weights and inference dependencies are included.",
412
  "PersonaMem bridge tensors are unchanged; training-only metadata is omitted.",
413
  "LoCoMo transfer used a shared Qwen3 writer and bounded-text fusion.",
414
+ "2026-09-29: PersonaMem-32K Qwen3-4B reader updated to validation-selected more-grpo epoch 5 (615 steps). Original paper release remains at revision d39842eeca8132288c5e2814fdd66733fa92fc3d. See OPTIMIZATION.json.",
415
+ "Post-release update of personamem-32k/qwen2.5-3b: current-ratio-epoch-4, selected by greedy validation mean. Previous weights: 10e71a32171ccde6f3a53bc8ee5b407264321c3d. See personamem-32k/qwen2.5-3b/optimization.json."
416
  ],
417
  "default_qwen3_4b_personamem32k": {
418
  "configuration": "more-grpo",
 
420
  "steps": 615,
421
  "source_revision": "01e46fd762eeff06efaba52d6170d7daea2f0da2",
422
  "original_release_revision": "d39842eeca8132288c5e2814fdd66733fa92fc3d"
423
+ },
424
+ "optimized_defaults": {
425
+ "personamem-32k/qwen2.5-3b": {
426
+ "checkpoint": "current-ratio-epoch-4",
427
+ "source_revision": "45b50a33bbcf55013f0c720a6980ea4fcc5e5db3",
428
+ "previous_revision": "10e71a32171ccde6f3a53bc8ee5b407264321c3d",
429
+ "validation_mean": 0.6759999999999999,
430
+ "test_mean": 0.6120000000000001
431
+ }
432
  }
433
  }
personamem-32k/qwen2.5-3b/optimization.json ADDED
@@ -0,0 +1,877 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "date": "2026-09-29",
3
+ "code_revision": "d049fd293e3232358594cba7cdae0f7a138bb58d",
4
+ "base_release": "10e71a32171ccde6f3a53bc8ee5b407264321c3d",
5
+ "bundle": "personamem-32k/qwen2.5-3b",
6
+ "phase": "on_policy_optimization",
7
+ "selection_criterion": "Maximum five-trial greedy validation mean; ties prefer earlier epoch then configuration order. Test only after selection.",
8
+ "selected_configuration": {
9
+ "name": "current-ratio",
10
+ "learning_rate": 3e-07,
11
+ "opd_weight": 0.002,
12
+ "grpo_weight": 0.3
13
+ },
14
+ "selected_epoch": 4,
15
+ "selected_optimizer_steps": 492,
16
+ "reference_kl_weight": 0,
17
+ "teacher_update_mode": "frozen_initial",
18
+ "epochs_per_configuration": 5,
19
+ "world_size": 4,
20
+ "questions_per_update": 4,
21
+ "rollouts_per_question": 8,
22
+ "rollouts_per_update": 32,
23
+ "seed": 42,
24
+ "sampling_temperature": 1.0,
25
+ "top_p": 0.98,
26
+ "max_response_tokens": 5,
27
+ "training_examples": 489,
28
+ "validation_examples": 50,
29
+ "test_examples": 50,
30
+ "evaluation": "Full-pipeline greedy memory generation and five fixed option permutations; test evaluated after validation selection.",
31
+ "validation_delta_vs_incumbent": 0.07599999999999996,
32
+ "validation_results": [
33
+ {
34
+ "name": "current-ratio-epoch-1",
35
+ "trials": [
36
+ {
37
+ "trial": 0,
38
+ "correct": 25,
39
+ "total": 50,
40
+ "accuracy": 0.5
41
+ },
42
+ {
43
+ "trial": 1,
44
+ "correct": 30,
45
+ "total": 50,
46
+ "accuracy": 0.6
47
+ },
48
+ {
49
+ "trial": 2,
50
+ "correct": 26,
51
+ "total": 50,
52
+ "accuracy": 0.52
53
+ },
54
+ {
55
+ "trial": 3,
56
+ "correct": 35,
57
+ "total": 50,
58
+ "accuracy": 0.7
59
+ },
60
+ {
61
+ "trial": 4,
62
+ "correct": 28,
63
+ "total": 50,
64
+ "accuracy": 0.56
65
+ }
66
+ ],
67
+ "mean_accuracy": 0.5760000000000001,
68
+ "best_accuracy": 0.7,
69
+ "seconds": 39.50442323461175
70
+ },
71
+ {
72
+ "name": "current-ratio-epoch-2",
73
+ "trials": [
74
+ {
75
+ "trial": 0,
76
+ "correct": 31,
77
+ "total": 50,
78
+ "accuracy": 0.62
79
+ },
80
+ {
81
+ "trial": 1,
82
+ "correct": 32,
83
+ "total": 50,
84
+ "accuracy": 0.64
85
+ },
86
+ {
87
+ "trial": 2,
88
+ "correct": 32,
89
+ "total": 50,
90
+ "accuracy": 0.64
91
+ },
92
+ {
93
+ "trial": 3,
94
+ "correct": 35,
95
+ "total": 50,
96
+ "accuracy": 0.7
97
+ },
98
+ {
99
+ "trial": 4,
100
+ "correct": 28,
101
+ "total": 50,
102
+ "accuracy": 0.56
103
+ }
104
+ ],
105
+ "mean_accuracy": 0.6319999999999999,
106
+ "best_accuracy": 0.7,
107
+ "seconds": 40.20348156709224
108
+ },
109
+ {
110
+ "name": "current-ratio-epoch-3",
111
+ "trials": [
112
+ {
113
+ "trial": 0,
114
+ "correct": 31,
115
+ "total": 50,
116
+ "accuracy": 0.62
117
+ },
118
+ {
119
+ "trial": 1,
120
+ "correct": 31,
121
+ "total": 50,
122
+ "accuracy": 0.62
123
+ },
124
+ {
125
+ "trial": 2,
126
+ "correct": 32,
127
+ "total": 50,
128
+ "accuracy": 0.64
129
+ },
130
+ {
131
+ "trial": 3,
132
+ "correct": 34,
133
+ "total": 50,
134
+ "accuracy": 0.68
135
+ },
136
+ {
137
+ "trial": 4,
138
+ "correct": 32,
139
+ "total": 50,
140
+ "accuracy": 0.64
141
+ }
142
+ ],
143
+ "mean_accuracy": 0.64,
144
+ "best_accuracy": 0.68,
145
+ "seconds": 39.52767578139901
146
+ },
147
+ {
148
+ "name": "current-ratio-epoch-4",
149
+ "trials": [
150
+ {
151
+ "trial": 0,
152
+ "correct": 32,
153
+ "total": 50,
154
+ "accuracy": 0.64
155
+ },
156
+ {
157
+ "trial": 1,
158
+ "correct": 34,
159
+ "total": 50,
160
+ "accuracy": 0.68
161
+ },
162
+ {
163
+ "trial": 2,
164
+ "correct": 34,
165
+ "total": 50,
166
+ "accuracy": 0.68
167
+ },
168
+ {
169
+ "trial": 3,
170
+ "correct": 38,
171
+ "total": 50,
172
+ "accuracy": 0.76
173
+ },
174
+ {
175
+ "trial": 4,
176
+ "correct": 31,
177
+ "total": 50,
178
+ "accuracy": 0.62
179
+ }
180
+ ],
181
+ "mean_accuracy": 0.6759999999999999,
182
+ "best_accuracy": 0.76,
183
+ "seconds": 39.636526766233146
184
+ },
185
+ {
186
+ "name": "current-ratio-epoch-5",
187
+ "trials": [
188
+ {
189
+ "trial": 0,
190
+ "correct": 30,
191
+ "total": 50,
192
+ "accuracy": 0.6
193
+ },
194
+ {
195
+ "trial": 1,
196
+ "correct": 35,
197
+ "total": 50,
198
+ "accuracy": 0.7
199
+ },
200
+ {
201
+ "trial": 2,
202
+ "correct": 31,
203
+ "total": 50,
204
+ "accuracy": 0.62
205
+ },
206
+ {
207
+ "trial": 3,
208
+ "correct": 35,
209
+ "total": 50,
210
+ "accuracy": 0.7
211
+ },
212
+ {
213
+ "trial": 4,
214
+ "correct": 30,
215
+ "total": 50,
216
+ "accuracy": 0.6
217
+ }
218
+ ],
219
+ "mean_accuracy": 0.644,
220
+ "best_accuracy": 0.7,
221
+ "seconds": 35.50570226460695
222
+ },
223
+ {
224
+ "name": "incumbent",
225
+ "trials": [
226
+ {
227
+ "trial": 0,
228
+ "correct": 26,
229
+ "total": 50,
230
+ "accuracy": 0.52
231
+ },
232
+ {
233
+ "trial": 1,
234
+ "correct": 29,
235
+ "total": 50,
236
+ "accuracy": 0.58
237
+ },
238
+ {
239
+ "trial": 2,
240
+ "correct": 31,
241
+ "total": 50,
242
+ "accuracy": 0.62
243
+ },
244
+ {
245
+ "trial": 3,
246
+ "correct": 35,
247
+ "total": 50,
248
+ "accuracy": 0.7
249
+ },
250
+ {
251
+ "trial": 4,
252
+ "correct": 29,
253
+ "total": 50,
254
+ "accuracy": 0.58
255
+ }
256
+ ],
257
+ "mean_accuracy": 0.6,
258
+ "best_accuracy": 0.7,
259
+ "seconds": 38.32996504101902
260
+ },
261
+ {
262
+ "name": "lower-lr-epoch-1",
263
+ "trials": [
264
+ {
265
+ "trial": 0,
266
+ "correct": 28,
267
+ "total": 50,
268
+ "accuracy": 0.56
269
+ },
270
+ {
271
+ "trial": 1,
272
+ "correct": 35,
273
+ "total": 50,
274
+ "accuracy": 0.7
275
+ },
276
+ {
277
+ "trial": 2,
278
+ "correct": 32,
279
+ "total": 50,
280
+ "accuracy": 0.64
281
+ },
282
+ {
283
+ "trial": 3,
284
+ "correct": 34,
285
+ "total": 50,
286
+ "accuracy": 0.68
287
+ },
288
+ {
289
+ "trial": 4,
290
+ "correct": 29,
291
+ "total": 50,
292
+ "accuracy": 0.58
293
+ }
294
+ ],
295
+ "mean_accuracy": 0.632,
296
+ "best_accuracy": 0.7,
297
+ "seconds": 39.12612637318671
298
+ },
299
+ {
300
+ "name": "lower-lr-epoch-2",
301
+ "trials": [
302
+ {
303
+ "trial": 0,
304
+ "correct": 27,
305
+ "total": 50,
306
+ "accuracy": 0.54
307
+ },
308
+ {
309
+ "trial": 1,
310
+ "correct": 32,
311
+ "total": 50,
312
+ "accuracy": 0.64
313
+ },
314
+ {
315
+ "trial": 2,
316
+ "correct": 31,
317
+ "total": 50,
318
+ "accuracy": 0.62
319
+ },
320
+ {
321
+ "trial": 3,
322
+ "correct": 34,
323
+ "total": 50,
324
+ "accuracy": 0.68
325
+ },
326
+ {
327
+ "trial": 4,
328
+ "correct": 29,
329
+ "total": 50,
330
+ "accuracy": 0.58
331
+ }
332
+ ],
333
+ "mean_accuracy": 0.6120000000000001,
334
+ "best_accuracy": 0.68,
335
+ "seconds": 31.27948883548379
336
+ },
337
+ {
338
+ "name": "lower-lr-epoch-3",
339
+ "trials": [
340
+ {
341
+ "trial": 0,
342
+ "correct": 28,
343
+ "total": 50,
344
+ "accuracy": 0.56
345
+ },
346
+ {
347
+ "trial": 1,
348
+ "correct": 32,
349
+ "total": 50,
350
+ "accuracy": 0.64
351
+ },
352
+ {
353
+ "trial": 2,
354
+ "correct": 32,
355
+ "total": 50,
356
+ "accuracy": 0.64
357
+ },
358
+ {
359
+ "trial": 3,
360
+ "correct": 35,
361
+ "total": 50,
362
+ "accuracy": 0.7
363
+ },
364
+ {
365
+ "trial": 4,
366
+ "correct": 32,
367
+ "total": 50,
368
+ "accuracy": 0.64
369
+ }
370
+ ],
371
+ "mean_accuracy": 0.636,
372
+ "best_accuracy": 0.7,
373
+ "seconds": 34.68806564621627
374
+ },
375
+ {
376
+ "name": "lower-lr-epoch-4",
377
+ "trials": [
378
+ {
379
+ "trial": 0,
380
+ "correct": 29,
381
+ "total": 50,
382
+ "accuracy": 0.58
383
+ },
384
+ {
385
+ "trial": 1,
386
+ "correct": 32,
387
+ "total": 50,
388
+ "accuracy": 0.64
389
+ },
390
+ {
391
+ "trial": 2,
392
+ "correct": 32,
393
+ "total": 50,
394
+ "accuracy": 0.64
395
+ },
396
+ {
397
+ "trial": 3,
398
+ "correct": 35,
399
+ "total": 50,
400
+ "accuracy": 0.7
401
+ },
402
+ {
403
+ "trial": 4,
404
+ "correct": 29,
405
+ "total": 50,
406
+ "accuracy": 0.58
407
+ }
408
+ ],
409
+ "mean_accuracy": 0.6279999999999999,
410
+ "best_accuracy": 0.7,
411
+ "seconds": 29.452694308012724
412
+ },
413
+ {
414
+ "name": "lower-lr-epoch-5",
415
+ "trials": [
416
+ {
417
+ "trial": 0,
418
+ "correct": 29,
419
+ "total": 50,
420
+ "accuracy": 0.58
421
+ },
422
+ {
423
+ "trial": 1,
424
+ "correct": 31,
425
+ "total": 50,
426
+ "accuracy": 0.62
427
+ },
428
+ {
429
+ "trial": 2,
430
+ "correct": 32,
431
+ "total": 50,
432
+ "accuracy": 0.64
433
+ },
434
+ {
435
+ "trial": 3,
436
+ "correct": 34,
437
+ "total": 50,
438
+ "accuracy": 0.68
439
+ },
440
+ {
441
+ "trial": 4,
442
+ "correct": 32,
443
+ "total": 50,
444
+ "accuracy": 0.64
445
+ }
446
+ ],
447
+ "mean_accuracy": 0.632,
448
+ "best_accuracy": 0.68,
449
+ "seconds": 28.85530649498105
450
+ },
451
+ {
452
+ "name": "more-grpo-epoch-1",
453
+ "trials": [
454
+ {
455
+ "trial": 0,
456
+ "correct": 23,
457
+ "total": 50,
458
+ "accuracy": 0.46
459
+ },
460
+ {
461
+ "trial": 1,
462
+ "correct": 29,
463
+ "total": 50,
464
+ "accuracy": 0.58
465
+ },
466
+ {
467
+ "trial": 2,
468
+ "correct": 28,
469
+ "total": 50,
470
+ "accuracy": 0.56
471
+ },
472
+ {
473
+ "trial": 3,
474
+ "correct": 34,
475
+ "total": 50,
476
+ "accuracy": 0.68
477
+ },
478
+ {
479
+ "trial": 4,
480
+ "correct": 28,
481
+ "total": 50,
482
+ "accuracy": 0.56
483
+ }
484
+ ],
485
+ "mean_accuracy": 0.5680000000000001,
486
+ "best_accuracy": 0.68,
487
+ "seconds": 29.11749249882996
488
+ },
489
+ {
490
+ "name": "more-grpo-epoch-2",
491
+ "trials": [
492
+ {
493
+ "trial": 0,
494
+ "correct": 29,
495
+ "total": 50,
496
+ "accuracy": 0.58
497
+ },
498
+ {
499
+ "trial": 1,
500
+ "correct": 32,
501
+ "total": 50,
502
+ "accuracy": 0.64
503
+ },
504
+ {
505
+ "trial": 2,
506
+ "correct": 29,
507
+ "total": 50,
508
+ "accuracy": 0.58
509
+ },
510
+ {
511
+ "trial": 3,
512
+ "correct": 35,
513
+ "total": 50,
514
+ "accuracy": 0.7
515
+ },
516
+ {
517
+ "trial": 4,
518
+ "correct": 31,
519
+ "total": 50,
520
+ "accuracy": 0.62
521
+ }
522
+ ],
523
+ "mean_accuracy": 0.624,
524
+ "best_accuracy": 0.7,
525
+ "seconds": 25.52914315275848
526
+ },
527
+ {
528
+ "name": "more-grpo-epoch-3",
529
+ "trials": [
530
+ {
531
+ "trial": 0,
532
+ "correct": 28,
533
+ "total": 50,
534
+ "accuracy": 0.56
535
+ },
536
+ {
537
+ "trial": 1,
538
+ "correct": 28,
539
+ "total": 50,
540
+ "accuracy": 0.56
541
+ },
542
+ {
543
+ "trial": 2,
544
+ "correct": 29,
545
+ "total": 50,
546
+ "accuracy": 0.58
547
+ },
548
+ {
549
+ "trial": 3,
550
+ "correct": 35,
551
+ "total": 50,
552
+ "accuracy": 0.7
553
+ },
554
+ {
555
+ "trial": 4,
556
+ "correct": 29,
557
+ "total": 50,
558
+ "accuracy": 0.58
559
+ }
560
+ ],
561
+ "mean_accuracy": 0.5960000000000001,
562
+ "best_accuracy": 0.7,
563
+ "seconds": 28.040696807205677
564
+ },
565
+ {
566
+ "name": "more-grpo-epoch-4",
567
+ "trials": [
568
+ {
569
+ "trial": 0,
570
+ "correct": 29,
571
+ "total": 50,
572
+ "accuracy": 0.58
573
+ },
574
+ {
575
+ "trial": 1,
576
+ "correct": 33,
577
+ "total": 50,
578
+ "accuracy": 0.66
579
+ },
580
+ {
581
+ "trial": 2,
582
+ "correct": 33,
583
+ "total": 50,
584
+ "accuracy": 0.66
585
+ },
586
+ {
587
+ "trial": 3,
588
+ "correct": 33,
589
+ "total": 50,
590
+ "accuracy": 0.66
591
+ },
592
+ {
593
+ "trial": 4,
594
+ "correct": 31,
595
+ "total": 50,
596
+ "accuracy": 0.62
597
+ }
598
+ ],
599
+ "mean_accuracy": 0.636,
600
+ "best_accuracy": 0.66,
601
+ "seconds": 27.200643570162356
602
+ },
603
+ {
604
+ "name": "more-grpo-epoch-5",
605
+ "trials": [
606
+ {
607
+ "trial": 0,
608
+ "correct": 29,
609
+ "total": 50,
610
+ "accuracy": 0.58
611
+ },
612
+ {
613
+ "trial": 1,
614
+ "correct": 33,
615
+ "total": 50,
616
+ "accuracy": 0.66
617
+ },
618
+ {
619
+ "trial": 2,
620
+ "correct": 29,
621
+ "total": 50,
622
+ "accuracy": 0.58
623
+ },
624
+ {
625
+ "trial": 3,
626
+ "correct": 36,
627
+ "total": 50,
628
+ "accuracy": 0.72
629
+ },
630
+ {
631
+ "trial": 4,
632
+ "correct": 28,
633
+ "total": 50,
634
+ "accuracy": 0.56
635
+ }
636
+ ],
637
+ "mean_accuracy": 0.62,
638
+ "best_accuracy": 0.72,
639
+ "seconds": 37.71861216891557
640
+ },
641
+ {
642
+ "name": "more-opd-epoch-1",
643
+ "trials": [
644
+ {
645
+ "trial": 0,
646
+ "correct": 26,
647
+ "total": 50,
648
+ "accuracy": 0.52
649
+ },
650
+ {
651
+ "trial": 1,
652
+ "correct": 33,
653
+ "total": 50,
654
+ "accuracy": 0.66
655
+ },
656
+ {
657
+ "trial": 2,
658
+ "correct": 30,
659
+ "total": 50,
660
+ "accuracy": 0.6
661
+ },
662
+ {
663
+ "trial": 3,
664
+ "correct": 34,
665
+ "total": 50,
666
+ "accuracy": 0.68
667
+ },
668
+ {
669
+ "trial": 4,
670
+ "correct": 26,
671
+ "total": 50,
672
+ "accuracy": 0.52
673
+ }
674
+ ],
675
+ "mean_accuracy": 0.5960000000000001,
676
+ "best_accuracy": 0.68,
677
+ "seconds": 34.81251425947994
678
+ },
679
+ {
680
+ "name": "more-opd-epoch-2",
681
+ "trials": [
682
+ {
683
+ "trial": 0,
684
+ "correct": 26,
685
+ "total": 50,
686
+ "accuracy": 0.52
687
+ },
688
+ {
689
+ "trial": 1,
690
+ "correct": 32,
691
+ "total": 50,
692
+ "accuracy": 0.64
693
+ },
694
+ {
695
+ "trial": 2,
696
+ "correct": 32,
697
+ "total": 50,
698
+ "accuracy": 0.64
699
+ },
700
+ {
701
+ "trial": 3,
702
+ "correct": 33,
703
+ "total": 50,
704
+ "accuracy": 0.66
705
+ },
706
+ {
707
+ "trial": 4,
708
+ "correct": 28,
709
+ "total": 50,
710
+ "accuracy": 0.56
711
+ }
712
+ ],
713
+ "mean_accuracy": 0.6040000000000001,
714
+ "best_accuracy": 0.66,
715
+ "seconds": 37.19308333192021
716
+ },
717
+ {
718
+ "name": "more-opd-epoch-3",
719
+ "trials": [
720
+ {
721
+ "trial": 0,
722
+ "correct": 27,
723
+ "total": 50,
724
+ "accuracy": 0.54
725
+ },
726
+ {
727
+ "trial": 1,
728
+ "correct": 32,
729
+ "total": 50,
730
+ "accuracy": 0.64
731
+ },
732
+ {
733
+ "trial": 2,
734
+ "correct": 30,
735
+ "total": 50,
736
+ "accuracy": 0.6
737
+ },
738
+ {
739
+ "trial": 3,
740
+ "correct": 35,
741
+ "total": 50,
742
+ "accuracy": 0.7
743
+ },
744
+ {
745
+ "trial": 4,
746
+ "correct": 31,
747
+ "total": 50,
748
+ "accuracy": 0.62
749
+ }
750
+ ],
751
+ "mean_accuracy": 0.6200000000000001,
752
+ "best_accuracy": 0.7,
753
+ "seconds": 25.204970656894147
754
+ },
755
+ {
756
+ "name": "more-opd-epoch-4",
757
+ "trials": [
758
+ {
759
+ "trial": 0,
760
+ "correct": 31,
761
+ "total": 50,
762
+ "accuracy": 0.62
763
+ },
764
+ {
765
+ "trial": 1,
766
+ "correct": 31,
767
+ "total": 50,
768
+ "accuracy": 0.62
769
+ },
770
+ {
771
+ "trial": 2,
772
+ "correct": 31,
773
+ "total": 50,
774
+ "accuracy": 0.62
775
+ },
776
+ {
777
+ "trial": 3,
778
+ "correct": 36,
779
+ "total": 50,
780
+ "accuracy": 0.72
781
+ },
782
+ {
783
+ "trial": 4,
784
+ "correct": 33,
785
+ "total": 50,
786
+ "accuracy": 0.66
787
+ }
788
+ ],
789
+ "mean_accuracy": 0.648,
790
+ "best_accuracy": 0.72,
791
+ "seconds": 27.25743418559432
792
+ },
793
+ {
794
+ "name": "more-opd-epoch-5",
795
+ "trials": [
796
+ {
797
+ "trial": 0,
798
+ "correct": 31,
799
+ "total": 50,
800
+ "accuracy": 0.62
801
+ },
802
+ {
803
+ "trial": 1,
804
+ "correct": 30,
805
+ "total": 50,
806
+ "accuracy": 0.6
807
+ },
808
+ {
809
+ "trial": 2,
810
+ "correct": 32,
811
+ "total": 50,
812
+ "accuracy": 0.64
813
+ },
814
+ {
815
+ "trial": 3,
816
+ "correct": 36,
817
+ "total": 50,
818
+ "accuracy": 0.72
819
+ },
820
+ {
821
+ "trial": 4,
822
+ "correct": 31,
823
+ "total": 50,
824
+ "accuracy": 0.62
825
+ }
826
+ ],
827
+ "mean_accuracy": 0.64,
828
+ "best_accuracy": 0.72,
829
+ "seconds": 28.25529603473842
830
+ }
831
+ ],
832
+ "selected_test": {
833
+ "name": "selected",
834
+ "trials": [
835
+ {
836
+ "trial": 0,
837
+ "correct": 32,
838
+ "total": 50,
839
+ "accuracy": 0.64
840
+ },
841
+ {
842
+ "trial": 1,
843
+ "correct": 28,
844
+ "total": 50,
845
+ "accuracy": 0.56
846
+ },
847
+ {
848
+ "trial": 2,
849
+ "correct": 30,
850
+ "total": 50,
851
+ "accuracy": 0.6
852
+ },
853
+ {
854
+ "trial": 3,
855
+ "correct": 30,
856
+ "total": 50,
857
+ "accuracy": 0.6
858
+ },
859
+ {
860
+ "trial": 4,
861
+ "correct": 33,
862
+ "total": 50,
863
+ "accuracy": 0.66
864
+ }
865
+ ],
866
+ "mean_accuracy": 0.6120000000000001,
867
+ "best_accuracy": 0.66,
868
+ "seconds": 31.559877825900912
869
+ },
870
+ "train_seconds": {
871
+ "current-ratio": 1146.5347720878199,
872
+ "lower-lr": 1171.0772221423686,
873
+ "more-grpo": 1160.5429365690798,
874
+ "more-opd": 1165.7883752603084
875
+ },
876
+ "scope": "Only the final mixed-loss phase was retrained. Compressor and initialization held fixed. The candidate is not a replacement for historical paper results."
877
+ }
personamem-32k/qwen2.5-3b/reader/adapter_config.json CHANGED
@@ -25,12 +25,12 @@
25
  "rank_pattern": {},
26
  "revision": "aa8e72537993ba99e69dfaafa59ed015b17504d1",
27
  "target_modules": [
28
- "q_proj",
29
- "v_proj",
30
  "k_proj",
31
- "up_proj",
32
- "gate_proj",
33
  "o_proj",
 
 
 
 
34
  "down_proj"
35
  ],
36
  "target_parameters": null,
 
25
  "rank_pattern": {},
26
  "revision": "aa8e72537993ba99e69dfaafa59ed015b17504d1",
27
  "target_modules": [
 
 
28
  "k_proj",
 
 
29
  "o_proj",
30
+ "v_proj",
31
+ "gate_proj",
32
+ "up_proj",
33
+ "q_proj",
34
  "down_proj"
35
  ],
36
  "target_parameters": null,
personamem-32k/qwen2.5-3b/reader/adapter_model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:ecf6da9d83cb82890aeffd4079d86880b29c3a62d7e8cf919329d47ae7197f71
3
  size 119801528
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:4fdf8e12b6211301883425c5d792843698d9d9b884ce391a9380d76e112d7a4d
3
  size 119801528