Caesarrr commited on
Commit
2c66bf8
·
verified ·
1 Parent(s): a020cd4

Publish atlantis qwenoft H8 evaluation artifacts

Browse files
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/README.md CHANGED
@@ -1,8 +1,8 @@
1
- # Atlantis qwenoft zero-latency H8 VLA
2
 
3
  Current RGB and eight binary decision-history windows; action horizon one. Initialized from the pinned Qwen3-VL-4B backbone and random action head.
4
 
5
  - Training: 5000 steps, seed 42, global batch 128
6
  - W&B: https://wandb.ai/saberrr-zju/starvla_tasks/runs/40be052b
7
- - Evaluation: 100 episodes, seeds 1000000–1000099, GPU simulator, parallel 32, capacity 1, 60/15 Hz, 7200 raw-frame cap, zero latency
8
- - Mean raw return: 39560.0
 
1
+ # Atlantis qwenoft zero-trained H8 VLA
2
 
3
  Current RGB and eight binary decision-history windows; action horizon one. Initialized from the pinned Qwen3-VL-4B backbone and random action head.
4
 
5
  - Training: 5000 steps, seed 42, global batch 128
6
  - W&B: https://wandb.ai/saberrr-zju/starvla_tasks/runs/40be052b
7
+ - Evaluation: 100 episodes, seeds 1000000–1000099, GPU simulator, parallel 32, capacity 1, 60/15 Hz, 7200 raw-frame cap, profile latency
8
+ - Mean raw return: 18588.0
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/eval_profile.json ADDED
@@ -0,0 +1,217 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "checkpoint_path": "/home/ubuntu/codex/atari-hist8-20260924/outputs/tasks/atari_hist8_v1/bundles/atlantis-qwenoft-hist8-zero-s42-v1/checkpoints/model.pt",
3
+ "experiment_name": "atlantis-qwenoft-hist8-zero-s42-v1-step5000-history01-profile-par32",
4
+ "latency": "profile_sample",
5
+ "latency_type": "profile_sample",
6
+ "lengths": [
7
+ 6502,
8
+ 4860,
9
+ 6607,
10
+ 5080,
11
+ 6437,
12
+ 3553,
13
+ 3278,
14
+ 5311,
15
+ 3748,
16
+ 5104,
17
+ 5768,
18
+ 3558,
19
+ 4806,
20
+ 6103,
21
+ 5411,
22
+ 5644,
23
+ 2602,
24
+ 4553,
25
+ 5281,
26
+ 4300,
27
+ 3326,
28
+ 2793,
29
+ 2433,
30
+ 5519,
31
+ 5334,
32
+ 3928,
33
+ 6605,
34
+ 3182,
35
+ 4062,
36
+ 7200,
37
+ 5238,
38
+ 4377,
39
+ 4530,
40
+ 4282,
41
+ 3908,
42
+ 6863,
43
+ 5245,
44
+ 6844,
45
+ 5974,
46
+ 4005,
47
+ 4792,
48
+ 4254,
49
+ 5061,
50
+ 7000,
51
+ 4033,
52
+ 5157,
53
+ 5771,
54
+ 3723,
55
+ 4875,
56
+ 3657,
57
+ 5048,
58
+ 3179,
59
+ 4804,
60
+ 3545,
61
+ 3753,
62
+ 5440,
63
+ 4275,
64
+ 4658,
65
+ 6598,
66
+ 4265,
67
+ 5644,
68
+ 3967,
69
+ 7200,
70
+ 3009,
71
+ 4001,
72
+ 7200,
73
+ 4702,
74
+ 5983,
75
+ 5245,
76
+ 5501,
77
+ 7200,
78
+ 4108,
79
+ 5188,
80
+ 7200,
81
+ 4230,
82
+ 7200,
83
+ 4982,
84
+ 4156,
85
+ 7132,
86
+ 5178,
87
+ 7200,
88
+ 5082,
89
+ 4799,
90
+ 4514,
91
+ 4587,
92
+ 6666,
93
+ 6847,
94
+ 5767,
95
+ 6339,
96
+ 5618,
97
+ 5613,
98
+ 3899,
99
+ 4318,
100
+ 4763,
101
+ 4798,
102
+ 4513,
103
+ 4716,
104
+ 4314,
105
+ 5162,
106
+ 5178
107
+ ],
108
+ "mean_length": 5017.61,
109
+ "mean_return": 18588.0,
110
+ "returns": [
111
+ 23600.0,
112
+ 18600.0,
113
+ 21700.0,
114
+ 21800.0,
115
+ 20500.0,
116
+ 12200.0,
117
+ 9500.0,
118
+ 16800.0,
119
+ 13800.0,
120
+ 19900.0,
121
+ 24500.0,
122
+ 13700.0,
123
+ 13700.0,
124
+ 17300.0,
125
+ 19000.0,
126
+ 27900.0,
127
+ 6900.0,
128
+ 21600.0,
129
+ 16000.0,
130
+ 17700.0,
131
+ 10500.0,
132
+ 9800.0,
133
+ 6300.0,
134
+ 17700.0,
135
+ 12400.0,
136
+ 8800.0,
137
+ 22900.0,
138
+ 9700.0,
139
+ 10300.0,
140
+ 26700.0,
141
+ 14400.0,
142
+ 22300.0,
143
+ 16300.0,
144
+ 13700.0,
145
+ 16800.0,
146
+ 29200.0,
147
+ 14400.0,
148
+ 28600.0,
149
+ 20100.0,
150
+ 14800.0,
151
+ 16500.0,
152
+ 14200.0,
153
+ 20700.0,
154
+ 27200.0,
155
+ 11400.0,
156
+ 19700.0,
157
+ 17400.0,
158
+ 7900.0,
159
+ 18900.0,
160
+ 17200.0,
161
+ 20300.0,
162
+ 9600.0,
163
+ 20400.0,
164
+ 13200.0,
165
+ 13300.0,
166
+ 21400.0,
167
+ 15400.0,
168
+ 16000.0,
169
+ 26000.0,
170
+ 12900.0,
171
+ 26400.0,
172
+ 9500.0,
173
+ 25500.0,
174
+ 9100.0,
175
+ 20000.0,
176
+ 25900.0,
177
+ 19200.0,
178
+ 23400.0,
179
+ 12600.0,
180
+ 15600.0,
181
+ 36300.0,
182
+ 13100.0,
183
+ 14700.0,
184
+ 38100.0,
185
+ 17300.0,
186
+ 29000.0,
187
+ 25900.0,
188
+ 13100.0,
189
+ 34900.0,
190
+ 19900.0,
191
+ 32600.0,
192
+ 14900.0,
193
+ 20900.0,
194
+ 14900.0,
195
+ 15100.0,
196
+ 25700.0,
197
+ 30300.0,
198
+ 20600.0,
199
+ 23500.0,
200
+ 27300.0,
201
+ 22800.0,
202
+ 12600.0,
203
+ 17400.0,
204
+ 15500.0,
205
+ 19500.0,
206
+ 17100.0,
207
+ 14800.0,
208
+ 20500.0,
209
+ 17100.0,
210
+ 24200.0
211
+ ],
212
+ "seed": 1000000,
213
+ "source_profile_path": "/home/ubuntu/codex/atari-hist8-20260924/outputs/tasks/atari_hist8_v1/teacher_configs/atlantis-qwenoft/profile/profile.json",
214
+ "std_return": 6544.543987169771,
215
+ "suite_name": "profile_sample",
216
+ "timestamp_utc": "2026-09-24T12:41:59.686946+00:00"
217
+ }
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/evaluation_profile/action_audit.json ADDED
@@ -0,0 +1,133 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "eval_protocol_version": "hist8-history01-sf-render-v4",
3
+ "history_encoding": "binary_0_1",
4
+ "episodes": 100,
5
+ "episode_seeds": [
6
+ 1000000,
7
+ 1000001,
8
+ 1000002,
9
+ 1000003,
10
+ 1000004,
11
+ 1000005,
12
+ 1000006,
13
+ 1000007,
14
+ 1000008,
15
+ 1000009,
16
+ 1000010,
17
+ 1000011,
18
+ 1000012,
19
+ 1000013,
20
+ 1000014,
21
+ 1000015,
22
+ 1000016,
23
+ 1000017,
24
+ 1000018,
25
+ 1000019,
26
+ 1000020,
27
+ 1000021,
28
+ 1000022,
29
+ 1000023,
30
+ 1000024,
31
+ 1000025,
32
+ 1000026,
33
+ 1000027,
34
+ 1000028,
35
+ 1000029,
36
+ 1000030,
37
+ 1000031,
38
+ 1000032,
39
+ 1000033,
40
+ 1000034,
41
+ 1000035,
42
+ 1000036,
43
+ 1000037,
44
+ 1000038,
45
+ 1000039,
46
+ 1000040,
47
+ 1000041,
48
+ 1000042,
49
+ 1000043,
50
+ 1000044,
51
+ 1000045,
52
+ 1000046,
53
+ 1000047,
54
+ 1000048,
55
+ 1000049,
56
+ 1000050,
57
+ 1000051,
58
+ 1000052,
59
+ 1000053,
60
+ 1000054,
61
+ 1000055,
62
+ 1000056,
63
+ 1000057,
64
+ 1000058,
65
+ 1000059,
66
+ 1000060,
67
+ 1000061,
68
+ 1000062,
69
+ 1000063,
70
+ 1000064,
71
+ 1000065,
72
+ 1000066,
73
+ 1000067,
74
+ 1000068,
75
+ 1000069,
76
+ 1000070,
77
+ 1000071,
78
+ 1000072,
79
+ 1000073,
80
+ 1000074,
81
+ 1000075,
82
+ 1000076,
83
+ 1000077,
84
+ 1000078,
85
+ 1000079,
86
+ 1000080,
87
+ 1000081,
88
+ 1000082,
89
+ 1000083,
90
+ 1000084,
91
+ 1000085,
92
+ 1000086,
93
+ 1000087,
94
+ 1000088,
95
+ 1000089,
96
+ 1000090,
97
+ 1000091,
98
+ 1000092,
99
+ 1000093,
100
+ 1000094,
101
+ 1000095,
102
+ 1000096,
103
+ 1000097,
104
+ 1000098,
105
+ 1000099
106
+ ],
107
+ "issued_actions": 62677,
108
+ "action_counts": {
109
+ "0": 15821,
110
+ "2": 29659,
111
+ "3": 12666,
112
+ "1": 4531
113
+ },
114
+ "largest_action_fraction": 0.47320388659316814,
115
+ "raw_score_min": [
116
+ -16.4777889251709,
117
+ -17.36240005493164,
118
+ -17.736751556396484,
119
+ -16.52406120300293
120
+ ],
121
+ "raw_score_max": [
122
+ 20.059650421142578,
123
+ 20.963884353637695,
124
+ 20.347089767456055,
125
+ 20.745391845703125
126
+ ],
127
+ "raw_score_mean": [
128
+ 1.5759939295484995,
129
+ -6.00103662508784,
130
+ 6.868458961622308,
131
+ -0.553409212165336
132
+ ]
133
+ }
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/evaluation_profile/episode_metrics.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/evaluation_profile/eval_config.yaml ADDED
@@ -0,0 +1,149 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ experiment:
2
+ name: atlantis-qwenoft-hist8-zero-s42-v1-step5000-history01-profile-par32
3
+ seed: 1000000
4
+ backend:
5
+ type: sample_factory
6
+ algo: APPO
7
+ device: cuda
8
+ train_dir: results/sample_factory
9
+ restart_behavior: resume
10
+ run_mode: eval
11
+ extra_args:
12
+ - --encoder_conv_architecture
13
+ - convnet_simple
14
+ executor:
15
+ mode: simulated
16
+ simulated_worker_capacity: 1
17
+ simulated_inference_pool: true
18
+ inference_devices:
19
+ - cuda:0
20
+ - cuda:1
21
+ inference_batch_size: 16
22
+ env:
23
+ name: gymnasium
24
+ task_name: atlantis
25
+ env_id: LatencyBench/Atlantis-v0
26
+ registration_imports:
27
+ - latency_bench.envs.gymnasium_atlantis
28
+ make_kwargs:
29
+ base_env_id: ALE/Atlantis-v5
30
+ render_mode: rgb_array
31
+ screen_size: 84
32
+ noop_max: 30
33
+ base_make_kwargs:
34
+ obs_type: rgb
35
+ frameskip: 1
36
+ repeat_action_probability: 0.0
37
+ full_action_space: false
38
+ mode: 0
39
+ difficulty: 0
40
+ max_num_frames_per_episode: 108000
41
+ env_fps: 60
42
+ obs_fps: 15
43
+ frame_stack: 1
44
+ simulator: gpu
45
+ terminal_on_life_loss: false
46
+ clip_reward: false
47
+ noop_action: noop
48
+ action_map:
49
+ noop: 0
50
+ fire: 1
51
+ rightfire: 2
52
+ leftfire: 3
53
+ action_order:
54
+ - noop
55
+ - fire
56
+ - rightfire
57
+ - leftfire
58
+ base_prompt: 'Defend Atlantis by firing at descending enemies. Choose exactly one
59
+ action from: noop, fire, rightfire, leftfire.'
60
+ obs_resize:
61
+ - 224
62
+ - 224
63
+ gpu_env_device: auto
64
+ noop_max: 30
65
+ action_history_decisions: 8
66
+ latency:
67
+ method: temporal
68
+ profile_path: /home/ubuntu/codex/atari-hist8-20260924/outputs/tasks/atari_hist8_v1/teacher_configs/atlantis-qwenoft/profile/profile.json
69
+ profile_worker_slot: 0
70
+ seed: 271828
71
+ add_latency_info: false
72
+ scheduler:
73
+ hold_policy: hold
74
+ ordering_policy: issue_order_fifo
75
+ policy:
76
+ type: starvla
77
+ actions:
78
+ - noop
79
+ - fire
80
+ - rightfire
81
+ - leftfire
82
+ checkpoint_path: /home/ubuntu/codex/atari-hist8-20260924/outputs/tasks/atari_hist8_v1/bundles/atlantis-qwenoft-hist8-zero-s42-v1/checkpoints/model.pt
83
+ model_config_path: /home/ubuntu/codex/atari-hist8-20260924/outputs/tasks/atari_hist8_v1/bundles/atlantis-qwenoft-hist8-zero-s42-v1/config.yaml
84
+ task_manifest_path: /home/ubuntu/codex/atari-hist8-20260924/outputs/tasks/atari_hist8_v1/bundles/atlantis-qwenoft-hist8-zero-s42-v1/manifest.json
85
+ device: cuda:0
86
+ unnorm_key: new_embodiment
87
+ prompt_mode: latency_neutral
88
+ backbone_path: /home/ubuntu/.cache/huggingface/hub/models--Qwen--Qwen3-VL-4B-Instruct/snapshots/ebb281ec70b05090aa6165b016eac8ec08e71b17
89
+ worker_python_executable: /home/ubuntu/miniconda3/envs/starvla_rl_games_openvla/bin/python3.10
90
+ state_source: transport
91
+ training:
92
+ train_for_env_steps: 1000000
93
+ num_workers: 4
94
+ num_envs_per_worker: 16
95
+ worker_num_splits: 1
96
+ num_policies: 1
97
+ batch_size: 1024
98
+ rollout: 64
99
+ recurrence: 32
100
+ num_epochs: 2
101
+ num_batches_per_epoch: 8
102
+ max_policy_lag: 300
103
+ learning_rate: 0.000303
104
+ lr_schedule: constant
105
+ nonlinearity: elu
106
+ gamma: 0.99
107
+ gae_lambda: 0.95
108
+ ppo_clip_ratio: 0.1
109
+ ppo_clip_value: 0.2
110
+ exploration_loss: entropy
111
+ exploration_loss_coeff: 0.003
112
+ value_loss_coeff: 0.5
113
+ max_grad_norm: 0.0
114
+ adam_eps: 1.0e-06
115
+ obs_scale: 1.0
116
+ adaptive_stddev: false
117
+ with_vtrace: false
118
+ async_rl: true
119
+ use_rnn: false
120
+ normalize_input: true
121
+ normalize_returns: true
122
+ stats_avg: 100
123
+ experiment_summaries_interval: 1
124
+ save_every_sec: 600
125
+ keep_checkpoints: 5
126
+ evaluation:
127
+ eval_interval_steps: null
128
+ eval_episodes: 100
129
+ eval_parallel_envs: 32
130
+ eval_max_steps: 7200
131
+ eval_deterministic: true
132
+ eval_suites:
133
+ fixed: []
134
+ normal: []
135
+ uniform: []
136
+ latency_bench_env_backend: atlantis_gpu_batched
137
+ eval_raw_reward: true
138
+ logging:
139
+ output_dir: /home/ubuntu/codex/atari-hist8-20260924/outputs/tasks/atari_hist8_v1/evals/atlantis-qwenoft-hist8-zero-s42-v1/zero_to_profile
140
+ video:
141
+ enabled: false
142
+ num_bins: 1
143
+ save_step_records: true
144
+ save_action_records: true
145
+ save_latency_records: true
146
+ wandb_project: null
147
+ wandb_group: null
148
+ wandb_job_type: null
149
+ wandb_tags: ''
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/evaluation_profile/provenance.json ADDED
@@ -0,0 +1,46 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "collection": {
3
+ "episodes": 200,
4
+ "filter_preset": "none",
5
+ "history": "decision_v1",
6
+ "history_decisions": 8,
7
+ "history_source": "teacher_transport",
8
+ "max_episode_raw_frames": 7200,
9
+ "rollout_overrides": {
10
+ "env_fps": 60,
11
+ "fixed_latency_ms": 0.0,
12
+ "latency_type": "fixed",
13
+ "obs_fps": 15,
14
+ "seed": 42
15
+ },
16
+ "student_image_frames": 1,
17
+ "teacher_checkpoint": "best_000026560_13598720_reward_675769.000.pth",
18
+ "teacher_path": "zero-latency/atlantis/small-policy/sample-factory-hist8-cpu15m-s0-v1",
19
+ "teacher_repo": "latency-sensitive-bench/benchmark-models",
20
+ "teacher_revision": "dfaeecb84cd6b5dfcc3adc0935d0ae86d0b504d2",
21
+ "teacher_sha256": "d41d728d8c2329d3623eb9e01bd9f3f2df4fc0d4f6f2c1f799c4badc24af78c5",
22
+ "train_labels": 324000,
23
+ "training_latency_condition": "zero"
24
+ },
25
+ "training_run_id": "atlantis-qwenoft-hist8-zero-s42-v1",
26
+ "training_step": 5000,
27
+ "global_batch": 128,
28
+ "base_model": "Qwen/Qwen3-VL-4B-Instruct",
29
+ "wandb_run_id": "40be052b",
30
+ "checkpoint_sha256": "f91d4079a94f779856220e4a954750624228a3da596b5ecdd70861e888a97877",
31
+ "eval_protocol_version": "hist8-history01-sf-render-v4",
32
+ "history_encoding": "binary_0_1",
33
+ "episodes": 100,
34
+ "env_condition": "profile",
35
+ "parallel_envs": 32,
36
+ "max_episode_raw_frames": 7200,
37
+ "inference_devices": [
38
+ "cuda:0",
39
+ "cuda:1"
40
+ ],
41
+ "episode_seeds": [
42
+ 1000000,
43
+ 1000099
44
+ ],
45
+ "profile_sha256": "2c8c1f14845e89c6171e1502ff5af6d61d12049e6445d98d63f31e4ffca384f4"
46
+ }
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/profile_publication.json ADDED
@@ -0,0 +1,28 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "model": "qwenoft",
3
+ "model_publication": {
4
+ "checkpoint_sha256": "f91d4079a94f779856220e4a954750624228a3da596b5ecdd70861e888a97877",
5
+ "episodes": 100,
6
+ "mean_return": 39560.0,
7
+ "model": "qwenoft",
8
+ "path": "zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1",
9
+ "repo_id": "latency-sensitive-bench/benchmark-models",
10
+ "revision": "15c9cd4ab12215ccbcef74a8cd4d051d8e3aae3f",
11
+ "run_id": "atlantis-qwenoft-hist8-zero-s42-v1",
12
+ "task": "atlantis",
13
+ "training_status": "complete",
14
+ "zero_evaluation_status": "complete"
15
+ },
16
+ "profile": {
17
+ "GPU_CLASS": "1x-rtx3090",
18
+ "INSTANCE_ID": "instance_8dd832253b2c3d69",
19
+ "PROFILE_PATH": "results/profiling/profiles/qwenoft/1x-rtx3090/atlantis/instance_8dd832253b2c3d69/profile.json",
20
+ "PROFILE_REF": "qwenoft/1x-rtx3090/atlantis/instance_8dd832253b2c3d69",
21
+ "RUN_ID": "20260924T110210282159Z"
22
+ },
23
+ "profile_path": "profiles/qwenoft/1x-rtx3090/atlantis/instance_8dd832253b2c3d69/profile.json",
24
+ "profile_repo": "latency-sensitive-bench/profiles",
25
+ "profile_revision": "7efc9af422864b9a1cb569aa81f1b2b1e216b89e",
26
+ "profile_sha256": "2c8c1f14845e89c6171e1502ff5af6d61d12049e6445d98d63f31e4ffca384f4",
27
+ "task": "atlantis"
28
+ }
zero-latency/atlantis/vla/starvla-qwenoft-h1/atlantis-qwenoft-hist8-zero-s42-v1/publication_status.json CHANGED
@@ -1,6 +1,8 @@
1
  {
2
  "episodes": 100,
3
- "mean_return": 39560.0,
 
 
4
  "training_status": "complete",
5
  "zero_evaluation_status": "complete"
6
  }
 
1
  {
2
  "episodes": 100,
3
+ "mean_return": 18588.0,
4
+ "profile_evaluation_status": "complete",
5
+ "profile_mean_return": 18588.0,
6
  "training_status": "complete",
7
  "zero_evaluation_status": "complete"
8
  }