File size: 2,198 Bytes
6de848e
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
{
  "schema_version": 1,
  "kind": "sa_wm_pretrained_model_release",
  "model_id": "sa_wm_pretrained_v1",
  "status": "selected_for_release_and_post_training",
  "selected_at": "2026-07-25",
  "selected_checkpoint": {
    "training_stage": "sa_wm_pretrain_v6_512x768",
    "micro_step": 55000,
    "path": "/mnt/cfs/wmcprg/RoboCoach/checkpoints/sa_wm_pretrain_v6_512x768/checkpoint-55000",
    "training_config": "configs/training/sa_wm_pretrain_v6.yaml",
    "weights": "dit_model.safetensors",
    "checkpoint_format": "safetensors+optim_shards",
    "post_training_load_mode": "model_only_init_from"
  },
  "selection_evidence": {
    "monitoring_selection": "/mnt/cfs/wmcprg/RoboCoach/checkpoints/sa_wm_pretrain_v6_512x768/validation/selection_closed_loop.json",
    "samples": 30,
    "rollout_chunks": 4,
    "future_frames": 1440,
    "metrics": {
      "psnr": 21.503524390835977,
      "ssim": 0.9163767135968223,
      "lpips_alex": 0.07815770965373506
    },
    "selection_reason": [
      "Human review preferred the motion and physical behavior of checkpoint-55000.",
      "DROID exterior_1 and exterior_2 PSNR, SSIM, and LPIPS all outperform checkpoint-70000.",
      "The fourth closed-loop chunk is slightly better than checkpoint-70000 in PSNR and LPIPS.",
      "The aggregate metric difference from checkpoint-70000 is too small to override the qualitative review."
    ]
  },
  "historical_comparison": {
    "checkpoint": "/mnt/cfs/wmcprg/RoboCoach/checkpoints/sa_wm_pretrain_v6_512x768/checkpoint-70000",
    "micro_step": 70000,
    "status": "retained_late_training_reference_not_selected",
    "metrics": {
      "psnr": 21.552308871443916,
      "ssim": 0.9175402801300698,
      "lpips_alex": 0.07684179975799957
    }
  },
  "usage_contract": {
    "inference": "load selected_checkpoint/dit_model.safetensors",
    "post_training": "initialize model weights from the selected checkpoint and start a new optimizer, scheduler, step counter, output directory, and run id",
    "exact_v6_resume": "allowed only for reproducing the historical v6 stage; do not use it for domain post-training",
    "immutability": "never overwrite the selected checkpoint directory"
  }
}