hackhackhack66666 commited on
Commit
502cb21
·
verified ·
1 Parent(s): 6b58a9d

Add train config seed0

Browse files
configs/bar_calvin_production_s0_h100_1gpu.yaml ADDED
@@ -0,0 +1,88 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ # Stage-3 production BAR on aicenter4: SmolVLM2-2.2B, single H100 GPU1.
2
+ # Exact twin of production_s1: same data/hparams/split; seed 0 + codec s0.
3
+ # Global batch 128 = microbatch 1 × 1 GPU × accumulate 128.
4
+ strategy: bar
5
+ seed: 0
6
+ DATA_SPLIT_SEED: 0
7
+
8
+ DATASET:
9
+ root: /datasets/askhabaliev_g/calvin_full/task_D_D
10
+ chunk_length: 30
11
+
12
+ MODEL:
13
+ model_name: HuggingFaceTB/SmolVLM2-2.2B-Instruct
14
+ answer_start_token: "Assistant:"
15
+
16
+ action_processor:
17
+ vocab_size: 2048
18
+ token_len: 48
19
+ _target_: transformers.AutoModel
20
+ kwargs:
21
+ pretrained_model_name_or_path: /datasets/askhabaliev_g/checkpoints/codec/rvq_posttrain_s0/best/model
22
+ trust_remote_code: true
23
+
24
+ vl_processor:
25
+ _target_: transformers.AutoProcessor
26
+ kwargs:
27
+ pretrained_model_name_or_path: ${MODEL.model_name}
28
+ trust_remote_code: true
29
+
30
+ vla_processor:
31
+ kwargs:
32
+ mode: "discrete"
33
+ vocab_shift: 0
34
+
35
+ vlm:
36
+ _target_: smolvla.bar.SmolVLABlockwiseAR
37
+ kwargs:
38
+ pretrained_model_name_or_path: ${MODEL.model_name}
39
+ trust_remote_code: true
40
+ torch_dtype: float32
41
+ token_budget: ${MODEL.action_processor.token_len}
42
+ num_blocks: 3
43
+ action_vocab_size: ${MODEL.action_processor.vocab_size}
44
+ _attn_implementation: sdpa
45
+
46
+ TRAINING:
47
+ ckpt_dir: /datasets/askhabaliev_g/checkpoints/vla/production_s0
48
+ training_steps: 30000
49
+ warmup_steps: 1000
50
+ batch_size: 1
51
+ num_workers: 4
52
+ val_check_interval: 1000
53
+ checkpoint_interval: 1000
54
+ checkpoint_monitor: val/l1_dist
55
+ checkpoint_mode: min
56
+ seed: 0
57
+
58
+ use_lora: false
59
+ lora_kwargs:
60
+ lora_alpha: 32
61
+ r: 16
62
+ task_type: "CAUSAL_LM"
63
+ target_modules:
64
+ - "q_proj"
65
+ - "k_proj"
66
+ - "v_proj"
67
+ - "o_proj"
68
+ - "up_proj"
69
+ - "down_proj"
70
+ - "gate_proj"
71
+
72
+ optimizer:
73
+ _target_: torch.optim.AdamW
74
+ kwargs:
75
+ lr: 1e-4
76
+ weight_decay: 1e-10
77
+ eps: 1e-8
78
+ betas: [0.9, 0.999]
79
+ # HAMI fragmentation: foreach=True OOMs on opt/restore with "free" cache.
80
+ foreach: false
81
+
82
+ trainer:
83
+ devices: [0]
84
+ accumulate_grad_batches: 128
85
+ precision: 16-mixed
86
+ gradient_clip_val: 1.0
87
+ log_every_n_steps: 10
88
+ strategy: auto