File size: 2,488 Bytes
954544e
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
{
  "format_version": 1,
  "name": "Dot",
  "release": "v0.4-thinking",
  "released_by": "MTEnt",
  "license": "Apache-2.0",
  "artifact": "complete BF16 text backbone plus recurrent-depth core",
  "architecture": {
    "class": "DotRecurrentDepthModel",
    "backbone_parameters": 8953803264,
    "recurrent_core_parameters": 864945224,
    "total_instantiated_parameters": 9818748488,
    "insertion_after_layer": 15,
    "copied_source_layers": [12, 13, 14, 15],
    "maximum_loops": 8,
    "active_loops": 4,
    "cache_supported": false,
    "precision": "bfloat16"
  },
  "lineage": {
    "upstream_repository": "Qwen/Qwen3.5-9B",
    "upstream_exact_revision": null,
    "upstream_revision_note": "The original training manifest did not record the exact source commit.",
    "semantic_stage": "Dot-9B-Semantic-v0.2 merged LoRA rank 32",
    "semantic_train_records": 11763,
    "recurrent_stage_source": "Dot-9B-Recurrent-v0.3",
    "thinking_repair_source_step": 938
  },
  "training": {
    "thinking_records": 60000,
    "thinking_record_sha256": "ff84822f7cf85be9c1a1e392e9c358b01b56d68f9aff821bcbc41c99534853e5",
    "validation_records": 4096,
    "validation_record_sha256": "e46a96db4c70398f7aacf02b7b900d21302f9560926b5cd0b133c2089891ef6c",
    "optimizer_steps": 938,
    "tokens_seen": 8642015,
    "backbone_frozen": true,
    "final_loss": 0.01007067202590406,
    "loop_scales": [
      0.15769609808921814,
      0.07985208928585052,
      0.05903945118188858,
      0.059150855988264084,
      0.0,
      0.0,
      0.0,
      0.0
    ]
  },
  "weights": {
    "model-00001-of-00004.safetensors": "330fa6fed85e394ba9dec36f986e017e51b8f3c588613bc0dc7f0ea8521aac14",
    "model-00002-of-00004.safetensors": "fa04e853ff08aef3b8ef51e2f2e602b8b3da7d1d9072925602812c052502e011",
    "model-00003-of-00004.safetensors": "6aa1824b165e566f6bcb6038c6dab6abd1c4b392d5e8cc0173d1f859c009f477",
    "model-00004-of-00004.safetensors": "dbb205aa4fab083cb7f409b337a34616a559d76fe7c834f48396da38d6de1a3a",
    "reasoning_core.safetensors": "7c897b83176c044a17840c90f4123ce7bcce437139aad1e84f76aa1555db6152"
  },
  "known_limits": [
    "custom loader required",
    "cache-backed decoding unsupported",
    "text only",
    "H200 BF16 runtime is the only verified hardware path",
    "targeted synthetic reasoning evaluation is not a broad capability benchmark",
    "new spatial generalization probe scored 21.875 percent exact match",
    "no independent safety evaluation"
  ]
}