AMOR-GatedDeltaNet-440M / export_metadata.json
FlyinGodzilla's picture
Provide verified sharded weights and optional inference routers
de79722 verified
Raw History Blame Contribute Delete
4.38 kB
{
"repo_id": "FlyinGodzilla/AMOR-GatedDeltaNet-440M",
"source_filename": "amor_gdn_440m_step_0018255.pt",
"source_sha256": "ea5c2b88561dd7eb4527ba25c868a7ccea675baa1907e7a5156eb793cd95cb9a",
"step": 18255,
"tokens_seen": 8973017088,
"n_params": 442133056,
"tensor_dtypes": {
"torch.float32": 458
},
"gate_buffers": {
"0": {
"running_threshold": 0.8792613744735718,
"running_std": 0.043029818683862686
},
"1": {
"running_threshold": 0.854272186756134,
"running_std": 0.057880498468875885
},
"2": {
"running_threshold": 0.8691568970680237,
"running_std": 0.04676362872123718
}
},
"optimizer_included": false,
"precision_changed": false,
"validation": {
"exported_tensors_equal": true,
"strict_model_load": true,
"generation": true,
"cuda": {
"versions": {
"torch": "2.6.0+cu124",
"transformers": "5.0.0",
"mamba-ssm": "2.3.0",
"causal-conv1d": "1.6.0",
"flash-linear-attention": "0.4.2",
"triton": "3.2.0",
"safetensors": "0.7.0"
},
"tokenizer_revision": "5590cc63198b0f111f4289daf006da0a12044173",
"strict_model_load": true,
"generation": true,
"device": "NVIDIA H100 NVL",
"torch": "2.6.0+cu124",
"autocast": "bfloat16",
"prefill_max_abs_diff": 0.0,
"decode_max_abs_diff": 0.0625,
"generated_token_ids": [
264,
4382,
7855,
6576,
430,
649,
387,
1511
],
"published_cli_passed": true,
"test_scope": "Short CUDA smoke test; not a benchmark rerun",
"comparison_tolerance": {
"rtol": 0.02,
"atol": 0.1
}
},
"router_release": {
"exact_router_export": true,
"base_model_association": true,
"default_forward_max_abs_diff": 0.0,
"modes": [
{
"router": false,
"cases": [
{
"batch": 1,
"length": 16,
"prefill_max_abs_diff": 0.0,
"decode_max_abs_diff": 0.0625
},
{
"batch": 1,
"length": 128,
"prefill_max_abs_diff": 0.0625,
"decode_max_abs_diff": 0.125
},
{
"batch": 2,
"length": 64,
"prefill_max_abs_diff": 0.0625
}
],
"canonical_router_exact": null
},
{
"router": true,
"cases": [
{
"batch": 1,
"length": 16,
"prefill_max_abs_diff": 0.0625,
"decode_max_abs_diff": 0.0625
},
{
"batch": 1,
"length": 128,
"prefill_max_abs_diff": 0.0,
"decode_max_abs_diff": 0.125
},
{
"batch": 2,
"length": 64,
"prefill_max_abs_diff": 0.0
}
],
"canonical_router_exact": true
}
],
"frozen_thresholds_preserved": true,
"training_router_rejected": true,
"gpu": "NVIDIA H100 NVL",
"torch": "2.6.0+cu124",
"generation_cli_modes_passed": [
"entropy",
"router"
]
},
"current_release": {
"strict_model_load": true,
"entropy_and_router_generation_passed": true,
"unit_tests_passed": 17,
"device": "NVIDIA H100 NVL"
}
},
"implementation": "AMOR shared release: native entropy gate and optional fitted inference router",
"router": {
"variant": "distill",
"hidden_dim": 512,
"activation": "silu",
"weights": "router.safetensors",
"config": "router_config.json",
"sha256": "41a61cd94e307aebea14a430d99cb7e95f2bb5e1ddc3b2b1f7ef4beb87ab5081",
"base_weight_sha256": "d2958bda89a6ecf036d00c545e9b36fb0d8562cb50d41b835397616a3d87530d",
"precision": "float32",
"inference_only": true
},
"weight_layout": {
"format": "safetensors-sharded-v1",
"index": "model.safetensors.index.json",
"shards": 2,
"weight_identity_sha256": "d2958bda89a6ecf036d00c545e9b36fb0d8562cb50d41b835397616a3d87530d",
"original_single_file_sha256": "577c139d54e9e4e353fd729f62e246f0b012d356cd2fcdc650ddb6d1fae9331b",
"all_tensor_bytes_identical": true
}
}