AMOR-GatedDeltaNet-1.5B / export_metadata.json
FlyinGodzilla's picture
Provide verified sharded weights and optional inference routers
9203c4e verified
Raw History Blame Contribute Delete
4.37 kB
{
"repo_id": "FlyinGodzilla/AMOR-GatedDeltaNet-1.5B",
"source_filename": "amor_gdn_1p5b_step_0062558.pt",
"source_sha256": "a635e9fbf742f17f71b607565440ed8c338a90a00173b1f0813e79806e45f0f7",
"step": 62558,
"tokens_seen": 30748520448,
"n_params": 1524025472,
"tensor_dtypes": {
"torch.float32": 458
},
"gate_buffers": {
"0": {
"running_threshold": 0.9387496709823608,
"running_std": 0.028903765603899956
},
"1": {
"running_threshold": 0.9252517819404602,
"running_std": 0.0452539287507534
},
"2": {
"running_threshold": 0.9380253553390503,
"running_std": 0.027104970067739487
}
},
"optimizer_included": false,
"precision_changed": false,
"validation": {
"exported_tensors_equal": true,
"strict_model_load": true,
"generation": true,
"cuda": {
"versions": {
"torch": "2.6.0+cu124",
"transformers": "5.0.0",
"mamba-ssm": "2.3.0",
"causal-conv1d": "1.6.0",
"flash-linear-attention": "0.4.2",
"triton": "3.2.0",
"safetensors": "0.7.0"
},
"tokenizer_revision": "5590cc63198b0f111f4289daf006da0a12044173",
"strict_model_load": true,
"generation": true,
"device": "NVIDIA H100 NVL",
"torch": "2.6.0+cu124",
"autocast": "bfloat16",
"prefill_max_abs_diff": 0.0,
"decode_max_abs_diff": 0.0625,
"generated_token_ids": [
264,
502,
4668,
304,
279,
469,
1950,
430
],
"published_cli_passed": true,
"test_scope": "Short CUDA smoke test; not a benchmark rerun",
"comparison_tolerance": {
"rtol": 0.02,
"atol": 0.1
}
},
"router_release": {
"exact_router_export": true,
"base_model_association": true,
"default_forward_max_abs_diff": 0.0,
"modes": [
{
"router": false,
"cases": [
{
"batch": 1,
"length": 16,
"prefill_max_abs_diff": 0.0,
"decode_max_abs_diff": 0.0625
},
{
"batch": 1,
"length": 128,
"prefill_max_abs_diff": 0.03125,
"decode_max_abs_diff": 0.125
},
{
"batch": 2,
"length": 64,
"prefill_max_abs_diff": 0.0
}
],
"canonical_router_exact": null
},
{
"router": true,
"cases": [
{
"batch": 1,
"length": 16,
"prefill_max_abs_diff": 0.0,
"decode_max_abs_diff": 0.0625
},
{
"batch": 1,
"length": 128,
"prefill_max_abs_diff": 0.0,
"decode_max_abs_diff": 0.125
},
{
"batch": 2,
"length": 64,
"prefill_max_abs_diff": 0.0
}
],
"canonical_router_exact": true
}
],
"frozen_thresholds_preserved": true,
"training_router_rejected": true,
"gpu": "NVIDIA H100 NVL",
"torch": "2.6.0+cu124",
"generation_cli_modes_passed": [
"entropy",
"router"
]
},
"current_release": {
"strict_model_load": true,
"entropy_and_router_generation_passed": true,
"unit_tests_passed": 17,
"device": "NVIDIA H100 NVL"
}
},
"implementation": "AMOR shared release: native entropy gate and optional fitted inference router",
"router": {
"variant": "distill",
"hidden_dim": 512,
"activation": "silu",
"weights": "router.safetensors",
"config": "router_config.json",
"sha256": "f67c3a65f7b92cb884d55cedaf49d39dca9a85b7a205c2492c9f8d2fc5b0fb9b",
"base_weight_sha256": "eec1f7d2e406b4df732714cb66ec968cd2bfa5c3f24888f3141e51ae1f001bdd",
"precision": "float32",
"inference_only": true
},
"weight_layout": {
"format": "safetensors-sharded-v1",
"index": "model.safetensors.index.json",
"shards": 8,
"weight_identity_sha256": "eec1f7d2e406b4df732714cb66ec968cd2bfa5c3f24888f3141e51ae1f001bdd",
"original_single_file_sha256": "87f2506b6d1f4ac36936b715920e80e348091fa9620391267a67ef27d17c7be5",
"all_tensor_bytes_identical": true
}
}