Download manifest.json from Axym-Labs/AxoSim: direct link, hf CLI and curl.
- Browser
- Download file 24.5 kB
-
https://huggingface.co/Axym-Labs/AxoSim/resolve/main/manifest.json
- Command line
-
hf download hf://Axym-Labs/AxoSim/manifest.json
-
curl -L -o manifest.json https://huggingface.co/Axym-Labs/AxoSim/resolve/main/manifest.json
24.5 kB
| { | |
| "artifacts": [ | |
| { | |
| "bytes": 624761, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "cumulative_train_presentations", | |
| "device", | |
| "inference_latency_median_ms", | |
| "last_components", | |
| "learning_rate", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "seed", | |
| "selection_assessments", | |
| "selection_deltas_vs_full", | |
| "selection_iteration", | |
| "selection_metrics", | |
| "selection_practical_gate", | |
| "selection_reason", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "source_checkpoint", | |
| "source_checkpoint_sha256", | |
| "train_presentations", | |
| "trainable_parameter_count", | |
| "training_step_latency_median_ms", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "parameter_count": 88858, | |
| "path": "released/axomamba.pt", | |
| "sha256": "c806ba612a828d72f89871650ded40052ab81a3b726baee6b4eab804bf9195cc", | |
| "status": "retained", | |
| "training": "original dual-objective release" | |
| }, | |
| { | |
| "bytes": 443317, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "device", | |
| "last_components", | |
| "learning_rate", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "seed", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "train_presentations", | |
| "trainable_parameter_count", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "parameter_count": 42748, | |
| "path": "released/axomamba-structured-compact.pt", | |
| "sha256": "a5805f84540a2391b26adecd03e609e612e74a5c1a91fcfd8598527a7c703804", | |
| "status": "retained", | |
| "training": "structured compact release" | |
| }, | |
| { | |
| "bytes": 401397, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "device", | |
| "estimated_adamw_training_state_bytes", | |
| "frozen_parameter_bytes", | |
| "frozen_parameter_count", | |
| "last_components", | |
| "learning_rate", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "per_neuron_streaming_state_bytes", | |
| "per_neuron_trainable_plus_streaming_bytes", | |
| "prefetch_batches", | |
| "resident_buffer_bytes", | |
| "resident_buffer_count", | |
| "resident_parameter_bytes", | |
| "resident_parameter_count", | |
| "resident_state_bytes", | |
| "seed", | |
| "shared_buffer_bytes", | |
| "shared_fixed_state_bytes", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "streaming_state_components_bytes", | |
| "streaming_state_dtype", | |
| "train_presentations", | |
| "trainable_parameter_bytes", | |
| "trainable_parameter_count", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "parameter_count": 31756, | |
| "path": "released/axomamba-regression.pt", | |
| "sha256": "c48c1d8a97552bba9f38f4871d3515cc936b463b576c4006400e0df698c98a27", | |
| "status": "retained", | |
| "training": "regression-specialized release" | |
| }, | |
| { | |
| "bytes": 426485, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "device", | |
| "estimated_adamw_training_state_bytes", | |
| "frozen_parameter_bytes", | |
| "frozen_parameter_count", | |
| "last_components", | |
| "learning_rate", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "per_neuron_streaming_state_bytes", | |
| "per_neuron_trainable_plus_streaming_bytes", | |
| "prefetch_batches", | |
| "resident_buffer_bytes", | |
| "resident_buffer_count", | |
| "resident_parameter_bytes", | |
| "resident_parameter_count", | |
| "resident_state_bytes", | |
| "seed", | |
| "shared_buffer_bytes", | |
| "shared_fixed_state_bytes", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "streaming_state_components_bytes", | |
| "streaming_state_dtype", | |
| "train_presentations", | |
| "trainable_parameter_bytes", | |
| "trainable_parameter_count", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "parameter_count": 38257, | |
| "path": "released/axomamba-spike.pt", | |
| "sha256": "50cb87986ea13ab467405ef45fef9727069405a4aed5d9d6392447f9a157db50", | |
| "status": "retained", | |
| "training": "spike-specialized release" | |
| }, | |
| { | |
| "bytes": 676643, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "connected_acceptance_iterations", | |
| "device", | |
| "estimated_adamw_training_state_bytes", | |
| "frozen_parameter_bytes", | |
| "frozen_parameter_count", | |
| "last_components", | |
| "learning_rate", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "per_neuron_adamw_training_state_bytes", | |
| "per_neuron_simulation_state_bytes", | |
| "per_neuron_streaming_state_bytes", | |
| "per_neuron_trainable_plus_streaming_bytes", | |
| "population_adamw_training_state_bytes_at_reference_scale", | |
| "population_memory_reference_neurons", | |
| "population_simulation_state_bytes_at_reference_scale", | |
| "prefetch_batches", | |
| "public_name", | |
| "resident_buffer_bytes", | |
| "resident_buffer_count", | |
| "resident_parameter_bytes", | |
| "resident_parameter_count", | |
| "resident_state_bytes", | |
| "seed", | |
| "shared_buffer_bytes", | |
| "shared_fixed_state_bytes", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "source_checkpoint_sha256", | |
| "source_iterations", | |
| "streaming_state_components_bytes", | |
| "streaming_state_dtype", | |
| "train_presentations", | |
| "trainable_parameter_bytes", | |
| "trainable_parameter_count", | |
| "training_distribution", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "parameter_count": 90214, | |
| "path": "released/axomamba-population.pt", | |
| "sha256": "f1c02362f335e6a40337a18219d0fb4537489984b00ea5f12c1fef85cb62f7f5", | |
| "status": "retained", | |
| "training": "7,500-trace development population corpus" | |
| }, | |
| { | |
| "bytes": 327738, | |
| "checkpoint_metadata_fields": [ | |
| "model_kind", | |
| "optimizer_steps", | |
| "processed_shared_batches", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "training_snapshot" | |
| ], | |
| "class": "AxoTemporalModel", | |
| "family": "AxoSim-GRU", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 20.71521737746768, | |
| "dynamics_sera_mv2_per_ms2": 429.12023099573895, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.6043874784323391, | |
| "spike_f1_area_0_5ms": 0.4991865910771507, | |
| "spike_f1_exact": 0.23169829923588858, | |
| "spike_mean_f1_0_5ms": 0.48566264070331117, | |
| "voltage_root_sera_mv": 15.617826051098955, | |
| "voltage_sera_mv2": 243.91649056238518 | |
| }, | |
| "parameter_count": 17432, | |
| "path": "ordinary/axosim-gru-small-32m.pt", | |
| "sha256": "e2225ab7df191986d4921b15f248ba7965729700b12b90fef6c0d69c7071d744", | |
| "status": "retained", | |
| "training": "ordinary AxoBench, 32M presentations" | |
| }, | |
| { | |
| "bytes": 749015, | |
| "checkpoint_metadata_fields": [ | |
| "model_kind", | |
| "optimizer_steps", | |
| "processed_shared_batches", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "training_snapshot" | |
| ], | |
| "class": "AxoTemporalModel", | |
| "family": "AxoSim-GRU", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 20.715431633179396, | |
| "dynamics_sera_mv2_per_ms2": 429.1291077489296, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.6405832320777642, | |
| "spike_f1_area_0_5ms": 0.5359902794653706, | |
| "spike_f1_exact": 0.26148238153098424, | |
| "spike_mean_f1_0_5ms": 0.5218307006885379, | |
| "voltage_root_sera_mv": 15.543812757698129, | |
| "voltage_sera_mv2": 241.61011504637912 | |
| }, | |
| "parameter_count": 122744, | |
| "path": "ordinary/axosim-gru-medium-40m.pt", | |
| "sha256": "767bbe91e7a3403d8f2364ef42ed531a44e48b8e6b6b986d7d500d38b52af5c5", | |
| "status": "retained", | |
| "training": "ordinary AxoBench, 40M presentations" | |
| }, | |
| { | |
| "bytes": 399135, | |
| "checkpoint_metadata_fields": [ | |
| "model_kind", | |
| "optimizer_steps", | |
| "processed_shared_batches", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "training_snapshot" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 20.70566130570553, | |
| "dynamics_sera_mv2_per_ms2": 428.7244101065912, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.5569999999999999, | |
| "spike_f1_area_0_5ms": 0.46449999999999997, | |
| "spike_f1_exact": 0.24100000000000002, | |
| "spike_mean_f1_0_5ms": 0.4535833333333334, | |
| "voltage_root_sera_mv": 15.629929894146915, | |
| "voltage_sera_mv2": 244.29470849594742 | |
| }, | |
| "parameter_count": 32536, | |
| "path": "ordinary/axosim-mamba-small-40m.pt", | |
| "sha256": "69a4b68c05b214588abeb535887b902c74681f2c3990b1fd6edf412a99cc56bd", | |
| "status": "retained", | |
| "training": "ordinary AxoBench, 40M presentations" | |
| }, | |
| { | |
| "bytes": 4790653, | |
| "checkpoint_metadata_fields": [ | |
| "model_kind", | |
| "optimizer_steps", | |
| "processed_shared_batches", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "training_snapshot" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 20.952361125905952, | |
| "dynamics_sera_mv2_per_ms2": 439.001436750375, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.6406528946697272, | |
| "spike_f1_area_0_5ms": 0.535067584799796, | |
| "spike_f1_exact": 0.2509563886763581, | |
| "spike_mean_f1_0_5ms": 0.5201904276120038, | |
| "voltage_root_sera_mv": 15.635292502976363, | |
| "voltage_sera_mv2": 244.46237165362885 | |
| }, | |
| "parameter_count": 1128200, | |
| "path": "ordinary/axosim-mamba-medium-40m.pt", | |
| "sha256": "27111580e9a9b11aa7bd7cf77654081eedd0354016a670a029f9a2b5b4fe3b8a", | |
| "status": "retained", | |
| "training": "ordinary AxoBench, 40M presentations" | |
| }, | |
| { | |
| "bytes": 381900, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "device", | |
| "estimated_adamw_training_state_bytes", | |
| "frozen_parameter_bytes", | |
| "frozen_parameter_count", | |
| "gradient_clipped_steps", | |
| "gradient_norm_observations", | |
| "last_components", | |
| "learning_rate", | |
| "loss_observations", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "per_neuron_adamw_training_state_bytes", | |
| "per_neuron_simulation_state_bytes", | |
| "per_neuron_streaming_state_bytes", | |
| "per_neuron_trainable_plus_streaming_bytes", | |
| "population_adamw_training_state_bytes_at_reference_scale", | |
| "population_memory_reference_neurons", | |
| "population_simulation_state_bytes_at_reference_scale", | |
| "preclip_gradient_norm_max", | |
| "preclip_gradient_norm_mean", | |
| "prefetch_batches", | |
| "processed_shared_batches", | |
| "resident_buffer_bytes", | |
| "resident_buffer_count", | |
| "resident_parameter_bytes", | |
| "resident_parameter_count", | |
| "resident_state_bytes", | |
| "resume_path", | |
| "seed", | |
| "shared_buffer_bytes", | |
| "shared_fixed_state_bytes", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "streaming_state_components_bytes", | |
| "streaming_state_dtype", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "trainable_parameter_bytes", | |
| "trainable_parameter_count", | |
| "training_complete", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoTemporalModel", | |
| "family": "AxoSim-GRU", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 4.4988664832787695, | |
| "dynamics_sera_mv2_per_ms2": 20.239799634369085, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.4700268789198425, | |
| "spike_f1_area_0_5ms": 0.38951930489863107, | |
| "spike_f1_exact": 0.18469360115016772, | |
| "spike_mean_f1_0_5ms": 0.37915946075469337, | |
| "voltage_root_sera_mv": 6.935555825886408, | |
| "voltage_sera_mv2": 48.10193461398689 | |
| }, | |
| "parameter_count": 18762, | |
| "path": "population-finetuned/axosim-gru-small.pt", | |
| "sha256": "fe6e8dd091f91ff6d112b7c4ef07f978e5eaaa218193f1d1e70b58296ae52ac8", | |
| "status": "retained", | |
| "training": "axobench-population, 30K presentations" | |
| }, | |
| { | |
| "bytes": 803596, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "device", | |
| "estimated_adamw_training_state_bytes", | |
| "frozen_parameter_bytes", | |
| "frozen_parameter_count", | |
| "gradient_clipped_steps", | |
| "gradient_norm_observations", | |
| "last_components", | |
| "learning_rate", | |
| "loss_observations", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "per_neuron_adamw_training_state_bytes", | |
| "per_neuron_simulation_state_bytes", | |
| "per_neuron_streaming_state_bytes", | |
| "per_neuron_trainable_plus_streaming_bytes", | |
| "population_adamw_training_state_bytes_at_reference_scale", | |
| "population_memory_reference_neurons", | |
| "population_simulation_state_bytes_at_reference_scale", | |
| "preclip_gradient_norm_max", | |
| "preclip_gradient_norm_mean", | |
| "prefetch_batches", | |
| "processed_shared_batches", | |
| "resident_buffer_bytes", | |
| "resident_buffer_count", | |
| "resident_parameter_bytes", | |
| "resident_parameter_count", | |
| "resident_state_bytes", | |
| "resume_path", | |
| "seed", | |
| "shared_buffer_bytes", | |
| "shared_fixed_state_bytes", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "streaming_state_components_bytes", | |
| "streaming_state_dtype", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "trainable_parameter_bytes", | |
| "trainable_parameter_count", | |
| "training_complete", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoTemporalModel", | |
| "family": "AxoSim-GRU", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 4.846691910811208, | |
| "dynamics_sera_mv2_per_ms2": 23.490422478322795, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.5044272452458033, | |
| "spike_f1_area_0_5ms": 0.41325386445983325, | |
| "spike_f1_exact": 0.18785241086551038, | |
| "spike_mean_f1_0_5ms": 0.40206819172580377, | |
| "voltage_root_sera_mv": 7.094813570716471, | |
| "voltage_sera_mv2": 50.3363796032226 | |
| }, | |
| "parameter_count": 124170, | |
| "path": "population-finetuned/axosim-gru-medium.pt", | |
| "sha256": "30e7e81aea7a329bec8ba7bb109a08ce8dca3d899ec0fc66d055c1d0c59a54cc", | |
| "status": "retained", | |
| "training": "axobench-population, 30K presentations" | |
| }, | |
| { | |
| "bytes": 453039, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "device", | |
| "estimated_adamw_training_state_bytes", | |
| "frozen_parameter_bytes", | |
| "frozen_parameter_count", | |
| "gradient_clipped_steps", | |
| "gradient_norm_observations", | |
| "last_components", | |
| "learning_rate", | |
| "loss_observations", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "per_neuron_adamw_training_state_bytes", | |
| "per_neuron_simulation_state_bytes", | |
| "per_neuron_streaming_state_bytes", | |
| "per_neuron_trainable_plus_streaming_bytes", | |
| "population_adamw_training_state_bytes_at_reference_scale", | |
| "population_memory_reference_neurons", | |
| "population_simulation_state_bytes_at_reference_scale", | |
| "preclip_gradient_norm_max", | |
| "preclip_gradient_norm_mean", | |
| "prefetch_batches", | |
| "processed_shared_batches", | |
| "resident_buffer_bytes", | |
| "resident_buffer_count", | |
| "resident_parameter_bytes", | |
| "resident_parameter_count", | |
| "resident_state_bytes", | |
| "resume_path", | |
| "seed", | |
| "shared_buffer_bytes", | |
| "shared_fixed_state_bytes", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "streaming_state_components_bytes", | |
| "streaming_state_dtype", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "trainable_parameter_bytes", | |
| "trainable_parameter_count", | |
| "training_complete", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 4.619637103264988, | |
| "dynamics_sera_mv2_per_ms2": 21.341046965862525, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.4503503566694022, | |
| "spike_f1_area_0_5ms": 0.37376849104644067, | |
| "spike_f1_exact": 0.18391094838288827, | |
| "spike_mean_f1_0_5ms": 0.3643288512930581, | |
| "voltage_root_sera_mv": 6.920495757974897, | |
| "voltage_sera_mv2": 47.89326153614854 | |
| }, | |
| "parameter_count": 33866, | |
| "path": "population-finetuned/axosim-mamba-small.pt", | |
| "sha256": "b835fdf5e6b20bbdfef4a621bf81b2e1bf03ee593716826c9bc1604bfbae83a2", | |
| "status": "retained", | |
| "training": "axobench-population, 30K presentations" | |
| }, | |
| { | |
| "bytes": 4844791, | |
| "checkpoint_metadata_fields": [ | |
| "batch_size", | |
| "burn_in", | |
| "candidate_compute_presentations_per_second", | |
| "candidate_compute_seconds", | |
| "candidate_id", | |
| "device", | |
| "estimated_adamw_training_state_bytes", | |
| "frozen_parameter_bytes", | |
| "frozen_parameter_count", | |
| "gradient_clipped_steps", | |
| "gradient_norm_observations", | |
| "last_components", | |
| "learning_rate", | |
| "loss_observations", | |
| "lr_schedule", | |
| "mean_training_loss", | |
| "model_kind", | |
| "morphology_conditioning", | |
| "objective", | |
| "optimizer", | |
| "optimizer_steps", | |
| "output_distillation_weight", | |
| "panel_candidates", | |
| "panel_presentations_per_second", | |
| "panel_wall_seconds", | |
| "parameter_count", | |
| "per_neuron_adamw_training_state_bytes", | |
| "per_neuron_simulation_state_bytes", | |
| "per_neuron_streaming_state_bytes", | |
| "per_neuron_trainable_plus_streaming_bytes", | |
| "population_adamw_training_state_bytes_at_reference_scale", | |
| "population_memory_reference_neurons", | |
| "population_simulation_state_bytes_at_reference_scale", | |
| "preclip_gradient_norm_max", | |
| "preclip_gradient_norm_mean", | |
| "prefetch_batches", | |
| "processed_shared_batches", | |
| "resident_buffer_bytes", | |
| "resident_buffer_count", | |
| "resident_parameter_bytes", | |
| "resident_parameter_count", | |
| "resident_state_bytes", | |
| "resume_path", | |
| "seed", | |
| "shared_buffer_bytes", | |
| "shared_fixed_state_bytes", | |
| "shared_super_batch_size", | |
| "shared_teacher_compute_seconds", | |
| "streaming_state_components_bytes", | |
| "streaming_state_dtype", | |
| "total_shared_batches", | |
| "train_presentations", | |
| "trainable_parameter_bytes", | |
| "trainable_parameter_count", | |
| "training_complete", | |
| "warmup_seconds", | |
| "weight_decay" | |
| ], | |
| "class": "AxoMamba", | |
| "family": "AxoSim-Mamba", | |
| "metrics": { | |
| "dynamics_root_sera_mv_per_ms": 4.922780051120665, | |
| "dynamics_sera_mv2_per_ms2": 24.23376343171158, | |
| "response_root_sera_mv": null, | |
| "response_sera_mv2": null, | |
| "spike_f1_5ms": 0.5115619739523345, | |
| "spike_f1_area_0_5ms": 0.4234074599096306, | |
| "spike_f1_exact": 0.1967750509435634, | |
| "spike_mean_f1_0_5ms": 0.4118676353326836, | |
| "voltage_root_sera_mv": 7.016426800370153, | |
| "voltage_sera_mv2": 49.23024504495255 | |
| }, | |
| "parameter_count": 1129674, | |
| "path": "population-finetuned/axosim-mamba-medium.pt", | |
| "sha256": "37b37b06f88e7276b92d302706f442a89654024e6c66ceca05c0faf121048707", | |
| "status": "retained", | |
| "training": "axobench-population, 30K presentations" | |
| } | |
| ], | |
| "excluded": "rejected trials, superseded snapshots, optimizer/resume state, caches, and third-party Branch-ELM weights", | |
| "format_version": 1, | |
| "repository": "Axym-Labs/AxoSim", | |
| "scope": "retained release, publication, and population-specialized AxoSim checkpoints" | |
| } | |