File size: 1,203 Bytes
a9a7638 | 1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18 19 20 21 22 23 24 25 26 27 28 29 30 31 32 33 34 | {
"model_name": "FlashRender",
"description": "Few-step generative rendering via camera-controlled video MeanFlow.",
"task": "video-to-video",
"license": "apache-2.0",
"code_repository": "https://github.com/byeongjun-park/FlashRender",
"base_model": "alibaba-pai/Wan2.1-Fun-V1.1-1.3B-Control-Camera",
"torch_dtype": "bfloat16",
"height": 480,
"width": 832,
"num_frames": 81,
"state_dict_format": "partial (trainable subset only); load with strict=False onto the base Wan DiT patched by flashrender_utils.model_utils.adjust_to_FlashRender",
"checkpoints": {
"epoch_20.safetensors": {
"stage": 1,
"training": "multi-step flow matching (t = r) with RETA feature alignment",
"num_inference_steps": 50,
"inference_mode": "mul"
},
"meanflow_epoch_20.safetensors": {
"stage": 2,
"training": "few-step MeanFlow distillation, resumed from stage 1",
"num_inference_steps": 4,
"inference_mode": "any"
},
"onpolicy_epoch_5.safetensors": {
"stage": 3,
"training": "on-policy flow-map distillation (DMD/VSD + GAN) of the stage-2 student",
"num_inference_steps": 4,
"inference_mode": "any"
}
}
}
|