Video-to-Video
PyTorch
File size: 1,203 Bytes
a9a7638
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
{
  "model_name": "FlashRender",
  "description": "Few-step generative rendering via camera-controlled video MeanFlow.",
  "task": "video-to-video",
  "license": "apache-2.0",
  "code_repository": "https://github.com/byeongjun-park/FlashRender",
  "base_model": "alibaba-pai/Wan2.1-Fun-V1.1-1.3B-Control-Camera",
  "torch_dtype": "bfloat16",
  "height": 480,
  "width": 832,
  "num_frames": 81,
  "state_dict_format": "partial (trainable subset only); load with strict=False onto the base Wan DiT patched by flashrender_utils.model_utils.adjust_to_FlashRender",
  "checkpoints": {
    "epoch_20.safetensors": {
      "stage": 1,
      "training": "multi-step flow matching (t = r) with RETA feature alignment",
      "num_inference_steps": 50,
      "inference_mode": "mul"
    },
    "meanflow_epoch_20.safetensors": {
      "stage": 2,
      "training": "few-step MeanFlow distillation, resumed from stage 1",
      "num_inference_steps": 4,
      "inference_mode": "any"
    },
    "onpolicy_epoch_5.safetensors": {
      "stage": 3,
      "training": "on-policy flow-map distillation (DMD/VSD + GAN) of the stage-2 student",
      "num_inference_steps": 4,
      "inference_mode": "any"
    }
  }
}