File size: 937 Bytes
59630ba
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
# Important: Be careful when modifying this file! The fields in file will be overridden by the dataset dependent config file in the `configurations/dataset_experiment` folder so consider making changes there instead! 

defaults:
  - dfot_video
  - override backbone: u_vit3d_pose


camera_pose_conditioning:
  normalize_by: first # first, mean
  bound: null # float to scale the camera positions to [-bound, bound]^3
  type: ray_encoding # global (flattened extrinsics), ray (per-pixel origin and direction), ray_encoding (ray mapped to high-dimensional space), plucker (ray in Plücker coordinates)

# repa loss , added by haoyu 
alignment: 
  alignment_coeff: 0.5 # loss weight 
  encoder_type: vggt
  apply_unnormalize_recon: True 
  alignment_context_length: 16 # number of frames to use for repa loss
  # vggt_layer_index: -1 
  # encoder_info: [1, 1374, 2048]
  encoder_info: [24, 512, 512]
  mid_channels: 128
  latents_info: null