Stable-Video-Infinity
[ICLR 26 Oral] Stable Video Infinity: Infinite-Length Video Generation with Error Recycling
File Explorer
Download Latest Version (.zip)- intro.png
- logo.png
- logo_white.png
- youtube1.png
- youtube2.png
- youtube3.png
- SVI-Shot-StreamingPrompt-1028.json
- SVI-Shot-StreamingPrompt.mp4
- test.jpg
- image.png
- pose.mp4
- frame.jpg
- prompt.txt
- frame.jpg
- prompt.txt
- frame.jpg
- prompt.txt
- obama.png
- obama_5min.wav
- frame.png
- prompt.txt
- dw_pose_with_foot_with_face.pkl
- dw_pose_with_foot_wo_face.pkl
- frame_data.pkl
- dw_pose_with_foot_with_face.pkl
- dw_pose_with_foot_wo_face.pkl
- frame_data.pkl
- dw_pose_with_foot_with_face.pkl
- dw_pose_with_foot_wo_face.pkl
- frame_data.pkl
- Cats.csv
- mixkit-a-white-cat-sits-in-front-of-a-white-wall-1535.mp4
- mixkit-a-woman-sits-on-a-couch-and-pets-a-cat-1536.mp4
- mixkit-a-boat-sailing-in-the-bay-1941.mp4
- mixkit-a-couple-at-the-beach-at-sunset-1040.mp4
- sea.csv
- audio_embedding.pkl
- frame_data.pkl
- audio_embedding.pkl
- frame_data.pkl
- audio_embedding.pkl
- frame_data.pkl
- 0a2bcd9380136373e4431e3f2a8b015c.wav
- 0a2c496c2246a88230f800052298612e.wav
- 0a2cda4f185353339eba0ee02ff536b0.wav
- __init__.py
- model_config.py
- model_config_talk.py
- __init__.py
- controlnet_unit.py
- processors.py
- __init__.py
- simple_text_image.py
- video.py
- __init__.py
- xdit_context_parallel.py
- __init__.py
- __init__.py
- accurate.py
- balanced.py
- fast.py
- interpolation.py
- __init__.py
- api.py
- cupy_kernels.py
- data.py
- patch_match.py
- __init__.py
- blip.py
- blip_pretrain.py
- med.py
- vit.py
- ViT-H-14.json
- __init__.py
- coca_model.py
- constants.py
- factory.py
- generation_utils.py
- hf_configs.py
- hf_model.py
- loss.py
- model.py
- modified_resnet.py
- openai.py
- pretrained.py
- push_to_hf_hub.py
- timm_model.py
- tokenizer.py
- transform.py
- transformer.py
- utils.py
- version.py
- __init__.py
- base_model.py
- clip_model.py
- cross_modeling.py
- __init__.py
- __init__.py
- aesthetic.py
- clip.py
- config.py
- hps.py
- imagereward.py
- mps.py
- pickscore.py
- __init__.py
- __init__.py
- __init__.py
- attention.py
- cog_dit.py
- cog_vae.py
- downloader.py
- flux_controlnet.py
- flux_dit.py
- flux_ipadapter.py
- flux_text_encoder.py
- flux_vae.py
- hunyuan_dit.py
- hunyuan_dit_text_encoder.py
- hunyuan_video_dit.py
- hunyuan_video_text_encoder.py
- hunyuan_video_vae_decoder.py
- hunyuan_video_vae_encoder.py
- kolors_text_encoder.py
- lora.py
- model_manager.py
- omnigen.py
- sd3_dit.py
- sd3_text_encoder.py
- sd3_vae_decoder.py
- sd3_vae_encoder.py
- sd_controlnet.py
- sd_ipadapter.py
- sd_motion.py
- sd_text_encoder.py
- sd_unet.py
- sd_vae_decoder.py
- sd_vae_encoder.py
- sdxl_controlnet.py
- sdxl_ipadapter.py
- sdxl_motion.py
- sdxl_text_encoder.py
- sdxl_unet.py
- sdxl_vae_decoder.py
- sdxl_vae_encoder.py
- stepvideo_dit.py
- stepvideo_text_encoder.py
- stepvideo_vae.py
- svd_image_encoder.py
- svd_unet.py
- svd_vae_decoder.py
- svd_vae_encoder.py
- tiler.py
- utils.py
- wan_video_dit.py
- wan_video_dit_talk.py
- wan_video_image_encoder.py
- wan_video_text_encoder.py
- wan_video_vae.py
- __init__.py
- base.py
- cog_video.py
- dancer.py
- flux_image.py
- hunyuan_image.py
- hunyuan_video.py
- omnigen_image.py
- pipeline_runner.py
- sd3_image.py
- sd_image.py
- sd_video.py
- sdxl_image.py
- sdxl_video.py
- step_video.py
- svd_video.py
- svi_video.py
- svi_video_dance.py
- svi_video_talk.py
- wan_video.py
- __init__.py
- base.py
- FastBlend.py
- PILEditor.py
- RIFE.py
- sequencial_processor.py
- __init__.py
- base_prompter.py
- cog_prompter.py
- flux_prompter.py
- hunyuan_dit_prompter.py
- hunyuan_video_prompter.py
- kolors_prompter.py
- omnigen_prompter.py
- omost.py
- prompt_refiners.py
- sd3_prompter.py
- sd_prompter.py
- sdxl_prompter.py
- stepvideo_prompter.py
- wan_prompter.py
- __init__.py
- continuous_ode.py
- ddim.py
- flow_match.py
- added_tokens.json
- special_tokens_map.json
- spiece.model
- tokenizer_config.json
- merges.txt
- special_tokens_map.json
- tokenizer_config.json
- vocab.json
- special_tokens_map.json
- spiece.model
- tokenizer.json
- tokenizer_config.json
- special_tokens_map.json
- tokenizer_config.json
- vocab.txt
- vocab_org.txt
- config.json
- special_tokens_map.json
- spiece.model
- tokenizer_config.json
- merges.txt
- special_tokens_map.json
- tokenizer_config.json
- vocab.json
- preprocessor_config.json
- special_tokens_map.json
- tokenizer.json
- tokenizer_config.json
- tokenizer.model
- tokenizer_config.json
- vocab.txt
- merges.txt
- special_tokens_map.json
- tokenizer_config.json
- vocab.json
- merges.txt
- special_tokens_map.json
- tokenizer_config.json
- vocab.json
- merges.txt
- special_tokens_map.json
- tokenizer_config.json
- vocab.json
- special_tokens_map.json
- spiece.model
- tokenizer.json
- tokenizer_config.json
- merges.txt
- special_tokens_map.json
- tokenizer_config.json
- vocab.json
- __init__.py
- __init__.py
- text_to_image.py
- __init__.py
- fm_solvers.py
- fm_solvers_unipc.py
- multitalk_utils.py
- prompt_extend.py
- qwen_vl_utils.py
- utils.py
- vace_processor.py
- __init__.py
- layers.py
- __init__.py
- causal.png
- DevLog.md
- FAQ.md
- talk.png
- wan22_preview.png
- __init__.py
- onnxdet.py
- onnxpose.py
- util.py
- wholebody.py
- prepare_video_audio.py
- prepare_video_pose.py
- process_mixkit.py
- svi_2.0.sh
- svi_dance.sh
- svi_film.sh
- svi_shot.sh
- svi_talk.sh
- svi_tom.sh
- svi_dance.sh
- svi_film.sh
- svi_shot.sh
- svi_talk.sh
- torch_utils.py
- wav2vec2.py
- __init__.py
- layers.py
- utils.py
- audio_process.py
- extract_lora.py
- image_process.py
- metadata_gen.py
- process.py
- project_utils.py
- run_align_pose.py
- text_utils.py
- video_process.py
- .gitignore
- gradio_demo.py
- gradio_demo.sh
- LICENSE
- README.md
- requirements.txt
- setup.py
- test_svi.py
- test_svi_dance.py
- test_svi_talk.py
- train_svi.py
- train_svi_dance.py
- train_svi_talk.py
// repository documentation
Was this content helpful?
(0 ratings)
