{"version":"1","model":{"id":"fal-ai/ltx-2.3-22b/image-to-video","name":"LTX-2.3 22B (I2V)","kind":"video","availability":"available","available":true,"restricted":false,"api_key_only":false,"description":"LTX-2.3 22B image-to-video — animate a start image with optional end-frame using the full 22B model, with native audio and advanced scheduler controls.","categories":["i2v"],"prompt_required":true,"prompt_supported":true,"supports":{"startImage":true,"endImage":true,"audioGeneration":true},"requires":{"startImage":true},"aspect_ratios":["16:9","9:16","1:1"],"resolutions":["1080p","1440p","2160p"],"durations":[3,4,5,6,7,8,10],"base_credits":30,"api_gross_margin_percent":20,"added_at":"2026-05-09","popularity_rank":288,"example_prompt":"A red London double-decker bus crossing Tower Bridge at blue hour, city lights reflecting on the Thames, camera slowly pulling back to reveal the full skyline","preview":{"video":"https://cdn.artemotion.ai/file/artemotion-ai-bckt/model-previews/ltx-2.3-22b-i2v-sm.mp4"},"schema_url":"/api/v1/models?id=fal-ai%2Fltx-2.3-22b%2Fimage-to-video","input_schema":{"$schema":"https://json-schema.org/draft/2020-12/schema","$id":"https://www.artemotion.ai/api/v1/models?id=fal-ai%2Fltx-2.3-22b%2Fimage-to-video","title":"LTX-2.3 22B (I2V) generation request","description":"Request body accepted by POST /api/v1/generate for fal-ai/ltx-2.3-22b/image-to-video.","type":"object","properties":{"model_id":{"type":"string","const":"fal-ai/ltx-2.3-22b/image-to-video","description":"ArtEmotion model identifier."},"extra":{"type":"object","additionalProperties":true,"description":"Model-specific settings may also be nested here."},"max_credits":{"type":"number","minimum":1,"description":"Reject before submission if the estimated list price exceeds this cap."},"webhook_url":{"type":"string","format":"uri","maxLength":2048},"webhook_secret":{"type":"string","maxLength":512},"folder_id":{"type":"string"},"prompt":{"type":"string"},"audio_cfg_scale":{"title":"Audio Cfg Scale","description":"The Classifier-Free Guidance (CFG) scale for the audio. Higher values result in more consistent and focused audio content.","default":7,"type":"number","minimum":1,"maximum":20},"generate_audio":{"title":"Generate Audio","description":"Whether to generate audio for the video.","default":true,"type":"boolean"},"audio_stg_scale":{"title":"Audio Stg Scale","description":"The Spatiotemporal Guidance (STG) scale for the audio. Higher values result in more consistent and focused audio content.","default":0,"type":"number","minimum":0,"maximum":20},"seed":{"title":"Seed","description":"The seed for the random number generator.","type":"integer"},"use_multiscale":{"title":"Use Multiscale","description":"Whether to use multi-scale generation. If True, the model will generate the video at a smaller scale first, then use the smaller video to guide the ge","default":true,"type":"boolean"},"audio_modality_scale":{"title":"Audio Modality Scale","description":"The modality scale for the audio. Controls the ratio between video and audio modalities.","default":3,"type":"number","minimum":0,"maximum":10},"use_restart_sampling":{"title":"Use Restart Sampling","description":"Whether to use restart sampling. This will inject a small amount of noise during each denoising step, which can help improve the quality of the genera","default":false,"type":"boolean"},"end_image_strength":{"title":"End Image Strength","description":"The strength of the end image to use for the video generation.","default":1,"type":"number","minimum":0,"maximum":1},"scheduler":{"title":"Scheduler","description":"The scheduler to use.","default":"ltx2","type":"string","enum":["ltx2","linear_quadratic","beta"]},"distill_lora_first_pass_scale":{"title":"Distill Lora First Pass Scale","description":"The scale of the distill LoRA to use for the first pass. Set to 0 to disable.","default":0.2,"type":"number","minimum":0,"maximum":1},"gradient_estimation_gamma":{"title":"Gradient Estimation Gamma","description":"The gamma of gradient estimation during denoising. Set to 0 to disable.","default":2,"type":"number","minimum":0,"maximum":10},"distill_lora_second_pass_scale":{"title":"Distill Lora Second Pass Scale","description":"The scale of the distill LoRA to use for the second and subsequent passes.","default":0.5,"type":"number","minimum":0,"maximum":1},"num_frames":{"title":"Num Frames","description":"The number of frames to generate.","default":121,"type":"number","minimum":9,"maximum":481,"multipleOf":1},"acceleration":{"title":"Acceleration","description":"The acceleration level to use.","default":"regular","type":"string","enum":["none","regular","high","full"]},"video_output_type":{"title":"Video Output Type","description":"The output type of the generated video.","default":"X264 (.mp4)","type":"string","enum":["X264 (.mp4)","VP9 (.webm)","PRORES4444 (.mov)","GIF (.gif)"]},"video_rescaling_scale":{"title":"Video Rescaling Scale","description":"The rescaling scale for the video. Controls the ratio between classifier-free guidance and spatiotemporal guidance.","default":0.7,"type":"number","minimum":0,"maximum":1},"video_quality":{"title":"Video Quality","description":"The quality of the generated video.","default":"high","type":"string","enum":["low","medium","high","maximum"]},"camera_lora_scale":{"title":"Camera LoRA Scale","description":"The scale of the camera LoRA to use. This allows you to control the camera movement of the generated video more accurately than just prompting the mod","default":1,"type":"number","minimum":0,"maximum":1},"audio_rescaling_scale":{"title":"Audio Rescaling Scale","description":"The rescaling scale for the audio. Controls the ratio between classifier-free guidance and spatiotemporal guidance.","default":0.7,"type":"number","minimum":0,"maximum":1},"video_modality_scale":{"title":"Video Modality Scale","description":"The modality scale for the video. Controls the ratio between video and audio modalities.","default":3,"type":"number","minimum":0,"maximum":10},"video_write_mode":{"title":"Video Write Mode","description":"The write mode of the generated video.","default":"balanced","type":"string","enum":["fast","balanced","small"]},"camera_lora":{"title":"Camera LoRA","description":"The camera LoRA to use. This allows you to control the camera movement of the generated video more accurately than just prompting the model to move th","default":"none","type":"string","enum":["dolly_in","dolly_out","dolly_left","dolly_right","jib_up","jib_down","static","none"]},"image_strength":{"title":"Image Strength","description":"The strength of the image to use for the video generation.","default":1,"type":"number","minimum":0,"maximum":1},"video_stg_scale":{"title":"Video Stg Scale","description":"The Spatiotemporal Guidance (STG) scale for the video. Higher values result in more consistent and focused video content.","default":0,"type":"number","minimum":0,"maximum":20},"video_cfg_scale":{"title":"Video Cfg Scale","description":"The Classifier-Free Guidance (CFG) scale for the video. Higher values result in more consistent and focused video content.","default":3,"type":"number","minimum":1,"maximum":20},"num_inference_steps":{"title":"Num Inference Steps","description":"The number of inference steps to use.","default":40,"type":"number","minimum":8,"maximum":50,"multipleOf":1},"negative_prompt":{"title":"Negative Prompt","description":"The negative prompt to generate the video from.","default":"news broadcast, 3d animation, computer graphics, pc game, console game, video game, cartoon, childish, watermark, logo, text, on screen text, subtitles, titles, signature, slowmo, static","type":"string"},"enable_prompt_expansion":{"title":"Enable Prompt Expansion","description":"Whether to enable prompt expansion.","default":true,"type":"boolean"},"enable_safety_checker":{"title":"Enable Safety Checker","description":"Run the safety checker to block unsafe content.","default":true,"type":"boolean"},"fps":{"title":"Fps","description":"The frames per second of the generated video.","default":24,"type":"number","minimum":1,"maximum":60},"interpolation_direction":{"title":"Interpolation Direction","description":"The direction to interpolate the image sequence in. 'Forward' goes from the start image to the end image, 'Backward' goes from the end image to the st","default":"forward","type":"string","enum":["forward","backward"]},"aspect_ratio":{"type":"string","enum":["16:9","9:16","1:1"]},"resolution":{"type":"string","enum":["1080p","1440p","2160p"]},"duration":{"type":"number","enum":[3,4,5,6,7,8,10]},"image_url":{"type":"string","format":"uri"},"end_image_url":{"type":"string","format":"uri"}},"required":["model_id","prompt","image_url"],"additionalProperties":false},"parameters":[{"key":"audio_cfg_scale","label":"Audio Cfg Scale","type":"number","default":7,"min":1,"max":20,"description":"The Classifier-Free Guidance (CFG) scale for the audio. Higher values result in more consistent and focused audio content."},{"key":"generate_audio","label":"Generate Audio","type":"toggle","default":true,"description":"Whether to generate audio for the video."},{"key":"audio_stg_scale","label":"Audio Stg Scale","type":"number","default":0,"min":0,"max":20,"description":"The Spatiotemporal Guidance (STG) scale for the audio. Higher values result in more consistent and focused audio content."},{"key":"seed","label":"Seed","type":"seed","description":"The seed for the random number generator."},{"key":"use_multiscale","label":"Use Multiscale","type":"toggle","default":true,"description":"Whether to use multi-scale generation. If True, the model will generate the video at a smaller scale first, then use the smaller video to guide the ge"},{"key":"audio_modality_scale","label":"Audio Modality Scale","type":"number","default":3,"min":0,"max":10,"description":"The modality scale for the audio. Controls the ratio between video and audio modalities."},{"key":"use_restart_sampling","label":"Use Restart Sampling","type":"toggle","default":false,"description":"Whether to use restart sampling. This will inject a small amount of noise during each denoising step, which can help improve the quality of the genera"},{"key":"end_image_strength","label":"End Image Strength","type":"number","default":1,"min":0,"max":1,"description":"The strength of the end image to use for the video generation."},{"key":"scheduler","label":"Scheduler","type":"select","default":"ltx2","options":["ltx2","linear_quadratic","beta"],"description":"The scheduler to use."},{"key":"distill_lora_first_pass_scale","label":"Distill Lora First Pass Scale","type":"number","default":0.2,"min":0,"max":1,"description":"The scale of the distill LoRA to use for the first pass. Set to 0 to disable."},{"key":"gradient_estimation_gamma","label":"Gradient Estimation Gamma","type":"number","default":2,"min":0,"max":10,"description":"The gamma of gradient estimation during denoising. Set to 0 to disable."},{"key":"distill_lora_second_pass_scale","label":"Distill Lora Second Pass Scale","type":"number","default":0.5,"min":0,"max":1,"description":"The scale of the distill LoRA to use for the second and subsequent passes."},{"key":"num_frames","label":"Num Frames","type":"number","default":121,"min":9,"max":481,"step":1,"description":"The number of frames to generate."},{"key":"acceleration","label":"Acceleration","type":"select","default":"regular","options":["none","regular","high","full"],"description":"The acceleration level to use."},{"key":"video_output_type","label":"Video Output Type","type":"select","default":"X264 (.mp4)","options":["X264 (.mp4)","VP9 (.webm)","PRORES4444 (.mov)","GIF (.gif)"],"description":"The output type of the generated video."},{"key":"video_rescaling_scale","label":"Video Rescaling Scale","type":"number","default":0.7,"min":0,"max":1,"description":"The rescaling scale for the video. Controls the ratio between classifier-free guidance and spatiotemporal guidance."},{"key":"video_quality","label":"Video Quality","type":"select","default":"high","options":["low","medium","high","maximum"],"description":"The quality of the generated video."},{"key":"camera_lora_scale","label":"Camera LoRA Scale","type":"number","default":1,"min":0,"max":1,"description":"The scale of the camera LoRA to use. This allows you to control the camera movement of the generated video more accurately than just prompting the mod"},{"key":"audio_rescaling_scale","label":"Audio Rescaling Scale","type":"number","default":0.7,"min":0,"max":1,"description":"The rescaling scale for the audio. Controls the ratio between classifier-free guidance and spatiotemporal guidance."},{"key":"video_modality_scale","label":"Video Modality Scale","type":"number","default":3,"min":0,"max":10,"description":"The modality scale for the video. Controls the ratio between video and audio modalities."},{"key":"video_write_mode","label":"Video Write Mode","type":"select","default":"balanced","options":["fast","balanced","small"],"description":"The write mode of the generated video."},{"key":"camera_lora","label":"Camera LoRA","type":"select","default":"none","options":["dolly_in","dolly_out","dolly_left","dolly_right","jib_up","jib_down","static","none"],"description":"The camera LoRA to use. This allows you to control the camera movement of the generated video more accurately than just prompting the model to move th"},{"key":"image_strength","label":"Image Strength","type":"number","default":1,"min":0,"max":1,"description":"The strength of the image to use for the video generation."},{"key":"video_stg_scale","label":"Video Stg Scale","type":"number","default":0,"min":0,"max":20,"description":"The Spatiotemporal Guidance (STG) scale for the video. Higher values result in more consistent and focused video content."},{"key":"video_cfg_scale","label":"Video Cfg Scale","type":"number","default":3,"min":1,"max":20,"description":"The Classifier-Free Guidance (CFG) scale for the video. Higher values result in more consistent and focused video content."},{"key":"num_inference_steps","label":"Num Inference Steps","type":"number","default":40,"min":8,"max":50,"step":1,"description":"The number of inference steps to use."},{"key":"negative_prompt","label":"Negative Prompt","type":"text","default":"news broadcast, 3d animation, computer graphics, pc game, console game, video game, cartoon, childish, watermark, logo, text, on screen text, subtitles, titles, signature, slowmo, static","description":"The negative prompt to generate the video from."},{"key":"enable_prompt_expansion","label":"Enable Prompt Expansion","type":"toggle","default":true,"description":"Whether to enable prompt expansion."},{"key":"enable_safety_checker","label":"Enable Safety Checker","type":"toggle","default":true,"description":"Run the safety checker to block unsafe content."},{"key":"fps","label":"Fps","type":"number","default":24,"min":1,"max":60,"description":"The frames per second of the generated video."},{"key":"interpolation_direction","label":"Interpolation Direction","type":"select","default":"forward","options":["forward","backward"],"description":"The direction to interpolate the image sequence in. 'Forward' goes from the start image to the end image, 'Backward' goes from the end image to the st"}]}}