PageSourceSearch

https://vibedream.ai/_next/static/chunks/ed17cc19c73aa398.js

js vibedream.ai collected 2026-10-05 13:26:43 UTC 111,706 bytes, 1 lines download raw bytes

1(globalThis.TURBOPACK||(globalThis.TURBOPACK=[])).push(["object"==typeof document?document.currentScript:void 0,26010,e=>{e.v(JSON.parse('{"version":"1.0.0","categories":{"text-to-image":{"id":"text-to-image","name":"Text to Image","description":"Create stunning images from text descriptions","icon":"image","models":[{"id":"z-image-turbo","name":"Z-Image Turbo","description":"6 billion parameter text-to-image model that generates photorealistic images in sub-second time","category":"text-to-image","wavespeedEndpoint":"wavespeed-ai/z-image/turbo","outputType":"image","estimatedTime":"1-3s","costPerRun":0.005,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/a169c110-955d-43c8-b95a-ebba714229a2-u1_40f5f03d-4b3f-4a55-add3-bded93384e74.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A futuristic cityscape at sunset with flying cars...","maxLength":2000},{"name":"image","type":"file","label":"Reference Image","description":"Optional reference image to guide the generation (for image-to-image)","required":false,"accept":"image/*"},{"name":"size","type":"select","label":"Size","description":"The size of the generated image in pixels","required":false,"default":"1024*1024","options":[{"value":"512*512","label":"512x512"},{"value":"768*768","label":"768x768"},{"value":"1024*1024","label":"1024x1024"},{"value":"1024*768","label":"1024x768 (Landscape)"},{"value":"768*1024","label":"768x1024 (Portrait)"},{"value":"1280*720","label":"1280x720 (HD)"},{"value":"720*1280","label":"720x1280 (HD Portrait)"}]},{"name":"strength","type":"number","label":"Strength","description":"Controls the strength of the transformation when using a reference image (0 to 1)","required":false,"default":0.6,"min":0,"max":1,"step":0.1},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"default":-1,"min":-1,"max":2147483647},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"jpeg","options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"},{"value":"webp","label":"WebP"}]}]},{"id":"seedream-v4.5","name":"Seedream V4.5","description":"ByteDance\'s next-gen text-to-image model optimized for typography - crisper text rendering, stronger prompt adherence, and up to 4K output","category":"text-to-image","wavespeedEndpoint":"bytedance/seedream-v4.5","outputType":"image","estimatedTime":"10-30s","costPerRun":0.04,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/3b23aca9-5d0b-4bd1-a4f5-b4e703bdd433-u1_11f963e4-41cb-4fea-b395-b3a193195c21.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A beautiful poster with bold typography...","maxLength":2000},{"name":"size","type":"select","label":"Size","description":"The size of the generated image in pixels (range 1024-4096)","required":false,"default":"2048*2048","options":[{"value":"1024*1024","label":"1024x1024"},{"value":"1536*1536","label":"1536x1536"},{"value":"2048*2048","label":"2048x2048"},{"value":"2048*1024","label":"2048x1024 (Landscape)"},{"value":"1024*2048","label":"1024x2048 (Portrait)"},{"value":"2048*1536","label":"2048x1536 (4:3)"},{"value":"1536*2048","label":"1536x2048 (3:4)"},{"value":"4096*4096","label":"4096x4096 (4K)"}]}]},{"id":"nano-banana-pro","name":"Nano Banana Pro","description":"Google\'s Gemini 3.0 Pro Image - cutting-edge text-to-image model enabling high-res 4K image generation","category":"text-to-image","wavespeedEndpoint":"google/nano-banana-pro/text-to-image","outputType":"image","estimatedTime":"10-30s","costPerRun":0.14,"pricing":{"paramBasedCost":{"params":["resolution"],"costs":{"1k":0.14,"2k":0.14,"4k":0.24}}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/b8e57921-a21d-4d1a-a25d-daafcc850fe1-u2_a80be9bd-97a2-4d95-9b8f-05efb48bfce5.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A stunning landscape with mountains...","maxLength":2000},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the generated image","required":false,"default":"1:1","options":[{"value":"1:1","label":"Square (1:1)"},{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"},{"value":"4:3","label":"Classic (4:3)"},{"value":"3:4","label":"Portrait Classic (3:4)"},{"value":"3:2","label":"Photo (3:2)"},{"value":"2:3","label":"Portrait Photo (2:3)"},{"value":"4:5","label":"Instagram (4:5)"},{"value":"5:4","label":"Landscape Instagram (5:4)"},{"value":"21:9","label":"Ultrawide (21:9)"}]},{"name":"resolution","type":"select","label":"Resolution","description":"The resolution of the output image","required":false,"default":"1k","options":[{"value":"1k","label":"1K"},{"value":"2k","label":"2K"},{"value":"4k","label":"4K"}]},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"png","options":[{"value":"png","label":"PNG"},{"value":"jpeg","label":"JPEG"}]}]},{"id":"nano-banana","name":"Nano Banana","description":"Google\'s cutting-edge text-to-image model that generates images from natural language prompts","category":"text-to-image","wavespeedEndpoint":"google/nano-banana/text-to-image","outputType":"image","estimatedTime":"5-15s","costPerRun":0.038,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/1e122a05-827a-4a20-adf3-7e6af4b6680f-u2_af63e952-1e4c-4185-a55d-4261962c4a95.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A vibrant digital artwork of a forest...","maxLength":2000},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the generated image","required":false,"default":"1:1","options":[{"value":"1:1","label":"Square (1:1)"},{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"},{"value":"4:3","label":"Classic (4:3)"},{"value":"3:4","label":"Portrait Classic (3:4)"},{"value":"3:2","label":"Photo (3:2)"},{"value":"2:3","label":"Portrait Photo (2:3)"},{"value":"4:5","label":"Instagram (4:5)"},{"value":"5:4","label":"Landscape Instagram (5:4)"},{"value":"21:9","label":"Ultrawide (21:9)"}]},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"png","options":[{"value":"png","label":"PNG"},{"value":"jpeg","label":"JPEG"}]}]},{"id":"flux-dev","name":"FLUX.1 Dev","description":"12-billion-
1parameter rectified-flow transformer for text-to-image generation with inpainting support","category":"text-to-image","wavespeedEndpoint":"wavespeed-ai/flux-dev","outputType":"image","estimatedTime":"5-15s","costPerRun":0.012,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/b4156692-e7c6-46ce-a601-59fe4a3f2b23-u2_c126b17b-d83d-4198-b8a6-08c1adb19ee8.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A beautiful sunset over mountains...","maxLength":2000},{"name":"image","type":"file","label":"Input Image","description":"Optional image for image-to-image generation","required":false,"accept":"image/*"},{"name":"mask_image","type":"file","label":"Mask Image","description":"Mask for inpainting - white areas will be regenerated, black areas preserved","required":false,"accept":"image/*"},{"name":"strength","type":"number","label":"Strength","description":"How much to transform the reference image (0 to 1)","required":false,"default":0.8,"min":0,"max":1,"step":0.1},{"name":"size","type":"select","label":"Size","description":"The size of the generated image in pixels","required":false,"default":"1024*1024","options":[{"value":"512*512","label":"512x512"},{"value":"768*768","label":"768x768"},{"value":"1024*1024","label":"1024x1024"},{"value":"1024*768","label":"1024x768 (Landscape)"},{"value":"768*1024","label":"768x1024 (Portrait)"},{"value":"1280*720","label":"1280x720 (HD)"},{"value":"720*1280","label":"720x1280 (HD Portrait)"}]},{"name":"num_inference_steps","type":"number","label":"Inference Steps","description":"Number of inference steps (higher = better quality, slower)","required":false,"default":28,"min":1,"max":50},{"name":"guidance_scale","type":"number","label":"Guidance Scale","description":"How closely to follow the prompt (higher = more faithful)","required":false,"default":3.5,"min":1,"max":20,"step":0.5},{"name":"num_images","type":"number","label":"Number of Images","description":"How many images to generate (1-4)","required":false,"default":1,"min":1,"max":4},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"default":-1,"min":-1,"max":2147483647},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"jpeg","options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"},{"value":"webp","label":"WebP"}]}]},{"id":"grok-imagine","name":"Grok Imagine","description":"X-AI\'s Grok Imagine model for precise text-to-image generation with AI-powered quality","category":"text-to-image","wavespeedEndpoint":"x-ai/grok-imagine-image/text-to-image","outputType":"image","estimatedTime":"5-15s","costPerRun":0.022,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/cb10b80e-9930-42f7-9c8f-e3b07a91e1ba-u2_62346208-2c54-4ba6-b86a-30ba61219f65.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A photorealistic portrait in dramatic lighting...","maxLength":2000},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Aspect ratio of the generated image","required":false,"default":"1:1","options":[{"value":"1:1","label":"Square (1:1)"},{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"},{"value":"4:3","label":"Classic (4:3)"},{"value":"3:4","label":"Portrait Classic (3:4)"},{"value":"3:2","label":"Photo (3:2)"},{"value":"2:3","label":"Portrait Photo (2:3)"},{"value":"2:1","label":"Wide (2:1)"},{"value":"1:2","label":"Tall (1:2)"},{"value":"20:9","label":"Ultrawide (20:9)"},{"value":"9:20","label":"Ultra Tall (9:20)"}]},{"name":"output_format","type":"select","label":"Output Format","description":"Output image format","required":false,"default":"jpeg","options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"}]}]},{"id":"flux-2-turbo","name":"FLUX.2 Turbo Text-to-Image","description":"FLUX 2 turbo from Black Forest Labs is the speed-optimized text-to-image model for real-time workflows. Generate photoreal images and clean typography with strong prompt adherence and consistent style—ideal for ads, posters, social posts, and rapid iteration. Built for low-latency, high-throughput use.","category":"text-to-image","wavespeedEndpoint":"wavespeed-ai/flux-2-turbo/text-to-image","outputType":"image","estimatedTime":"Low-latency","costPerRun":0.01,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of the desired image (the more specific the prompt, the more consistent the result).","required":true,"placeholder":"Describe the image you want to generate"},{"name":"size","type":"select","label":"Size","description":"Output image dimensions","required":false,"default":"1024x1024","options":[{"value":"1024x1024","label":"1024 x 1024 (Square)"},{"value":"1152x896","label":"1152 x 896 (Landscape)"},{"value":"896x1152","label":"896 x 1152 (Portrait)"},{"value":"1344x768","label":"1344 x 768 (Wide)"},{"value":"768x1344","label":"768 x 1344 (Tall)"},{"value":"1536x640","label":"1536 x 640 (Ultra Wide)"},{"value":"640x1536","label":"640 x 1536 (Ultra Tall)"}]},{"name":"seed","type":"number","label":"Seed","description":"Use -1 for random results, or set a fixed value for reproducible generations.","required":false,"min":-1,"max":1000000,"step":1,"placeholder":"-1 for random"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/273d8f5e-b765-4892-946f-fc1a0535df3e-u2_efe10e1d-2e78-41a3-8135-2712c18e9e81.png"},{"id":"seedream-v3-1","name":"ByteDance Seedream V3.1","description":"Seedream V3.1 by ByteDance is a text-to-image model with upgraded visuals, stronger style fidelity, and rich detail from text prompts. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"text-to-image","wavespeedEndpoint":"bytedance/seedream-v3.1","outputType":"image","estimatedTime":"5-15s","costPerRun":0.027,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text prompt describing the desired image","required":true,"placeholder":"Enter your image description here"},{"name":"size","type":"select","label":"Size","description":"Select the aspect ratio for the image","required":true,"options":[{"value":"1:1","label":"1:1"},{"value":"16:9","label":"16:9"},{"value":"9:16","label":"9:16"},{"value":"4:3","label":"4:3"},{"value":"3:4","label":"3:4"},{"value":"3:2","label":"3:2"},{"value":"2:3","label":"2:3"}]}
1,{"name":"width","type":"number","label":"Width","description":"Width of the output image in pixels","required":false,"min":512,"max":2048,"step":1,"placeholder":"Enter width (512-2048)"},{"name":"height","type":"number","label":"Height","description":"Height of the output image in pixels","required":false,"min":512,"max":2048,"step":1,"placeholder":"Enter height (512-2048)"},{"name":"seed","type":"number","label":"Seed","description":"Random seed for image generation","required":false,"max":10000,"step":1,"placeholder":"Enter a seed value (optional)"},{"name":"enable_prompt_expansion","type":"boolean","label":"Enable Prompt Expansion","description":"If set to true, the function will wait for the image to be generated and uploaded before returning the response.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/be027a99-62ca-4f54-a6ae-420ec39fbe9b-u1_29ddb041-41bb-48b2-8533-15a3ad5483ad.png"},{"id":"wan-2.6-t2i","name":"Alibaba WAN 2.6 Text-to-Image","description":"Generates high-quality images from natural-language prompts with strong prompt adherence and clean composition. It supports multiple aspect ratios and size control, seed-based reproducibility, and flexible styles (photorealistic to illustrative) for ads, product shots, and social visuals. Built for stable production use with a ready-to-use REST API, no cold starts, and predictable pricing.","category":"text-to-image","wavespeedEndpoint":"alibaba/wan-2.6/text-to-image","outputType":"image","estimatedTime":"5-15s","costPerRun":0.03,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of the image you want to generate.","required":true,"placeholder":"A modern tea shop interior, warm afternoon light, minimalist wood design, cinematic photography."},{"name":"width","type":"number","label":"Width","description":"Output width (within allowed limits).","required":true,"min":768,"max":1440,"step":1},{"name":"height","type":"number","label":"Height","description":"Output height (within allowed limits).","required":true,"min":768,"max":1440,"step":1},{"name":"enable_prompt_expansion","type":"boolean","label":"Enable Prompt Expansion","description":"Toggle prompt expansion to enrich short prompts.","required":false},{"name":"seed","type":"number","label":"Seed","description":"Set a fixed seed for more repeatable iterations (-1 for random).","required":false,"min":-1,"max":10000,"step":1}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/1f61186b-e866-4198-86c5-23e607585b97-u1_a7fbe3c3-19f4-4bcd-93e4-726b2eadfb89.png"},{"id":"gpt-image-1","name":"OpenAI GPT Image 1","description":"OpenAI GPT Image-1 generates images from text prompts, ideal for creating visual assets. It combines the reasoning power of GPT-4-Turbo with DALL·E-class visual synthesis, allowing for high-quality, creative, and context-aware image generation across various styles and purposes.","category":"text-to-image","wavespeedEndpoint":"openai/gpt-image-1/text-to-image","outputType":"image","estimatedTime":"5-15s","costPerRun":0.063,"pricing":{"paramBasedCost":{"params":["quality","size"],"costs":{"low|1024*1024":0.011,"low|1024*1536":0.016,"low|1536*1024":0.016,"medium|1024*1024":0.042,"medium|1024*1536":0.063,"medium|1536*1024":0.063,"high|1024*1024":0.167,"high|1024*1536":0.25,"high|1536*1024":0.25}}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Required text description of the desired image.","required":true,"placeholder":"Describe the image you want to create..."},{"name":"size","type":"select","label":"Size","description":"Select the dimensions of the image.","required":false,"options":[{"value":"1024*1024","label":"1024×1024"},{"value":"1024*1536","label":"1024×1536"},{"value":"1536*1024","label":"1536×1024"}]},{"name":"quality","type":"select","label":"Quality","description":"Choose the desired quality of the image output.","required":false,"options":[{"value":"low","label":"Low"},{"value":"medium","label":"Medium"},{"value":"high","label":"High"}
1]},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/ff954613-cc25-47ee-8ef4-b79d4af62eed-u2_2ac8843a-0f16-4dfb-a791-9da727646ae3.png"},{"id":"dall-e-3","name":"OpenAI DALL-E 3 Text-To-Image Generation API","description":"OpenAI DALL·E 3 is OpenAI\'s most advanced text-to-image system, capable of generating highly detailed, realistic, and creative visuals directly from natural language descriptions. It builds upon OpenAI\'s extensive world knowledge and artistic training to create images that are accurate, expressive, and aligned with your intent.","category":"text-to-image","wavespeedEndpoint":"openai/dall-e-3","outputType":"image","estimatedTime":"5-15s","costPerRun":0.08,"pricing":{"paramBasedCost":{"params":["quality","size"],"costs":{"standard|1024x1024":0.04,"standard|1024x1792":0.08,"standard|1792x1024":0.08,"hd|1024x1024":0.08,"hd|1024x1792":0.12,"hd|1792x1024":0.12}}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"The text description based on which the image will be generated.","required":true,"placeholder":"Enter your prompt here"},{"name":"size","type":"select","label":"Image Size","description":"Select the desired resolution for the generated image.","required":true,"options":[{"value":"1024x1024","label":"1024x1024"},{"value":"1024x1792","label":"1024x1792"},{"value":"1792x1024","label":"1792x1024"}]},{"name":"quality","type":"select","label":"Image Quality","description":"Select the quality of the resulting image.","required":true,"options":[{"value":"standard","label":"Standard"},{"value":"hd","label":"HD"}]},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/121c01d6-2c37-47ad-aaae-de0ac5e3d284-u1_8d972084-9e04-4c81-8cf1-40fda6b049ee.png"}]},"image-to-image":{"id":"image-to-image","name":"Image to Image","description":"Edit, enhance, and transform images with AI","icon":"image","models":[{"id":"z-image-turbo-i2i","name":"Z-Image Turbo I2I","description":"6 billion parameter model that enhances image quality and applies style transformations in sub-second time","category":"image-to-image","wavespeedEndpoint":"wavespeed-ai/z-image-turbo/image-to-image","outputType":"image","estimatedTime":"1-3s","costPerRun":0.005,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/757a0651-4e57-42d5-b2c6-ba77a97ce37a-u2_79331ca1-7b4b-4245-802c-89e35a4ca09f.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"Reference image to transform or enhance","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the desired transformation or enhancement","required":true,"placeholder":"Enhance this image with vibrant colors and sharp details...","maxLength":2000},{"name":"size","type":"select","label":"Size","description":"The size of the output image in pixels","required":false,"default":"1024*1024","options":[{"value":"512*512","label":"512x512"},{"value":"768*768","label":"768x768"},{"value":"1024*1024","label":"1024x1024"},{"value":"1024*768","label":"1024x768 (Landscape)"},{"value":"768*1024","label":"768x1024 (Portrait)"}]},{"name":"strength","type":"number","label":"Strength","description":"Controls the strength of the transformation (0 = subtle, 1 = dramatic)","required":false,"default":0.6,"min":0,"max":1,"step":0.1},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"default":-1,"min":-1,"max":2147483647},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"jpeg","options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"},{"value":"webp","label":"WebP"}]}]},{"id":"qwen-image-edit","name":"Qwen Image Edit 2511","description":"Advanced image editor with strong edit consistency, multi-person identity preservation, built-in LoRA styles, and geometric reasoning","category":"image-to-image","wavespeedEndpoint":"wavespeed-ai/qwen-image/edit-2511","outputType":"image","estimatedTime":"5-15s","costPerRun":0.03,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/0b5dc15e-9430-4c70-a0aa-77269e725dae-u2_0f119699-aff3-4bc0-823c-669e09ea2f5c.png","parameters":[{"name":"images","type":"file","label":"Input Image","description":"The image to edit (up to 3 reference images supported via API)","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the edit you want to apply to the image","required":true,"placeholder":"Change the background to a 
1sunset beach scene...","maxLength":2000},{"name":"size","type":"select","label":"Size","description":"The size of the output image in pixels","required":false,"default":"1024*1024","options":[{"value":"512*512","label":"512x512"},{"value":"768*768","label":"768x768"},{"value":"1024*1024","label":"1024x1024"},{"value":"1024*768","label":"1024x768 (Landscape)"},{"value":"768*1024","label":"768x1024 (Portrait)"}]},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"default":-1,"min":-1,"max":2147483647},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"jpeg","options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"},{"value":"webp","label":"WebP"}]}]},{"id":"gpt-image-edit","name":"GPT Image 1.5 Edit","description":"OpenAI\'s image editor for precise natural-language edits - add/remove objects, swap backgrounds, retouch, adjust colors, and edit text","category":"image-to-image","wavespeedEndpoint":"openai/gpt-image-1.5/edit","outputType":"image","estimatedTime":"5-20s","costPerRun":0.02,"pricing":{"paramBasedCost":{"params":["quality","size"],"costs":{"low|auto":0.009,"low|1024*1024":0.009,"low|1024*1536":0.034,"low|1536*1024":0.013,"medium|auto":0.034,"medium|1024*1024":0.034,"medium|1024*1536":0.051,"medium|1536*1024":0.051,"high|auto":0.133,"high|1024*1024":0.133,"high|1024*1536":0.2,"high|1536*1024":0.2}}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/ee579264-b657-4a4b-8eeb-3fe88fb6dc65-u2_fd6dc2e1-3bb9-45ac-9272-ef6942a9dcf6.png","parameters":[{"name":"images","type":"file","label":"Input Image","description":"The image to edit","required":false,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the edit you want (e.g. \'remove the person in the background\')","required":true,"placeholder":"Remove the background and replace with a studio setting...","maxLength":2000},{"name":"size","type":"select","label":"Size","description":"Output image size","required":false,"default":"1024*1024","options":[{"value":"auto","label":"Auto"},{"value":"1024*1024","label":"1024x1024"},{"value":"1024*1536","label":"1024x1536 (Portrait)"},{"value":"1536*1024","label":"1536x1024 (Landscape)"}]},{"name":"quality","type":"select","label":"Quality","description":"Image quality level (higher quality costs more)","required":false,"default":"medium","options":[{"value":"low","label":"Low"},{"value":"medium","label":"Medium"},{"value":"high","label":"High"}]},{"name":"input_fidelity","type":"select","label":"Input Fidelity","description":"How closely to preserve details from input (faces, logos)","required":false,"default":"high","options":[{"value":"low","label":"Low"},{"value":"high","label":"High"}]},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"jpeg","options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"}]}]},{"id":"grok-imagine-edit","name":"Grok Imagine Edit","description":"X-AI\'s Grok Imagine model for precise image editing - transform and modify images using text prompts","category":"image-to-image","wavespeedEndpoint":"x-ai/grok-imagine-image/edit","outputType":"image","estimatedTime":"5-15s","costPerRun":0.022,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/21439dd1-5f69-4941-a796-65587e72bfeb-u1_8c7704e3-7898-4d43-a1b2-732aaadffc42.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"The source image to edit","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the edit you want to apply","required":true,"placeholder":"Make the sky a dramatic sunset with warm orange tones...","maxLength":2000},{"name":"output_format","type":"select","label":"Output Format","description":"Output image format","required":false,"default":"jpeg","options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"}]}]},{"id":"nano-banana-edit","name":"Nano Banana Edit","description":"Google\'s image editing model for precise inpainting, outpainting, background replacement, and stylized transformations","category":"image-to-image","wavespeedEndpoint":"google/nano-banana/edit","outputType":"image","estimatedTime":"5-15s","costPerRun":0.038,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/e119c0f1-27ae-4666-9557-70ce10b06f86-u2_2fe1dc6c-a534-46a8-ba77-4a7087aedb5d.png","parameters":[{"name":"images","type":"file","label":"Input Image","description":"The image to edit (up to 10 images supported via API)","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the edit or transformation","required":true,"placeholder":"Replace the background with a tropical beach...","maxLength":2000},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the output image","required":false,"default":"1:1","options":[{"value":"1:1","label":"Square (1:1)"},{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"},{"value":"4:3","label":"Classic (4:3)"},{"value":"3:4","label":"Portrait Classic (3:4)"},{"value":"3:2","label":"Photo (3:2)"},{"value":"2:3","label":"Portrait Photo (2:3)"},{"value":"4:5","label":"Instagram (4:5)"},{"value":"5:4","label":"Landscape Instagram (5:4)"},{"value":"21:9","label":"Ultrawide (21:9)"}]},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"png","options":[{"value":"png","label":"PNG"},{"value":"jpeg","label":"JPEG"}]}]},{"id":"nano-banana-pro-edit","name":"Nano Banana Pro Edit","description":"Google Gemini 3.0 Pro image editor with 4K-capable output for high-resolution image editing and manipulation","category":"image-to-image","wavespeedEndpoint":"google/nano-banana-pro/edit","outputType":"image","estimatedTime":"10-30s","costPerRun":0.14,"pricing":{"paramBasedCost":{"params":["resolution"],"costs":{"1k":0.14,"2k":0.14,"4k":0.24}}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/d6acae1a-4834-4c2d-81e2-37f7ba446d27-u2_f9e16ab4-1c9a-456b-a023-bd3ae9e64f35.png","parameters":[{"name":"images","type":"file","label":"Input Image","description":"The image to edit (up to 14 images supported via API)","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the edit or transformation","required":true,"placeholder":"Upscale and enhance the details of this photo...","maxLength":2000},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the output image","required":false,"default":"1:1","options":[{"value":"1:1","label":"Square (1:1)"},{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"},{"value":"4:3","label":"Classic (4:3)"},{"value":"3:4","label":"Portrait Classic (3:4)"},{"value":"3:2","label":"Photo (3:2)"},{"value":"2:3","label":"Portrait Photo (2:3)"},{"value":"4:5","label":"Instagram (4:5)"},{"value":"5:4","label":"Landscape Instagram (5:4)"},{"value":"21:9","label":"Ultrawide (21:9)"}]},{"name":"resolution","type":"select","label":"Resolution","description":"Output resolution (4K costs more: ~$0.24/run)","required":false,"default":"1k","options":[{"value":"1k","label":"1K"},{"value":"2k","label":"2K"},{"value":"4k","label":"4K"}]},{"name":"output_format","type":"select","label":"Output Format","description":"The format of the output image","required":false,"default":"png","options":[{"value":"png","label":"PNG"},{"value":"jpeg","label":"JPEG"}]}]},{"id":"flux-2-klein-9b-edit","name":"FLUX.2 Klein 9B Edit","description":"FLUX.2 Klein 9B Edit is a high-quality image editing model with 9B parameters, offering precise modifications using natural language instructions. Ready-to-use REST inference API, best performance, no cold starts, affordable pricing.","category":"image-to-image","wavespeedEndpoint":"wavespeed-ai/flux-2-klein-9b/edit","outputType":"image","estimatedTime":"
160-120s","costPerRun":0.016,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of the desired edit.","required":true,"placeholder":"Describe your edit here..."},{"name":"images","type":"file","label":"Images","description":"Source images to edit (can add multiple).","required":true,"accept":"image/*"},{"name":"size","type":"number","label":"Size","description":"Output dimensions (empty = same as input image).","required":false,"step":1,"placeholder":"Enter output size"},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random).","required":false,"min":-1,"max":10000,"step":1,"placeholder":"Enter seed value"},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/e98629eb-6268-48a8-88f4-4effe4062f17-u1_14f38a11-9539-43b3-9d52-e37c1b23f659.png"},{"id":"seedream-v4-edit","name":"Seedream V4 Edit","description":"Seedream V4 Edit is ByteDance\'s state-of-the-art image editing model that outperforms Nano Banana in fidelity and edit quality. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"image-to-image","wavespeedEndpoint":"bytedance/seedream-v4/edit","outputType":"image","estimatedTime":"Processing time not specified","costPerRun":0.027,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"Clear textual prompt for image details","required":true,"placeholder":"Enter your prompt here"},{"name":"images","type":"file","label":"Upload Images","description":"Upload up to 10 base images for editing","required":true,"accept":"image/*"},{"name":"size","type":"select","label":"Size","description":"Specify the size of the output image","required":false,"options":[{"value":"4096x4096","label":"4096 x 4096"},{"value":"1920x1080","label":"1920 x 1080"},{"value":"1280x720","label":"1280 x 720"}]},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"Wait for the result to be generated before returning","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"Encode output into a BASE64 string instead of a URL","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/11fc05c0-511b-4c87-b7d2-15520eaf2884-u1_e9a7f0a6-7955-4374-acdb-be5021145d4d.png"},{"id":"image-edit","name":"Alibaba WAN 2.6 Image-Edit","description":"Alibaba WAN 2.6 Image-Edit turns prompts into precise photo edits—adjusting color and lighting, restyling aesthetics, replacing backgrounds, removing objects, and refining details while preserving subject identity. Built for stable, repeatable image-to-image pipelines. Ready-to-use REST API, best performance, no cold starts, affordable pricing.","category":"image-to-image","wavespeedEndpoint":"alibaba/wan-2.6/image-edit","outputType":"image","estimatedTime":"5-15s","costPerRun":0.035,"parameters":[{"name":"prompt","type":"textarea","label":"Image Edit Instruction","description":"The edit instruction describing what to change and what to keep (e.g., \\"change the jacket to leather, keep face and pose unchanged\\").","required":true,"placeholder":"Describe the edit..."},{"name":"images","type":"file","label":"Input Images","description":"One or more input images to edit (uploaded files or public URLs).","required":true,"placeholder":"Drag and drop or click to upload images","accept":"image/*"},{"name":"seed","type":"number","label":"Seed","description":"Optional integer for reproducibility; use a fixed seed to iterate with smaller prompt changes.","required":false,"max":10000,"step":1,"placeholder":"Enter seed number"},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Optional list of things you don’t want (e.g., \\"text, watermark, extra fingers, blurry face\\").","required":false,"placeholder":"List unwanted elements..."}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/421c9fe4-0002-4e85-8e0e-0f53583ba38a-u2_2908618a-4f01-474f-a7a7-3dec3e0570c7.png"},{"id":"seedream-4.5-edit","name":"ByteDance Seedream 4.5 Edit","description":"Seedream 4.5 Edit preserves facial features, lighting, and color tone from reference images, delivering professional, high-fidelity edits up to 4K with strong prompt adherence. Ready-to-use REST inference API, best performance, no cold starts, affordable pricing.","category":"image-to-image","wavespeedEndpoint":"bytedance/seedream-v4.5/edit","outputType":"image","estimatedTime":"5-15s","costPerRun":0.04,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"Describe what should change and what must stay the same.","required":true,"placeholder":"e.g. Change jacket to red leather, keep pose and background, cinematic lighting."},{"name":"images","type":"file","label":"Image Upload","description":"Upload 1–10 images that will be edited.","required":true,"accept":"image/*"},{"name":"size","type":"select","label":"Size","description":"Choose width and height for output image (if left empty, default resolution is used).","required":false,"options":[{"value":"1024x1024","label":"1024x1024"},{"value":"2048x2048","label":"2048x2048"},{"value":"4096x4096","label":"4096x4096"}]},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL.","required":false},{"name":"enable_safety_checker","type":"boolean","label":"Enable Safety Checker","description":"Toggle to enable safety checks during processing.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/30cedbd5-6c6f-47a5-b9fe-5ab1861d84fe-u1_f0dd3e76-24b2-4db7-ab46-f0bc2c7eacbb.png"},{"id":"flux-2-pro-edit","name":"FLUX.2 [pro] Edit","description":"FLUX.2 [pro] Edit delivers production-grade image-to-image editing from Black Forest Labs—apply natural-language instructions and exact hex color control for consistent, studio-quality results. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"image-to-image","wavespeedEndpoint":"wavespeed-ai/flux-2-pro/edit","outputType":"image","estimatedTime":"Variable, depending on task","costPerRun":0.06,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"Detailed instructions on how to edit the image.","required":true,"placeholder":"Enter your prompt here"}
1,{"name":"images","type":"file","label":"Images","description":"Upload the image(s) to be edited. You can drag and drop a file or click to upload.","required":true,"accept":"image/*"},{"name":"size","type":"select","label":"Image Size","description":"Select the desired output size for the edited image.","required":false,"options":[{"value":"small","label":"Small"},{"value":"medium","label":"Medium"},{"value":"large","label":"Large"}]},{"name":"seed","type":"number","label":"Random Seed","description":"Set a seed for the random number generator to ensure reproducibility.","required":false,"max":10000,"step":1,"placeholder":"Enter a seed value"},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL.","required":false},{"name":"enable_safety_checker","type":"boolean","label":"Enable Safety Checker","description":"Enable additional checks to ensure content safety and compliance.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/edee1b53-2764-4089-9020-0576aca27c0e-u1_a40d798c-4221-4218-b51e-d29104926489.png"},{"id":"edit-sequential","name":"ByteDance Seedream 4.5 Edit Sequential","description":"Seedream 4.5 Edit Sequential performs multi-image editing while locking character and object identity across shots. It detects main subjects, preserves continuity, and applies controlled edits with up to 4K output. Ready-to-use REST inference API, best performance, no cold starts, affordable pricing.","category":"image-to-image","wavespeedEndpoint":"bytedance/seedream-v4.5/edit-sequential","outputType":"image","estimatedTime":"5-15s","costPerRun":0.08,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"The edit prompt describing the shared change you want across the whole set.","required":true,"placeholder":"e.g. Change outfit to a black suit, add soft studio lighting, keep poses the same."},{"name":"images","type":"file","label":"Images","description":"Upload source images to edit sequentially, containing the same main subject or product.","required":true,"accept":"image/*"},{"name":"size","type":"select","label":"Size","description":"Select the target resolution. Supports sizes up to 4096 × 4096 for maximum detail.","required":false,"options":[{"value":"4096x4096","label":"4096 × 4096"},{"value":"2048x2048","label":"2048 × 2048"},{"value":"1024x1024","label":"1024 × 1024"}]},{"name":"max_images","type":"number","label":"Max Images","description":"Specify how many edited outputs you want the model to generate from your input set.","required":true,"min":1,"max":10,"step":1,"placeholder":"1-10"},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL. This property is only available through the API.","required":false},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response.","required":false},{"name":"Enable Safety Checker","type":"boolean","label":"Enable Safety Checker","description":"Enable safety precautions for the outputs.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/11a89b3f-ecf8-4961-a828-ff3b09b747c3-u2_e0a70325-57d2-4271-b09b-0261de058831.png"},{"id":"edit-ultra","name":"Google Nano Banana Pro Edit Ultra","description":"Nano Banana Pro Edit (Gemini 3.0 Pro Image) is Google’s advanced AI-powered image editing and generation model, designed to make visual transformation as intuitive as describing it in words. Built on Google’s cutting-edge computer vision and generative research, it combines precision, flexibility, and semantic awareness for professional-grade editing.","category":"image-to-image","wavespeedEndpoint":"google/nano-banana-pro/edit-ultra","outputType":"image","estimatedTime":"5-15s","costPerRun":0.15,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"Input text prompt for image editing.","required":true,"placeholder":"Describe the desired image changes."},{"name":"images","type":"file","label":"Image Upload","description":"Upload the existing image to be edited.","required":true,"accept":"image/*"},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Choose the aspect ratio for the output image.","required":false,"options":[{"value":"1:1","label":"1:1"},{"value":"4:3","label":"4:3"},{"value":"16:9","label":"16:9"},{"value":"21:9","label":"21:9"}]}
1,{"name":"resolution","type":"select","label":"Resolution","description":"Select the resolution for the output image.","required":false,"options":[{"value":"4k","label":"4k"},{"value":"8k","label":"8k"}]},{"name":"output_format","type":"select","label":"Output Format","description":"Choose the output format of the image.","required":false,"options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"}]},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"Wait for the result to be generated before returning the response.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"Encode the output into a BASE64 string instead of a URL.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/44ed1a98-8a45-4e24-8b8b-56d1255b2932-u1_481168b8-9b24-49a9-972f-f783e7565ad1.png"}]},"text-to-video":{"id":"text-to-video","name":"Text to Video","description":"Generate videos from text descriptions","icon":"video","models":[{"id":"minimax-video-01","name":"MiniMax Video 01","description":"High-compression text-to-video with cinematic quality, smooth motion, and prompt-faithful generation","category":"text-to-video","wavespeedEndpoint":"minimax/video-01","outputType":"video","estimatedTime":"60-120s","costPerRun":0.5,"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/f67655c5-7a7e-4968-88b4-4b2aa57db292-u1_bc8f6216-3570-4b3d-9957-293e070cf2d6.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the video you want to generate","required":true,"placeholder":"A serene lake at sunset with gentle waves...","maxLength":2000},{"name":"prompt_optimizer","type":"boolean","label":"Prompt Optimizer","description":"Automatically enhance your prompt for better results","required":false,"default":true}]},{"id":"kling-v2.6-pro-t2v","name":"Kling 2.6 Pro","description":"Top-tier text-to-video with smooth motion, cinematic visuals, strong prompt adherence, and optional native audio","category":"text-to-video","wavespeedEndpoint":"kwaivgi/kling-v2.6-pro/text-to-vi
1deo","outputType":"video","estimatedTime":"60-180s","costPerRun":0.35,"pricing":{"durationPricing":{"5":0.35,"10":0.7},"costMultipliers":{"sound":2}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/1dc3a309-f777-4439-8929-08ad849af949-u1_fb072f7b-b6b9-44f4-be3d-8d3d94a0dd7c.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A cinematic aerial shot of a city at night...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Describe what you don\'t want in the output","required":false,"placeholder":"blurry, low quality, distorted...","maxLength":1000},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"5","options":[{"value":"5","label":"5 seconds"},{"value":"10","label":"10 seconds"}]},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the generated video","required":false,"default":"16:9","options":[{"value":"1:1","label":"Square (1:1)"},{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"}]},{"name":"sound","type":"boolean","label":"Generate Audio","description":"Generate native audio with the video (doubles cost)","required":false,"default":false},{"name":"cfg_scale","type":"number","label":"CFG Scale","description":"Prompt adherence strength (higher = stricter prompt following)","required":false,"default":0.5,"min":0,"max":1,"step":0.1}]},{"id":"grok-imagine-video","name":"Grok Imagine Video","description":"X-AI\'s Grok Imagine Video generates high-quality videos from text with customizable duration, aspect ratio, and resolution","category":"text-to-video","wavespeedEndpoint":"x-ai/grok-imagine-video/text-to-vi
1deo","outputType":"video","estimatedTime":"30-120s","costPerRun":0.275,"pricing":{"costPerSecond":0.055},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/8cc271c5-caea-47ba-ac32-a6551c95fe9e-u1_3142e8df-635b-4884-a183-a42a33b56ea8.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of the desired video","required":true,"placeholder":"A slow-motion shot of a hummingbird in flight...","maxLength":2000},{"name":"duration","type":"number","label":"Duration","description":"Video duration in seconds (1-15)","required":false,"default":6,"min":1,"max":15},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Aspect ratio of the generated video","required":false,"default":"16:9","options":[{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"},{"value":"4:3","label":"Classic (4:3)"},{"value":"3:4","label":"Portrait Classic (3:4)"},{"value":"3:2","label":"Photo (3:2)"},{"value":"2:3","label":"Portrait Photo (2:3)"},{"value":"1:1","label":"Square (1:1)"}]},{"name":"resolution","type":"select","label":"Resolution","description":"Resolution of the output video","required":false,"default":"720p","options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"}]}]},{"id":"wan-2.6-t2v","name":"WAN 2.6","description":"Alibaba\'s text-to-video model with coherent cinematic clips, stable motion, and strong instruction-following for ads and explainers","category":"text-to-video","wavespeedEndpoint":"alibaba/wan-2.6/text-to-video","outputType":"video","estimatedTime":"60-180s","costPerRun":0.5,"pricing":{"perSecondByParam"
1:{"param":"size","rates":{"1280*720":0.1,"720*1280":0.1,"1920*1080":0.15,"1080*1920":0.15}}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/91d291c3-eb1b-4868-85e2-f0590c2144e5-u1_4c99bea2-30fe-460f-b6bb-670fff577a8a.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A product showcase video with smooth camera movements...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Describe what you don\'t want in the output","required":false,"placeholder":"blurry, low quality, distorted...","maxLength":1000},{"name":"size","type":"select","label":"Resolution","description":"The size of the generated video","required":false,"default":"1280*720","options":[{"value":"1280*720","label":"720p Landscape"},{"value":"720*1280","label":"720p Portrait"},{"value":"1920*1080","label":"1080p Landscape"},{"value":"1080*1920","label":"1080p Portrait"}]},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"5","options":[{"value":"5","label":"5 seconds"},{"value":"10","label":"10 seconds"},{"value":"15","label":"15 seconds"}]},{"name":"shot_type","type":"select","label":"Shot Type","description":"Single continuous shot or multi-shot video","required":false,"default":"single","options":[{"value":"single","label":"Single Shot"},{"value":"multi","label":"Multi Shot"}]},{"name":"enable_prompt_expansion","type":"boolean","label":"Prompt Optimizer","description":"Automatically enhance your prompt for better results","required":false,"default":false},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"default":-1,"min":-1,"max":2147483647}]},{"id":"sora-2-t2v","name":"OpenAI Sora 2","description":"State-of-the-art text-to-video with realistic visuals, accurate physics, synchronized audio, and strong steerability","category":"text-to-video","wavespeedEndpoint":"openai/sora-2/text-to-video","outputType":"video","estimatedTime":"60-180s","costPerRun":0.4,"pricing":{"durationPricing":{"4":0.4,"8":0.8,"12":1.2}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/7c6021a6-a578-423f-aec8-f348a1c8a9d9-u1_b8b8d2a6-e23b-4bba-ada3-8c285c656ab4.png","parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"A golden retriever running through autumn leaves in slow motion...","maxLength":2000},{"name":"size","type":"select","label":"Orientation","description":"The orientation of the generated video","required":false,"default":"1280*720","options":[{"value":"1280*720","label":"Landscape (1280x720)"},{"value":"720*1280","label":"Portrait (720x1280)"}]},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"4","options":[{"value":"4","label":"4 seconds"},{"value":"8","label":"8 seconds"},{"value":"12","label":"12 seconds"}]}]},{"id":"veo3","name":"Google Veo3","description":"Google Veo3 is Google\'s flagship text-to-video model with built-in audio, producing synchronized video and sound from text prompts. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"text-to-video","wavespeedEndpoint":"google/veo3","outputType":"video","estimatedTime":"60-120s","costPerRun":1.2,"pricing":{"costMultipliers":{"generate_audio":2.6667}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Description of the scene you want to create.","required":true,"placeholder":"A close-up shot of melting icicles..."},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the output video.","required":false,"options":[{"value":"16:9","label":"16:9"},{"value":"16:10","label":"16:10"}]},{"name":"duration","type":"number","label":"Duration","description":"Length of the video in seconds.","required":true,"min":1,"max":8,"step":1,"placeholder":"8"},{"name":"resolution","type":"select","label":"Resolution","description":"Resolution of the output video.","required":false,"options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"generate_audio","type":"boolean","label":"Generate Audio","de
1scription":"Whether to generate audio.","required":false},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Elements to avoid in the video.","required":false,"placeholder":"Avoid bright colors..."},{"name":"seed","type":"number","label":"Seed","description":"Random seed for generation.","required":false,"max":10000,"step":1,"placeholder":"1234"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/ed78f7a6-b01c-4b5e-9f9d-fd520030f160-u2_d52bd780-b4de-45cc-98d5-256613c4454b.png"},{"id":"seedance-v1-lite-t2v-720p","name":"ByteDance Seedance V1 Lite","description":"Seedance V1 Lite produces coherent multi-shot 720p videos with smooth motion and accurate following of detailed text prompts. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"text-to-video","wavespeedEndpoint":"bytedance/seedance-v1-lite-t2v-720p","outputType":"video","estimatedTime":"5-15s","costPerRun":0.16,"pricing":{"costPerSecond":0.032},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe what happens in the video (subject, action, scene, mood)","required":true},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Output aspect ratio (e.g., 16:9, 9:16, 1:1)","required":false,"options":[{"value":"16:9","label":"16:9"},{"value":"9:16","label":"9:16"},{"value":"1:1","label":"1:1"}]},{"name":"duration","type":"number","label":"Duration","description":"Video length in seconds","required":true,"min":5,"max":20,"step":1},{"name":"camera_fixed","type":"boolean","label":"Camera Fixed","description":"Whether to fix the camera position.","required":false},{"name":"seed","type":"number","label":"Random Seed","description":"Random seed (-1 for random; fixed value for reproducible results)","required":false,"min":-1,"max":99999,"step":1}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/b51a1221-4b04-4a3b-a00b-5e2875320fc3-u2_d54a9860-e810-46ef-9e22-f6a71942ecde.png"},{"id":"seedance-1.5-pro-t2v","name":"Seedance 1.5 Pro Text to Video","description":"Seedance 1.5 Pro generates cinematic, live-action–leaning clips from text with strong prompt adherence, expressive motion, and stable aesthetics. It supports 4–12s duration control, multiple aspect ratios, and reproducible generation via seeds.","category":"text-to-video","wavespeedEndpoint":"bytedance/seedance-v1.5-pro/text-to-vi
1deo","outputType":"video","estimatedTime":"60-120s","costPerRun":0.26,"pricing":{"perSecondByParam":{"param":"resolution","rates":{"480p":0.012,"720p":0.026,"1080p":0.052}},"costMultipliers":{"generate_audio":2}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of the scene, style, subject actions, camera motion, and overall mood.","required":true},{"name":"duration","type":"number","label":"Duration","description":"Integer seconds in [4, 12]. Use -1 for Smart Duration (model decides within [4, 12]).","required":false,"min":4,"max":12,"step":1},{"name":"resolution","type":"select","label":"Resolution","description":"Select the video resolution.","required":false,"options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"}]},{"name":"fixed_camera","type":"boolean","label":"Fixed Camera","description":"If true, the camera remains fixed; if false, camera motion is driven by the prompt.","required":false},{"name":"seed","type":"number","label":"Seed","description":"Controls randomness; the same seed tends to produce more similar outputs. Set -1 to cancel random seed.","required":false,"min":-1,"max":2147483647,"step":1},{"name":"fps","type":"number","label":"Frames Per Second","description":"Fixed at 24.","required":false,"min":24,"max":24,"step":1},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Choose the aspect ratio for the video.","required":false,"options":[{"value":"adaptive","label":"Adaptive"},{"value":"16:9","label":"16:9"},{"value":"9:16","label":"9:16"},{"value":"4:3","label":"4:3"},{"value":"3:4","label":"3:4"},{"value":"1:1","label":"1:1"},{"value":"21:9","label":"21:9"}]}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/b8d86761-ad2f-40f0-b032-e7ecb446c7f7-u1_6e33b0de-531e-4c22-80fa-f961c2f32af6.png"},{"id":"t2v-720p","name":"WAN 2.1 T2V 720P","description":"WAN 2.1 T2V 720P offers text-to-video 720p generation from prompts, enabling unlimited AI video creation for social and marketing. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"text-to-video","wavespeedEndpoint":"wavespeed-ai/wan-2.1/t2v-720p","outputType":"video","estimatedTime":"5-15s","costPerRun":0.3,"pricing":{"durationPricing":{"5":0.3,"10":0.6}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of the video you want to generate.","required":true,"placeholder":"Describe the scene, action, and style you want."},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Elements to avoid in the generated video.","required":false,"placeholder":"Specify unwanted elements."},{"name":"size","type":"select","label":"Output Resolution","description":"Output resolution (default: 1280×720).","required":false,"options":[{"value":"1280x720","label":"1280x720"},{"value":"720x1280","label":"720x1280"}]},{"name":"num_inference_steps","type":"number","label":"Number of Inference Steps","description":"Quality/speed trade-off (default: 30).","required":false,"min":1,"max":100,"step":1,"placeholder":"30"},{"name":"duration","type":"select","label":"Video Duration","description":"Video length in seconds: 5 or 10 (default: 5).","required":false,"options":[{"value":"5","label":"5 seconds"},{"value":"10","label":"10 seconds"}]},{"name":"guidance_scale","type":"number","label":"Guidance Scale","description":"Prompt adherence strength (default: 5).","required":false,"min":1,"max":10,"step":1,"placeholder":"5"},{"name":"flow_shift","type":"number","label":"Flow Shift","description":"Motion intensity control (default: 5).","required":false,"min":1,"max":10,"step":1,"placeholder":"5"},{"name":"seed","type":"number","label":"Seed","description":"Set for reproducibility; -1 for random.","required":false,"min":-1,"max":100,"step":1,"placeholder":"-1"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/09fe3d8a-b728-47db-858d-429c610880c6-u2_03a38cc1-489f-4ddc-b057-f177cf8d2870.png"},{"id":"vidu-q3-t2v","name":"Vidu Q3 Text-To-Video","description":"Vidu Q3 Text-to-Video turns text prompts into high-quality videos with exceptional visual fidelity and diverse motion. It features multiple styles, resolutions up to 1080p, flexible duration, audio generation, and motion control.","category":"text-to-vi
1deo","wavespeedEndpoint":"vidu/q3/text-to-video","outputType":"video","estimatedTime":"5-15s","costPerRun":0.375,"pricing":{"perSecondByParam":{"param":"resolution","rates":{"540p":0.07,"720p":0.15,"1080p":0.16}}},"parameters":[{"name":"prompt","type":"textarea","label":"Text Description","description":"Text description of the video scene and action","required":true,"placeholder":"Describe your scene here"},{"name":"style","type":"select","label":"Visual Style","description":"Choose between general realistic style or anime aesthetic.","required":false,"options":[{"value":"general","label":"General"},{"value":"anime","label":"Anime"}]},{"name":"resolution","type":"select","label":"Output Quality","description":"Output quality for the video.","required":false,"options":[{"value":"540p","label":"540p"},{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"duration","type":"number","label":"Video Length","description":"Video length in seconds (1-16).","required":false,"min":1,"max":16,"step":1,"placeholder":"5"},{"name":"aspect_ratio","type":"select","label":"Output Ratio","description":"Output ratio for the video.","required":false,"options":[{"value":"16:9","label":"16:9"},{"value":"4:3","label":"4:3"},{"value":"9:16","label":"9:16"}]},{"name":"movement_amplitude","type":"select","label":"Motion Intensity","description":"Motion intensity level.","required":false,"options":[{"value":"auto","label":"Auto"},{"value":"small","label":"Small"},{"value":"medium","label":"Medium"},{"value":"large","label":"Large"}]},{"name":"generate_audio","type":"boolean","label":"Generate Audio","description":"Whether to generate synchronized audio.","required":false},{"name":"bgm","type":"boolean","label":"Background Music","description":"Add background music to the output.","required":false},{"name":"seed","type":"number","label":"Random Seed","description":"Random seed for reproducibility (-1 for random).","required":false,"min":-1,"step":1}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/2291da3c-dc26-4a56-8b8e-4101a759e8b5-u1_1bd281e1-3f43-400a-991e-04ba89a5f6b9.png"},{"id":"wan-2.5-t2v","name":"Alibaba WAN 2.5 Text-to-Video Model","description":"Alibaba WAN 2.5 makes 480p-1080p text/image-to-video with synced audio and is faster, more affordable than Google Veo3.","category":"text-to-video","wavespeedEndpoint":"alibaba/wan-2.5/text-to-vi
1deo","outputType":"video","estimatedTime":"5-15s","costPerRun":0.5,"pricing":{"perSecondByParam":{"param":"size","rates":{"832x480":0.05,"480x832":0.05,"1280x720":0.1,"720x1280":0.1,"1920x1080":0.15,"1080x1920":0.15}}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Input text prompt for the video generation.","required":true,"placeholder":"Enter your prompt here"},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Input text prompts to avoid certain content.","required":false,"placeholder":"Enter prompts to avoid here"},{"name":"audio","type":"file","label":"Audio","description":"Optional audio file for voice/music.","required":false,"placeholder":"Upload an audio file","accept":"audio/wav,audio/mp3"},{"name":"size","type":"select","label":"Video Size","description":"Select the resolution/aspect of the video.","required":true,"options":[{"value":"1280x720","label":"1280 x 720"},{"value":"832x480","label":"832 x 480"},{"value":"480x832","label":"480 x 832"},{"value":"720x1280","label":"720 x 1280"},{"value":"1920x1080","label":"1920 x 1080"},{"value":"1080x1920","label":"1080 x 1920"}]},{"name":"duration","type":"number","label":"Duration","description":"Select the duration of the video in seconds (5 or 10).","required":true,"min":5,"max":10,"step":5},{"name":"enable_prompt_expansion","type":"boolean","label":"Enable Prompt Expansion","description":"If set to true, the prompt optimizer will be enabled.","required":false},{"name":"seed","type":"number","label":"Seed","description":"Random seed for generation.","required":false,"max":10000,"step":1,"placeholder":"Enter a seed value"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/1ea32f46-b238-4ebe-b5c4-3811bc9d79e4-u2_6457b48b-5f6e-449a-a4a9-455e3b9b9e41.png"}]},"image-to-video":{"id":"image-to-video","name":"Image to Video","description":"Animate images into videos","icon":"film","models":[{"id":"wan-2.2-i2v","name":"WAN 2.2","description":"Turns a single image into smooth, cinematic motion - ideal for storyboards, mood shots, and product demos","category":"image-to-video","wavespeedEndpoint":"wavespeed-ai/wan-2.2/image-to-video","outputType":"video","estimatedTime":"60-120s","costPerRun":0.15,"pricing":{"perSecondByParam"
1:{"param":"resolution","rates":{"480p":0.03,"720p":0.06}}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/f0a6704c-b128-45ba-9489-d3c19059d56a-u1_8344cf65-bd6b-4bf4-b1c2-dc1904da2ced.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"The image to animate into a video","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Describe the motion and scene you want","required":true,"placeholder":"Camera slowly pans across the scene with gentle motion...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Describe what you don\'t want in the output","required":false,"placeholder":"blurry, distorted, low quality...","maxLength":1000},{"name":"resolution","type":"select","label":"Resolution","description":"The resolution of the generated video","required":false,"default":"480p","options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"}]},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"5","options":[{"value":"5","label":"5 seconds"},{"value":"8","label":"8 seconds"}]},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"default":-1,"min":-1,"max":2147483647}]},{"id":"grok-imagine-video-i2v","name":"Grok Imagine Video","description":"X-AI\'s model that transforms images into videos with natural motion, scene continuity, and synchronized audio","category":"image-to-video","wavespeedEndpoint":"x-ai/grok-imagine-video/image-to-video","outputType":"video","estimatedTime":"30-120s","costPerRun":0.275,"pricing":{"costPerSecond":0.055},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/16274318-ff2c-4fe5-9292-c7b6ef51c4dd-u1_e2f0c62c-f151-4993-9f7a-f076d7a2b76d.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"The image to animate into a video","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of desired motion or changes","required":true,"placeholder":"The flowers gently sway in the breeze...","maxLength":2000},{"name":"duration","type":"number","label":"Duration","description":"Video duration in seconds (1-15)","required":false,"default":6,"min":1,"max":15},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Aspect ratio of the generated video","required":false,"default":"16:9","options":[{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"},{"value":"4:3","label":"Classic (4:3)"},{"value":"3:4","label":"Portrait Classic (3:4)"},{"value":"3:2","label":"Photo (3:2)"},{"value":"2:3","label":"Portrait Photo (2:3)"},{"value":"1:1","label":"Square (1:1)"}]},{"name":"resolution","type":"select","label":"Resolution","description":"Resolution of the output video","required":false,"default":"720p","options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"}]}]},{"id":"kling-v2.6-pro-i2v","name":"Kling 2.6 Pro","description":"Top-tier image-to-video with smooth motion, cinematic visuals, accurate prompt adherence, and optional native audio","category":"image-to-video","wavespeedEndpoint":"kwaivgi/kling-v2.6-pro/image-to-video","outputType":"video","estimatedTime":"60-180s","costPerRun":0.35,"pricing":{"durationPricing":{"5":0.35,"10":0.7},"costMultipliers":{"sound":2}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/b51d115f-b929-40e3-bda7-79573688fd89-u2_ff5a8806-d47b-4aa7-ad24-ab8d98d0e4e0.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"The image to animate (JPG/PNG, max 10MB, min 300px)","required":true,"accept":"image/jpeg,image/png"},{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"The character walks forward with fluid motion...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Describe what you don\'t want in the output","required":false,"placeholder":"blurry, distorted, low quality...","maxLength":1000},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"5","options":[{"value":"5","label":"5 seconds"},{"value":"10","label":"10 seconds"}]}
1,{"name":"sound","type":"boolean","label":"Generate Audio","description":"Generate native audio with the video (doubles cost)","required":false,"default":false},{"name":"cfg_scale","type":"number","label":"CFG Scale","description":"Prompt adherence strength (higher = stricter prompt following)","required":false,"default":0.5,"min":0,"max":1,"step":0.1}]},{"id":"veo3.1-fast","name":"Google Veo 3.1 Fast","description":"Google\'s fast Image-to-Video model with native 1080p output and audio generation","category":"image-to-video","wavespeedEndpoint":"google/veo3.1-fast/image-to-video","outputType":"video","estimatedTime":"60-180s","costPerRun":0.8,"pricing":{"costMultipliers":{"generate_audio":1.5}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/a2d2d809-9c11-4d29-8b6c-7add58f289a3-u1_a8d298cc-08bc-44ae-9dc4-c30cc9c5b184.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"The image to use as the starting frame","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the video generation","required":true,"placeholder":"A cinematic shot of waves crashing on the shore...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Things to avoid in the generation","required":false,"placeholder":"blurry, low quality, distorted...","maxLength":1000},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the generated video","required":false,"default":"16:9","options":[{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"}]},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"8","options":[{"value":"4","label":"4 seconds"},{"value":"6","label":"6 seconds"},{"value":"8","label":"8 seconds"}]},{"name":"resolution","type":"select","label":"Resolution","description":"Video resolution","required":false,"default":"1080p","options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"},{"value":"4k","label":"4K"}]},{"name":"generate_audio","type":"boolean","label":"Generate Audio","description":"Generate audio for the video ($1.20 with audio, $0.80 without)","required":false,"default":true},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility","required":false,"min":0,"max":2147483647}]},{"id":"sora-2-i2v","name":"OpenAI Sora 2 Pro","description":"Creates physics-aware, realistic videos from images with synchronized audio and greater steerability","category":"image-to-video","wavespeedEndpoint":"openai/sora-2/image-to-video-pro","outputType":"video","estimatedTime":"60-180s","costPerRun":1.2,"pricing":{"durationPricing":{"4":1.2,"8":2.4,"12":3.6}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/9b7bb961-59e7-4b1b-ad65-04d721da765e-u1_2eeeb6c5-446c-4b97-bc98-398e8a39b0c1.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"The image to animate into a video","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the generation","required":true,"placeholder":"The scene comes alive with natural motion and lighting...","maxLength":2000},{"name":"resolution","type":"select","label":"Resolution","description":"Resolution of the output video","required":false,"default":"720p","options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"4","options":[{"value":"4","label":"4 seconds"},{"value":"8","label":"8 seconds"},{"value":"12","label":"12 seconds"}]}]},{"id":"veo3.1","name":"Google Veo 3.1","description":"Google\'s premium Image-to-Video model with native 1080P output for highest quality videos with audio","category":"image-to-video","wavespeedEndpoint":"google/veo3.1/image-to-video","outputType":"video","estimatedTime":"120-300s","costPerRun":0.8,"pricing":{"durationPricing":{"4":0.8,"6":1.2,"8":1.6},"costMultipliers":{"generate_audio":2}},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/723c45f7-6758-4c8f-aef9-c0cb34213536-u1_32e2300f-26d1-40ab-91a8-c02e9aa3a928.png","parameters":[{"name":"image","type":"file","label":"Input Image","description":"The image to use as the starting frame","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"The positive prompt for the video generation","required":true,"placeholder":"An epic cinematic sequence with dramatic camera movement...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Things to avoid in the generation","required":false,"placeholder":"blurry, low quality, distorted...","maxLength":1000},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"The aspect ratio of the generated video","required":false,"default":"16:9","options":[{"value":"16:9","label":"Landscape (16:9)"},{"value":"9:16","label":"Portrait (9:16)"}]},{"name":"duration","type":"select","label":"Duration","description":"The duration of the generated video in seconds","required":false,"default":"8","options":[{"value":"4","label":"4 seconds"},{"value":"6","label":"6 seconds"},{"value":"8","label":"8 seconds"}]},{"name":"resolution","type":"select","label":"Resolution","description":"Video resolution","required":false,"default":"1080p","options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"},{"value":"4k","label":"4K"}]},{"name":"generate_audio","type":"boolean","label":"Generate Audio","de
1scription":"Generate audio ($0.40/sec with audio, $0.20/sec without)","required":false,"default":true},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility","required":false,"min":0,"max":2147483647}]},{"id":"gen4-turbo","name":"RunwayML Gen4 Turbo","description":"RunwayML Gen4 Turbo is an image-to-video model that generates high-quality videos from images. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"image-to-video","wavespeedEndpoint":"runwayml/gen4-turbo","outputType":"video","estimatedTime":"5-15s","costPerRun":0.5,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of the motion and action you want (e.g., \\"The model walks forward, fabric flowing in the wind\\").","required":true,"placeholder":"Enter your prompt here"},{"name":"image","type":"file","label":"Image","description":"Source image to animate (upload or public URL).","required":true,"accept":"image/*"},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Output aspect ratio: 16:9, 4:3, 1:1, 3:4, or 9:16. Leave empty to match source.","required":false,"options":[{"value":"16:9","label":"16:9"},{"value":"4:3","label":"4:3"},{"value":"1:1","label":"1:1"},{"value":"3:4","label":"3:4"},{"value":"9:16","label":"9:16"}]}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/082383b4-b67f-4277-a3df-52c7d5bf7d85-u2_d9f218c5-da78-4938-9353-0e15cbdd4100.png"},{"id":"infinitetalk","name":"InfiniteTalk","description":"InfiniteTalk converts one photo + audio into audio-driven talking or singing avatar videos (Image-to-Video), up to 10 minutes, 720p tier.","category":"image-to-video","wavespeedEndpoint":"wavespeed-ai/infinitetalk","outputType":"video","estimatedTime":"10-30s per second of video","costPerRun":0.15,"parameters":[{"name":"audio","type":"file","label":"Audio File","description":"Upload audio file for the avatar to sync with.","required":true,"accept":"audio/*"},{"name":"image","type":"file","label":"Image File","description":"Upload an image of the avatar to animate.","required":true,"accept":"image/*"},{"name":"mask_image","type":"file","label":"Mask Image","description":"Optional mask image to specify regions that can move during animation.","required":false,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"Optional prompt to guide expression, style, or pose.","required":false,"placeholder":"Enter your prompt here..."},{"name":"resolution","type":"select","label":"Resolution","description":"Select the output resolution of the video.","required":true,"options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"}]},{"name":"seed","type":"number","label":"Seed","description":"Set a fixed number for reproducibility.","required":false,"max":10000,"step":1,"placeholder":"Enter a seed value..."}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/571bf5ad-2d36-4114-a53c-8eefb3d9386e-u2_8a6f39ac-9313-4914-be6a-5646af19342e.png"},{"id":"ltx-2-19b","name":"LTX-2 19B","description":"LTX-2 19B is the first DiT-based audio-video foundation model with synchronized audio and video, high fidelity, multiple performance modes, and production-ready outputs in one model. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"image-to-video","wavespeedEndpoint":"wavespeed-ai/ltx-2-19b/image-to-video","outputType":"video","estimatedTime":"5-20s","costPerRun":0.08,"pricing":{"perSecondByParam"
1:{"param":"resolution","rates":{"480p":0.012,"720p":0.016,"1080p":0.024}}},"parameters":[{"name":"image","type":"file","label":"Reference Image","description":"Reference image to animate (JPG or PNG)","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of motion, action, and audio cues","required":true},{"name":"resolution","type":"select","label":"Output Resolution","description":"Output resolution: 480p, 720p (default), or 1080p","required":false,"options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"duration","type":"number","label":"Video Length","description":"Video length in seconds (5-20)","required":false,"min":5,"max":20,"step":1},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"min":-1,"max":999999,"step":1}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/e47cdddd-d80c-4083-afe3-533e0c918d20-u2_795982d8-6d11-454f-a74b-5746825e3bcf.png"},{"id":"hunyuan-1.5-i2v","name":"HunyuanVideo-1.5","description":"HunyuanVideo-1.5 (i2v) is a lightweight 8.3B parameter image-to-video model that generates high-quality videos from images with top-tier visual quality and motion coherence. Optimized for fast inference on consumer-grade GPUs. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"image-to-video","wavespeedEndpoint":"wavespeed-ai/hunyuan-video-1.5/image-to-video","outputType":"video","estimatedTime":"5-10s","costPerRun":0.4,"parameters":[{"name":"image","type":"file","label":"Input Image","description":"Upload your input image (this becomes the starting frame of the video).","required":true,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt","description":"Enter a prompt describing the motion, camera movement, environment changes, and overall mood.","required":true,"placeholder":"Describe the scene..."},{"name":"resolution","type":"select","label":"Resolution","description":"Choose the resolution of the video.","required":true,"options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"}]},{"name":"duration","type":"select","label":"Duration","description":"Select the duration of the video.","required":true,"options":[{"value":"5","label":"5 seconds"},{"value":"8","label":"8 seconds"},{"value":"10","label":"10 seconds"}]},{"name":"seed","type":"number","label":"Seed","description":"Set a seed number for reproducibility.","required":false,"max":10000,"step":1,"placeholder":"Enter a seed number"},{"name":"safetyChecker","type":"boolean","label":"Enable Safety Checker","description":"Optionally enable a safety checker for the generation.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/f848f796-4d3e-4ccb-9864-75b81ba03f8e-u2_5eb4f928-0f98-4151-bfdf-7dcb16c93bfb.png"},{"id":"vidu-i2v","name":"Vidu Image-to-Video","description":"Vidu Image-to-Video converts images into smooth-transition videos with high visual quality and diverse motion for cinematic results. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"image-to-video","wavespeedEndpoint":"vidu/image-to-video","outputType":"video","estimatedTime":"5-15s","costPerRun":0.2,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text description of desired motion and style.","required":true,"placeholder":"Describe the motion, mood, and style you want"},{"name":"image","type":"file","label":"Image","description":"Source image (upload or public URL).","required":true,"accept":"image/*"},{"name":"movement_amplitude","type":"select","label":"Movement Amplitude","description":"Motion intensity: auto, small, medium, or large (default: auto).","required":false,"options":[{"value":"auto","label":"Auto"},{"value":"small","label":"Small"},{"value":"medium","label":"Medium"},{"value":"large","label":"Large"}]},{"name":"seed","type":"number","label":"Seed","description":"Set for reproducibility; -1 for random.","required":false,"min":-1,"max":10000,"step":1,"placeholder":"Enter seed value"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/04561669-891c-4ab1-ace9-696200e1a2dd-u2_cacfe2c3-a7d4-4e62-ba41-ada53ebf1e01.png"},{"id":"seedance-1.5-pro-i2v","name":"Seedance 1.5 Pro Image to Video","description":"Seedance 1.5 Pro Image-to-Video generates cinematic, live-action–leaning clips from a text prompt plus a first-frame image, preserving the image\'s subject and composition while adding expressive motion and stable aesthetics. It supports 4–
112s duration control, adaptive aspect ratio, and reproducible outputs via seeds, ideal for ad creatives and short-drama shots that need a strong visual anchor.","category":"image-to-video","wavespeedEndpoint":"bytedance/seedance-v1.5-pro/image-to-video","outputType":"video","estimatedTime":"60-120s","costPerRun":0.26,"pricing":{"perSecondByParam":{"param":"resolution","rates":{"480p":0.012,"720p":0.026,"1080p":0.052}},"costMultipliers":{"generate_audio":2}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"The instruction describing what should happen in the video (action + camera + style).","required":true,"placeholder":"Enter your prompt here"},{"name":"image","type":"file","label":"Reference Image","description":"The reference image that anchors composition, subject identity, and lighting.","required":true,"accept":"image/*"},{"name":"last_image","type":"file","label":"Last Image (optional)","description":"An optional ending frame to steer the final composition (if supported).","required":false,"accept":"image/*"},{"name":"duration","type":"number","label":"Duration","description":"Video length in seconds.","required":false,"min":4,"max":12,"step":1,"placeholder":"Enter duration"},{"name":"resolution","type":"select","label":"Resolution","description":"Output resolution.","required":false,"options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Output aspect ratio.","required":false,"options":[{"value":"16:9","label":"16:9"},{"value":"9:16","label":"9:16"},{"value":"1:1","label":"1:1"},{"value":"4:3","label":"4:3"},{"value":"3:4","label":"3:4"},{"value":"21:9","label":"21:9"},{"value":"auto","label":"Auto"}]},{"name":"camera_fixed","type":"boolean","label":"Camera Fixed","description":"Whether to keep the camera position fixed.","required":false},{"name":"seed","type":"number","label":"Seed","description":"Random seed for reproducibility.","required":false,"min":-1,"max":1000,"step":1,"placeholder":"Enter seed value"},{"name":"generate_audio","type":"boolean","label":"Generate Audio","description":"To decide whether to generate videos with audio.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/b40cfa51-5bbe-4c19-9c29-a94009971174-u2_33b6c84c-95a2-4517-a36b-c192da42b189.png"},{"id":"vidu-q3-i2v","name":"Vidu Q3 Image-to-Video","description":"Vidu Q3 Image-to-Video turns text prompts into high-quality videos with exceptional visual fidelity and diverse motion. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"image-to-video","wavespeedEndpoint":"vidu/q3/image-to-video","outputType":"video","estimatedTime":"60-120s","costPerRun":0.375,"pricing":{"perSecondByParam"
1:{"param":"resolution","rates":{"540p":0.07,"720p":0.15,"1080p":0.16}}},"parameters":[{"name":"prompt","type":"textarea","label":"Text Prompt","description":"Text description of the desired motion and action","required":true,"placeholder":"Describe the motion you want..."},{"name":"image","type":"file","label":"Reference Image","description":"Reference image to animate (URL or upload)","required":true,"accept":"image/*"},{"name":"resolution","type":"select","label":"Output Resolution","description":"Output quality: 540p, 720p, 1080p","required":false,"options":[{"value":"540p","label":"540p"},{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"duration","type":"number","label":"Video Duration","description":"Video length in seconds (1-16, default: 5)","required":false,"min":1,"max":16,"step":1,"placeholder":"Set duration"},{"name":"movement_amplitude","type":"select","label":"Movement Amplitude","description":"Motion intensity: auto, small, medium, large","required":false,"options":[{"value":"auto","label":"Auto"},{"value":"small","label":"Small"},{"value":"medium","label":"Medium"},{"value":"large","label":"Large"}]},{"name":"generate_audio","type":"boolean","label":"Generate Audio","description":"Generate synchronized audio","required":false},{"name":"bgm","type":"boolean","label":"Background Music","description":"Add background music","required":false},{"name":"seed","type":"number","label":"Random Seed","description":"Random seed for reproducibility","required":false,"step":1,"placeholder":"Enter seed value"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/90feff51-5dfb-4695-bb63-28a5084be65b-u1_2093a477-2b61-4ed4-bdc2-d283b2b9f71e.png"},{"id":"i2v-pro","name":"MiniMax Hailuo 2.3 Pro","description":"An image-to-video model for ultra-clear 1080P output and physics-aware scenes with responsive rendering. Ready-to-use REST inference API, best performance, no cold starts, affordable pricing.","category":"image-to-video","wavespeedEndpoint":"minimax/hailuo-2.3/i2v-pro","outputType":"video","estimatedTime":"5 seconds","costPerRun":0.49,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"A text prompt to guide motion, camera style, or lighting.","required":false,"placeholder":"Describe motion, camera angle, or lighting"},{"name":"image","type":"file","label":"Input Image","description":"Upload a reference image to define your scene or subject. Accepts JPEG or PNG files.","required":true,"accept":"image/*"},{"name":"enable_prompt_expansion","type":"boolean","label":"Enable Safety Checker","description":"Automatically optimizes incoming prompts to enhance output quality, also activates the safety checker.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/3be463bd-d63f-45e8-b521-8fada17c3820-u1_3b38b823-7422-4ed9-9e5b-a249887e068c.png"},{"id":"veo3-i2v","name":"Google Veo3","description":"Google Veo 3 is Google\'s flagship image-to-video model that creates audio-enabled videos from images. It transforms still images into cinematic 1080p videos with smooth, realistic motion, consistent lighting, and synchronized native audio.","category":"image-to-video","wavespeedEndpoint":"google/veo3/image-to-video","outputType":"video","estimatedTime":"60-120s","costPerRun":1.2,"pricing":{"costMultipliers":{"generate_audio":2.6667}},"parameters":[{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"Describe the desired motion, mood, and camera movement.","required":true,"placeholder":"Example: \'Slow cinematic zoom out as wind moves through the trees.\'"},{"name":"image","type":"file","label":"Upload Image","description":"Choose a clear, high-quality still image to define the subject.","required":true,"accept":"image/*"},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Select the aspect ratio for the video.","required":false,"options":[{"value":"16:9","label":"16:9"},{"value":"16:9","label":"16:9"}]},{"name":"duration","type":"number","label":"Duration","description":"Select the video duration (maximum 8 seconds).","required":false,"min":1,"max":8,"step":1},{"name":"resolution","type":"select","label":"Resolution","description":"Choose video resolution.","required":false,"options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"generate_audio","type":"boolean","label":"Generate Audio","de
1scription":"Whether to generate audio for the video.","required":false},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Things you want to avoid in the video output.","required":false},{"name":"seed","type":"number","label":"Seed","description":"Random seed for generation consistency.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/ed78f7a6-b01c-4b5e-9f9d-fd520030f160-u2_d52bd780-b4de-45cc-98d5-256613c4454b.png"}]},"motion-control":{"id":"motion-control","name":"Motion Control","description":"Transfer motion from reference videos to images","icon":"move","models":[{"id":"kling-v2.6-std-motion","name":"Kling 2.6 Standard Motion Control","description":"Transfer motion from reference videos to animate still images with smooth, realistic results at an affordable price","category":"motion-control","wavespeedEndpoint":"kwaivgi/kling-v2.6-std/motion-control","outputType":"video","estimatedTime":"60-120s","costPerRun":0.21,"pricing":{"costPerSecond":0.07,"billingBlockSeconds":3},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/c67cc62d-a45f-4bd6-8cfc-4996b4a00093-u1_8868bea7-3036-4fc7-ba6a-ad8dc3d22871.png","parameters":[{"name":"image","type":"file","label":"Character Image","description":"Upload the character image (JPG/PNG, max 10MB, min 300px, aspect ratio 1:2.5 to 2.5:1)","required":true,"accept":"image/jpeg,image/png"},{"name":"video","type":"file","label":"Motion Reference Video","description":"Upload the motion reference video (MP4/MOV, max 10MB, min 300px)","required":true,"accept":"video/mp4,video/quicktime"},{"name":"character_orientation","type":"select","label":"Character Source","description":"Whether the character orientation matches the image or video","required":true,"default":"image","options":[{"value":"image","label":"Image"},{"value":"video","label":"Video"}]},{"name":"prompt","type":"textarea","label":"Prompt","description":"Optional description to guide the generation","required":false,"placeholder":"A person dancing gracefully...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Describe what you don\'t want in the output","required":false,"placeholder":"blurry, distorted, low quality...","maxLength":1000},{"name":"keep_original_sound","type":"boolean","label":"Keep Original Sound","description":"Whether to retain the original video sound","required":false,"default":true}]},{"id":"kling-v2.6-pro-motion","name":"Kling 2.6 Pro Motion Control","description":"Premium motion transfer with superior identity preservation, temporal consistency, and native-audio option for dance, action, and gesture animations","category":"motion-control","wavespeedEndpoint":"kwaivgi/kling-v2.6-pro/motion-control","outputType":"video","estimatedTime":"60-180s","costPerRun":0.336,"pricing":{"costPerSecond":0.112,"billingBlockSeconds":3},"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/d11c543e-65e4-414f-b593-76d0d4288760-u2_5b7b1c8b-80ee-4ccf-bb96-5c5c01c7ca12.png","parameters":[{"name":"image","type":"file","label":"Character Image","description":"Upload the character image (JPG/PNG, max 10MB, min 300px, aspect ratio 1:2.5 to 2.5:1)","required":true,"accept":"image/jpeg,image/png"},{"name":"video","type":"file","label":"Motion Reference Video","description":"Upload the motion reference video (MP4/MOV, max 10MB, min 300px)","required":true,"accept":"video/mp4,video/quicktime"},{"name":"character_orientation","type":"select","label":"Character Source","description":"Whether the character orientation matches the image or video","required":true,"default":"image","options":[{"value":"image","label":"Image"},{"value":"video","label":"Video"}]},{"name":"prompt","type":"textarea","label":"Prompt","description":"Optional description to guide the generation","required":false,"placeholder":"A person dancing gracefully...","maxLength":2000},{"name":"negative_prompt","type":"textarea","label":"Negative Prompt","description":"Describe what you don\'t want in the output","required":false,"placeholder":"blurry, distorted, low quality...","maxLength":1000},{"name":"keep_original_sound","type":"boolean","label":"Keep Original Sound","description":"Whether to retain the original video sound","required":false,"default":true}]},{"id":"fun-control","name":"Wan 2.2 Fun Control","description":"Wan2.2-Fun-Control is an advanced video generation and control model developed by the Alibaba PAI team, designed for precise and creative video synthesis. It uses Control Codes and multi-modal inputs to generate preset-controlled videos up to 120s at 720p, offering features like multi-modal control, high-quality video generation, and intelligent composition.","category":"motion-control","wavespeedEndpoint":"wavespeed-ai/wan-2.2/fun-control","outputType":"video","estimatedTime":"
160-120s","costPerRun":0.2,"parameters":[{"name":"image","type":"file","label":"Input Image","description":"Upload an image to guide the video generation.","required":true,"accept":"image/*"},{"name":"video","type":"file","label":"Input Video","description":"Upload a video to guide the video generation.","required":true,"accept":"video/*"},{"name":"prompt","type":"textarea","label":"Prompt Enhancer","description":"Describe the video you want to generate.","required":false,"placeholder":"Enter your prompt here"},{"name":"resolution","type":"select","label":"Video Resolution","description":"Choose the resolution for video output.","required":true,"options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"}]},{"name":"seed","type":"number","label":"Random Seed","description":"Set a random seed for generation control.","required":false,"step":1},{"name":"Enable Safety Checker","type":"boolean","label":"Enable Safety Checker","description":"Toggle to enable/disable the safety checker for generated content.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/b8aee257-4c45-41ec-bc1e-374ee42c6e23-u2_d258499a-d39b-4951-b3ac-a2c6bafaf57e.png"},{"id":"ltx-2-19b-controlnet","name":"LTX-2 19B ControlNet","description":"LTX-2 19B ControlNet generates synchronized audio-video (up to 20s) from video input with pose, depth, or canny edge guidance. Supports audio preservation, generation, or removal for flexible video transformation. Ready-to-use REST inference API, best performance, no cold starts, affordable pricing.","category":"motion-control","wavespeedEndpoint":"wavespeed-ai/ltx-2-19b/control","outputType":"video","estimatedTime":"60-120s","costPerRun":0.2,"parameters":[{"name":"video","type":"file","label":"Input Video","description":"Input video providing motion and structure","required":true,"accept":"video/*"},{"name":"image","type":"file","label":"Reference Image","description":"Reference image for appearance guidance","required":false,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt Description","description":"Text description of desired output","required":false,"placeholder":"Describe the desired transformation"},{"name":"mode","type":"select","label":"Control Mode","description":"Control mode: pose (default), depth, or canny","required":false,"options":[{"value":"pose","label":"Pose"},{"value":"depth","label":"Depth"},{"value":"canny","label":"Canny"}]},{"name":"audio_mode","type":"select","label":"Audio Handling","description":"Audio handling: preserve (default), generate, or none","required":false,"options":[{"value":"preserve","label":"Preserve"},{"value":"generate","label":"Generate"},{"value":"none","label":"None"}]},{"name":"resolution","type":"select","label":"Output Resolution","description":"Output resolution: 480p, 720p (default), or 1080p","required":false,"options":[{"value":"480p","label":"480p"},{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"}]},{"name":"seed","type":"number","label":"Random Seed","description":"Random seed for reproducibility (-1 for random)","required":false,"min":-1,"max":2147483647,"step":1,"placeholder":"-1"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/8d31a9c8-e192-47b3-adde-349800bc37bb-u2_3723e1cb-3dfe-43f5-a646-170cab1b9b4e.png"}]},"video-to-video":{"id":"video-to-video","name":"Video to Video","description":"Transform and enhance videos","icon":"film","models":[{"id":"gen4-aleph","name":"RunwayML Gen4 Aleph","description":"RunwayML Gen4 Aleph is a Video-to-Video model for editing, transforming, and generating video at $0.18 per second. It utilizes natural language instructions to edit and modify footage—removing objects, changing environments, and applying styles. The API supports context-aware transformations and flexible aspect ratios, ensuring high-quality outputs.","category":"video-to-video","wavespeedEndpoint":"runwayml/gen4-aleph","outputType":"video","estimatedTime":"5-15s","costPerRun":0.9,"parameters":[{"name":"prompt","type":"textarea","label":"Prompt","description":"Text instruction describing the edit (e.g., \'Remove people from the video\').","required":true,"placeholder":"Describe the transformation you want."},{"name":"video","type":"file","label":"Video","description":"Source video file (upload or public URL).","required":true,"accept":"video/*"},{"name":"aspect_ratio","type":"select","label":"Aspect Ratio","description":"Output aspect ratio: 16:9, 4:3, 1:1, 3:4, or 9:16. Default: 16:9.","required":false,"options":[{"value":"16:9","label":"16:9"},{"value":"4:3","label":"4:3"},{"value":"1:1","label":"1:1"},{"value":"3:4","label":"3:4"},{"value":"9:16","label":"9:16"}]},{"name":"reference_image","type":"file","label":"Reference Image","description":"Reference image to guide style or appearance (upload or URL).","required":false,"accept":"image/*"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/0f6e03a2-6b54-4314-9402-a5245f3bad06-u2_4c60b85c-8f67-4921-96f6-e495e1d194f5.png"},{"id":"video-to-video","name":"InfiniteTalk Fast Video-To-Video","description":"Audio-driven infinitetalk-fast turns one video plus audio into realistic talking or singing videos with lip-sync. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"video-to-video","wavespeedEndpoint":"wavespeed-ai/infinitetalk-fast/video-to-video","outputType":"video","estimatedTime":"10–30s per 1 second of video","costPerRun":0.15,"parameters":[{"name":"audio","type":"file","label":"Audio File","description":"Upload the audio file to sync with the video.","required":true,"accept":"audio/*"},{"name":"video","type":"file","label":"Video File","description":"Upload the base video for the transformation.","required":true,"accept":"video/*"},{"name":"mask_image","type":"file","label":"Mask Image (Optional)","description":"Upload a mask image to control which regions can move (optional).","required":false,"accept":"image/*"},{"name":"prompt","type":"textarea","label":"Prompt Enhancer (Optional)","description":"Specify the style, pose or expressions you want to guide your video (optional).","required":false,"placeholder":"Write your prompt here..."},{"name":"seed","type":"number","label":"Seed (Optional)","description":"Set the seed for reproducibility of the results (optional).","required":false,"max":1000,"step":1,"placeholder":"Enter seed value"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/cc4539a1-b83b-4dca-9c99-bb4e5beaa45b-u2_f8d62070-d798-4586-87bb-32ce4bc0dbf3.png"}]},"upscaler":{"id":"upscaler","name":"Upscaler","description":"Upscale and enhance image/video quality","icon":"maximize","models":[{"id":"real-esrgan","name":"Real-ESRGAN","description":"Real-ESRGAN delivers high-quality image super-resolution with optional face correction and adjustable upscale factors. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"upscaler","wavespeedEndpoint":"wavespeed-ai/real-esrgan","outputType":"image","estimatedTime":"5-15s","costPerRun":0.0024,"parameters":[{"name":"image","type":"file","label":"Image","description":"The image to upscale/enhance","required":true,"accept":"image/*"}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/dd3fe698-2e86-47ae-ba3f-f1771842cb87-u2_ca67c4cf-dbc1-4244-9412-ed24e63184e0.png"},{"id":"increase-resolution","name":"Bria Increase Resolution","description":"Bria Increase Resolution upscales images with a method that preserves original content without regeneration, producing sharper, higher-quality output. Ready-to-use REST API, best performance, no coldstarts, affordable pricing.","category":"upscaler","wavespeedEndpoint":"bria/increase-resolution","outputType":"image","estimatedTime":"5-15s","costPerRun":0.04,"parameters":[{"name":"image","type":"file","label":"Source Image","description":"Source image (URL or upload).","required":true,"placeholder":"Drag and drop a file or click to upload","accept":"image/*"},{"name":"desired_increase","type":"select","label":"Desired Increase","description":"Resolution multiplier: 2 or 4.","required":true,"options":[{"value":"2","label":"2x"},{"value":"4","label":"4x"}]},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response. This property is only available through the API.","required":false},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL. This property is only available through the API.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/7dc9838d-980b-4086-93c5-2ee20e450007-u1_93a57619-10d7-4f0c-a028-61598bc89028.png"},{"id":"flashvsr","name":"FlashVSR Video Upscaler","description":"FlashVSR is a fast, high-quality video upscaler that boosts resolution and restores clarity for low-resolution or blurry footage. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"upscaler","wavespeedEndpoint":"wavespeed-ai/flashvsr","outputType":"video","estimatedTime":"3–20 seconds per second of video","costPerRun":0.06,"parameters":[{"name":"video","type":"file","label":"Upload Video","description":"The video file you want to upscale.","required":true,"placeholder":"Drag and drop a file or click to upload","accept":"video/*"},{"name":"target_resolution","type":"select","label":"Target Resolution","description":"Choose the desired resolution for the output video.","required":true,"options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"},{"value":"2k","label":"2K"},{"value":"4k","label":"4K"}]}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/658979bd-dbcb-4312-ab65-f9
1612526e61c-u2_a904fd10-0296-4b2e-923f-ab31e89b44fe.png"},{"id":"image-upscaler","name":"WaveSpeed AI Image Upscaler","description":"AI Image Upscaler that enhances image resolution to 4K or 8K while improving detail and clarity for photos and graphics. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"upscaler","wavespeedEndpoint":"wavespeed-ai/image-upscaler","outputType":"image","estimatedTime":"Fast & Efficient","costPerRun":0.01,"parameters":[{"name":"image","type":"file","label":"Input Image","description":"The image you want to upscale.","required":true,"placeholder":"Drag and drop a file or click to upload.","accept":"image/*"},{"name":"target_resolution","type":"select","label":"Target Resolution","description":"Choose the resolution you want to upscale to.","required":true,"options":[{"value":"2k","label":"2K"},{"value":"4k","label":"4K"},{"value":"8k","label":"8K"}]},{"name":"output_format","type":"select","label":"Output Format","description":"Select the format of the output image.","required":true,"options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"},{"value":"webp","label":"WEBP"}]},{"name":"enable_base64_output","type":"boolean","label":"Enable Base64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL.","required":false},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated before returning the response.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/f43b3777-ac8a-4b60-a9f2-e30633819521-u1_a1a816e3-b22f-4208-98a8-b0c4ebf89dbe.png"},{"id":"video-upscaler","name":"Video Upscaler","description":"AI Video Upscaler enhances resolution and clarity to fix blurry outputs and improve low-resolution footage with ML upscaling. It features a ready-to-use REST inference API, optimized performance, and affordable pricing.","category":"upscaler","wavespeedEndpoint":"wavespeed-ai/video-upscaler","outputType":"video","estimatedTime":"5–10 seconds per second of video","costPerRun":0.025,"parameters":[{"name":"video","type":"file","label":"Upload Video","description":"Video file to be upscaled.","required":true,"placeholder":"Drag and drop or click to upload","accept":"video/*"},{"name":"target_resolution","type":"select","label":"Target Resolution","description":"Choose the resolution for the output video.","required":true,"options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"},{"value":"2k","label":"2K"},{"value":"4k","label":"4K"}]},{"name":"Enable Safety Checker","type":"boolean","label":"Enable Safety Checker","description":"Toggle to enable safety checks during processing.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/a7ac7f26-4ee0-4b68-99e4-b52f87c1d1ad-u1_b26945e3-e829-4a91-aaf7-5772c78e7ab6.png"},{"id":"ultimate-image-upscaler","name":"Ultimate Image Upscaler","description":"The Ultimate Image Upscaler is a high-performance enhancement model that intelligently increases image resolution while preserving details, sharpness, and natural texture. It uses advanced deep learning upscaling to restore fine features and eliminate blur or compression artifacts.","category":"upscaler","wavespeedEndpoint":"wavespeed-ai/ultimate-image-upscaler","outputType":"image","estimatedTime":"5-15s","costPerRun":0.06,"parameters":[{"name":"image","type":"file","label":"Input Image","description":"Upload the image you want to upscale.","required":true,"accept":"image/*"},{"name":"target_resolution","type":"select","label":"Target Resolution","description":"Select the desired output resolution.","required":true,"options":[{"value":"2k","label":"2K"},{"value":"4k","label":"4K"},{"value":"8k","label":"8K"}]},{"name":"output_format","type":"select","label":"Output Format","description":"Choose the output format for the upscaled image.","required":true,"options":[{"value":"jpeg","label":"JPEG"},{"value":"png","label":"PNG"},{"value":"webp","label":"WEBP"}]}
1,{"name":"enable_base64_output","type":"boolean","label":"Enable BASE64 Output","description":"If enabled, the output will be encoded into a BASE64 string instead of a URL. This property is only available through the API.","required":false},{"name":"enable_sync_mode","type":"boolean","label":"Enable Sync Mode","description":"If set to true, the function will wait for the result to be generated and uploaded before returning the response. This property is only available through the API.","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/ef7e76d2-583e-4bb0-99fa-eb5cea5de04d-u2_ea350ed9-9e19-469d-ba22-49bf7c054412.png"},{"id":"video-upscaler-pro","name":"AI Video Upscaler Pro","description":"AI Video Upscaler Pro converts low-resolution videos into crisp 4K footage with seamless motion dynamics and frame consistency. Ready-to-use REST inference API, best performance, no coldstarts, affordable pricing.","category":"upscaler","wavespeedEndpoint":"wavespeed-ai/video-upscaler-pro","outputType":"video","estimatedTime":"10-30s per second of video","costPerRun":0.1,"parameters":[{"name":"video","type":"file","label":"Upload Video","description":"The video file to upscale","required":true,"placeholder":"Drag and drop or click to upload","accept":"video/*"},{"name":"target_resolution","type":"select","label":"Target Resolution","description":"Select the desired output resolution for upscaling","required":true,"options":[{"value":"720p","label":"720p"},{"value":"1080p","label":"1080p"},{"value":"2k","label":"2K"},{"value":"4k","label":"4K"}]},{"name":"Enable Safety Checker","type":"boolean","label":"Enable Safety Checker","description":"Enables safety check during processing","required":false}],"previewImage":"https://d2p7pge43lyniu.cloudfront.net/output/cac6ae48-340f-41d2-9501-40eef396a3e3-u1_4a18cbf5-e393-4417-91fd-3f3431b7e2d9.png"}]}}}'))}]);

Line numbers count LF bytes from the start of the resource, as the search results do. Vendor segments are library code the classifier recognised; they are stored but not indexed. Bytes are shown as Latin1 characters, one per byte.