wan2-1-VACE-fast

Running on Zero

App Files Files Community

linoyts HF Staff commited on 14 days ago

Commit

33f10c8

verified ·

1 Parent(s): b15e441

Update app.py

Browse files

Files changed (1) hide show

app.py +11 -11

app.py CHANGED Viewed

@@ -49,9 +49,9 @@ MAX_FRAMES_MODEL = 81
 # Default prompts for different modes - Updated with new mode names
 MODE_PROMPTS = {
-    "reference": "the playful penguin picks up the green cat eye sunglasses and puts them on",
-    "first - last frame": "CG animation style, a small blue bird takes off from the ground, flapping its wings. The bird's feathers are delicate, with a unique pattern on its chest. The background shows a blue sky with white clouds under bright sunshine. The camera follows the bird upward, capturing its flight and the vastness of the sky from a close-up, low-angle perspective.",
-    "random transitions": "Various different characters appear and disappear in a fast transition video showcasting their unique features and personalities. The video is about showcasing different dance styles, with each character performing a distinct dance move. The background is a vibrant, colorful stage with dynamic lighting that changes with each dance style. The camera captures close-ups of the dancers' expressions and movements. Highly dynamic, fast-paced music video, with quick cuts and transitions between characters, cinematic, vibrant colors"
 }
 default_negative_prompt = "Bright tones, overexposed, static, blurred details, subtitles, style, works, paintings, images, static, overall gray, worst quality, low quality, JPEG compression residue, ugly, incomplete, extra fingers, poorly drawn hands, poorly drawn faces, deformed, disfigured, misshapen limbs, fused fingers, still picture, messy background, three legs, many people in the background, walking backwards, watermark, text, signature"
@@ -335,7 +335,7 @@ def get_duration(gallery_images, mode, prompt, height, width,
         base_duration = 75
     # Add extra time for background removal processing
-    if mode == "reference" and remove_bg:  # Updated to use new mode name
         base_duration += 30
     return base_duration
@@ -351,7 +351,7 @@ def generate_video(gallery_images, mode, prompt, height, width,
     Args:
         gallery_images (list): List of PIL images from the gallery
-        mode (str): Processing mode - "reference", "first - last frame", or "random transitions"
         prompt (str): Text prompt describing the desired animation
         height (int): Target height for the output video
         width (int): Target width for the output video
@@ -376,7 +376,7 @@ def generate_video(gallery_images, mode, prompt, height, width,
             image = img[0]  # Extract PIL image from gallery format
             # Apply background removal only for reference mode if checkbox is checked
-            if mode == "reference" and remove_bg:  # Updated to use new mode name
                 image = remove_background_from_image(image)
             # Always remove alpha channels to ensure RGB format
@@ -385,9 +385,9 @@ def generate_video(gallery_images, mode, prompt, height, width,
         gallery_images = processed_images
-    if mode == "first - last frame" and len(gallery_images) >= 2:  # Updated mode name
         gallery_images = gallery_images[:2]
-    elif mode == "first - last frame" and len(gallery_images) < 2:  # Updated mode name
         raise gr.Error("First - Last Frame mode requires at least 2 images, but only {} were supplied.".format(len(gallery_images)))
     target_h = max(MOD_VALUE, (int(height) // MOD_VALUE) * MOD_VALUE)
@@ -398,7 +398,7 @@ def generate_video(gallery_images, mode, prompt, height, width,
     current_seed = random.randint(0, MAX_SEED) if randomize_seed else int(seed)
     # Process images based on the selected mode
-    if mode == "first - last frame":  # Updated mode name
         frames, mask = prepare_video_and_mask_FLF2V(
             first_img=gallery_images[0],
             last_img=gallery_images[1],
@@ -407,7 +407,7 @@ def generate_video(gallery_images, mode, prompt, height, width,
             num_frames=num_frames
         )
         reference_images = None
-    elif mode == "reference":  # Updated mode name
         frames, mask = prepare_video_and_mask_Ref2V(height=target_h, width=target_w, num_frames=num_frames)
         reference_images = gallery_images
     else:  # mode == "random transitions"  # Updated mode name
@@ -460,7 +460,7 @@ with gr.Blocks() as demo:
             with gr.Group():
                  # Radio button for mode selection with updated names
                 mode_radio = gr.Radio(
-                    choices=["Reference", "First - Last frame", "Random Transitions"],
                     value="reference",
                     label="Control Mode",
                     #info="Reference: upload reference images to take elements from | First - Last Frame: upload 1st and last frames| Random Transitions: upload images to be used as frame anchors"

 # Default prompts for different modes - Updated with new mode names
 MODE_PROMPTS = {
+    "Reference": "the playful penguin picks up the green cat eye sunglasses and puts them on",
+    "First - Last Frame": "CG animation style, a small blue bird takes off from the ground, flapping its wings. The bird's feathers are delicate, with a unique pattern on its chest. The background shows a blue sky with white clouds under bright sunshine. The camera follows the bird upward, capturing its flight and the vastness of the sky from a close-up, low-angle perspective.",
+    "Random Transitions": "Various different characters appear and disappear in a fast transition video showcasting their unique features and personalities. The video is about showcasing different dance styles, with each character performing a distinct dance move. The background is a vibrant, colorful stage with dynamic lighting that changes with each dance style. The camera captures close-ups of the dancers' expressions and movements. Highly dynamic, fast-paced music video, with quick cuts and transitions between characters, cinematic, vibrant colors"
 }
 default_negative_prompt = "Bright tones, overexposed, static, blurred details, subtitles, style, works, paintings, images, static, overall gray, worst quality, low quality, JPEG compression residue, ugly, incomplete, extra fingers, poorly drawn hands, poorly drawn faces, deformed, disfigured, misshapen limbs, fused fingers, still picture, messy background, three legs, many people in the background, walking backwards, watermark, text, signature"
         base_duration = 75
     # Add extra time for background removal processing
+    if mode == "Reference" and remove_bg:  # Updated to use new mode name
         base_duration += 30
     return base_duration
     Args:
         gallery_images (list): List of PIL images from the gallery
+        mode (str): Processing mode - "Reference", "first - last frame", or "random transitions"
         prompt (str): Text prompt describing the desired animation
         height (int): Target height for the output video
         width (int): Target width for the output video
             image = img[0]  # Extract PIL image from gallery format
             # Apply background removal only for reference mode if checkbox is checked
+            if mode == "Reference" and remove_bg:  # Updated to use new mode name
                 image = remove_background_from_image(image)
             # Always remove alpha channels to ensure RGB format
         gallery_images = processed_images
+    if mode == "First - Last Frame" and len(gallery_images) >= 2:  # Updated mode name
         gallery_images = gallery_images[:2]
+    elif mode == "First - Last Frame" and len(gallery_images) < 2:  # Updated mode name
         raise gr.Error("First - Last Frame mode requires at least 2 images, but only {} were supplied.".format(len(gallery_images)))
     target_h = max(MOD_VALUE, (int(height) // MOD_VALUE) * MOD_VALUE)
     current_seed = random.randint(0, MAX_SEED) if randomize_seed else int(seed)
     # Process images based on the selected mode
+    if mode == "First - Last Frame":  # Updated mode name
         frames, mask = prepare_video_and_mask_FLF2V(
             first_img=gallery_images[0],
             last_img=gallery_images[1],
             num_frames=num_frames
         )
         reference_images = None
+    elif mode == "Reference":  # Updated mode name
         frames, mask = prepare_video_and_mask_Ref2V(height=target_h, width=target_w, num_frames=num_frames)
         reference_images = gallery_images
     else:  # mode == "random transitions"  # Updated mode name
             with gr.Group():
                  # Radio button for mode selection with updated names
                 mode_radio = gr.Radio(
+                    choices=["Reference", "First - Last Frame", "Random Transitions"],
                     value="reference",
                     label="Control Mode",
                     #info="Reference: upload reference images to take elements from | First - Last Frame: upload 1st and last frames| Random Transitions: upload images to be used as frame anchors"