{
  "status": "complete",
  "samples": [
    {
      "id": 1,
      "slug": "via-mlx-serve-by-561487",
      "title": {
        "en": "via MLX Serve by",
        "zh": "小恐龙打喷嚏变出蝴蝶"
      },
      "category": "animation",
      "source": {
        "kind": "x",
        "name": "@albertgao",
        "url": "https://x.com/albertgao/status/2085137397134561487"
      },
      "source_prompt": "\"A cute little girl walks through a sunny forest when a small, friendly dinosaur steps out, waves, and says hi. Suddenly, the dinosaur sneezes—and a shower of colorful butterflies bursts from its nose. They both laugh. Cute, colorful animated style.\"\n\n- time: 15min ish\n- 768x768",
      "source_duration": "5s",
      "prompt": "integrated_multimodal_description: [Shot 1] A cute little girl walks through a sunny forest when a small, friendly dinosaur with a light, cheerful English voice (S1) steps out, waves, and says: <d>[English] hi</d>. Suddenly, the dinosaur sneezes—and a shower of colorful butterflies bursts from its nose. They both laugh. Cute, colorful animated style. Pace the meeting, sneeze and shared laughter over a continuous 15-second scene, allowing time to see the butterflies settle around both characters.\n\noverall_soundscape: Light forest wind, birds and small footsteps accompany the meeting. A comic sneeze releases a fluttering rush of butterfly wings, followed by the girl and dinosaur laughing together.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "954337465d95f9ab7b1fd512ba3879fa925f8e668f493e314c9d84a91e99fad3",
      "duration": 15,
      "seed": 42,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源条目元数据为 5s，本次明确扩展为 15 秒。"
      ],
      "source_title": {
        "en": "via MLX Serve by",
        "zh": "via MLX Serve by"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-01",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] A cute little girl walks through a sunny forest when a small, friendly dinosaur with a light, cheerful English voice (S1) steps out, waves, and says: <d>[English] hi</d>. Suddenly, the dinosaur sneezes—and a shower of colorful butterflies bursts from its nose. They both laugh. Cute, colorful animated style. Pace the meeting, sneeze and shared laughter over a continuous 15-second scene, allowing time to see the butterflies settle around both characters.\n\noverall_soundscape: Light forest wind, birds and small footsteps accompany the meeting. A comic sneeze releases a fluttering rush of butterfly wings, followed by the girl and dinosaur laughing together.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "954337465d95f9ab7b1fd512ba3879fa925f8e668f493e314c9d84a91e99fad3"
    },
    {
      "id": 2,
      "slug": "the-last-thing-you-see-in-your-first-and-last-290296",
      "title": {
        "en": "The last thing you see in your first and last",
        "zh": "太空旅客窗外的黑洞"
      },
      "category": "cinematic-travel",
      "source": {
        "kind": "x",
        "name": "@ivanfioravanti",
        "url": "https://x.com/ivanfioravanti/status/2086553101029290296"
      },
      "source_prompt": "amateur handheld pov footage of a tourist in their plush and comfortable room of a space cruise, carpeted floor and comfy bed, large window shows a view of Gargantua blackhole, you can see their reflection in the window, they turn off the light half way through so we can see outside better and say: \"Wow look at that!\", then they show back the room",
      "source_duration": "10s",
      "prompt": "integrated_multimodal_description: [Shot 1] amateur handheld pov footage of a tourist with an excited, conversational English voice (S1) in their plush and comfortable room of a space cruise, carpeted floor and comfy bed, large window shows a view of Gargantua blackhole, you can see their reflection in the window, they turn off the light half way through so we can see outside better and say: <d>[English] Wow look at that!</d>, then they show back the room Use one uninterrupted 15-second handheld take. Switch off the room light around halfway through, hold the window view for the spoken reaction, then turn back to show the room.\n\noverall_soundscape: A steady enclosed-room ventilation hum and soft carpeted footsteps continue under slight handheld handling noise. A light-switch click marks the change to the darkened room.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "712f02d91f39b4d5a029cf4c3396609ac32d151511bf04949501f697139405e5",
      "duration": 15,
      "seed": 43,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源条目元数据为 10s，本次明确扩展为 15 秒。"
      ],
      "source_title": {
        "en": "The last thing you see in your first and last",
        "zh": "The last thing you see in your first and last"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-02",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] amateur handheld pov footage of a tourist with an excited, conversational English voice (S1) in their plush and comfortable room of a space cruise, carpeted floor and comfy bed, large window shows a view of Gargantua blackhole, you can see their reflection in the window, they turn off the light half way through so we can see outside better and say: <d>[English] Wow look at that!</d>, then they show back the room Use one uninterrupted 15-second handheld take. Switch off the room light around halfway through, hold the window view for the spoken reaction, then turn back to show the room.\n\noverall_soundscape: A steady enclosed-room ventilation hum and soft carpeted footsteps continue under slight handheld handling noise. A light-switch click marks the change to the darkened room.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "712f02d91f39b4d5a029cf4c3396609ac32d151511bf04949501f697139405e5"
    },
    {
      "id": 3,
      "slug": "ultra-cinematic-macro-shot-a-calm-mountain-lake-039920",
      "title": {
        "en": "Ultra cinematic macro shot: A calm mountain lake at dawn",
        "zh": "水滴汇成 MORNING"
      },
      "category": "cinematic-travel",
      "source": {
        "kind": "x",
        "name": "@CharaspowerAI",
        "url": "https://x.com/CharaspowerAI/status/2085622021812039920"
      },
      "source_prompt": "Ultra cinematic macro shot: A calm mountain lake at dawn reflects the first golden rays of sunlight. Tiny ripples spread across the mirror-like surface as morning mist drifts gracefully above the water. Billions of sparkling droplets slowly rise into the air, suspended as if gravity has stopped. They spiral together with incredible fluid realism, naturally sculpting the word \"MORNING\" entirely from crystal-clear water. The liquid letters shimmer with golden reflections before gently collapsing into a spectacular explosion of droplets illuminated by the sunrise. Hollywood title sequence, ultra-realistic fluid simulation, volumetric lighting, IMAX quality.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Ultra cinematic macro shot: A calm mountain lake at dawn reflects the first golden rays of sunlight. Tiny ripples spread across the mirror-like surface as morning mist drifts gracefully above the water. Billions of sparkling droplets slowly rise into the air, suspended as if gravity has stopped. They spiral together with incredible fluid realism, naturally sculpting the word \"MORNING\" entirely from crystal-clear water. The liquid letters shimmer with golden reflections before gently collapsing into a spectacular explosion of droplets illuminated by the sunrise. Hollywood title sequence, ultra-realistic fluid simulation, volumetric lighting, IMAX quality. In one continuous 15-second camera move, establish the dawn lake first, raise and spiral the droplets in the middle, hold the completed readable word before its final liquid collapse.\n\noverall_soundscape: Quiet mountain wind, small lake ripples and faint distant dawn birds establish the lakeside. Delicate water movement grows into the rushing liquid-letter collapse and a crisp shower of droplets.\n\nnon_diegetic_music: Slow, sustained string notes and sparse high piano tones rise gently with the water lettering and resolve as the droplets fall.",
      "prompt_sha256": "7abb51ba6a7327cb9266b13da9a8fa218928aeb0d97f444c82dc95ea21c245c2",
      "duration": 15,
      "seed": 44,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Ultra cinematic macro shot: A calm mountain lake at dawn",
        "zh": "Ultra cinematic macro shot: A calm mountain lake at dawn"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-03",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Ultra cinematic macro shot: A calm mountain lake at dawn reflects the first golden rays of sunlight. Tiny ripples spread across the mirror-like surface as morning mist drifts gracefully above the water. Billions of sparkling droplets slowly rise into the air, suspended as if gravity has stopped. They spiral together with incredible fluid realism, naturally sculpting the word \"MORNING\" entirely from crystal-clear water. The liquid letters shimmer with golden reflections before gently collapsing into a spectacular explosion of droplets illuminated by the sunrise. Hollywood title sequence, ultra-realistic fluid simulation, volumetric lighting, IMAX quality. In one continuous 15-second camera move, establish the dawn lake first, raise and spiral the droplets in the middle, hold the completed readable word before its final liquid collapse.\n\noverall_soundscape: Quiet mountain wind, small lake ripples and faint distant dawn birds establish the lakeside. Delicate water movement grows into the rushing liquid-letter collapse and a crisp shower of droplets.\n\nnon_diegetic_music: Slow, sustained string notes and sparse high piano tones rise gently with the water lettering and resolve as the droplets fall.",
      "effective_prompt_sha256": "7abb51ba6a7327cb9266b13da9a8fa218928aeb0d97f444c82dc95ea21c245c2"
    },
    {
      "id": 4,
      "slug": "strawberry-seasonal-match-cut-food-commercial-117943",
      "title": {
        "en": "Strawberry seasonal match-cut food commercial",
        "zh": "草莓穿过四季落入奶油"
      },
      "category": "product-commercial",
      "source": {
        "kind": "x",
        "name": "@HBCoop_",
        "url": "https://x.com/HBCoop_/status/2082827829172117943"
      },
      "source_prompt": "Premium food commercial. Extreme macro shot of a single ripe strawberry falling toward a white ceramic bowl. Before it lands, the environment changes through four seamless seasonal match cuts: a sunlit spring greenhouse, a hot summer picnic, an autumn kitchen and a candlelit winter dining room. The strawberry stays in exactly the same screen position, size, orientation and downward motion through every transition. It lands in fresh cream during the winter scene, sending one graceful crown-shaped splash upward. Photorealistic food texture, high-speed product cinematography, clean commercial lighting. Each season has distinct environmental sound, but one continuous musical phrase connects the entire sequence. No hands, labels or additional fruit.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Photorealistic premium food commercial, extreme macro high-speed product cinematography. A single ripe strawberry falls toward a white ceramic bowl in a sunlit spring greenhouse. Its red seeds and moist surface are sharply resolved against soft green glasshouse bokeh. Keep the strawberry's screen position, size, orientation and downward motion consistent through every seasonal match cut. No hands, labels or additional fruit.\n[Shot 2] At 00:03.000, the shot transitions to a hot summer picnic in a seamless seasonal match cut around the same falling strawberry and bowl. Warm sunlight replaces the greenhouse light; distant insects accompany its uninterrupted descent.\n[Shot 3] At 00:06.000, the shot transitions to an autumn kitchen with amber window light, preserving the strawberry and bowl alignment and the same downward motion. Soft indoor room tone replaces the outdoor insects.\n[Shot 4] At 00:09.000, the shot transitions to a candlelit winter dining room. The strawberry completes its fall into fresh cream in the white bowl, sending one graceful crown-shaped splash upward. Hold the detailed splash and its settling ripples through the final seconds. One continuous musical phrase connects all four seasons.\n\noverall_soundscape: Spring birds and a gentle greenhouse breeze change to summer insects, quiet autumn kitchen ambience, then the still winter dining room. The final cream splash and falling droplets remain clearly audible.\n\nnon_diegetic_music: One uninterrupted light piano-and-plucked-string phrase runs at a moderate tempo across all four seasons, resolving with the strawberry splash.",
      "prompt_sha256": "0c4614d6b1c88a86974b1393eab7a574817abc11c37bf72e3cb53e2d9bbf718f",
      "duration": 15,
      "seed": 45,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Strawberry seasonal match-cut food commercial",
        "zh": "草莓四季匹配剪辑食品广告"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-04",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Photorealistic premium food commercial, extreme macro high-speed product cinematography. A single ripe strawberry falls toward a white ceramic bowl in a sunlit spring greenhouse. Its red seeds and moist surface are sharply resolved against soft green glasshouse bokeh. Keep the strawberry's screen position, size, orientation and downward motion consistent through every seasonal match cut. No hands, labels or additional fruit.\n[Shot 2] At 00:03.000, the shot transitions to a hot summer picnic in a seamless seasonal match cut around the same falling strawberry and bowl. Warm sunlight replaces the greenhouse light; distant insects accompany its uninterrupted descent.\n[Shot 3] At 00:06.000, the shot transitions to an autumn kitchen with amber window light, preserving the strawberry and bowl alignment and the same downward motion. Soft indoor room tone replaces the outdoor insects.\n[Shot 4] At 00:09.000, the shot transitions to a candlelit winter dining room. The strawberry completes its fall into fresh cream in the white bowl, sending one graceful crown-shaped splash upward. Hold the detailed splash and its settling ripples through the final seconds. One continuous musical phrase connects all four seasons.\n\noverall_soundscape: Spring birds and a gentle greenhouse breeze change to summer insects, quiet autumn kitchen ambience, then the still winter dining room. The final cream splash and falling droplets remain clearly audible.\n\nnon_diegetic_music: One uninterrupted light piano-and-plucked-string phrase runs at a moderate tempo across all four seasons, resolving with the strawberry splash.",
      "effective_prompt_sha256": "0c4614d6b1c88a86974b1393eab7a574817abc11c37bf72e3cb53e2d9bbf718f"
    },
    {
      "id": 5,
      "slug": "stunning-action-packed-scene-following-a-little-robot-r-288619",
      "title": {
        "en": "Little Robot Cyberpunk Escape",
        "zh": "赛博城市里的机器人追逐"
      },
      "category": "cinematic-story",
      "source": {
        "kind": "x",
        "name": "@AIandDesign",
        "url": "https://x.com/AIandDesign/status/2082522979339288619"
      },
      "source_prompt": "A stunning action packed scene following a little robot running away from soldiers, drones and other things that are trying to get him. Chase cam, dystopian futuristic cyberpunk environment, cinematic lighting. The camera follows the robot from behind, chase cam style. Every now and then the robot looks behind him and the entities chasing it are shown.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] A stunning action packed scene following a little robot running away from soldiers, drones and other things that are trying to get him. Chase cam, dystopian futuristic cyberpunk environment, cinematic lighting. The camera follows the robot from behind, chase cam style. Every now and then the robot looks behind him and the entities chasing it are shown. Maintain one continuous 15-second rear chase shot with brief arcs that reveal the pursuers when the robot looks back; preserve the robot’s identity and forward motion.\n\noverall_soundscape: Rapid metallic footfalls and small robot servo movements run through the city ambience. Heavy pursuing footsteps, drone rotors and passing machinery track the chase without intelligible dialogue.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "be75f61112b0d84692fb5ce91ced32ae50c3f1a3a1ad6b115fba2089944d4073",
      "duration": 15,
      "seed": 46,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Little Robot Cyberpunk Escape",
        "zh": "小机器人赛博朋克逃亡"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-05",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] A stunning action packed scene following a little robot running away from soldiers, drones and other things that are trying to get him. Chase cam, dystopian futuristic cyberpunk environment, cinematic lighting. The camera follows the robot from behind, chase cam style. Every now and then the robot looks behind him and the entities chasing it are shown. Maintain one continuous 15-second rear chase shot with brief arcs that reveal the pursuers when the robot looks back; preserve the robot’s identity and forward motion.\n\noverall_soundscape: Rapid metallic footfalls and small robot servo movements run through the city ambience. Heavy pursuing footsteps, drone rotors and passing machinery track the chase without intelligible dialogue.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "be75f61112b0d84692fb5ce91ced32ae50c3f1a3a1ad6b115fba2089944d4073"
    },
    {
      "id": 6,
      "slug": "created-with-minimax-h3-max-120666",
      "title": {
        "en": "Created with MiniMax H3 Max",
        "zh": "湖畔花束与风"
      },
      "category": "cinematic-story",
      "source": {
        "kind": "x",
        "name": "@TaliaAariz",
        "url": "https://x.com/TaliaAariz/status/2095857062987120666"
      },
      "source_prompt": "A cinematic medium shot of a gentle East Asian young woman with dark hair styled in a bun adorned with a blue-and-white checkered bow, standing on a sunny pebbled lakefront. She wears a sleeveless light blue polka-dot tulle dress with long, translucent white ribbons flowing from her shoulder that flutter gracefully in a gentle breeze. Holding a lush bouquet of pastel pink, baby blue, and white flowers, she smiles warmly at the camera, softly touches the flower petals, and tilts her head back slightly to close her eyes and feel the breeze. In the background, glistening blue lake water sparkles under the sun, bordered by tall green reeds along the shore and serene blue mountains under a bright, clear sky, captured in soft natural daylight with a shallow depth of field.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] A cinematic medium shot of a gentle East Asian young woman with dark hair styled in a bun adorned with a blue-and-white checkered bow, standing on a sunny pebbled lakefront. She wears a sleeveless light blue polka-dot tulle dress with long, translucent white ribbons flowing from her shoulder that flutter gracefully in a gentle breeze. Holding a lush bouquet of pastel pink, baby blue, and white flowers, she smiles warmly at the camera, softly touches the flower petals, and tilts her head back slightly to close her eyes and feel the breeze. In the background, glistening blue lake water sparkles under the sun, bordered by tall green reeds along the shore and serene blue mountains under a bright, clear sky, captured in soft natural daylight with a shallow depth of field. Keep one continuous 15-second medium shot, with a very slow small push-in as she touches the flowers, smiles and finally closes her eyes to feel the breeze.\n\noverall_soundscape: Gentle shore wind moves the ribbons and reeds while small waves wash over the pebbles. Quiet bouquet and fabric rustles accompany her hand movements.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "7e05bce3aa444d99e753af9bcc4865bd2934e2c9102f5a1ce63f952302516d15",
      "duration": 15,
      "seed": 47,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Created with MiniMax H3 Max",
        "zh": "Created with MiniMax H3 Max"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-06",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] A cinematic medium shot of a gentle East Asian young woman with dark hair styled in a bun adorned with a blue-and-white checkered bow, standing on a sunny pebbled lakefront. She wears a sleeveless light blue polka-dot tulle dress with long, translucent white ribbons flowing from her shoulder that flutter gracefully in a gentle breeze. Holding a lush bouquet of pastel pink, baby blue, and white flowers, she smiles warmly at the camera, softly touches the flower petals, and tilts her head back slightly to close her eyes and feel the breeze. In the background, glistening blue lake water sparkles under the sun, bordered by tall green reeds along the shore and serene blue mountains under a bright, clear sky, captured in soft natural daylight with a shallow depth of field. Keep one continuous 15-second medium shot, with a very slow small push-in as she touches the flowers, smiles and finally closes her eyes to feel the breeze.\n\noverall_soundscape: Gentle shore wind moves the ribbons and reeds while small waves wash over the pebbles. Quiet bouquet and fabric rustles accompany her hand movements.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "7e05bce3aa444d99e753af9bcc4865bd2934e2c9102f5a1ce63f952302516d15"
    },
    {
      "id": 7,
      "slug": "created-with-minimax-h3-on-989121",
      "title": {
        "en": "Created with MiniMax H3 on .",
        "zh": "画家与湖畔城市天际线"
      },
      "category": "cinematic-travel",
      "source": {
        "kind": "x",
        "name": "@aiwithaly",
        "url": "https://x.com/aiwithaly/status/2087408671101989121"
      },
      "source_prompt": "A young Western artist sets up an easel in a peaceful city park and paints the skyline overlooking a lake. As the artwork comes to life, she compares the finished painting with the real view, holding the canvas in front of the skyline and smiling proudly. Ultra-photorealistic visuals, cinematic documentary style, realistic brush strokes, authentic hand movements, natural English dialogue, handheld camera movement, detailed cityscape, premium lighting, seamless continuity, and immersive outdoor ambience.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] A young Western artist with a relaxed conversational English voice (S1) sets up an easel in a peaceful city park and paints the skyline overlooking a lake. As the artwork comes to life, she compares the finished painting with the real view, holding the canvas in front of the skyline and smiling proudly. Ultra-photorealistic visuals, cinematic documentary style, realistic brush strokes, authentic hand movements, natural English dialogue, handheld camera movement, detailed cityscape, premium lighting, seamless continuity, and immersive outdoor ambience. Her brief natural English remarks follow her painting actions. Keep the artist, easel and skyline consistent over the 15-second take. Begin with her positioning the easel, follow several visible finishing brush strokes, then give the final canvas-to-skyline comparison time to read.\n\noverall_soundscape: Soft city-park wind, distant traffic and lakeside birds surround the easel. Light wood movements, brush-on-canvas friction and cloth handling follow the artist’s actions.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "f39be188be70b922d6337aa1436e276c98592d5adf01475408b43bbd227e6cbc",
      "duration": 15,
      "seed": 48,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源文未提供固定台词；仅保留原文要求的语言与说话动作。"
      ],
      "source_title": {
        "en": "Created with MiniMax H3 on .",
        "zh": "Created with MiniMax H3 on ."
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-07",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] A young Western artist with a relaxed conversational English voice (S1) sets up an easel in a peaceful city park and paints the skyline overlooking a lake. As the artwork comes to life, she compares the finished painting with the real view, holding the canvas in front of the skyline and smiling proudly. Ultra-photorealistic visuals, cinematic documentary style, realistic brush strokes, authentic hand movements, natural English dialogue, handheld camera movement, detailed cityscape, premium lighting, seamless continuity, and immersive outdoor ambience. Her brief natural English remarks follow her painting actions. Keep the artist, easel and skyline consistent over the 15-second take. Begin with her positioning the easel, follow several visible finishing brush strokes, then give the final canvas-to-skyline comparison time to read.\n\noverall_soundscape: Soft city-park wind, distant traffic and lakeside birds surround the easel. Light wood movements, brush-on-canvas friction and cloth handling follow the artist’s actions.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "f39be188be70b922d6337aa1436e276c98592d5adf01475408b43bbd227e6cbc"
    },
    {
      "id": 8,
      "slug": "created-with-minimax-h3-522089",
      "title": {
        "en": "Created with MiniMax H3.",
        "zh": "街头摄影师、老人和小狗"
      },
      "category": "music-video",
      "source": {
        "kind": "x",
        "name": "@aiwithaly",
        "url": "https://x.com/aiwithaly/status/2087102541146522089"
      },
      "source_prompt": "A young Western female street photographer walks through a lively downtown street and notices an elderly man sitting outside a café with his small dog. She carefully composes the candid moment through her camera, captures the photo, then turns the camera toward the viewer to proudly show the shot she just took. She smiles, says “Look at that,” then continues walking through the city. Ultra-photorealistic visuals, natural handheld documentary movement, realistic camera interaction, authentic facial expressions, accurate hand movements, realistic dog behavior, natural daylight, cinematic depth of field, continuous character consistency, immersive city ambience, premium documentary realism.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] A young Western female street photographer with a warm conversational English voice (S1) walks through a lively downtown street and notices an elderly man sitting outside a café with his small dog. She carefully composes the candid moment through her camera, captures the photo, then turns the camera toward the viewer to proudly show the shot she just took. She smiles, says: <d>[English] Look at that,</d> then continues walking through the city. Ultra-photorealistic visuals, natural handheld documentary movement, realistic camera interaction, authentic facial expressions, accurate hand movements, realistic dog behavior, natural daylight, cinematic depth of field, continuous character consistency, immersive city ambience, premium documentary realism. Over one continuous 15-second handheld shot, follow her noticing and photographing the café scene, then turn with her as she shows the camera display, speaks and continues walking.\n\noverall_soundscape: Layered downtown footsteps, distant traffic and soft café room spill continue beneath a camera shutter click. The dog shifts and lightly pants beside the seated man.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "be91933cc24647e7105ff883517384c8c83c2ccef42801f72555f509c467e319",
      "duration": 15,
      "seed": 49,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Created with MiniMax H3.",
        "zh": "Created with MiniMax H3."
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-08",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] A young Western female street photographer with a warm conversational English voice (S1) walks through a lively downtown street and notices an elderly man sitting outside a café with his small dog. She carefully composes the candid moment through her camera, captures the photo, then turns the camera toward the viewer to proudly show the shot she just took. She smiles, says: <d>[English] Look at that,</d> then continues walking through the city. Ultra-photorealistic visuals, natural handheld documentary movement, realistic camera interaction, authentic facial expressions, accurate hand movements, realistic dog behavior, natural daylight, cinematic depth of field, continuous character consistency, immersive city ambience, premium documentary realism. Over one continuous 15-second handheld shot, follow her noticing and photographing the café scene, then turn with her as she shows the camera display, speaks and continues walking.\n\noverall_soundscape: Layered downtown footsteps, distant traffic and soft café room spill continue beneath a camera shutter click. The dog shifts and lightly pants beside the seated man.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "be91933cc24647e7105ff883517384c8c83c2ccef42801f72555f509c467e319"
    },
    {
      "id": 9,
      "slug": "low-angle-fashion-tracking-film-062019",
      "title": {
        "en": "Low-Angle Fashion Tracking Film",
        "zh": "小巷里的时装与歌声"
      },
      "category": "fashion",
      "source": {
        "kind": "x",
        "name": "@Kiber_Alla",
        "url": "https://x.com/Kiber_Alla/status/2083583963512062019"
      },
      "source_prompt": "Ultra realistic cinematic video, shot on ARRI Alexa 35 with spherical lens, natural daylight, photorealistic, looks like real footage filmed on location, high-end fashion film aesthetic, extreme detail, realistic skin texture, real fabric physics, subtle film grain. Continuous low angle tracking shot, camera smoothly moving backward as a beautiful young woman with long dark wind-blown hair walks and performs down a narrow sun-drenched urban alley between old textured stone buildings under deep blue sky. She wears a flowing milky white asymmetrical off-shoulder top and high-waist wide-leg technical cotton culottes with a thin rope belt. The wide pants and fabric move naturally and dramatically with the wind and her body. She sings with genuine joy, smiling and softly laughing, maintaining strong eye contact with the camera, while performing free, light dance elements — fluid arm gestures, body swaying, and elegant full turns around herself as she continues walking forward. Her movement is energetic yet effortless, full of lightness and positive energy. Cinematic natural lighting with strong sunlight and deep shadows, realistic motion blur, high dynamic range, masterpiece, best quality, no artificial look.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Ultra realistic cinematic video, shot on ARRI Alexa 35 with spherical lens, natural daylight, photorealistic, looks like real footage filmed on location, high-end fashion film aesthetic, extreme detail, realistic skin texture, real fabric physics, subtle film grain. Continuous low angle tracking shot, camera smoothly moving backward as a beautiful young woman with long dark wind-blown hair walks and performs down a narrow sun-drenched urban alley between old textured stone buildings under deep blue sky. She wears a flowing milky white asymmetrical off-shoulder top and high-waist wide-leg technical cotton culottes with a thin rope belt. The wide pants and fabric move naturally and dramatically with the wind and her body. The woman, with a light melodic singing voice (S1), sings with genuine joy, smiling and softly laughing, maintaining strong eye contact with the camera, while performing free, light dance elements — fluid arm gestures, body swaying, and elegant full turns around herself as she continues walking forward. Her movement is energetic yet effortless, full of lightness and positive energy. Cinematic natural lighting with strong sunlight and deep shadows, realistic motion blur, high dynamic range, masterpiece, best quality, no artificial look. She performs a wordless vocal melody. Keep the low-angle backward tracking shot continuous for the full 15 seconds. Leave enough space for a full turn and return to eye contact without cutting.\n\noverall_soundscape: Light footsteps on stone, wind through the alley and natural fabric movement accompany the performance. Her soft laughter and breaths are synchronized to the visible movement.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "7c2058cc5dc84a848b29169d61ecc607158e3ce4e68dd07d43622c10bfd34d91",
      "duration": 15,
      "seed": 50,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源文未提供固定台词或歌词；未编造逐字台词，歌唱使用无词旋律。"
      ],
      "source_title": {
        "en": "Low-Angle Fashion Tracking Film",
        "zh": "低机位时尚跟拍短片"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-09",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Ultra realistic cinematic video, shot on ARRI Alexa 35 with spherical lens, natural daylight, photorealistic, looks like real footage filmed on location, high-end fashion film aesthetic, extreme detail, realistic skin texture, real fabric physics, subtle film grain. Continuous low angle tracking shot, camera smoothly moving backward as a beautiful young woman with long dark wind-blown hair walks and performs down a narrow sun-drenched urban alley between old textured stone buildings under deep blue sky. She wears a flowing milky white asymmetrical off-shoulder top and high-waist wide-leg technical cotton culottes with a thin rope belt. The wide pants and fabric move naturally and dramatically with the wind and her body. The woman, with a light melodic singing voice (S1), sings with genuine joy, smiling and softly laughing, maintaining strong eye contact with the camera, while performing free, light dance elements — fluid arm gestures, body swaying, and elegant full turns around herself as she continues walking forward. Her movement is energetic yet effortless, full of lightness and positive energy. Cinematic natural lighting with strong sunlight and deep shadows, realistic motion blur, high dynamic range, masterpiece, best quality, no artificial look. She performs a wordless vocal melody. Keep the low-angle backward tracking shot continuous for the full 15 seconds. Leave enough space for a full turn and return to eye contact without cutting.\n\noverall_soundscape: Light footsteps on stone, wind through the alley and natural fabric movement accompany the performance. Her soft laughter and breaths are synchronized to the visible movement.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "7c2058cc5dc84a848b29169d61ecc607158e3ce4e68dd07d43622c10bfd34d91"
    },
    {
      "id": 10,
      "slug": "generated-a-french-bulldog-with-minimax-h3-182507",
      "title": {
        "en": "Generated a french bulldog with MiniMax H3",
        "zh": "法国斗牛犬闯障碍"
      },
      "category": "cinematic-story",
      "source": {
        "kind": "x",
        "name": "@cocktailpeanut",
        "url": "https://x.com/cocktailpeanut/status/2085448660612182507"
      },
      "source_prompt": "A casual handheld smartphone video filmed by the dog's owner in an ordinary suburban backyard. A cute, stocky French bulldog excitedly runs through a small homemade agility course, awkwardly weaving between plastic poles, hopping over two low hurdles, scrambling through a short fabric tunnel, and finally bumping a little bell with its nose.",
      "source_duration": "14s",
      "prompt": "integrated_multimodal_description: [Shot 1] A casual handheld smartphone video filmed by the dog's owner in an ordinary suburban backyard. A cute, stocky French bulldog excitedly runs through a small homemade agility course, awkwardly weaving between plastic poles, hopping over two low hurdles, scrambling through a short fabric tunnel, and finally bumping a little bell with its nose. Follow the complete obstacle sequence in one 15-second handheld take, with the bell contact visible in the final seconds.\n\noverall_soundscape: Backyard birds and distant neighborhood ambience surround quick paw taps and excited dog panting. Plastic poles rattle, the fabric tunnel scrapes, and a small bell rings once at the finish.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "b441f7d70cc0c7a54efccb096492e34839955225bc31781700cf560af438bc8d",
      "duration": 15,
      "seed": 51,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源条目元数据为 14s，本次明确扩展为 15 秒。"
      ],
      "source_title": {
        "en": "Generated a french bulldog with MiniMax H3",
        "zh": "Generated a french bulldog with MiniMax H3"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-10",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] A casual handheld smartphone video filmed by the dog's owner in an ordinary suburban backyard. A cute, stocky French bulldog excitedly runs through a small homemade agility course, awkwardly weaving between plastic poles, hopping over two low hurdles, scrambling through a short fabric tunnel, and finally bumping a little bell with its nose. Follow the complete obstacle sequence in one 15-second handheld take, with the bell contact visible in the final seconds.\n\noverall_soundscape: Backyard birds and distant neighborhood ambience surround quick paw taps and excited dog panting. Plastic poles rattle, the fabric tunnel scrapes, and a small bell rings once at the finish.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "b441f7d70cc0c7a54efccb096492e34839955225bc31781700cf560af438bc8d"
    },
    {
      "id": 11,
      "slug": "created-with-hailuo-h3-182619",
      "title": {
        "en": "Created with Hailuo H3.",
        "zh": "冰川崩塌组成 EXTINCTION"
      },
      "category": "title-sequence",
      "source": {
        "kind": "x",
        "name": "@CharaspowerAI",
        "url": "https://x.com/CharaspowerAI/status/2090740724803182619"
      },
      "source_prompt": "Ultra cinematic title sequence: A colossal glacier fractures across an endless frozen ocean. Deep blue cracks spread for miles beneath translucent ice before the entire shelf collapses. Mountains of ice are launched into the sky, suspended in breathtaking slow motion. Every crystal catches the sunlight while frozen debris rotates around the camera. The crystalline storm assembles the word \"EXTINCTION\", carved from pure glacial ice with intricate frozen textures. Hairline fractures slowly spread through every letter until the title violently shatters into millions of glittering ice fragments. Photorealistic ice simulation, IMAX-scale environmental VFX.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Ultra cinematic title sequence: A colossal glacier fractures across an endless frozen ocean. Deep blue cracks spread for miles beneath translucent ice before the entire shelf collapses. Mountains of ice are launched into the sky, suspended in breathtaking slow motion. Every crystal catches the sunlight while frozen debris rotates around the camera. The crystalline storm assembles the word \"EXTINCTION\", carved from pure glacial ice with intricate frozen textures. Hairline fractures slowly spread through every letter until the title violently shatters into millions of glittering ice fragments. Photorealistic ice simulation, IMAX-scale environmental VFX. Over 15 seconds, establish the glacier, follow the expanding fractures and airborne ice, hold the completed readable ice title, then finish on its violent shatter.\n\noverall_soundscape: Deep frozen-ocean wind surrounds spreading ice creaks and resonant fractures. Heavy ice collapse becomes a crystalline debris storm, followed by the sharp final title shatter.\n\nnon_diegetic_music: Sustained low strings swell with the ice fracture, joined by deep percussion at the shelf collapse and final title shatter.",
      "prompt_sha256": "4450f406537c05f4eaf4263bdbb5e1e4d60f9b8e253dbe76aac8d0be5e2b24bd",
      "duration": 15,
      "seed": 52,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Created with Hailuo H3.",
        "zh": "Created with Hailuo H3."
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-11",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Ultra cinematic title sequence: A colossal glacier fractures across an endless frozen ocean. Deep blue cracks spread for miles beneath translucent ice before the entire shelf collapses. Mountains of ice are launched into the sky, suspended in breathtaking slow motion. Every crystal catches the sunlight while frozen debris rotates around the camera. The crystalline storm assembles the word \"EXTINCTION\", carved from pure glacial ice with intricate frozen textures. Hairline fractures slowly spread through every letter until the title violently shatters into millions of glittering ice fragments. Photorealistic ice simulation, IMAX-scale environmental VFX. Over 15 seconds, establish the glacier, follow the expanding fractures and airborne ice, hold the completed readable ice title, then finish on its violent shatter.\n\noverall_soundscape: Deep frozen-ocean wind surrounds spreading ice creaks and resonant fractures. Heavy ice collapse becomes a crystalline debris storm, followed by the sharp final title shatter.\n\nnon_diegetic_music: Sustained low strings swell with the ice fracture, joined by deep percussion at the shelf collapse and final title shatter.",
      "effective_prompt_sha256": "4450f406537c05f4eaf4263bdbb5e1e4d60f9b8e253dbe76aac8d0be5e2b24bd"
    },
    {
      "id": 12,
      "slug": "shatter-695817",
      "title": {
        "en": "SHATTER.",
        "zh": "玻璃碎片组成 SHATTER"
      },
      "category": "music-video",
      "source": {
        "kind": "x",
        "name": "@CharaspowerAI",
        "url": "https://x.com/CharaspowerAI/status/2096554023809695817"
      },
      "source_prompt": "Premium action-film title sequence beginning with a gigantic transparent glass monolith floating against complete blackness. A tiny projectile impacts the exact center and an intricate network of cracks instantly races across the entire surface. The camera pushes directly into one expanding fracture, traveling through the microscopic crystalline structure as tension builds everywhere. Suddenly the monolith violently detonates into millions of glass fragments. The camera performs a fast 360-degree orbital move through the suspended crystal storm as every shard reflects sharp beams of cinematic light. An invisible force begins pulling selected fragments together, constructing the enormous word \"SHATTER\" from overlapping razor-thin glass shards. The completed transparent typography refracts the entire environment like a giant prism before another pressure wave passes through it, producing thousands of secondary fractures. The letters explode outward directly toward the lens in spectacular ultra slow motion. Extremely detailed glass physics, realistic refraction and caustics, elegant brutality, premium Hollywood VFX.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Premium action-film title sequence beginning with a gigantic transparent glass monolith floating against complete blackness. A tiny projectile impacts the exact center and an intricate network of cracks instantly races across the entire surface. The camera pushes directly into one expanding fracture, traveling through the microscopic crystalline structure as tension builds everywhere. Suddenly the monolith violently detonates into millions of glass fragments. The camera performs a fast 360-degree orbital move through the suspended crystal storm as every shard reflects sharp beams of cinematic light. An invisible force begins pulling selected fragments together, constructing the enormous word \"SHATTER\" from overlapping razor-thin glass shards. The completed transparent typography refracts the entire environment like a giant prism before another pressure wave passes through it, producing thousands of secondary fractures. The letters explode outward directly toward the lens in spectacular ultra slow motion. Extremely detailed glass physics, realistic refraction and caustics, elegant brutality, premium Hollywood VFX. Keep the camera move continuous over 15 seconds: establish the floating monolith, enter the spreading fracture, orbit the suspended fragments, hold the completed title briefly, then finish on fragments expanding toward the lens.\n\noverall_soundscape: A projectile impact triggers fine traveling glass cracks and a swelling structural groan. Dense crystalline shattering and passing shards follow the orbit, with a second pressure-wave crack and outward burst at the end.\n\nnon_diegetic_music: A sparse low electronic pulse builds beneath bowed-metal tones, with two deep percussion hits at the main and final glass explosions.",
      "prompt_sha256": "b198e1ddb0d1d5c79cab984ce06e7944ecd4efbca41f105960c7312c540dd572",
      "duration": 15,
      "seed": 53,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "SHATTER.",
        "zh": "SHATTER."
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-12",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Premium action-film title sequence beginning with a gigantic transparent glass monolith floating against complete blackness. A tiny projectile impacts the exact center and an intricate network of cracks instantly races across the entire surface. The camera pushes directly into one expanding fracture, traveling through the microscopic crystalline structure as tension builds everywhere. Suddenly the monolith violently detonates into millions of glass fragments. The camera performs a fast 360-degree orbital move through the suspended crystal storm as every shard reflects sharp beams of cinematic light. An invisible force begins pulling selected fragments together, constructing the enormous word \"SHATTER\" from overlapping razor-thin glass shards. The completed transparent typography refracts the entire environment like a giant prism before another pressure wave passes through it, producing thousands of secondary fractures. The letters explode outward directly toward the lens in spectacular ultra slow motion. Extremely detailed glass physics, realistic refraction and caustics, elegant brutality, premium Hollywood VFX. Keep the camera move continuous over 15 seconds: establish the floating monolith, enter the spreading fracture, orbit the suspended fragments, hold the completed title briefly, then finish on fragments expanding toward the lens.\n\noverall_soundscape: A projectile impact triggers fine traveling glass cracks and a swelling structural groan. Dense crystalline shattering and passing shards follow the orbit, with a second pressure-wave crack and outward burst at the end.\n\nnon_diegetic_music: A sparse low electronic pulse builds beneath bowed-metal tones, with two deep percussion hits at the main and final glass explosions.",
      "effective_prompt_sha256": "b198e1ddb0d1d5c79cab984ce06e7944ecd4efbca41f105960c7312c540dd572"
    },
    {
      "id": 13,
      "slug": "shot-opens-tight-on-a-beautiful-woman-face-wind-452380",
      "title": {
        "en": "Cliffside Descent to the Ocean",
        "zh": "从悬崖人像到俯拍海浪"
      },
      "category": "cinematic-story",
      "source": {
        "kind": "x",
        "name": "@umesh_ai",
        "url": "https://x.com/umesh_ai/status/2082700637444452380"
      },
      "source_prompt": "The shot opens tight on a beautiful woman face, wind tearing at her coat as she stands on a jagged cliff. The camera eases over her shoulder and tilts down, following her line of sight. We descend past the crumbling lip into the void, the frame widening to reveal cliffs dropping away. The ocean roars into view, violent surf detonating against rock. The move continues into a high overhead, locking into an aerial top-down that frames the lone figure at the precipice as waves explode below.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] The shot opens tight on a beautiful woman face, wind tearing at her coat as she stands on a jagged cliff. The camera eases over her shoulder and tilts down, following her line of sight. We descend past the crumbling lip into the void, the frame widening to reveal cliffs dropping away. The ocean roars into view, violent surf detonating against rock. The move continues into a high overhead, locking into an aerial top-down that frames the lone figure at the precipice as waves explode below. Perform the described camera path as one continuous 15-second move, ending on the high overhead view of the lone figure and exploding surf.\n\noverall_soundscape: Strong coastal wind drives the coat fabric. Deep ocean roar becomes clearer as the camera descends, with violent waves striking rock beneath the final overhead view.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "8a9fb882646cdb521b55ea777eea72c255e9b3863f257602d34d5bda4da3b879",
      "duration": 15,
      "seed": 54,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Cliffside Descent to the Ocean",
        "zh": "悬崖下行至海边"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-13",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] The shot opens tight on a beautiful woman face, wind tearing at her coat as she stands on a jagged cliff. The camera eases over her shoulder and tilts down, following her line of sight. We descend past the crumbling lip into the void, the frame widening to reveal cliffs dropping away. The ocean roars into view, violent surf detonating against rock. The move continues into a high overhead, locking into an aerial top-down that frames the lone figure at the precipice as waves explode below. Perform the described camera path as one continuous 15-second move, ending on the high overhead view of the lone figure and exploding surf.\n\noverall_soundscape: Strong coastal wind drives the coat fabric. Deep ocean roar becomes clearer as the camera descends, with violent waves striking rock beneath the final overhead view.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "8a9fb882646cdb521b55ea777eea72c255e9b3863f257602d34d5bda4da3b879"
    },
    {
      "id": 14,
      "slug": "ugc-style-skincare-video-featuring-a-realistic-young-wo-242986",
      "title": {
        "en": "Bathroom Mirror Skincare UGC Ad",
        "zh": "镜前护肤产品口播"
      },
      "category": "product-commercial",
      "source": {
        "kind": "x",
        "name": "@ZaraIrahh",
        "url": "https://x.com/ZaraIrahh/status/2083011066800242986"
      },
      "source_prompt": "UGC-style skincare video featuring a realistic young woman speaking directly to the camera in front of a bathroom mirror while holding a premium facial serum bottle. She smiles naturally, talks about her skincare routine, and demonstrates applying the serum to her face. Authentic, conversational presentation with genuine facial expressions and hand gestures, natural morning light, casual home bathroom setting, clean skincare aesthetic, healthy glowing skin, minimal makeup, smartphone-shot look, candid lifestyle beauty content, soft depth of field, luxury beauty brand feel, vertical 9:16, social media ad quality, highly realistic natural skin texture, not overly retouched.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] UGC-style skincare video featuring a realistic young woman with a clear conversational English voice (S1) speaking directly to the camera in front of a bathroom mirror while holding a premium facial serum bottle. She smiles naturally, talks about her skincare routine, and demonstrates applying the serum to her face. Authentic, conversational presentation with genuine facial expressions and hand gestures, natural morning light, casual home bathroom setting, clean skincare aesthetic, healthy glowing skin, minimal makeup, smartphone-shot look, candid lifestyle beauty content, soft depth of field, luxury beauty brand feel, landscape 16:9, social media ad quality, highly realistic natural skin texture, not overly retouched. Her short natural English presentation stays about her skincare routine. Use a continuous 15-second landscape smartphone shot framed to keep both her face and the bottle clearly visible. She presents the serum, applies it to her face, and returns her attention to the camera.\n\noverall_soundscape: Quiet bathroom room tone, small glass-bottle handling sounds and soft hand-on-skin movements accompany the presentation. Keep breaths and clothing movement natural and close.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "61f30fc7218593117f92f6c7e46d55e03de10b31642a4122a69a9eba43d40b4c",
      "duration": 15,
      "seed": 55,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源文未提供固定台词；仅保留原文要求的语言与说话动作。"
      ],
      "source_title": {
        "en": "Bathroom Mirror Skincare UGC Ad",
        "zh": "浴室镜前护肤 UGC 广告"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-14",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] UGC-style skincare video featuring a realistic young woman with a clear conversational English voice (S1) speaking directly to the camera in front of a bathroom mirror while holding a premium facial serum bottle. She smiles naturally, talks about her skincare routine, and demonstrates applying the serum to her face. Authentic, conversational presentation with genuine facial expressions and hand gestures, natural morning light, casual home bathroom setting, clean skincare aesthetic, healthy glowing skin, minimal makeup, smartphone-shot look, candid lifestyle beauty content, soft depth of field, luxury beauty brand feel, landscape 16:9, social media ad quality, highly realistic natural skin texture, not overly retouched. Her short natural English presentation stays about her skincare routine. Use a continuous 15-second landscape smartphone shot framed to keep both her face and the bottle clearly visible. She presents the serum, applies it to her face, and returns her attention to the camera.\n\noverall_soundscape: Quiet bathroom room tone, small glass-bottle handling sounds and soft hand-on-skin movements accompany the presentation. Keep breaths and clothing movement natural and close.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "61f30fc7218593117f92f6c7e46d55e03de10b31642a4122a69a9eba43d40b4c"
    },
    {
      "id": 15,
      "slug": "found-footage-film-with-335284",
      "title": {
        "en": "Horror Film Study 335284",
        "zh": "渔夫接到诡异来电"
      },
      "category": "horror",
      "source": {
        "kind": "x",
        "name": "@cocktailpeanut",
        "url": "https://x.com/cocktailpeanut/status/2084842607411335284"
      },
      "source_prompt": "\"Hyper-realistic found-footage horror video, filmed by a cheap action camera mounted inside a small fishing boat on a quiet lake. A lone fisherman violently reels in something heavy. He pulls up a muddy smartphone tangled in weeds. He answers, breathing hard. A distorted voice says, \"Finally. I’ve been trying to switch places with you.\" The phone flashes blinding white and the footage glitches. When the image returns, the phone shows the terrified fisherman trapped inside its screen, silently pounding on the glass. Behind him, an identical soaking-wet fisherman slowly climbs out of the lake into the boat, notices the camera, and reaches toward the lens. Abrupt cut to black. Raw handheld realism, imperfect autofocus, clipped audio, no music, no stylization.\"",
      "source_duration": "14s",
      "prompt": "integrated_multimodal_description: [Shot 1] \"Hyper-realistic found-footage horror video, filmed by a cheap action camera mounted inside a small fishing boat on a quiet lake. A lone fisherman violently reels in something heavy. He pulls up a muddy smartphone tangled in weeds. He answers, breathing hard. A low distorted English voice from the phone speaker (S1) says: <d>[English] Finally. I’ve been trying to switch places with you.</d> The phone flashes blinding white and the footage glitches. When the image returns, the phone shows the terrified fisherman trapped inside its screen, silently pounding on the glass. Behind him, an identical soaking-wet fisherman slowly climbs out of the lake into the boat, notices the camera, and reaches toward the lens. Abrupt cut to black. Raw handheld realism, imperfect autofocus, clipped audio, no music, no stylization.\" Keep the cheap boat-mounted camera viewpoint continuous for the 15-second scene, with only the specified white-flash glitch interrupting the image. Establish the fishing action, allow the phone line to finish, show the trapped image and then the wet double before the abrupt final black frame.\n\noverall_soundscape: Low lake water lapping, boat creaks and fishing-reel strain accompany the fisherman’s heavy breathing. Wet weeds scrape the hull, the phone glitches with clipped electronic noise, then water splashes as the soaked double enters the boat.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "778e3b3e9a6ca16d59312279338db624dd6281b72e5ade48402f63ed4b25f617",
      "duration": 15,
      "seed": 56,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源条目元数据为 14s，本次明确扩展为 15 秒。"
      ],
      "source_title": {
        "en": "Horror Film Study 335284",
        "zh": "恐怖短片案例 335284"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-15",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] \"Hyper-realistic found-footage horror video, filmed by a cheap action camera mounted inside a small fishing boat on a quiet lake. A lone fisherman violently reels in something heavy. He pulls up a muddy smartphone tangled in weeds. He answers, breathing hard. A low distorted English voice from the phone speaker (S1) says: <d>[English] Finally. I’ve been trying to switch places with you.</d> The phone flashes blinding white and the footage glitches. When the image returns, the phone shows the terrified fisherman trapped inside its screen, silently pounding on the glass. Behind him, an identical soaking-wet fisherman slowly climbs out of the lake into the boat, notices the camera, and reaches toward the lens. Abrupt cut to black. Raw handheld realism, imperfect autofocus, clipped audio, no music, no stylization.\" Keep the cheap boat-mounted camera viewpoint continuous for the 15-second scene, with only the specified white-flash glitch interrupting the image. Establish the fishing action, allow the phone line to finish, show the trapped image and then the wet double before the abrupt final black frame.\n\noverall_soundscape: Low lake water lapping, boat creaks and fishing-reel strain accompany the fisherman’s heavy breathing. Wet weeds scrape the hull, the phone glitches with clipped electronic noise, then water splashes as the soaked double enters the boat.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "778e3b3e9a6ca16d59312279338db624dd6281b72e5ade48402f63ed4b25f617"
    },
    {
      "id": 16,
      "slug": "cinematic-sci-fi-suspense-video-a-realistic-spaceship-crew-in-024495",
      "title": {
        "en": "cinematic sci-fi suspense video. A realistic spaceship crew in full futuristic",
        "zh": "太空船门后的小学教室"
      },
      "category": "cinematic-story",
      "source": {
        "kind": "x",
        "name": "@cocktailpeanut",
        "url": "https://x.com/cocktailpeanut/status/2085498501384024495"
      },
      "source_prompt": "cinematic sci-fi suspense video. A realistic spaceship crew in full futuristic uniforms follows a mysterious distress signal and boards a dark, abandoned spacecraft drifting in space. Flashlights cut through darkness as they move cautiously through narrow metallic corridors. Warning lights flicker, metal creaks, and their boots echo on the floor. One crew member says, \"Signal's coming from deeper inside.\" They enter a large chamber filled with hundreds of empty cryogenic pods, all open and abandoned. Another crew member says nervously, \"Where did everyone go?\" They approach a final closed door at the far end. The captain slowly opens it, and instead of another ship compartment, it reveals a completely normal present-day elementary school classroom in bright daylight. Children sit at desks, backpacks on the floor, colorful posters on the walls, and a teacher at the front turns to the crew and scolds them: \"You're late again. Take your seats.\" The astronauts stand frozen in confusion at the doorway. Realistic tone, smooth camera movement, strong suspense buildup, with the final reveal clear, sudden, and absurd. Include realistic audio: radio static, low ship hum, footsteps, pod machinery ambience, door hiss, then ordinary classroom sounds and the teacher's voice.",
      "source_duration": "14s",
      "prompt": "integrated_multimodal_description: [Shot 1] Live-action cinematic science-fiction suspense, realistic spaceship crew in full futuristic uniforms boarding a dark abandoned spacecraft in response to a distress signal. A smooth tracking shot follows them through a narrow metallic corridor; flashlight beams cross flickering warning lights. One crew member with a restrained mid-pitched English voice (S1) says: <d>[English] Signal's coming from deeper inside.</d> Radio static, a low ship hum, creaking metal and hard boot echoes accompany their advance.\n[Shot 2] At 00:04.500, the camera cuts to a wide view of a large chamber with hundreds of open, empty cryogenic pods. A second crew member with an uneasy, slightly higher English voice (S2) looks from the abandoned pods toward a final closed door and says nervously: <d>[English] Where did everyone go?</d> The captain approaches that door as the camera pushes forward with the crew. Pod machinery resonates beneath their footsteps.\n[Shot 3] At 00:09.000, the camera cuts behind the captain as the door hisses open. Bright daylight reveals an ordinary present-day elementary school classroom, children at desks, backpacks on the floor and colorful posters on the walls. A teacher at the front with a clear, matter-of-fact English voice (S3) turns toward the astronauts and scolds them: <d>[English] You're late again. Take your seats.</d> The astronauts remain frozen in confusion at the doorway. Hold the absurd reveal until the end while ordinary classroom room tone replaces the ship ambience.\n\noverall_soundscape: Radio static, a low ship hum, creaking metal, boot echoes and empty pod machinery fill the spacecraft. A door hiss gives way abruptly to ordinary classroom chair movement and quiet paper rustling.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "515b6f9c02a9aff5ebdd48c593434b27848f0a398afc1c0651103dfdd475c0da",
      "duration": 15,
      "seed": 57,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源条目元数据为 14s，本次明确扩展为 15 秒。"
      ],
      "source_title": {
        "en": "cinematic sci-fi suspense video. A realistic spaceship crew in full futuristic",
        "zh": "cinematic sci-fi suspense video. A realistic spaceship crew in full futuristic"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-16",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Live-action cinematic science-fiction suspense, realistic spaceship crew in full futuristic uniforms boarding a dark abandoned spacecraft in response to a distress signal. A smooth tracking shot follows them through a narrow metallic corridor; flashlight beams cross flickering warning lights. One crew member with a restrained mid-pitched English voice (S1) says: <d>[English] Signal's coming from deeper inside.</d> Radio static, a low ship hum, creaking metal and hard boot echoes accompany their advance.\n[Shot 2] At 00:04.500, the camera cuts to a wide view of a large chamber with hundreds of open, empty cryogenic pods. A second crew member with an uneasy, slightly higher English voice (S2) looks from the abandoned pods toward a final closed door and says nervously: <d>[English] Where did everyone go?</d> The captain approaches that door as the camera pushes forward with the crew. Pod machinery resonates beneath their footsteps.\n[Shot 3] At 00:09.000, the camera cuts behind the captain as the door hisses open. Bright daylight reveals an ordinary present-day elementary school classroom, children at desks, backpacks on the floor and colorful posters on the walls. A teacher at the front with a clear, matter-of-fact English voice (S3) turns toward the astronauts and scolds them: <d>[English] You're late again. Take your seats.</d> The astronauts remain frozen in confusion at the doorway. Hold the absurd reveal until the end while ordinary classroom room tone replaces the ship ambience.\n\noverall_soundscape: Radio static, a low ship hum, creaking metal, boot echoes and empty pod machinery fill the spacecraft. A door hiss gives way abruptly to ordinary classroom chair movement and quiet paper rustling.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "515b6f9c02a9aff5ebdd48c593434b27848f0a398afc1c0651103dfdd475c0da"
    },
    {
      "id": 17,
      "slug": "skyship-flight-across-a-floating-kingdom-single-continu-344891",
      "title": {
        "en": "Skyship Through a Floating Kingdom",
        "zh": "穿越浮岛王国的飞艇"
      },
      "category": "cinematic-story",
      "source": {
        "kind": "x",
        "name": "@Strength04_X",
        "url": "https://x.com/Strength04_X/status/2082712692159344891"
      },
      "source_prompt": "Skyship flight across a floating kingdom (single continuous shot) From above a sea of endless clouds, the camera dives toward a majestic skyship soaring between gigantic floating islands. Lock-on: the airship accelerating into a violent thunderstorm suspended in the sky. The camera slingshots beneath the hull, whips across the towering sails, then drops tight behind the glowing engines: lightning tears across the clouds, rain lashes the deck, shattered rock fragments tumble through the air. A collapsing floating island sends massive boulders crashing into the storm; the captain steers beneath a breaking stone arch before threading between waterfalls pouring into the sky, hanging bridges and drifting islands in one uninterrupted line. The camera mirrors every impossible maneuver before bursting through the storm into brilliant sunlight, revealing an endless kingdom of floating continents stretching beyond the horizon above a sea of clouds.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Skyship flight across a floating kingdom (single continuous shot) From above a sea of endless clouds, the camera dives toward a majestic skyship soaring between gigantic floating islands. Lock-on: the airship accelerating into a violent thunderstorm suspended in the sky. The camera slingshots beneath the hull, whips across the towering sails, then drops tight behind the glowing engines: lightning tears across the clouds, rain lashes the deck, shattered rock fragments tumble through the air. A collapsing floating island sends massive boulders crashing into the storm; the captain steers beneath a breaking stone arch before threading between waterfalls pouring into the sky, hanging bridges and drifting islands in one uninterrupted line. The camera mirrors every impossible maneuver before bursting through the storm into brilliant sunlight, revealing an endless kingdom of floating continents stretching beyond the horizon above a sea of clouds. Complete the described path in one uninterrupted 15-second flight, establishing the skyship first and reserving the final seconds for the sunlight reveal beyond the storm.\n\noverall_soundscape: High-altitude wind and engine thrust deepen as the skyship enters thunder, lashing rain and cracking rock. Waterfalls rush past the hull before the storm sound falls behind the ship in open sunlight.\n\nnon_diegetic_music: Broad sustained strings and steady low orchestral percussion build through the storm, then open into a long brass-and-string chord in sunlight.",
      "prompt_sha256": "15ed03a4ff9ad4d170616c8220d15f8c7ddadb16a059a522c567a1cede541838",
      "duration": 15,
      "seed": 58,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Skyship Through a Floating Kingdom",
        "zh": "飞空艇穿越浮空王国"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-17",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Skyship flight across a floating kingdom (single continuous shot) From above a sea of endless clouds, the camera dives toward a majestic skyship soaring between gigantic floating islands. Lock-on: the airship accelerating into a violent thunderstorm suspended in the sky. The camera slingshots beneath the hull, whips across the towering sails, then drops tight behind the glowing engines: lightning tears across the clouds, rain lashes the deck, shattered rock fragments tumble through the air. A collapsing floating island sends massive boulders crashing into the storm; the captain steers beneath a breaking stone arch before threading between waterfalls pouring into the sky, hanging bridges and drifting islands in one uninterrupted line. The camera mirrors every impossible maneuver before bursting through the storm into brilliant sunlight, revealing an endless kingdom of floating continents stretching beyond the horizon above a sea of clouds. Complete the described path in one uninterrupted 15-second flight, establishing the skyship first and reserving the final seconds for the sunlight reveal beyond the storm.\n\noverall_soundscape: High-altitude wind and engine thrust deepen as the skyship enters thunder, lashing rain and cracking rock. Waterfalls rush past the hull before the storm sound falls behind the ship in open sunlight.\n\nnon_diegetic_music: Broad sustained strings and steady low orchestral percussion build through the storm, then open into a long brass-and-string chord in sunlight.",
      "effective_prompt_sha256": "15ed03a4ff9ad4d170616c8220d15f8c7ddadb16a059a522c567a1cede541838"
    },
    {
      "id": 18,
      "slug": "a-hyper-realistic-handheld-phone-video-of-a-quiet-suburban-495564",
      "title": {
        "en": "\"A hyper-realistic handheld phone video of a quiet suburban backyard on a bright",
        "zh": "后院上空突然坠落的天空"
      },
      "category": "music-video",
      "source": {
        "kind": "x",
        "name": "@cocktailpeanut",
        "url": "https://x.com/cocktailpeanut/status/2086879654116495564"
      },
      "source_prompt": "\"A hyper-realistic handheld phone video of a quiet suburban backyard on a bright normal afternoon. A person casually waters plants near a fence while birds chirp, distant traffic hums, and everything feels completely ordinary and mundane. The camera drifts naturally like a real home video, capturing the lawn, patio furniture, and blue sky with soft white clouds. After several seconds of normal peaceful activity, a deep cracking sound suddenly comes from above. Without warning, the entire sky begins dropping straight downward like a gigantic solid ceiling, with the blue sky and clouds moving as one physical surface descending toward the yard. The person looks up in shock just as the sky rapidly fills the frame, swallowing the scene in a terrifying instant. The camera jerks and falls to the ground at the last moment. No cinematic buildup, no surreal visual style, it should feel like a totally normal real-life recording interrupted by one sudden impossible and visceral event.\"",
      "source_duration": "14s",
      "prompt": "integrated_multimodal_description: [Shot 1] \"A hyper-realistic handheld phone video of a quiet suburban backyard on a bright normal afternoon. A person casually waters plants near a fence while birds chirp, distant traffic hums, and everything feels completely ordinary and mundane. The camera drifts naturally like a real home video, capturing the lawn, patio furniture, and blue sky with soft white clouds. After several seconds of normal peaceful activity, a deep cracking sound suddenly comes from above. Without warning, the entire sky begins dropping straight downward like a gigantic solid ceiling, with the blue sky and clouds moving as one physical surface descending toward the yard. The person looks up in shock just as the sky rapidly fills the frame, swallowing the scene in a terrifying instant. The camera jerks and falls to the ground at the last moment. No cinematic buildup, no surreal visual style, it should feel like a totally normal real-life recording interrupted by one sudden impossible and visceral event.\" Use a single 15-second home-video take. Preserve the first several seconds of ordinary watering, begin the overhead crack after that peaceful lead-in, and end with the sky descending and camera falling.\n\noverall_soundscape: Ordinary backyard birds, distant traffic and a watering stream continue for several seconds. A sudden deep overhead crack becomes a heavy descending rumble, followed by startled breathing and the phone hitting the ground.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "42f8695e30223aeca06566bf0476f7882751c30d3db90ad0136585d517c6ba26",
      "duration": 15,
      "seed": 59,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源条目元数据为 14s，本次明确扩展为 15 秒。"
      ],
      "source_title": {
        "en": "\"A hyper-realistic handheld phone video of a quiet suburban backyard on a bright",
        "zh": "\"A hyper-realistic handheld phone video of a quiet suburban backyard on a bright"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-18",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] \"A hyper-realistic handheld phone video of a quiet suburban backyard on a bright normal afternoon. A person casually waters plants near a fence while birds chirp, distant traffic hums, and everything feels completely ordinary and mundane. The camera drifts naturally like a real home video, capturing the lawn, patio furniture, and blue sky with soft white clouds. After several seconds of normal peaceful activity, a deep cracking sound suddenly comes from above. Without warning, the entire sky begins dropping straight downward like a gigantic solid ceiling, with the blue sky and clouds moving as one physical surface descending toward the yard. The person looks up in shock just as the sky rapidly fills the frame, swallowing the scene in a terrifying instant. The camera jerks and falls to the ground at the last moment. No cinematic buildup, no surreal visual style, it should feel like a totally normal real-life recording interrupted by one sudden impossible and visceral event.\" Use a single 15-second home-video take. Preserve the first several seconds of ordinary watering, begin the overhead crack after that peaceful lead-in, and end with the sky descending and camera falling.\n\noverall_soundscape: Ordinary backyard birds, distant traffic and a watering stream continue for several seconds. A sudden deep overhead crack becomes a heavy descending rumble, followed by startled breathing and the phone hitting the ground.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "42f8695e30223aeca06566bf0476f7882751c30d3db90ad0136585d517c6ba26"
    },
    {
      "id": 19,
      "slug": "also-has-an-impressive-understanding-of-differen-911311",
      "title": {
        "en": "also has an impressive understanding of different cultures",
        "zh": "穿行爪哇市场的肉丸摊贩"
      },
      "category": "cinematic-story",
      "source": {
        "kind": "x",
        "name": "@junwatu",
        "url": "https://x.com/junwatu/status/2084840152715911311"
      },
      "source_prompt": "Traditional Javanese market in Yogyakarta, early morning, present day, ultra-realistic handheld one-take. A bakso vendor hurries through impossibly narrow, crowded market aisles carrying a large steaming bowl of freshly prepared bakso. Shoppers and vendors step aside to let him pass. He weaves between vegetable stalls overflowing with chilies, shallots, and leafy greens, ducks beneath low-hanging tarps, squeezes past batik merchants and baskets of traditional snacks, and carefully avoids children running through the market. The market feels dense, humid, noisy, and full of life. He passes porters carrying heavy sacks on their shoulders, women bargaining in Javanese, and motorcycles slowly pushing through the narrow lane. He finally reaches a tiny food stall at the back of the market and serves the steaming bowl just as the power briefly flickers, leaving the market dim for a moment before the bustle continues. Oppressive realism, natural morning light mixed with fluorescent bulbs, worn concrete floors, damp walls, wooden stalls, hanging plastic signs, Javanese conversations, vendors calling out prices, clattering bowls, footsteps echoing through the covered market, immersive cinematic pacing. No cyberpunk, no stylization, documentary-like realism.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Traditional Javanese market in Yogyakarta, early morning, present day, ultra-realistic handheld one-take. A bakso vendor hurries through impossibly narrow, crowded market aisles carrying a large steaming bowl of freshly prepared bakso. Shoppers and vendors step aside to let him pass. He weaves between vegetable stalls overflowing with chilies, shallots, and leafy greens, ducks beneath low-hanging tarps, squeezes past batik merchants and baskets of traditional snacks, and carefully avoids children running through the market. The market feels dense, humid, noisy, and full of life. He passes porters carrying heavy sacks on their shoulders, women bargaining in Javanese, and motorcycles slowly pushing through the narrow lane. He finally reaches a tiny food stall at the back of the market and serves the steaming bowl just as the power briefly flickers, leaving the market dim for a moment before the bustle continues. Oppressive realism, natural morning light mixed with fluorescent bulbs, worn concrete floors, damp walls, wooden stalls, hanging plastic signs, Javanese conversations, vendors calling out prices, clattering bowls, footsteps echoing through the covered market, immersive cinematic pacing. No cyberpunk, no stylization, documentary-like realism. Nearby women bargaining in Javanese (S1) and vendors calling prices in Javanese (S2) remain overlapping background voices. Follow the vendor in one continuous 15-second handheld take, keeping the steaming bowl and his careful steps legible; finish with serving the bowl and the brief power flicker.\n\noverall_soundscape: Dense covered-market bustle includes footsteps, bowl clatter, steaming food and sacks brushing clothing. Slow motorcycle engines pass among the stalls, followed by a brief fluorescent power buzz and the uninterrupted market activity.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "644ee708e360b5247f1f86b9510bc54be46b1ede81be1dfd8e283227dfab466f",
      "duration": 15,
      "seed": 60,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。",
        "源文未提供固定台词；仅保留原文要求的语言与说话动作。"
      ],
      "source_title": {
        "en": "also has an impressive understanding of different cultures",
        "zh": "also has an impressive understanding of different cultures"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-19",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Traditional Javanese market in Yogyakarta, early morning, present day, ultra-realistic handheld one-take. A bakso vendor hurries through impossibly narrow, crowded market aisles carrying a large steaming bowl of freshly prepared bakso. Shoppers and vendors step aside to let him pass. He weaves between vegetable stalls overflowing with chilies, shallots, and leafy greens, ducks beneath low-hanging tarps, squeezes past batik merchants and baskets of traditional snacks, and carefully avoids children running through the market. The market feels dense, humid, noisy, and full of life. He passes porters carrying heavy sacks on their shoulders, women bargaining in Javanese, and motorcycles slowly pushing through the narrow lane. He finally reaches a tiny food stall at the back of the market and serves the steaming bowl just as the power briefly flickers, leaving the market dim for a moment before the bustle continues. Oppressive realism, natural morning light mixed with fluorescent bulbs, worn concrete floors, damp walls, wooden stalls, hanging plastic signs, Javanese conversations, vendors calling out prices, clattering bowls, footsteps echoing through the covered market, immersive cinematic pacing. No cyberpunk, no stylization, documentary-like realism. Nearby women bargaining in Javanese (S1) and vendors calling prices in Javanese (S2) remain overlapping background voices. Follow the vendor in one continuous 15-second handheld take, keeping the steaming bowl and his careful steps legible; finish with serving the bowl and the brief power flicker.\n\noverall_soundscape: Dense covered-market bustle includes footsteps, bowl clatter, steaming food and sacks brushing clothing. Slow motorcycle engines pass among the stalls, followed by a brief fluorescent power buzz and the uninterrupted market activity.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "644ee708e360b5247f1f86b9510bc54be46b1ede81be1dfd8e283227dfab466f"
    },
    {
      "id": 20,
      "slug": "16-9-cinematic-wuxia-mystery-set-in-a-bamboo-571041",
      "title": {
        "en": "Bamboo Forest Wuxia Mystery",
        "zh": "夜雪竹林里的武侠悬疑"
      },
      "category": "product-commercial",
      "source": {
        "kind": "x",
        "name": "@AIwithAliya",
        "url": "https://x.com/AIwithAliya/status/2083132770650571041"
      },
      "source_prompt": "A 16:9 cinematic wuxia mystery set in a bamboo forest at night. Use a low-saturation palette of cold blue, ink green, charcoal, and gray. Thin mist fills the forest and fine snow drifts through the air. The mood is austere, lethal, and controlled, with the tension of a martial-arts sect investigating a case and exchanging secret intelligence.\n\nDeep in the forest, dense vertical bamboo fills the background while cold white mist-light glows in the distance. Soft, out-of-focus leaves partially obscure the foreground, creating the sense of watching from within the grove. A cool, soft front-side key lights the actors’ faces; backlight keeps the foreground dark and the distance luminous. Use shallow depth of field so leaves, snow, and bamboo dissolve into soft bokeh.\nPrioritize facial close-ups and measured shot/reverse-shot coverage. Keep the rhythm restrained but tense. Photoreal period-drama production value, cinematic lighting, and no modern elements. No subtitles, on-screen text, watermarks, modern clothing or architecture, animation styling, over-smoothed skin, bright daylight, or comic performance.",
      "source_duration": "15s",
      "prompt": "integrated_multimodal_description: [Shot 1] Live-action cinematic wuxia mystery in a bamboo forest at night, in a restrained cold-blue, ink-green, charcoal and gray palette. Thin mist and fine drifting snow surround two period-clothed martial-arts investigators. Dense vertical bamboo fills the background; soft foreground leaves partially obscure a tight facial close-up of the first investigator. Cool, soft front-side light reveals natural skin texture while cold white mist-light glows in the distance. The investigator looks toward the other person with measured, wary attention. No modern objects, text, subtitles, watermarks or comic performance.\n[Shot 2] At 00:05.000, the camera cuts to the other investigator in matching close-up for a restrained reverse angle. Keep the same cold lighting, shallow depth of field and softly blurred falling snow. A small shift of the eyes toward the dark grove conveys the exchange of secret intelligence without adding invented dialogue. Clothing rustles quietly under the bamboo wind.\n[Shot 3] At 00:10.000, the camera cuts back to the first investigator seen through out-of-focus leaves. A slow, small push-in follows a controlled glance into the luminous mist. Hold the austere, tense expression until the end, with the distant forest remaining still except for snow and moving leaves. Photoreal period-drama lighting, no animation styling, bright daylight or over-smoothed skin.\n\noverall_soundscape: Thin wind passes through bamboo leaves with occasional dry stalk creaks. Soft period-clothing movement and faint footsteps remain close while falling snow and distant forest space stay quiet.\n\nnon_diegetic_music: N/A",
      "prompt_sha256": "11628027721bd267f57680bb10ff37aff3623b6ea6bbfb064178144173e3a828",
      "duration": 15,
      "seed": 61,
      "adaptations": [
        "保留社区原文及作者链接；本次按 H3 T2VA 三字段、镜头和语音规则整理。",
        "按用户要求生成 15 秒（362 帧 / 24 FPS），画幅统一为 1344×768；新增环境音和镜头节奏均属于本次整理。"
      ],
      "source_title": {
        "en": "Bamboo Forest Wuxia Mystery",
        "zh": "竹林武侠谜案"
      },
      "frames": 362,
      "prompt_version": "v2-h3-15s",
      "prompt_origin": "H3-format adaptation of attributed community source; not claimed to be the original author’s structured prompt",
      "task": "t2va",
      "key": "t2va-20",
      "effective_prompt": "integrated_multimodal_description: [Shot 1] Live-action cinematic wuxia mystery in a bamboo forest at night, in a restrained cold-blue, ink-green, charcoal and gray palette. Thin mist and fine drifting snow surround two period-clothed martial-arts investigators. Dense vertical bamboo fills the background; soft foreground leaves partially obscure a tight facial close-up of the first investigator. Cool, soft front-side light reveals natural skin texture while cold white mist-light glows in the distance. The investigator looks toward the other person with measured, wary attention. No modern objects, text, subtitles, watermarks or comic performance.\n[Shot 2] At 00:05.000, the camera cuts to the other investigator in matching close-up for a restrained reverse angle. Keep the same cold lighting, shallow depth of field and softly blurred falling snow. A small shift of the eyes toward the dark grove conveys the exchange of secret intelligence without adding invented dialogue. Clothing rustles quietly under the bamboo wind.\n[Shot 3] At 00:10.000, the camera cuts back to the first investigator seen through out-of-focus leaves. A slow, small push-in follows a controlled glance into the luminous mist. Hold the austere, tense expression until the end, with the distant forest remaining still except for snow and moving leaves. Photoreal period-drama lighting, no animation styling, bright daylight or over-smoothed skin.\n\noverall_soundscape: Thin wind passes through bamboo leaves with occasional dry stalk creaks. Soft period-clothing movement and faint footsteps remain close while falling snow and distant forest space stay quiet.\n\nnon_diegetic_music: N/A",
      "effective_prompt_sha256": "11628027721bd267f57680bb10ff37aff3623b6ea6bbfb064178144173e3a828"
    },
    {
      "id": 1,
      "source_case": 1,
      "slug": "reference-case-1",
      "title": {
        "zh": "雨夜双角色交锋 · 人物与环境参考",
        "en": "Public Ref2VA case 1"
      },
      "source": {
        "name": "LightX2V public Ref2VA testset",
        "url": "https://github.com/ModelTC/Minimax-H3-Turbo/blob/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/prompts_ref2va_test.json"
      },
      "source_prompt": "subject_definitions:\n<Subject 1> is the long-haired East Asian female hunter in <Picture 1>, preserving her face, long wet black hair, black leather combat suit, tattered black cloak, black gloves, and single silver short blade.\n<Subject 2> is the East Asian female tech hunter in <Picture 2>, preserving her face, short black hair, white fox mask with red markings, black trench coat with cyan circuit lines, black boots, and single blue energy dagger.\n<Subject 3> is the rainy cyberpunk intersection in <Picture 3>, preserving its skyscrapers, giant blue-purple holographic billboard, neon signs, wet reflective asphalt, rain, and drifting fog.\n\nsummary:\n[reference generation] In one continuous 5.1667-second cinematic shot, <Subject 1> and <Subject 2> begin in a tense standoff inside <Subject 3>, charge across the wet street, exchange one fast blade combination, and end locked blade-to-blade in a burst of blue energy.\n\nretention_analysis:\n<Subject 1> (appears throughout): fully_preserved - retain the same face, long hair, black combat suit, cloak, body proportions, and one silver blade.\n<Subject 2> (appears throughout): fully_preserved - retain the same face, short hair, fox mask, cyan-lit trench coat, body proportions, and one energy dagger.\n<Subject 3> (appears throughout): fully_preserved - retain the wet intersection, billboard, neon palette, rain, fog, and reflective pavement.\n\ndetailed_description:\nPhotorealistic live-action cyberpunk action cinema, 16:9, high contrast purple-and-cyan night lighting, heavy rain, physically coherent motion, and stable identities. One continuous shot with no cuts, no duplicate people, no duplicate weapons, no costume changes, no gore, and no readable text.\n\n[Shot 1, 00:00.000-00:01.200] A low wide camera frames both fighters five meters apart on the rain-soaked street. <Subject 1> crouches on frame left with her silver blade held low; <Subject 2> stands on frame right with the blue energy dagger raised. Rain streaks through the holographic light and fog crosses the reflective road.\n\n[Shot 1, 00:01.200-00:03.600] Both women burst forward. The camera tracks sideways at waist height as <Subject 1>'s cloak opens behind her. She makes one upward diagonal slash; <Subject 2> sidesteps and blocks with the energy dagger. A compact blue-white spark and a ring of displaced droplets mark the impact. <Subject 2> answers with one horizontal counter, which <Subject 1> deflects with the flat of her blade.\n\n[Shot 1, 00:03.600-00:05.166] They rotate once around their joined weapons and stop in a close blade lock, faces and costumes still distinct. Cyan electricity crawls briefly across the crossed blades while rain runs from the fox mask and black cloak. The camera settles into a medium two-shot as the billboard flickers behind them and both hold the unresolved confrontation.\n\noverall_soundscape:\nHeavy rain, wet footfalls, two sharp blade impacts, controlled fabric movement, electrical crackle, distant traffic, and a low thunder roll. No dialogue.\n\nnon_diegetic_music:\nA dark electronic pulse rises through the charge and stops on the final blade lock.\n",
      "source_duration": 5.1666666667,
      "prompt": "subject_definitions:\n<Subject 1> is the long-haired East Asian female hunter in <Picture 1>, preserving her face, long wet black hair, black leather combat suit, tattered black cloak, black gloves, and single silver short blade.\n<Subject 2> is the East Asian female tech hunter in <Picture 2>, preserving her face, short black hair, white fox mask with red markings, black trench coat with cyan circuit lines, black boots, and single blue energy dagger.\n<Subject 3> is the rainy cyberpunk intersection in <Picture 3>, preserving its skyscrapers, giant blue-purple holographic billboard, neon signs, wet reflective asphalt, rain, and drifting fog.\n\nsummary:\n[reference generation] In one continuous 15-second cinematic shot, <Subject 1> and <Subject 2> begin in a tense standoff inside <Subject 3>, charge across the wet street, exchange one fast blade combination, and end locked blade-to-blade in a burst of blue energy.\n\nretention_analysis:\n<Subject 1> (appears throughout): fully_preserved - retain the same face, long hair, black combat suit, cloak, body proportions, and one silver blade.\n<Subject 2> (appears throughout): fully_preserved - retain the same face, short hair, fox mask, cyan-lit trench coat, body proportions, and one energy dagger.\n<Subject 3> (appears throughout): fully_preserved - retain the wet intersection, billboard, neon palette, rain, fog, and reflective pavement.\n\ndetailed_description:\nPhotorealistic live-action cyberpunk action cinema, 16:9, high contrast purple-and-cyan night lighting, heavy rain, physically coherent motion, and stable identities. One continuous shot with no cuts, no duplicate people, no duplicate weapons, no costume changes, no gore, and no readable text.\n\n[Shot 1] A low wide camera frames both fighters five meters apart on the rain-soaked street. <Subject 1> crouches on frame left with her silver blade held low; <Subject 2> stands on frame right with the blue energy dagger raised. Rain streaks through the holographic light and fog crosses the reflective road.\n\nDuring the middle of the same continuous shot, Both women burst forward. The camera tracks sideways at waist height as <Subject 1>'s cloak opens behind her. She makes one upward diagonal slash; <Subject 2> sidesteps and blocks with the energy dagger. A compact blue-white spark and a ring of displaced droplets mark the impact. <Subject 2> answers with one horizontal counter, which <Subject 1> deflects with the flat of her blade.\n\nDuring the final phase of the same continuous shot, They rotate once around their joined weapons and stop in a close blade lock, faces and costumes still distinct. Cyan electricity crawls briefly across the crossed blades while rain runs from the fox mask and black cloak. The camera settles into a medium two-shot as the billboard flickers behind them and both hold the unresolved confrontation.\n\nThe character sheets provide identity and clothing, not a collage to reproduce. Show one physical version of each fighter in the environment, with their separate silhouettes legible against the billboard. The silver blade belongs only to <Subject 1>; the blue dagger belongs only to <Subject 2>. Keep their hands connected naturally to the weapon handles and show planted feet disturbing shallow puddles. During the opening four seconds, let the rain establish depth while both women watch each other without changing sides. Across the middle six seconds, preserve the contact order of approach, slash, block and counter; the sideways camera travels only far enough to keep both bodies in view. Reserve the last five seconds for the rotation, final lock and held reaction. The final two-shot keeps both referenced faces readable, with rain, breathing and small cloak movements continuing around their otherwise stable pose.\n\noverall_soundscape:\nHeavy rain, wet footfalls, two sharp blade impacts, controlled fabric movement, electrical crackle, distant traffic, and a low thunder roll. No dialogue.\n\nnon_diegetic_music:\nA dark electronic pulse rises through the charge and stops on the final blade lock.\n",
      "prompt_sha256": "5d2a08b4360b9080603ddd8fb85c73c3386a5dda88600dd557e1626eaaabedb9",
      "references": [
        {
          "type": "image",
          "path": "ref-inputs/ref2va_testset/ref2va_test_1_1.jpg",
          "label": "Picture 1",
          "sha256": "1665c52d9e20ace593cbc30d1e17b5cf996ad25a59c7f1890d636711da93dbd3",
          "bytes": 344933,
          "source_url": "https://raw.githubusercontent.com/ModelTC/Minimax-H3-Turbo/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/ref2va_testset/ref2va_test_1_1.jpg",
          "public_path": "assets/ref2va-inputs/ref2va_test_1_1.jpg"
        },
        {
          "type": "image",
          "path": "ref-inputs/ref2va_testset/ref2va_test_1_2.jpg",
          "label": "Picture 2",
          "sha256": "df77cc201252e90cb828132863d4126c4d62d512f79299ec2651bcc9be8c5b1f",
          "bytes": 396459,
          "source_url": "https://raw.githubusercontent.com/ModelTC/Minimax-H3-Turbo/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/ref2va_testset/ref2va_test_1_2.jpg",
          "public_path": "assets/ref2va-inputs/ref2va_test_1_2.jpg"
        },
        {
          "type": "image",
          "path": "ref-inputs/ref2va_testset/ref2va_test_1_3.jpg",
          "label": "Picture 3",
          "sha256": "82cdb7a20d6ec0d7c6c396cc8e99d3f5fe49a2f67f7c3e5e75d45102f4a91062",
          "bytes": 390923,
          "source_url": "https://raw.githubusercontent.com/ModelTC/Minimax-H3-Turbo/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/ref2va_testset/ref2va_test_1_3.jpg",
          "public_path": "assets/ref2va-inputs/ref2va_test_1_3.jpg"
        }
      ],
      "duration": 15,
      "frames": 362,
      "seed": 7301,
      "adaptations": [
        "保留上游六字段提示词与原始参考图片；时间安排由约 5.17 秒扩展为 15 秒。",
        "将上游重复的 Shot 1 时间段整理为一个连续镜头；明确对白说话人；原文台词未改。",
        "新增动作衔接与参考保持说明属于本次整理；本组仅测试图片参考，声音由模型生成。"
      ],
      "task": "ref2va",
      "key": "ref2va-01",
      "effective_prompt": "subject_definitions:\n<Subject 1> is the long-haired East Asian female hunter in <Picture 1>, preserving her face, long wet black hair, black leather combat suit, tattered black cloak, black gloves, and single silver short blade.\n<Subject 2> is the East Asian female tech hunter in <Picture 2>, preserving her face, short black hair, white fox mask with red markings, black trench coat with cyan circuit lines, black boots, and single blue energy dagger.\n<Subject 3> is the rainy cyberpunk intersection in <Picture 3>, preserving its skyscrapers, giant blue-purple holographic billboard, neon signs, wet reflective asphalt, rain, and drifting fog.\n\nsummary:\n[reference generation] In one continuous 15-second cinematic shot, <Subject 1> and <Subject 2> begin in a tense standoff inside <Subject 3>, charge across the wet street, exchange one fast blade combination, and end locked blade-to-blade in a burst of blue energy.\n\nretention_analysis:\n<Subject 1> (appears throughout): fully_preserved - retain the same face, long hair, black combat suit, cloak, body proportions, and one silver blade.\n<Subject 2> (appears throughout): fully_preserved - retain the same face, short hair, fox mask, cyan-lit trench coat, body proportions, and one energy dagger.\n<Subject 3> (appears throughout): fully_preserved - retain the wet intersection, billboard, neon palette, rain, fog, and reflective pavement.\n\ndetailed_description:\nPhotorealistic live-action cyberpunk action cinema, 16:9, high contrast purple-and-cyan night lighting, heavy rain, physically coherent motion, and stable identities. One continuous shot with no cuts, no duplicate people, no duplicate weapons, no costume changes, no gore, and no readable text.\n\n[Shot 1] A low wide camera frames both fighters five meters apart on the rain-soaked street. <Subject 1> crouches on frame left with her silver blade held low; <Subject 2> stands on frame right with the blue energy dagger raised. Rain streaks through the holographic light and fog crosses the reflective road.\n\nDuring the middle of the same continuous shot, Both women burst forward. The camera tracks sideways at waist height as <Subject 1>'s cloak opens behind her. She makes one upward diagonal slash; <Subject 2> sidesteps and blocks with the energy dagger. A compact blue-white spark and a ring of displaced droplets mark the impact. <Subject 2> answers with one horizontal counter, which <Subject 1> deflects with the flat of her blade.\n\nDuring the final phase of the same continuous shot, They rotate once around their joined weapons and stop in a close blade lock, faces and costumes still distinct. Cyan electricity crawls briefly across the crossed blades while rain runs from the fox mask and black cloak. The camera settles into a medium two-shot as the billboard flickers behind them and both hold the unresolved confrontation.\n\nThe character sheets provide identity and clothing, not a collage to reproduce. Show one physical version of each fighter in the environment, with their separate silhouettes legible against the billboard. The silver blade belongs only to <Subject 1>; the blue dagger belongs only to <Subject 2>. Keep their hands connected naturally to the weapon handles and show planted feet disturbing shallow puddles. During the opening four seconds, let the rain establish depth while both women watch each other without changing sides. Across the middle six seconds, preserve the contact order of approach, slash, block and counter; the sideways camera travels only far enough to keep both bodies in view. Reserve the last five seconds for the rotation, final lock and held reaction. The final two-shot keeps both referenced faces readable, with rain, breathing and small cloak movements continuing around their otherwise stable pose.\n\noverall_soundscape:\nHeavy rain, wet footfalls, two sharp blade impacts, controlled fabric movement, electrical crackle, distant traffic, and a low thunder roll. No dialogue.\n\nnon_diegetic_music:\nA dark electronic pulse rises through the charge and stops on the final blade lock.",
      "effective_prompt_sha256": "18a9df9d2b582543496c29a76643782650be317d3ae63c83662f1d9da53107c9"
    },
    {
      "id": 2,
      "source_case": 2,
      "slug": "reference-case-2",
      "title": {
        "zh": "维多利亚书房对白 · 双角色身份",
        "en": "Public Ref2VA case 2"
      },
      "source": {
        "name": "LightX2V public Ref2VA testset",
        "url": "https://github.com/ModelTC/Minimax-H3-Turbo/blob/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/prompts_ref2va_test.json"
      },
      "source_prompt": "subject_definitions:\n<Subject 1> is the fictional Victorian detective in <Picture 3>, preserving his angular clean-shaven face, combed-back dark hair, long dark overcoat, brown waistcoat, cream cravat, watch chain, and briar pipe.\n<Subject 2> is the fictional Victorian doctor in <Picture 2>, preserving his sturdy build, cropped brown hair, heavy moustache, brown tweed suit, dark tie, and wooden cane.\n<Subject 3> is the Baker Street study in <Picture 1>, preserving the wood paneling, lit fireplace, bookshelves, rain-streaked windows, leather armchair, desk, chandelier, rug, and warm-firelight-versus-cool-window-light composition.\n\nsummary:\n[reference generation] In one continuous 5.1667-second locked-off Victorian dialogue shot, <Subject 2> asks whether <Subject 1> could have become a criminal; <Subject 1> answers with dry confidence while both remain naturally alive in <Subject 3>.\n\nretention_analysis:\n<Subject 1> (appears throughout): fully_preserved - retain his face, hair, clothing, pipe, slim proportions, and fireplace-side placement.\n<Subject 2> (appears throughout): fully_preserved - retain his face, moustache, tweed suit, cane, sturdy proportions, and armchair placement.\n<Subject 3> (appears throughout): fully_preserved - retain the room layout, fireplace, shelves, windows, desk, furniture, rain, and mixed warm-cool lighting.\n\ndetailed_description:\nPhotorealistic live-action Victorian feature-film style, restrained 35 mm texture, natural skin, subtle performance, and consistent eyelines. One continuous medium-wide two-shot from a fully static locked-off camera. No cuts, camera movement, identity swapping, wardrobe changes, extra people, readable text, or exaggerated gestures.\n\n[Shot 1, 00:00.000-00:01.300] <Subject 1> stands beside the fireplace on frame left, his pipe held near his chest, while <Subject 2> sits in the leather armchair on frame right with one hand resting on his cane. Wind-driven rain traces the cool windows behind them. A coal shifts in the grate, briefly sharpening warm firelight across <Subject 1>'s angular profile and <Subject 2>'s moustached face.\n\n[Shot 1, 00:01.300-00:03.200] <Subject 2> leans forward just enough to make the leather chair creak, studies the detective, and asks with familiar curiosity, <d>[English] What if you'd chosen crime?</d> His fingers tighten once around the cane handle. <Subject 1> remains still for a beat, draws quietly from the pipe, and turns only his eyes toward the doctor as smoke curls through the warm light.\n\n[Shot 1, 00:03.200-00:05.166] <Subject 1> meets his eyeline and replies in a low, dry voice, <d>[English] I'd have been remarkably successful.</d> <Subject 2>'s amused expression tightens into thoughtful concern. <Subject 1> allows the faintest private smile as the fire gives one crisp crackle and a distant lightning reflection passes across the rain-streaked glass. Both settle into the charged silence without changing position.\n\noverall_soundscape:\nSteady rain on glass, close coal-fire crackle, faint room tone, one subtle leather creak, and distant carriage wheels. Dialogue stays clear and intimate.\n\nnon_diegetic_music:\nN/A\n",
      "source_duration": 5.1666666667,
      "prompt": "subject_definitions:\n<Subject 1> is the fictional Victorian detective in <Picture 3>, preserving his angular clean-shaven face, combed-back dark hair, long dark overcoat, brown waistcoat, cream cravat, watch chain, and briar pipe.\n<Subject 2> is the fictional Victorian doctor in <Picture 2>, preserving his sturdy build, cropped brown hair, heavy moustache, brown tweed suit, dark tie, and wooden cane.\n<Subject 3> is the Baker Street study in <Picture 1>, preserving the wood paneling, lit fireplace, bookshelves, rain-streaked windows, leather armchair, desk, chandelier, rug, and warm-firelight-versus-cool-window-light composition.\n\nsummary:\n[reference generation] In one continuous 15-second locked-off Victorian dialogue shot, <Subject 2> asks whether <Subject 1> could have become a criminal; <Subject 1> answers with dry confidence while both remain naturally alive in <Subject 3>.\n\nretention_analysis:\n<Subject 1> (appears throughout): fully_preserved - retain his face, hair, clothing, pipe, slim proportions, and fireplace-side placement.\n<Subject 2> (appears throughout): fully_preserved - retain his face, moustache, tweed suit, cane, sturdy proportions, and armchair placement.\n<Subject 3> (appears throughout): fully_preserved - retain the room layout, fireplace, shelves, windows, desk, furniture, rain, and mixed warm-cool lighting.\n\ndetailed_description:\nPhotorealistic live-action Victorian feature-film style, restrained 35 mm texture, natural skin, subtle performance, and consistent eyelines. One continuous medium-wide two-shot from a fully static locked-off camera. No cuts, camera movement, identity swapping, wardrobe changes, extra people, readable text, or exaggerated gestures.\n\n[Shot 1] <Subject 1> stands beside the fireplace on frame left, his pipe held near his chest, while <Subject 2> sits in the leather armchair on frame right with one hand resting on his cane. Wind-driven rain traces the cool windows behind them. A coal shifts in the grate, briefly sharpening warm firelight across <Subject 1>'s angular profile and <Subject 2>'s moustached face.\n\nDuring the middle of the same continuous shot, <Subject 2> (S1) leans forward just enough to make the leather chair creak, studies the detective, and asks in a warm, mid-pitched English voice with familiar curiosity, <d>[English] What if you'd chosen crime?</d> His fingers tighten once around the cane handle. <Subject 1> remains still for a beat, draws quietly from the pipe, and turns only his eyes toward the doctor as smoke curls through the warm light.\n\nDuring the final phase of the same continuous shot, <Subject 1> (S2) meets his eyeline and replies in a low, dry English voice, <d>[English] I'd have been remarkably successful.</d> <Subject 2>'s amused expression tightens into thoughtful concern. <Subject 1> allows the faintest private smile as the fire gives one crisp crackle and a distant lightning reflection passes across the rain-streaked glass. Both settle into the charged silence without changing position.\n\nUse the portrait and turnaround sheets to preserve each man's identity, not to reproduce a reference-board layout. The seated doctor's sturdy shoulders, moustache and cane distinguish him from the standing detective's narrow profile and pipe. Keep their hands and props separate and their eyelines directed toward each other. The first four seconds establish their positions and the room's depth: chair legs sit firmly on the rug, the desk remains behind them, and warm reflections move subtly across the wood paneling. Allow roughly five seconds for the doctor's complete question and the detective's silent consideration, then leave the remaining six seconds for the complete reply and both reactions. Do not add another line. The doctor listens with his lips closed during the reply; the detective keeps his lips closed during the question. Small breaths, a natural blink and a restrained finger adjustment prevent a frozen tableau. Preserve the fixed viewpoint through the final pause, keeping the fire visible and the cool rainy windows distinct behind the two warm-lit faces.\n\noverall_soundscape:\nSteady rain on glass, close coal-fire crackle, faint room tone, one subtle leather creak, and distant carriage wheels. Dialogue stays clear and intimate.\n\nnon_diegetic_music:\nN/A\n",
      "prompt_sha256": "593dde5ba7be9ded40395af7de00ed01fd5cc7e1ac778276bac5c61c09f7f278",
      "references": [
        {
          "type": "image",
          "path": "ref-inputs/ref2va_testset/ref2va_test_2_1.jpg",
          "label": "Picture 1",
          "sha256": "3808f06a4b21910dcc04a4f62f6fe95714d0c9773cd9f81fd67695d4eb924373",
          "bytes": 383601,
          "source_url": "https://raw.githubusercontent.com/ModelTC/Minimax-H3-Turbo/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/ref2va_testset/ref2va_test_2_1.jpg",
          "public_path": "assets/ref2va-inputs/ref2va_test_2_1.jpg"
        },
        {
          "type": "image",
          "path": "ref-inputs/ref2va_testset/ref2va_test_2_2.jpg",
          "label": "Picture 2",
          "sha256": "b62cf2f93c256975c80fd58869216a819749a8f3aab4bc03b396b6147f0848eb",
          "bytes": 278692,
          "source_url": "https://raw.githubusercontent.com/ModelTC/Minimax-H3-Turbo/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/ref2va_testset/ref2va_test_2_2.jpg",
          "public_path": "assets/ref2va-inputs/ref2va_test_2_2.jpg"
        },
        {
          "type": "image",
          "path": "ref-inputs/ref2va_testset/ref2va_test_2_3.jpg",
          "label": "Picture 3",
          "sha256": "d1aac91407587e4457ad191e3867e5f24dd878870ca4bbafc40ca24c605c86df",
          "bytes": 194821,
          "source_url": "https://raw.githubusercontent.com/ModelTC/Minimax-H3-Turbo/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/ref2va_testset/ref2va_test_2_3.jpg",
          "public_path": "assets/ref2va-inputs/ref2va_test_2_3.jpg"
        }
      ],
      "duration": 15,
      "frames": 362,
      "seed": 7302,
      "adaptations": [
        "保留上游六字段提示词与原始参考图片；时间安排由约 5.17 秒扩展为 15 秒。",
        "将上游重复的 Shot 1 时间段整理为一个连续镜头；明确对白说话人；原文台词未改。",
        "新增动作衔接与参考保持说明属于本次整理；本组仅测试图片参考，声音由模型生成。"
      ],
      "task": "ref2va",
      "key": "ref2va-02",
      "effective_prompt": "subject_definitions:\n<Subject 1> is the fictional Victorian detective in <Picture 3>, preserving his angular clean-shaven face, combed-back dark hair, long dark overcoat, brown waistcoat, cream cravat, watch chain, and briar pipe.\n<Subject 2> is the fictional Victorian doctor in <Picture 2>, preserving his sturdy build, cropped brown hair, heavy moustache, brown tweed suit, dark tie, and wooden cane.\n<Subject 3> is the Baker Street study in <Picture 1>, preserving the wood paneling, lit fireplace, bookshelves, rain-streaked windows, leather armchair, desk, chandelier, rug, and warm-firelight-versus-cool-window-light composition.\n\nsummary:\n[reference generation] In one continuous 15-second locked-off Victorian dialogue shot, <Subject 2> asks whether <Subject 1> could have become a criminal; <Subject 1> answers with dry confidence while both remain naturally alive in <Subject 3>.\n\nretention_analysis:\n<Subject 1> (appears throughout): fully_preserved - retain his face, hair, clothing, pipe, slim proportions, and fireplace-side placement.\n<Subject 2> (appears throughout): fully_preserved - retain his face, moustache, tweed suit, cane, sturdy proportions, and armchair placement.\n<Subject 3> (appears throughout): fully_preserved - retain the room layout, fireplace, shelves, windows, desk, furniture, rain, and mixed warm-cool lighting.\n\ndetailed_description:\nPhotorealistic live-action Victorian feature-film style, restrained 35 mm texture, natural skin, subtle performance, and consistent eyelines. One continuous medium-wide two-shot from a fully static locked-off camera. No cuts, camera movement, identity swapping, wardrobe changes, extra people, readable text, or exaggerated gestures.\n\n[Shot 1] <Subject 1> stands beside the fireplace on frame left, his pipe held near his chest, while <Subject 2> sits in the leather armchair on frame right with one hand resting on his cane. Wind-driven rain traces the cool windows behind them. A coal shifts in the grate, briefly sharpening warm firelight across <Subject 1>'s angular profile and <Subject 2>'s moustached face.\n\nDuring the middle of the same continuous shot, <Subject 2> (S1) leans forward just enough to make the leather chair creak, studies the detective, and asks in a warm, mid-pitched English voice with familiar curiosity, <d>[English] What if you'd chosen crime?</d> His fingers tighten once around the cane handle. <Subject 1> remains still for a beat, draws quietly from the pipe, and turns only his eyes toward the doctor as smoke curls through the warm light.\n\nDuring the final phase of the same continuous shot, <Subject 1> (S2) meets his eyeline and replies in a low, dry English voice, <d>[English] I'd have been remarkably successful.</d> <Subject 2>'s amused expression tightens into thoughtful concern. <Subject 1> allows the faintest private smile as the fire gives one crisp crackle and a distant lightning reflection passes across the rain-streaked glass. Both settle into the charged silence without changing position.\n\nUse the portrait and turnaround sheets to preserve each man's identity, not to reproduce a reference-board layout. The seated doctor's sturdy shoulders, moustache and cane distinguish him from the standing detective's narrow profile and pipe. Keep their hands and props separate and their eyelines directed toward each other. The first four seconds establish their positions and the room's depth: chair legs sit firmly on the rug, the desk remains behind them, and warm reflections move subtly across the wood paneling. Allow roughly five seconds for the doctor's complete question and the detective's silent consideration, then leave the remaining six seconds for the complete reply and both reactions. Do not add another line. The doctor listens with his lips closed during the reply; the detective keeps his lips closed during the question. Small breaths, a natural blink and a restrained finger adjustment prevent a frozen tableau. Preserve the fixed viewpoint through the final pause, keeping the fire visible and the cool rainy windows distinct behind the two warm-lit faces.\n\noverall_soundscape:\nSteady rain on glass, close coal-fire crackle, faint room tone, one subtle leather creak, and distant carriage wheels. Dialogue stays clear and intimate.\n\nnon_diegetic_music:\nN/A",
      "effective_prompt_sha256": "296e594352cad39ca8efd4a078fcf8c00b8c5ffc6c5b960f8cf51135a20a0a33"
    },
    {
      "id": 3,
      "source_case": 4,
      "slug": "reference-case-4",
      "title": {
        "zh": "沙漠观测者 · 单图人物与场景",
        "en": "Public Ref2VA case 4"
      },
      "source": {
        "name": "LightX2V public Ref2VA testset",
        "url": "https://github.com/ModelTC/Minimax-H3-Turbo/blob/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/prompts_ref2va_test.json"
      },
      "source_prompt": "subject_definitions:\n<Subject 1> is a tan-skinned adult man with close-cropped dark hair, light stubble, a sand-colored field jacket, and a burgundy scarf, as shown in <Picture 1>.\n<Subject 2> is a vintage brass telescope on a dark tripod, as shown in <Picture 1>.\n<Subject 3> is a small red-glowing metal lantern on a weathered table, as shown in <Picture 1>.\n<Subject 4> is a quiet desert observatory at dusk with soft mountain silhouettes, as shown in <Picture 1>.\n\nsummary:\n[reference generation] In one continuous 5.1667-second shot inside <Subject 4>, <Subject 1> notices a distant flash, turns toward <Subject 2>, adjusts its focus ring once, and reaches the eyepiece as <Subject 3> pulses red.\n\nretention_analysis:\n<Subject 1> (appears in [Shot 1]): fully_preserved - Preserve the same adult identity, facial structure, hair, stubble, jacket, scarf, and body proportions.\n<Subject 2> (appears in [Shot 1]): fully_preserved - Preserve the same brass tube, black fittings, tripod geometry, scale, and position.\n<Subject 3> (appears in [Shot 1]): fully_preserved - Preserve the same metal housing and red light.\n<Subject 4> (appears in [Shot 1]): fully_preserved - Preserve the dusk palette, mountain horizon, and sparse desert setting.\n\ndetailed_description:\nVisual style: cinematic naturalism with warm lantern light against a cool violet dusk. One locked medium shot preserves the composition from <Picture 1>. The camera stays fixed with no pan, tilt, zoom, or dolly.\n\n[Shot 1, 00:00.000-00:01.300] <Subject 1> stands beside <Subject 2> as the final dusk light outlines the brass tube and distant mountain ridge. A faint cold flash crosses his eyes from outside frame. He stops breathing for a beat and turns his gaze sharply toward the telescope while <Subject 3> holds a steady red glow on the table.\n\n[Shot 1, 00:01.300-00:03.600] He reaches across the brass tube, grips the focus ring, and makes exactly one deliberate adjustment. The mechanism gives a precise metal click. Reflected violet sky slides across the brass surface, his scarf lifts in the desert wind, and <Subject 3> flickers once brighter without moving.\n\n[Shot 1, 00:03.600-00:05.166] <Subject 1> leans quickly but naturally to the eyepiece, steadies the telescope with his free hand, and peers through it with sudden concentration. The lantern throws a brief red edge across his cheek and telescope fittings. He holds the tense viewing pose as the distant flash fades behind the mountains.\n\noverall_soundscape:\nA faint desert breeze, a soft metal adjustment click, fabric rustle, and a quiet lantern flame.\n\nnon_diegetic_music:\nN/A\n",
      "source_duration": 5.1666666667,
      "prompt": "subject_definitions:\n<Subject 1> is a tan-skinned adult man with close-cropped dark hair, light stubble, a sand-colored field jacket, and a burgundy scarf, as shown in <Picture 1>.\n<Subject 2> is a vintage brass telescope on a dark tripod, as shown in <Picture 1>.\n<Subject 3> is a small red-glowing metal lantern on a weathered table, as shown in <Picture 1>.\n<Subject 4> is a quiet desert observatory at dusk with soft mountain silhouettes, as shown in <Picture 1>.\n\nsummary:\n[reference generation] In one continuous 15-second shot inside <Subject 4>, <Subject 1> notices a distant flash, turns toward <Subject 2>, adjusts its focus ring once, and reaches the eyepiece as <Subject 3> pulses red.\n\nretention_analysis:\n<Subject 1> (appears in [Shot 1]): fully_preserved - Preserve the same adult identity, facial structure, hair, stubble, jacket, scarf, and body proportions.\n<Subject 2> (appears in [Shot 1]): fully_preserved - Preserve the same brass tube, black fittings, tripod geometry, scale, and position.\n<Subject 3> (appears in [Shot 1]): fully_preserved - Preserve the same metal housing and red light.\n<Subject 4> (appears in [Shot 1]): fully_preserved - Preserve the dusk palette, mountain horizon, and sparse desert setting.\n\ndetailed_description:\nVisual style: cinematic naturalism with warm lantern light against a cool violet dusk. One locked medium shot keeps the referenced man, telescope, lantern and desert setting spatially coherent. The camera stays fixed with no pan, tilt, zoom, or dolly.\n\n[Shot 1] <Subject 1> stands beside <Subject 2> as the final dusk light outlines the brass tube and distant mountain ridge. A faint cold flash crosses his eyes from outside frame. He stops breathing for a beat and turns his gaze sharply toward the telescope while <Subject 3> holds a steady red glow on the table.\n\nDuring the middle of the same continuous shot, He reaches across the brass tube, grips the focus ring, and makes exactly one deliberate adjustment. The mechanism gives a precise metal click. Reflected violet sky slides across the brass surface, his scarf lifts in the desert wind, and <Subject 3> flickers once brighter without moving.\n\nDuring the final phase of the same continuous shot, <Subject 1> leans quickly but naturally to the eyepiece, steadies the telescope with his free hand, and peers through it with sudden concentration. The lantern throws a brief red edge across his cheek and telescope fittings. He holds the tense viewing pose as the distant flash fades behind the mountains.\n\nTreat <Picture 1> as a source for the man, telescope, lantern and desert setting rather than as a required first frame. Keep a single adult man in the scene with the same cropped hair, stubble and burgundy scarf; preserve the jacket seams, brass tube proportions and dark tripod supports. His hands remain visibly attached to his arms as one hand finds the focus ring and the other steadies the instrument. The opening four seconds allow the distant light to pass across his eyes while the lantern remains steady. Over the next six seconds, his gaze follows the flash, his body turns toward the instrument, and his fingers make one deliberate focusing adjustment. During the final five seconds, he leans to the eyepiece and settles into concentrated observation. Let his scarf and jacket respond to the same gentle wind throughout. Keep the table and lantern stationary, with the red pulse changing only the light on nearby surfaces. The camera holds its viewpoint, maintaining a clear gap between the man, telescope tube and mountain skyline; no new people, additional telescopes, text or sudden camera cuts appear.\n\noverall_soundscape:\nA faint desert breeze, a soft metal adjustment click, fabric rustle, and a quiet lantern flame.\n\nnon_diegetic_music:\nN/A\n",
      "prompt_sha256": "e427f08a215a42180ea4e363c45c306819885e80729580173bb9a4f59f321da2",
      "references": [
        {
          "type": "image",
          "path": "ref-inputs/ref2va_testset/ref2va_test_4_1.jpg",
          "label": "Picture 1",
          "sha256": "3b54a682135c5e60c3584df15754c8d54177738686a0778ca89f0a326fe50bbd",
          "bytes": 290315,
          "source_url": "https://raw.githubusercontent.com/ModelTC/Minimax-H3-Turbo/02e26d591f7a04d5d1a074c9566d5dd4f22f6225/examples/ref2va_testset/ref2va_test_4_1.jpg",
          "public_path": "assets/ref2va-inputs/ref2va_test_4_1.jpg"
        }
      ],
      "duration": 15,
      "frames": 362,
      "seed": 7303,
      "adaptations": [
        "保留上游六字段提示词与原始参考图片；时间安排由约 5.17 秒扩展为 15 秒。",
        "将上游重复的 Shot 1 时间段整理为一个连续镜头；明确对白说话人；原文台词未改。",
        "新增动作衔接与参考保持说明属于本次整理；本组仅测试图片参考，声音由模型生成。"
      ],
      "task": "ref2va",
      "key": "ref2va-03",
      "effective_prompt": "subject_definitions:\n<Subject 1> is a tan-skinned adult man with close-cropped dark hair, light stubble, a sand-colored field jacket, and a burgundy scarf, as shown in <Picture 1>.\n<Subject 2> is a vintage brass telescope on a dark tripod, as shown in <Picture 1>.\n<Subject 3> is a small red-glowing metal lantern on a weathered table, as shown in <Picture 1>.\n<Subject 4> is a quiet desert observatory at dusk with soft mountain silhouettes, as shown in <Picture 1>.\n\nsummary:\n[reference generation] In one continuous 15-second shot inside <Subject 4>, <Subject 1> notices a distant flash, turns toward <Subject 2>, adjusts its focus ring once, and reaches the eyepiece as <Subject 3> pulses red.\n\nretention_analysis:\n<Subject 1> (appears in [Shot 1]): fully_preserved - Preserve the same adult identity, facial structure, hair, stubble, jacket, scarf, and body proportions.\n<Subject 2> (appears in [Shot 1]): fully_preserved - Preserve the same brass tube, black fittings, tripod geometry, scale, and position.\n<Subject 3> (appears in [Shot 1]): fully_preserved - Preserve the same metal housing and red light.\n<Subject 4> (appears in [Shot 1]): fully_preserved - Preserve the dusk palette, mountain horizon, and sparse desert setting.\n\ndetailed_description:\nVisual style: cinematic naturalism with warm lantern light against a cool violet dusk. One locked medium shot keeps the referenced man, telescope, lantern and desert setting spatially coherent. The camera stays fixed with no pan, tilt, zoom, or dolly.\n\n[Shot 1] <Subject 1> stands beside <Subject 2> as the final dusk light outlines the brass tube and distant mountain ridge. A faint cold flash crosses his eyes from outside frame. He stops breathing for a beat and turns his gaze sharply toward the telescope while <Subject 3> holds a steady red glow on the table.\n\nDuring the middle of the same continuous shot, He reaches across the brass tube, grips the focus ring, and makes exactly one deliberate adjustment. The mechanism gives a precise metal click. Reflected violet sky slides across the brass surface, his scarf lifts in the desert wind, and <Subject 3> flickers once brighter without moving.\n\nDuring the final phase of the same continuous shot, <Subject 1> leans quickly but naturally to the eyepiece, steadies the telescope with his free hand, and peers through it with sudden concentration. The lantern throws a brief red edge across his cheek and telescope fittings. He holds the tense viewing pose as the distant flash fades behind the mountains.\n\nTreat <Picture 1> as a source for the man, telescope, lantern and desert setting rather than as a required first frame. Keep a single adult man in the scene with the same cropped hair, stubble and burgundy scarf; preserve the jacket seams, brass tube proportions and dark tripod supports. His hands remain visibly attached to his arms as one hand finds the focus ring and the other steadies the instrument. The opening four seconds allow the distant light to pass across his eyes while the lantern remains steady. Over the next six seconds, his gaze follows the flash, his body turns toward the instrument, and his fingers make one deliberate focusing adjustment. During the final five seconds, he leans to the eyepiece and settles into concentrated observation. Let his scarf and jacket respond to the same gentle wind throughout. Keep the table and lantern stationary, with the red pulse changing only the light on nearby surfaces. The camera holds its viewpoint, maintaining a clear gap between the man, telescope tube and mountain skyline; no new people, additional telescopes, text or sudden camera cuts appear.\n\noverall_soundscape:\nA faint desert breeze, a soft metal adjustment click, fabric rustle, and a quiet lantern flame.\n\nnon_diegetic_music:\nN/A",
      "effective_prompt_sha256": "995ab1d031e19ff55039114038c9d44e3a64aa853b16e156421ecb211b509bed"
    }
  ],
  "pairs": [
    {
      "key": "t2va-01",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "611a20742f6896d6bd43b230dad1639f59926a509a3a0788c83a4b490b5b87f0",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "e2152306e5568d1f362c9625521c0c1dd575dfe4038281c7ebef160cb1d63b72",
          "generator_device": "cpu",
          "generator_seed": 42
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/01.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/01.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/01.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3171299,
            "sha256": "9e022aaa69247b9bef6802cec248049aed51d16863a24085dcc49f28de6efa18",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/01.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/01.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/01.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/01.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3951804,
            "sha256": "3362ecf93f99da4478839221365dc6e016b194cf84964461570241fb5d06ef88",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/01.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-02",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "f41d67dd2b5f4f5a4a868f3f77509fc3bec7213547f21f9a48113f57903851c2",
          "generator_device": "cpu",
          "generator_seed": 43
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "1e99678a9b11d1381737f3841ea25f9e7bcca8cdca19c43d81e0e284c27ac0ba",
          "generator_device": "cpu",
          "generator_seed": 43
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/02.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/02.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/02.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3896053,
            "sha256": "6ed1640523584e4c06e8217f9d6b227b56903399ba82375af77d9aebd2accdda",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/02.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/02.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/02.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/02.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 4967801,
            "sha256": "3e71ad96fc3cdcb1ae2dd891f0539fbd32882764cd7f9df9e625b1c1c9a9b3b3",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/02.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-03",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "bc03bdf7faaa51d126f06c2048373e8f7efd73c83d2414627a8deb25e44dd9e7",
          "generator_device": "cpu",
          "generator_seed": 44
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "5b9ecfb74648c5c756a68ecc39db2f32f57cf60471ee6a8c440c24111b80c666",
          "generator_device": "cpu",
          "generator_seed": 44
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/03.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/03.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/03.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3167457,
            "sha256": "afb0aea3c8b1cc3c3c97de117af9ccf0d82f9501652151d0ed4722f47dc60712",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/03.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/03.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/03.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/03.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 4835356,
            "sha256": "4b2a72dc78dfbccff841bdeb489c345da97bf176242b3c3d97fe1316e9b52222",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/03.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-04",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "b07dd1ef807e06b3f6a41d948626a9f82248c96d66756a058ca83730f46812b6",
          "generator_device": "cpu",
          "generator_seed": 45
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "63ccae9b397303d87bc8ebf48b292d1c3311a2fe6f9fa2dfb593ee14dd7501e9",
          "generator_device": "cpu",
          "generator_seed": 45
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/04.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/04.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/04.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 1557590,
            "sha256": "11ecdad03db6b9f6375146599545fd3f49632a551d1c371aa515392e27fb02c8",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/04.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/04.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/04.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/04.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 1920257,
            "sha256": "46a9cd618af6f9f34814a6684d37a6b4284d5090cc946d70923719d250712b3f",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/04.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-05",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "aaa77d0630e7d06662485c047522cbfae7ac96e3c650c94ca4b84de7ab5fa72d",
          "generator_device": "cpu",
          "generator_seed": 46
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "ef8c171dbd7d0f4b90253f01b8730e7431d2d3da9dd0ee62b364713b6e0d6717",
          "generator_device": "cpu",
          "generator_seed": 46
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/05.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/05.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/05.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 6953633,
            "sha256": "bdd7e4c6444260c00ce05242dba7ccf2621ed9a4909fb7270928ccbe36ab8291",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/05.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/05.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/05.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/05.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 8359725,
            "sha256": "173be10075153ae92d4127de38ce46826cd0b7cc09d951b0d01caccc0e35b300",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/05.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-06",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "baeb819db7875cb1a3d5f556ba559b1f41810605728d71095fc80f0da034510f",
          "generator_device": "cpu",
          "generator_seed": 47
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "f0f37f16ecc16f05694ee317eee4363a579fdf85df45ab8f04fc1e44c69fba3e",
          "generator_device": "cpu",
          "generator_seed": 47
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/06.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/06.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/06.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3378163,
            "sha256": "e6e782085803b289d5c47130692d828ab846e69917a45e552d0e69f9fa89c4a5",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/06.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/06.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/06.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/06.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3765970,
            "sha256": "e695befc10db7c488f0d9a46b4de53095741018d8f039a8a14243c6ed23b9c46",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/06.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-07",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "51b4f245600ed4e02e256f42e714bb1585c9835af9c7092acef47c927d7b048b",
          "generator_device": "cpu",
          "generator_seed": 48
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "8fa56a39b023c816908bdd15d75499f8044adbf46026e3a83efeb1535e964bd4",
          "generator_device": "cpu",
          "generator_seed": 48
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/07.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/07.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/07.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3400770,
            "sha256": "ea8894328cd77b0c90bcfe7da6c786effcd37fc5fa2795635c2ae0a140eaf8ab",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/07.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/07.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/07.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/07.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 4050643,
            "sha256": "c33000152838da80e1fc7b84bf1e16b0b19793d0b67b9cf1ce16b46105a213c7",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/07.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-08",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "32c26d4721965753d25723e2ff5181705d225cf1da89333f38e0964f2db51538",
          "generator_device": "cpu",
          "generator_seed": 49
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "8e706fa0ff83f00256fed9ca13f02b663bcc222a3d566043ecf997a2710723d0",
          "generator_device": "cpu",
          "generator_seed": 49
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/08.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/08.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/08.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 9554663,
            "sha256": "fc46a5bf800d415f57a26e673568238eeae679ea96eb088f825301b554af9619",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/08.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/08.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/08.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/08.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 11443462,
            "sha256": "0547998b404d1e6784444fae8da3d4a5a34129bc5c8853d3a1c267f762ba549a",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/08.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-09",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "540ff5a9c6338b427e74997e399ed0889c53850a33b4be840e3a5317299f8ecb",
          "generator_device": "cpu",
          "generator_seed": 50
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "d00f0f6ab5fa9da9f0886279d9a30944b4fa337860e09c9afa410361ed568c1b",
          "generator_device": "cpu",
          "generator_seed": 50
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/09.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/09.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/09.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083008,
            "bytes": 8293780,
            "sha256": "bf0310d67ffcab11e869374de1e147ed13f6bb2488b796c4b2e8d50dc0bda127",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/09.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/09.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/09.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/09.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 9382437,
            "sha256": "198859bd23a1d721459a6c037dd30606f0c53486fe8654ecf831dfb3eeff4cbf",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/09.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-10",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "a2670d58b295804d61fe5ac80ff018968621f57e57740a185b225b6274a20a1e",
          "generator_device": "cpu",
          "generator_seed": 51
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "7dc6657d948c0f74e009888963d536dfb2a03c6e0409b32ab31e819fd2c32d0a",
          "generator_device": "cpu",
          "generator_seed": 51
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/10.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/10.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/10.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 14421162,
            "sha256": "2392fc49e132bf0fd208f4dfb2f5e1954484d5a00d2088c6fe7bf29ca48b5608",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/10.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/10.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/10.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/10.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 16746211,
            "sha256": "30362fd1142f36b2106618814272d888baf18c44addce4710fa47a91ecb09f6c",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/10.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-11",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "63ad955d7cac58eeec66c61095703e137f2c286638111f2b256e2adc0dce23d7",
          "generator_device": "cpu",
          "generator_seed": 52
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "49b589c088f520d774db67cd51a2870d94c58c48c5ec6a183de0c4ab37501505",
          "generator_device": "cpu",
          "generator_seed": 52
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/11.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/11.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/11.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 5608773,
            "sha256": "dcbcbddca0c9674cbb22e790a7cfa21930a2b408380ab8d638e49600d424c83e",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/11.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/11.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/11.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/11.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 6986128,
            "sha256": "a5d7e761f50db7d1ecfbdc2a810fc3d6f578437de65ae345bda860d5e74e70ed",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/11.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-12",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "11ca4af9de2f51096152ef04cc38af3a48078ebf5b19e105105315e837407c7d",
          "generator_device": "cpu",
          "generator_seed": 53
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "cfc3663f23d3ac5be1a08764cbaa3d2f5f1689f35be8400c3cc8fe3b2147c438",
          "generator_device": "cpu",
          "generator_seed": 53
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/12.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/12.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/12.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 4170521,
            "sha256": "c5f60e69bad1f8def53a109dc94d8952009231b95dac653f0620f766865b33f8",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/12.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/12.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/12.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/12.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 7367797,
            "sha256": "3f2610fd345f1181aa594d934618d5276cbcc8b61fc51822391dae0f06363a9e",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/12.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-13",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "1acb7e05af638ab46e9abd89c9f3242265255275dda057ff5a6034b0870a170e",
          "generator_device": "cpu",
          "generator_seed": 54
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "b9d3a7efbcda171412b88b4d86735e0ffa337ec3ebd21274c1520f80911bddd3",
          "generator_device": "cpu",
          "generator_seed": 54
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/13.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/13.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/13.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 7142241,
            "sha256": "a0de8231483624747fd9783af414b5668bc906adb8407eb166f90bd3d832d6e4",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/13.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/13.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/13.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/13.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 8699789,
            "sha256": "137ba9b834155ab12d88af2bd1957889a0b9b69fb1b691d6aca8a18a9dbc2420",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/13.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-14",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "7137acbbc69dd7f6cb03a2eb16ed12279bb6072e3bbc349273825be8516bb298",
          "generator_device": "cpu",
          "generator_seed": 55
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "8dacb40df69d998b7d703dac54ecc48b6c2a7e3f921708349bc70a688d8fee65",
          "generator_device": "cpu",
          "generator_seed": 55
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/14.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/14.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/14.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3790233,
            "sha256": "432b97689878829adf82f0ec71654c2a5a8009f80dabf08ee54840b339e16e04",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/14.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/14.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/14.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/14.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3951027,
            "sha256": "adc8d10bbfc42385e1faa9f6792d10c550bf2040d28ffaf3f9a56192357c2c07",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/14.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-15",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "bccfd8f76b57a75de3a47e5cc5f5fe659e868fd04786d69401d78b61d38ad8e8",
          "generator_device": "cpu",
          "generator_seed": 56
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "57a1e4289fbdf052ee040aa6ce4ce836cbd9fc680efc255f06e6f6feb8c44191",
          "generator_device": "cpu",
          "generator_seed": 56
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/15.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/15.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/15.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 5344410,
            "sha256": "3e68ba6cd2b56975bad0a1fd28a7888849d87ed6c7526c024cebb9ccf87a7a23",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/15.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/15.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/15.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/15.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 5917405,
            "sha256": "2db681b75b9f595fdcd6d94d3db5f6600436461b6fa0e117a60b4989a6098102",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/15.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-16",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "2a21234f77a623be4fd2dea60cc8b1065a51333891da139986e8154f43501f64",
          "generator_device": "cpu",
          "generator_seed": 57
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "995c973788b7e124290b8c8e7fc72fc2cf80694e7da024641ec1eaf952dfe932",
          "generator_device": "cpu",
          "generator_seed": 57
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/16.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/16.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/16.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 4701340,
            "sha256": "231f418e9dffe2be40a966d02aed4eee91c7fc6023a7d0ae57b031adb9c3b17f",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/16.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/16.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/16.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/16.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 5031774,
            "sha256": "75a7afab19cceb6b5855bf534ac8365c725baabbd5c9bb4cff5272066843ea0f",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/16.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-17",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "aae74593566b92f58019f355520957468570be23342f819d1b14a5f82d0a8eed",
          "generator_device": "cpu",
          "generator_seed": 58
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "74fff792e1b2ba69394a5eaa0816aa565bacf3a26e4507a1d489479f3ff56715",
          "generator_device": "cpu",
          "generator_seed": 58
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/17.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/17.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/17.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 6570995,
            "sha256": "53d2a517114f0e3861e046c03eb0001171c870e35a4cc59f7c37eca6ab5c2687",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/17.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/17.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/17.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/17.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 7940195,
            "sha256": "919033b30cd6211b28a93700fcf07b0a843a5b9fcad8bf38beacc7c93e707e7c",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/17.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-18",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "93fa59f16d1a97403f55ea3502c9e1f808c8ec7a95f5a3dfdfd74b65011c393e",
          "generator_device": "cpu",
          "generator_seed": 59
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "81126727754844ee17b6c515cb16735d8ce75c0f544e04b1f46fd11b2bdc4b24",
          "generator_device": "cpu",
          "generator_seed": 59
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/18.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/18.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/18.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3144422,
            "sha256": "fac60c285b113d970594198fc74357e1f6da267e2c8a29b9d588db355b33fd1d",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/18.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/18.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/18.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/18.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 9017609,
            "sha256": "39731f1d0f5667bd9ad38f7f593fc64babd572024c13a721a45249f699b30fed",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/18.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-19",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "ab3ef539b89179cf26385d6ffbc145b4acec6da1f25a1ca5527265772aedb67f",
          "generator_device": "cpu",
          "generator_seed": 60
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "acb1caa9c9c40638b1262bca58d499f6b58b52147b61feaa9f3ca12f0c2978c7",
          "generator_device": "cpu",
          "generator_seed": 60
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/19.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/19.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/19.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 7548706,
            "sha256": "bfac59d63ee44e00714feebbfe5590fc5920ad7f080e4dfc48aac5ddae19f7d7",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/19.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/19.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/19.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/19.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 7431929,
            "sha256": "395a525a357b4ac8a7ff7176d100bcda0ce6374739629f67a02afb4e6c5c3b44",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/19.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "t2va-20",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "4df2305ab044463a6a65f19ec0f6546b8daa7ee44d79be0f76d615d811b1fe5d",
          "generator_device": "cpu",
          "generator_seed": 61
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "68b7900f691a970b6c08fc9591b78a1effbf23d163091f8a69cc4372322b4ca5",
          "generator_device": "cpu",
          "generator_seed": 61
        }
      ],
      "input_state_checks": {
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "text_indices": true,
        "audio_latents": true,
        "token_tags": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/t2va/20.mp4",
          "local_video": "assets/acceleration-20260914/authors/t2va/20.mp4",
          "poster": "assets/acceleration-20260914/authors/t2va/20.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 1969783,
            "sha256": "40cb8788fdafa0faaeab27d4b2c520692e19385b232360da2c9f5cd5f2fa3b91",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/t2va/20.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/t2va/20.mp4",
          "local_video": "assets/acceleration-20260914/sol/t2va/20.mp4",
          "poster": "assets/acceleration-20260914/sol/t2va/20.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 2218527,
            "sha256": "34f9df983ddf83e8211f36704f0a81d947d4f540ef8e269a17de825fe1c497f5",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/t2va/20.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 482400,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "ref2va-01",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "reference_posterior",
          "shape": [
            1,
            24,
            1,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "8cee66d94679f998af42be779f012a25ae98f082478275a8bf501280e83f020d",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "reference_posterior",
          "shape": [
            1,
            24,
            1,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "8cee66d94679f998af42be779f012a25ae98f082478275a8bf501280e83f020d",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "reference_posterior",
          "shape": [
            1,
            24,
            1,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "8cee66d94679f998af42be779f012a25ae98f082478275a8bf501280e83f020d",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "reference_condition",
          "shape": [
            1,
            24,
            1,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "ae86ca8a4bc18f95d0f1ca624073957724e7998b4efc505230cdebfa3133b266",
          "generator_device": "cpu",
          "generator_seed": 7301
        },
        {
          "stream": "reference_condition",
          "shape": [
            1,
            24,
            1,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "0120dc90a4fb9aeacb4ffad23f0e5eb6a877d12e710b1bcb14d28e7045d4ba45",
          "generator_device": "cpu",
          "generator_seed": 7301
        },
        {
          "stream": "reference_condition",
          "shape": [
            1,
            24,
            1,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "4eac4a0763fb9a4587d375f09fe0030e172bf4883c27b6a7381fad7d710dbff0",
          "generator_device": "cpu",
          "generator_seed": 7301
        },
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "8e2cda8f1b19338712a359d66706a8e8b2b2255ec1bcf111cb898d1711f78ccf",
          "generator_device": "cpu",
          "generator_seed": 7301
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "6349aebcf54623510d38b471c21dbd52f816bfdf54a06018a6b291916aabf91d",
          "generator_device": "cpu",
          "generator_seed": 7301
        }
      ],
      "input_state_checks": {
        "text_indices": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true,
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "condition_latents": true,
        "audio_latents": true,
        "token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/ref2va/01.mp4",
          "local_video": "assets/acceleration-20260914/authors/ref2va/01.mp4",
          "poster": "assets/acceleration-20260914/authors/ref2va/01.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 3946326,
            "sha256": "1b9d9093936b60d5db848a20be303e5fa97e253c004d06146d7d0a90ce5527a7",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/ref2va/01.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/ref2va/01.mp4",
          "local_video": "assets/acceleration-20260914/sol/ref2va/01.mp4",
          "poster": "assets/acceleration-20260914/sol/ref2va/01.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 5277826,
            "sha256": "78a37cd092cba95757d3560c58e328a3ff88d5a263a898bf85b7ce7759a8b857",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/ref2va/01.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "ref2va-02",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "reference_posterior",
          "shape": [
            1,
            24,
            1,
            44,
            90
          ],
          "dtype": "torch.float32",
          "sha256": "a6e65e03a36b1a4721ce5b1765010d16bd707ee5cdd9230405b5187c6de2dd7c",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "reference_posterior",
          "shape": [
            1,
            24,
            1,
            48,
            86
          ],
          "dtype": "torch.float32",
          "sha256": "b7db95f38428eaecdeed368b3bd9dac0c1a80df7448a6462f9d0058e1ef8105f",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "reference_posterior",
          "shape": [
            1,
            24,
            1,
            48,
            86
          ],
          "dtype": "torch.float32",
          "sha256": "b7db95f38428eaecdeed368b3bd9dac0c1a80df7448a6462f9d0058e1ef8105f",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "reference_condition",
          "shape": [
            1,
            24,
            1,
            44,
            90
          ],
          "dtype": "torch.float32",
          "sha256": "71767cc0a486284e238a760a21d596fd461e06ce2997594c4d7ae9d3ec24ab9b",
          "generator_device": "cpu",
          "generator_seed": 7302
        },
        {
          "stream": "reference_condition",
          "shape": [
            1,
            24,
            1,
            48,
            86
          ],
          "dtype": "torch.float32",
          "sha256": "37f3c5dcc78f1bbc766c1021bef007698aaf33b3331109e640b63630964f3c13",
          "generator_device": "cpu",
          "generator_seed": 7302
        },
        {
          "stream": "reference_condition",
          "shape": [
            1,
            24,
            1,
            48,
            86
          ],
          "dtype": "torch.float32",
          "sha256": "1521eb25f6f371063d64af0b4fe9a25baa771aa254624f0f55d7c0d530e6fbf1",
          "generator_device": "cpu",
          "generator_seed": 7302
        },
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "c9c965b3a8c8586e33406e0329a53832da0f5c2f7126e3b56f140eb0f56e9f27",
          "generator_device": "cpu",
          "generator_seed": 7302
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "8af4ce669369ad94c8615f89023472620675bbbe70583b75f5f5353024c893b9",
          "generator_device": "cpu",
          "generator_seed": 7302
        }
      ],
      "input_state_checks": {
        "text_indices": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true,
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "condition_latents": true,
        "audio_latents": true,
        "token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/ref2va/02.mp4",
          "local_video": "assets/acceleration-20260914/authors/ref2va/02.mp4",
          "poster": "assets/acceleration-20260914/authors/ref2va/02.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 2959945,
            "sha256": "830c87cac7b819e99d2f0541fbfce201a015328c0d7c392c2516345b92b89ed2",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/ref2va/02.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/ref2va/02.mp4",
          "local_video": "assets/acceleration-20260914/sol/ref2va/02.mp4",
          "poster": "assets/acceleration-20260914/sol/ref2va/02.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 2218510,
            "sha256": "92286c4d349b5403b3dc702eb41261a7e3c8afe3d9e9b3f28aacb538ef9794e3",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/ref2va/02.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    },
    {
      "key": "ref2va-03",
      "noise_exact": true,
      "noise_draws": [
        {
          "stream": "reference_posterior",
          "shape": [
            1,
            24,
            1,
            56,
            70
          ],
          "dtype": "torch.float32",
          "sha256": "04cfbeb376ea7c11a2694f2304a97fe07f2908b210b97f76188be2fc79d708ba",
          "generator_device": "cpu",
          "generator_seed": 42
        },
        {
          "stream": "reference_condition",
          "shape": [
            1,
            24,
            1,
            56,
            70
          ],
          "dtype": "torch.float32",
          "sha256": "a96e4e535e93d4ba58974eaf2e85b02edabdc8f57478033e8c789d2a6a9b53f0",
          "generator_device": "cpu",
          "generator_seed": 7303
        },
        {
          "stream": "target",
          "shape": [
            1,
            24,
            107,
            48,
            84
          ],
          "dtype": "torch.float32",
          "sha256": "877693a4be7d39c54fec8183771074af2cf6f1c523b4f4163a7928dee591ac48",
          "generator_device": "cpu",
          "generator_seed": 7303
        },
        {
          "stream": "target",
          "shape": [
            1206,
            32
          ],
          "dtype": "torch.float32",
          "sha256": "0f1626a806cd76f29bc0b4384aa416d75c264812c820c46a8df7a8d82a20a517",
          "generator_device": "cpu",
          "generator_seed": 7303
        }
      ],
      "input_state_checks": {
        "text_indices": true,
        "prompt_embeds": true,
        "latents": true,
        "text_token_tags": true,
        "audio_indices": true,
        "video_indices": true,
        "position_ids": true,
        "condition_latents": true,
        "audio_latents": true,
        "token_tags": true
      },
      "sigma_grid_equal": true,
      "video_sigmas": [
        1.0,
        0.9939097762107849,
        0.9842875003814697,
        0.9660636782646179,
        0.9230769276618958,
        0.8349423408508301,
        0.6968520283699036,
        0.4687533378601074,
        0.0
      ],
      "audio_sigmas": [
        1.0,
        0.9760763049125671,
        0.9399791955947876,
        0.8767979145050049,
        0.75,
        0.5584253072738647,
        0.3649502694606781,
        0.1807248592376709,
        0.0
      ],
      "outputs": {
        "authors": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/authors/ref2va/03.mp4",
          "local_video": "assets/acceleration-20260914/authors/ref2va/03.mp4",
          "poster": "assets/acceleration-20260914/authors/ref2va/03.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 2029456,
            "sha256": "5afec19cb6eaab69dec205e444d85d881e03e23d36b7cfdda5159f7d053d0ef0",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/authors/ref2va/03.mp4"
        },
        "sol": {
          "video": "https://huggingface.co/datasets/Lawrence-cj/sol-h3-hyperflow-20260913-media/resolve/3ecc22fa54e5ec24e820fd94eefac5803bc0e45a/videos/acceleration-20260914/sol/ref2va/03.mp4",
          "local_video": "assets/acceleration-20260914/sol/ref2va/03.mp4",
          "poster": "assets/acceleration-20260914/sol/ref2va/03.jpg",
          "media_validation": {
            "video_codec": "h264",
            "audio_codec": "aac",
            "frames": 362,
            "width": 1344,
            "height": 768,
            "fps": 24,
            "sample_rate": 32000,
            "channels": 2,
            "duration_s": 15.083333,
            "bytes": 2037716,
            "sha256": "d2ba473c6f438b1bd2738b51964b52f012dfb65ea8986583f345e64ce082f5f8",
            "verification": "ffprobe frame count and complete audio/video decode"
          },
          "dataset_path": "videos/acceleration-20260914/sol/ref2va/03.mp4"
        }
      },
      "previous_sol_reproduction": {
        "frames_compared": 362,
        "video_decoded_identical": true,
        "audio_identical_after_container_alignment": true,
        "audio_alignment": {
          "old_offset_samples": 1024,
          "samples_compared": 483328,
          "exact": true,
          "max_abs": 0.0
        }
      }
    }
  ],
  "checkpoint_sha256": "4d7dec1363ebcb9fd63117621b65f8bd19fecacf7ba41f38dd098be363d3972d",
  "base_revision": "83db0c0efe6ef9824e0e194be110346c0a9542ed",
  "historical_noise": "Historical noise was not saved; these are new paired runs with the previous prompts, seeds, and matching RNG behavior.",
  "updated_at": "2026-09-14T05:06:39.863064+00:00"
}
