From 229d3d63a6a19df026421b2bcbe64306051af5fb Mon Sep 17 00:00:00 2001 From: lin-bot23 Date: Wed, 12 Aug 2026 02:52:06 +0800 Subject: [PATCH] docs: add LTX-2.5 tutorial pages and update Lightricks pricing - Add open-source LTX-2.5 page (tutorials/video/ltx/ltx-2-5) covering T2V, I2V, and FLF2V native workflows with model downloads - Add LTX-2.5 API page (tutorials/partner-nodes/lightricks/ltx-2-5) covering the three cloud API workflows with Fast and Pro tiers - Add Lightricks provider group to Partner Nodes navigation in all 4 languages - Update Lightricks pricing table with LTX-2.5 Fast and Pro rates --- docs.json | 38 ++- .../partner-nodes/lightricks/ltx-2-5.mdx | 108 ++++++++ ja/tutorials/partner-nodes/pricing.mdx | 30 +- ja/tutorials/video/ltx/ltx-2-5.mdx | 258 +++++++++++++++++ .../partner-nodes/lightricks/ltx-2-5.mdx | 108 ++++++++ ko/tutorials/partner-nodes/pricing.mdx | 30 +- ko/tutorials/video/ltx/ltx-2-5.mdx | 259 ++++++++++++++++++ .../partner-nodes/lightricks/ltx-2-5.mdx | 106 +++++++ tutorials/partner-nodes/pricing.mdx | 26 +- tutorials/video/ltx/ltx-2-5.mdx | 251 +++++++++++++++++ .../partner-nodes/lightricks/ltx-2-5.mdx | 108 ++++++++ zh/tutorials/partner-nodes/pricing.mdx | 28 +- zh/tutorials/video/ltx/ltx-2-5.mdx | 258 +++++++++++++++++ 13 files changed, 1558 insertions(+), 50 deletions(-) create mode 100644 ja/tutorials/partner-nodes/lightricks/ltx-2-5.mdx create mode 100644 ja/tutorials/video/ltx/ltx-2-5.mdx create mode 100644 ko/tutorials/partner-nodes/lightricks/ltx-2-5.mdx create mode 100644 ko/tutorials/video/ltx/ltx-2-5.mdx create mode 100644 tutorials/partner-nodes/lightricks/ltx-2-5.mdx create mode 100644 tutorials/video/ltx/ltx-2-5.mdx create mode 100644 zh/tutorials/partner-nodes/lightricks/ltx-2-5.mdx create mode 100644 zh/tutorials/video/ltx/ltx-2-5.mdx diff --git a/docs.json b/docs.json index 934bd99b8..11be7abe6 100644 --- a/docs.json +++ b/docs.json @@ -307,7 +307,8 @@ "pages": [ "tutorials/video/ltxv", "tutorials/video/ltx/ltx-2", - "tutorials/video/ltx/ltx-2-3" + "tutorials/video/ltx/ltx-2-3", + "tutorials/video/ltx/ltx-2-5" ] }, { @@ -486,6 +487,12 @@ "tutorials/partner-nodes/luma/luma-image-to-video" ] }, + { + "group": "Lightricks", + "pages": [ + "tutorials/partner-nodes/lightricks/ltx-2-5" + ] + }, { "group": "MiniMax", "pages": [ @@ -3154,7 +3161,8 @@ "pages": [ "zh/tutorials/video/ltxv", "zh/tutorials/video/ltx/ltx-2", - "zh/tutorials/video/ltx/ltx-2-3" + "zh/tutorials/video/ltx/ltx-2-3", + "zh/tutorials/video/ltx/ltx-2-5" ] }, { @@ -3333,6 +3341,12 @@ "zh/tutorials/partner-nodes/luma/luma-image-to-video" ] }, + { + "group": "Lightricks", + "pages": [ + "zh/tutorials/partner-nodes/lightricks/ltx-2-5" + ] + }, { "group": "MiniMax", "pages": [ @@ -6016,7 +6030,8 @@ "pages": [ "ja/tutorials/video/ltxv", "ja/tutorials/video/ltx/ltx-2", - "ja/tutorials/video/ltx/ltx-2-3" + "ja/tutorials/video/ltx/ltx-2-3", + "ja/tutorials/video/ltx/ltx-2-5" ] }, { @@ -6196,6 +6211,12 @@ "ja/tutorials/partner-nodes/luma/luma-image-to-video" ] }, + { + "group": "Lightricks", + "pages": [ + "ja/tutorials/partner-nodes/lightricks/ltx-2-5" + ] + }, { "group": "MiniMax", "pages": [ @@ -8962,7 +8983,8 @@ "pages": [ "ko/tutorials/video/ltxv", "ko/tutorials/video/ltx/ltx-2", - "ko/tutorials/video/ltx/ltx-2-3" + "ko/tutorials/video/ltx/ltx-2-3", + "ko/tutorials/video/ltx/ltx-2-5" ] }, { @@ -9140,6 +9162,12 @@ "ko/tutorials/partner-nodes/luma/luma-image-to-video" ] }, + { + "group": "Lightricks", + "pages": [ + "ko/tutorials/partner-nodes/lightricks/ltx-2-5" + ] + }, { "group": "MiniMax", "pages": [ @@ -12054,4 +12082,4 @@ "destination": "/ko/basic-concepts/dependencies" } ] -} \ No newline at end of file +} diff --git a/ja/tutorials/partner-nodes/lightricks/ltx-2-5.mdx b/ja/tutorials/partner-nodes/lightricks/ltx-2-5.mdx new file mode 100644 index 000000000..470a8f0de --- /dev/null +++ b/ja/tutorials/partner-nodes/lightricks/ltx-2-5.mdx @@ -0,0 +1,108 @@ +--- +title: "ComfyUIでのLTX-2.5 APIビデオ生成" +description: "ComfyUIでLTX-2.5 APIノードをクラウドワークフローとともに使用して、テキストから動画、画像から動画、最初と最後のフレームから動画への生成方法を、FastとProモデルティアを含めて学びます。" +sidebarTitle: "LTX-2.5 API" +translationSourceHash: 131a7a07 +translationFrom: tutorials/partner-nodes/lightricks/ltx-2-5.mdx +--- + +import ReqHint from "/snippets/ja/tutorials/partner-nodes/req-hint.mdx"; +import UpdateReminder from "/snippets/ja/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://ltx.io/model/ltx-2-5) は Lightricks のオープンウェイトのビデオおよびワールドモデルであり、LTX-2.5 APIノードを通じて ComfyUI で利用できます。APIワークフローは Lightricks のサーバー上で生成を実行するため、モデルのダウンロードやローカルGPUは不要で、生成されたビデオの1秒ごとに Comfy API アカウントに課金されます。 + +LTX-2.5 API は2つのモデルティアを提供します。最大品質の **LTX-2.5 (Pro)** と、蒸留モデルによる高速生成の **LTX-2.5 (Fast)** です。両ティアともに、ネイティブなマルチショットシーンを備えた同期オーディオとビデオ生成をサポートし、Fastティアでは最大4K解像度で出力できます。 + + + + +## LTX-2.5 API ワークフロー + +テンプレートライブラリには、3つのAPIサンプルワークフローが用意されています。 + +- **テキストから動画へ**: テキストプロンプトからビデオを生成します +- **画像から動画へ**: 単一の開始画像からビデオを生成します +- **FLF2V**: 最初と最後のフレーム画像の間のビデオを生成します + +### LTX-2.5 テキストから動画へ + +LTX-2.5 API を通じて、テキストプロンプトから同期オーディオ付きのビデオを生成します。 + + + + + + Comfy Cloud で開く + + + JSONをダウンロードするか、テンプレートライブラリで「LTX-2.5 T2V」を検索してください + + + +#### プロンプト作成のヒント + +1. **モデル**: 最高品質の場合は `LTX-2.5 (Pro)` を、高速生成の場合は `LTX-2.5 (Fast)` を選択します +2. **再生時間**: Pro は2〜10秒、Fast は2〜20秒に対応しています。10秒を超える場合は、720p または 1080p の解像度と 24/25 FPS が必要です +3. **解像度と比率**: `1280x720`(横)や `720x1280`(縦)などの解像度プリセットを選択します。Fast では `2560x1440` と `3840x2160` のオプションが追加されます +4. **FPS**: Fast では 24、25、48、50 FPS から、Pro では 24、25、50 FPS から選択できます +5. **オーディオ**: オーディオ生成はデフォルトでオンになっています。無音のビデオのみが必要な場合は、ノードの高度な設定でオフにできます +6. **プロンプト**: ショット、アクション、キャラクター、カメラの動き、希望するオーディオを含め、シーン全体を1つの流れるような段落で説明します + +### LTX-2.5 画像から動画へ + +LTX-2.5 API を通じて、単一の開始画像から同期オーディオ付きのビデオを生成します。 + + + + + + Comfy Cloud で開く + + + JSONをダウンロードするか、テンプレートライブラリで「LTX-2.5 I2V」を検索してください + + + + + + ワークフローの開始フレームです。または、ご自身の画像を使用できます。 + + + +#### プロンプト作成のヒント + +1. **モデル**: このワークフローは高速なイテレーションのために `LTX-2.5 (Fast)` が選択済みです。最高品質の場合は `LTX-2.5 (Pro)` に切り替えてください +2. **画像は1枚のみ**: 画像入力は開始画像を正確に1枚受け付けます +3. **プロンプト**: その後に起こること、つまり開始画像から続く動き、カメラの動き、音を説明します。「提供された開始画像を最初のフレームとして使用してください」のような表現は生成を固定するのに役立ちます +4. **再生時間と解像度**: Fast ティアは最大20秒、最大4K解像度に対応しています。10秒を超える場合は、24/25 FPS での720p/1080pが必要です + +### LTX-2.5 FLF2V + +LTX-2.5 API を通じて、最初と最後のフレーム画像の間のビデオを生成し、両方のフレームを1つのアニメーションシーケンスにブレンドします。 + + + + + + Comfy Cloud で開く + + + JSONをダウンロードするか、テンプレートライブラリで「LTX-2.5 FLF2V」を検索してください + + + + + + ビデオの最初のフレームです。 + + + ビデオの最後のフレームで、ノードの `last_frame` 入力に接続されます。 + + + +#### プロンプト作成のヒント + +1. **フレーム**: ワークフローでは、最初の画像を `image` 入力に、2番目の画像を `last_frame` 入力に接続します。モデルはその間の動きを生成します +2. **フレームの整合性を保つ**: スムーズな補間のために、同じアスペクト比の画像を使用してください +3. **モデル**: このワークフローは `LTX-2.5 (Pro)` が選択済みです。高速生成の場合は `LTX-2.5 (Fast)` に切り替えてください +4. **プロンプト**: フレーム間の遷移と、希望するオーディオ(会話、SFX、音楽)を説明します \ No newline at end of file diff --git a/ja/tutorials/partner-nodes/pricing.mdx b/ja/tutorials/partner-nodes/pricing.mdx index 6fc441a50..34383a8a8 100644 --- a/ja/tutorials/partner-nodes/pricing.mdx +++ b/ja/tutorials/partner-nodes/pricing.mdx @@ -3,7 +3,7 @@ title: "料金" description: "この記事では、現在のパートナーノードの料金を一覧表示します。" sidebarTitle: "料金" mode: "wide" -translationSourceHash: 6ae911be +translationSourceHash: 320cca0f translationFrom: tutorials/partner-nodes/pricing.mdx translationBlockHashes: "_intro": 471c1619 @@ -21,7 +21,7 @@ translationBlockHashes: "Ideogram": a6ecfb20 "Krea": fb20c257 "Kling": d61f855f - "Lightricks": 4870964f + "Lightricks": 838c1ed7 "Luma": 655207a5 "Meshy": 6cd06881 "Minimax": bf30472e @@ -636,16 +636,22 @@ Seedream 5.0 Proは解像度ベースの課金方式を採用しています:* 合計クレジット = **(クレジット / 秒)× `duration`**。 -| モデル | 解像度 | クレジット | -| :----------- | :--------- | :----------- | -| LTX-2 (Pro) | 1080p | 12.66 / 秒 | -| LTX-2 (Pro) | 1440p | 25.32 / 秒 | -| LTX-2 (Pro) | 4K | 50.64 / 秒 | -| LTX-2 (Fast) | 1080p | 8.44 / 秒 | -| LTX-2 (Fast) | 1440p | 16.88 / 秒 | -| LTX-2 (Fast) | 4K | 33.76 / 秒 | - -解像度は `1920x1080`、`2560x1440`、`3840x2160` に対応します。10秒を超える時間は **LTX-2(Fast)** の1080pおよび25 FPSでのみ利用可能です。 +| モデル | 解像度 | クレジット | +| :------------ | :--------- | :---------- | +| LTX-2.5 (Pro) | 720p | 36.21 / 秒 | +| LTX-2.5 (Pro) | 1080p | 51.29 / 秒 | +| LTX-2.5 (Fast)| 720p | 27.16 / 秒 | +| LTX-2.5 (Fast)| 1080p | 39.22 / 秒 | +| LTX-2.5 (Fast)| 1440p | 57.33 / 秒 | +| LTX-2.5 (Fast)| 4K | 90.52 / 秒 | +| LTX-2 (Pro) | 1080p | 12.66 / 秒 | +| LTX-2 (Pro) | 1440p | 25.32 / 秒 | +| LTX-2 (Pro) | 4K | 50.64 / 秒 | +| LTX-2 (Fast) | 1080p | 8.44 / 秒 | +| LTX-2 (Fast) | 1440p | 16.88 / 秒 | +| LTX-2 (Fast) | 4K | 33.76 / 秒 | + +解像度は `1280x720` / `720x1280`(720p)、`1920x1080` / `1080x1920`(1080p)、`2560x1440` / `1440x2560`(1440p)、`3840x2160` / `2160x3840`(4K)に対応します。**LTX-2.5(Fast)** で10秒を超える時間は720pまたは1080p解像度と24/25 FPSが必要です。**LTX-2(Fast)** で10秒を超える時間は1080pおよび25 FPSでのみ利用可能です。 ## Luma diff --git a/ja/tutorials/video/ltx/ltx-2-5.mdx b/ja/tutorials/video/ltx/ltx-2-5.mdx new file mode 100644 index 000000000..81cdf5a7c --- /dev/null +++ b/ja/tutorials/video/ltx/ltx-2-5.mdx @@ -0,0 +1,258 @@ +--- +title: "LTX-2.5: ComfyUI ワークフロー例" +description: "ComfyUI で LTX-2.5 を使用する方法を説明します。テキストから動画へ、画像から動画へ、先頭・最終フレームから動画生成のためのネイティブワークフローを備え、ピクセルディフュージョン品質、ネイティブマルチショットシーン、同期オーディオに対応しています。" +sidebarTitle: "LTX-2.5" +translationSourceHash: 34280157 +translationFrom: tutorials/video/ltx/ltx-2-5.mdx +translationBlockHashes: + "_intro": b8e0c526 + "Key features": c81ccf83 + "Getting started": d9d37046 + "ComfyUI Native Workflows": 616316a6 + "Resources": f89b58b6 +--- +import UpdateReminder from "/snippets/ja/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5) は、Lightricks が開発したオープンウェイトのビデオ・ワールドモデルで、ComfyUI でネイティブにサポートされるようになりました。LTX-2.3 を基盤に、シーンの複雑さに応じてレンダリング計算を割り当てる Diffusion Fidelity Rendering パイプラインを搭載しています。このモデルは LTX-2.3 の同期オーディオ・ビデオ生成を維持しつつ、ネイティブのマルチショットを追加しています。1回の生成で複数の接続されたショットが生成され、カットをまたいでキャラクター、環境、照明、声が保持されます。 + +LTX-2.5 には、カスタムの Gemma 4 12B テキストエンコーダーが同梱されており、複雑なプロンプトにわたって複数の被写体、アクション、照明、カメラの方向を保持します。さらに、プロンプトエンハンサーも搭載され、短いプロンプトを詳細な指示に拡張します。Auto duration は diffusion が始まる前に、記述されたアクションからクリップの長さを予測します。蒸留モデルはより高速な生成を提供し、出力はネイティブ4K HDR、最大50 FPS、プロフェッショナルな仕上げパイプラインのための RAW ワークフローに対応しています。 + +ComfyUI には LTX-2.5 用のネイティブワークフローが3つ同梱されています: + +- **テキストから動画へ** (T2V):テキストプロンプトからビデオを生成 +- **画像から動画へ** (I2V):入力画像からビデオを生成 +- **FLF2V**:最初と最後のフレーム画像の間を補間 + + + +## 主な機能 + +- **Diffusion Fidelity Rendering**: キーフレーム優先の生成で、高忠実度キーフレームのグリッド上にすべてのシーンを構築し、ピクセル単位の品質を実現 +- **Diffusion Video Decoder**: よりシャープな顔、読みやすいテキスト、速い動きでのスミアの低減 +- **Native multishot**: 1回の生成で、カットをまたいでキャラクター、環境、照明、声、スタイルを保持した、接続された複数のショットを生成 +- **Gemma 4 12B text encoder**: 複雑なプロンプト全体で、複数の被写体、アクション、照明、ディテール、カメラの方向を保持 +- **Prompt enhancer**: 短いプロンプトを豊かな映画的な指示へと展開する軽量モデル +- **Auto duration**: 拡散の開始前に、記述されたアクションから適切なクリップの再生時間を予測 +- **Native 4K HDR**: プロフェッショナルな仕上げに適した高解像度HDR出力、最大50 FPSでの同期オーディオ・ビデオ対応 + +## はじめに + +LTX-2.5はComfyUIでネイティブにサポートされています。はじめるには: + +1. ComfyUIを最新バージョンにアップデート +2. **テンプレートライブラリ** > **ビデオ** > 任意のLTX-2.5ワークフローを選択 +3. Hugging Faceでモデルリポジトリへのアクセスを申請し(下記参照)、ポップアップに従ってモデルをダウンロードし、ワークフローを実行します + + + LTX-2.5 のモデルファイルは、Hugging Face のアクセス制限付き(gated)リポジトリ [Lightricks/LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5) でホストされています。リポジトリのページを開き、モデルライセンスに同意し、アクセス申請が承認されるまで待ってからモデルをダウンロードしてください。アクセスがないとモデルのダウンロードは失敗します。 + + +## ComfyUI ネイティブワークフロー + +### LTX-2.5 テキストからビデオ生成(T2V) + +テキストプロンプトからビデオを生成します。オプションの空間アップスケーリングを使用して、解像度を向上させることができます。 + + + + + + Comfy Cloud で開く + + + JSONをダウンロードするか、テンプレートライブラリで「LTX-2.5 T2V」を検索してください + + + +#### モデルのダウンロード + + + + ComfyUI/models/diffusion_models/ に配置してください + + + ComfyUI/models/text_encoders/ に配置してください + + + ComfyUI/models/text_encoders/ に配置してください + + + ComfyUI/models/vae/ に配置してください + + + ComfyUI/models/vae/ に配置してください + + + ComfyUI/models/latent_upscale_models/ に配置してください + + + +#### モデルの保存場所 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### プロンプト作成のヒント + +1. **シーン全体を説明する**:ショットタイプ、シーン、アクション、キャラクター、カメラの動きを1つの流れるような段落に含めます +2. **オーディオ**:シーンに必要な音とダイアログを説明します。モデルはビデオと同期したオーディオを生成します +3. **プロンプトエンハンサーを使用する**:ワークフローは短いプロンプトを自動的に詳細な映画的指示へと拡張します。簡単なシーンのアイデアを書き、残りはエンハンサーに任せましょう + +--- + +### LTX-2.5 画像からビデオ(I2V) + +入力画像からビデオを生成し、提供された最初のフレームからモデルがシーンをアニメーション化します。 + + + + + + Comfy Cloud で開く + + + JSONをダウンロードするか、テンプレートライブラリで「LTX-2.5 I2V」を検索してください。 + + + + + + デフォルトの入力画像をダウンロードするか、ご自身の画像を使用してください。 + + + +#### モデルのダウンロード + +I2Vワークフローは、テキストから動画へと同じモデルセットを使用します。 + + + + ComfyUI/models/diffusion_models/ に配置してください。 + + + ComfyUI/models/text_encoders/ に配置してください。 + + + ComfyUI/models/text_encoders/ に配置してください。 + + + ComfyUI/models/vae/ に配置してください。 + + + ComfyUI/models/vae/ に配置してください。 + + + ComfyUI/models/latent_upscale_models/ に配置してください。 + + + +#### モデルの保存場所 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### プロンプトのヒント + +1. **次に起こることを説明する**: 入力画像の後に続く動き、カメラの動き、サウンドを記述します。すでに表示されているものを再説明しないでください。 +2. **最初のフレームを固定する**: 続きを書くときは、「提供された開始画像を最初のフレームとして使用する」のような表現を使います。 +3. **オーディオ**: シーンの会話とサウンドを記述します。モデルが同期したオーディオを生成します。 + +--- + +### LTX-2.5 FLF2V + +最初のフレーム画像と最後のフレーム画像の間のビデオを生成し、両方のフレームを1つのアニメーションシーケンスにブレンドします。 + + + + + + Comfy Cloud で開く + + + JSONをダウンロード、またはテンプレートライブラリで「LTX-2.5 FLF2V」を検索してください + + + + + + ビデオの最初のフレームです。 + + + ビデオの最後のフレームです。 + + + +#### モデルのダウンロード + +FLF2Vワークフローは、テキストから動画へと同じモデルセットを使用しますが、空間アップスケーラーは含まれません。 + + + + ComfyUI/models/diffusion_models/ に配置します + + + ComfyUI/models/text_encoders/ に配置します + + + ComfyUI/models/text_encoders/ に配置します + + + ComfyUI/models/vae/ に配置します + + + ComfyUI/models/vae/ に配置します + + + +#### モデルの保存場所 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ └── 📂 vae/ +│ ├── ltx-2.5-video-vae-bf16.safetensors +│ └── ltx-2.5-audio-vae-bf16.safetensors +``` + +#### プロンプトのヒント + +1. **トランジションを説明する**: 2つのフレーム間で何が起こるかを記述します。カメラの動きやオーディオを含めます +2. **フレームを揃える**: スムーズな補間のために、同じアスペクト比の画像を使用します +3. **オーディオ**: シーンの会話とサウンドを記述します。モデルが同期したオーディオを生成します + +## リソース + +- [Hugging Face モデル](https://huggingface.co/Lightricks/LTX-2.5) +- [LTX モデルページ](https://ltx.io/model/ltx-2-5) +- [GitHub リポジトリ](https://github.com/Lightricks/LTX-Video) +- [ComfyUI-LTXVideo](https://github.com/Lightricks/ComfyUI-LTXVideo/) diff --git a/ko/tutorials/partner-nodes/lightricks/ltx-2-5.mdx b/ko/tutorials/partner-nodes/lightricks/ltx-2-5.mdx new file mode 100644 index 000000000..c00a05fad --- /dev/null +++ b/ko/tutorials/partner-nodes/lightricks/ltx-2-5.mdx @@ -0,0 +1,108 @@ +--- +title: "ComfyUI에서 LTX-2.5 API 비디오 생성" +description: "클라우드 워크플로에서 텍스트 기반 비디오 생성, 이미지 기반 비디오 생성, 첫 번째 및 마지막 프레임 기반 비디오 생성을 위한 LTX-2.5 API 노드 사용법과 Fast 및 Pro 모델 등급에 대해 알아보세요." +sidebarTitle: "LTX-2.5 API" +translationSourceHash: 131a7a07 +translationFrom: tutorials/partner-nodes/lightricks/ltx-2-5.mdx +--- + +import ReqHint from "/snippets/ko/tutorials/partner-nodes/req-hint.mdx"; +import UpdateReminder from "/snippets/ko/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://ltx.io/model/ltx-2-5)는 Lightricks의 오픈 가중치 비디오 및 월드 모델로, LTX-2.5 API 노드를 통해 ComfyUI에서 사용할 수 있습니다. API 워크플로는 Lightricks 서버에서 생성을 실행하므로 모델 다운로드나 로컬 GPU가 필요하지 않으며, 생성된 비디오 1초마다 Comfy API 계정으로 청구됩니다. + +LTX-2.5 API는 두 가지 모델 등급을 제공합니다. 최고 품질을 위한 **LTX-2.5 (Pro)** 와 증류된 모델로 더 빠른 생성을 위한 **LTX-2.5 (Fast)** 입니다. 두 등급 모두 기본 멀티샷 장면과 동기화된 오디오 및 비디오 생성을 지원하며, Fast 등급에서는 최대 4K 해상도로 출력할 수 있습니다. + + + + +## LTX-2.5 API 워크플로 + +템플릿 라이브러리에는 세 가지 API 예제 워크플로가 포함되어 있습니다. + +- **텍스트 기반 비디오 생성**: 텍스트 프롬프트에서 비디오 생성 +- **이미지 기반 비디오 생성**: 단일 시작 이미지에서 비디오 생성 +- **FLF2V**: 첫 번째 프레임과 마지막 프레임 이미지 사이의 비디오 생성 + +### LTX-2.5 텍스트 기반 비디오 생성 + +LTX-2.5 API를 통해 텍스트 프롬프트에서 동기화된 오디오와 함께 비디오를 생성합니다. + + + + + + Comfy Cloud에서 열기 + + + JSON을 다운로드하거나 템플릿 라이브러리에서 "LTX-2.5 T2V"를 검색하세요. + + + +#### 프롬프트 팁 + +1. **모델**: 최고 품질을 위해서는 `LTX-2.5 (Pro)`를, 더 빠른 생성을 위해서는 `LTX-2.5 (Fast)`를 선택하세요 +2. **재생 시간**: Pro는 2-10초, Fast는 2-20초를 지원합니다. 10초를 초과하는 재생 시간은 720p 또는 1080p 해상도와 24/25 FPS가 필요합니다 +3. **해상도 및 비율**: `1280x720`(가로) 또는 `720x1280`(세로) 같은 해상도 프리셋을 선택하세요. Fast는 `2560x1440` 및 `3840x2160` 옵션을 추가로 제공합니다 +4. **FPS**: Fast에서는 24, 25, 48, 50 FPS 중, Pro에서는 24, 25, 50 FPS 중에서 선택하세요 +5. **오디오**: 오디오 생성은 기본적으로 켜져 있습니다. 무성 비디오만 필요한 경우 노드의 고급 설정에서 비활성화하세요 +6. **프롬프트**: 장면, 동작, 캐릭터, 카메라 움직임, 원하는 오디오를 포함해 전체 장면을 하나의 자연스러운 문단으로 설명하세요 + +### LTX-2.5 이미지 기반 비디오 생성 + +LTX-2.5 API를 통해 단일 시작 이미지에서 동기화된 오디오와 함께 비디오를 생성합니다. + + + + + + Comfy Cloud에서 열기 + + + JSON을 다운로드하거나 템플릿 라이브러리에서 "LTX-2.5 I2V"를 검색하세요. + + + + + + 워크플로의 시작 프레임이거나, 직접 만든 이미지를 사용할 수 있습니다. + + + +#### 프롬프트 팁 + +1. **모델**: 이 워크플로는 더 빠른 반복을 위해 `LTX-2.5 (Fast)`가 선택된 상태로 제공됩니다. 최고 품질을 원한다면 `LTX-2.5 (Pro)`로 전환하세요 +2. **단일 이미지만 지원**: 이미지 입력에는 정확히 하나의 시작 이미지를 넣을 수 있습니다 +3. **프롬프트**: 시작 이미지에서 이어지는 동작, 카메라 움직임, 사운드 등 그 다음에 일어날 일을 설명하세요. "제공된 시작 이미지를 첫 번째 프레임으로 사용하세요"와 같은 표현은 생성을 안정적으로 이끄는 데 도움이 됩니다 +4. **재생 시간 및 해상도**: Fast 등급은 최대 20초와 최대 4K 해상도를 지원합니다. 10초를 초과하는 재생 시간에는 24/25 FPS에서 720p/1080p가 필요합니다 + +### LTX-2.5 FLF2V + +LTX-2.5 API를 통해 첫 번째 프레임과 마지막 프레임 이미지 사이의 비디오를 생성하고, 두 프레임을 하나의 애니메이션 시퀀스로 자연스럽게 연결합니다. + + + + + + Comfy Cloud에서 열기 + + + JSON을 다운로드하거나 템플릿 라이브러리에서 "LTX-2.5 FLF2V"를 검색하세요. + + + + + + 비디오의 첫 번째 프레임입니다. + + + 비디오의 마지막 프레임으로, 노드의 `last_frame` 입력에 연결됩니다. + + + +#### 프롬프트 팁 + +1. **프레임**: 워크플로는 첫 번째 이미지를 `image` 입력에, 두 번째 이미지를 `last_frame` 입력에 연결합니다. 모델은 두 프레임 사이의 움직임을 생성합니다 +2. **프레임 정렬 유지**: 부드러운 보간을 위해 동일한 화면 비율의 이미지를 사용하세요 +3. **모델**: 이 워크플로는 `LTX-2.5 (Pro)`가 선택된 상태로 제공됩니다. 더 빠른 생성을 위해 `LTX-2.5 (Fast)`로 전환하세요 +4. **프롬프트**: 프레임 사이의 전환과 원하는 오디오(대사, 효과음, 음악)를 설명하세요 \ No newline at end of file diff --git a/ko/tutorials/partner-nodes/pricing.mdx b/ko/tutorials/partner-nodes/pricing.mdx index 43e1202f6..8875408c3 100644 --- a/ko/tutorials/partner-nodes/pricing.mdx +++ b/ko/tutorials/partner-nodes/pricing.mdx @@ -3,7 +3,7 @@ title: "Pricing" description: "이 문서에는 현재 파트너 노드의 가격 정책이 나와 있습니다." sidebarTitle: "가격 정책" mode: "wide" -translationSourceHash: 6ae911be +translationSourceHash: 320cca0f translationFrom: tutorials/partner-nodes/pricing.mdx translationBlockHashes: "_intro": 471c1619 @@ -21,7 +21,7 @@ translationBlockHashes: "Ideogram": a6ecfb20 "Krea": fb20c257 "Kling": d61f855f - "Lightricks": 4870964f + "Lightricks": 838c1ed7 "Luma": 655207a5 "Meshy": 6cd06881 "Minimax": bf30472e @@ -636,16 +636,22 @@ Seedream 5.0 Pro는 해상도 기반 요금제를 사용합니다: **1K** (≤2. 총 크레딧 = **(크레딧 / 초) × `duration`**. -| 모델 | 해상도 | 크레딧 | -| :----------- | :--------- | :----------- | -| LTX-2 (Pro) | 1080p | 12.66 / 초 | -| LTX-2 (Pro) | 1440p | 25.32 / 초 | -| LTX-2 (Pro) | 4K | 50.64 / 초 | -| LTX-2 (Fast) | 1080p | 8.44 / 초 | -| LTX-2 (Fast) | 1440p | 16.88 / 초 | -| LTX-2 (Fast) | 4K | 33.76 / 초 | - -해상도는 `1920x1080`, `2560x1440`, `3840x2160`에 해당합니다. 10초를 초과하는 지속 시간은 **LTX-2 (Fast)** 에서 1080p 및 25 FPS로만 사용할 수 있습니다. +| 모델 | 해상도 | 크레딧 | +| :------------ | :--------- | :---------- | +| LTX-2.5 (Pro) | 720p | 36.21 / 초 | +| LTX-2.5 (Pro) | 1080p | 51.29 / 초 | +| LTX-2.5 (Fast)| 720p | 27.16 / 초 | +| LTX-2.5 (Fast)| 1080p | 39.22 / 초 | +| LTX-2.5 (Fast)| 1440p | 57.33 / 초 | +| LTX-2.5 (Fast)| 4K | 90.52 / 초 | +| LTX-2 (Pro) | 1080p | 12.66 / 초 | +| LTX-2 (Pro) | 1440p | 25.32 / 초 | +| LTX-2 (Pro) | 4K | 50.64 / 초 | +| LTX-2 (Fast) | 1080p | 8.44 / 초 | +| LTX-2 (Fast) | 1440p | 16.88 / 초 | +| LTX-2 (Fast) | 4K | 33.76 / 초 | + +해상도는 `1280x720` / `720x1280`(720p), `1920x1080` / `1080x1920`(1080p), `2560x1440` / `1440x2560`(1440p), `3840x2160` / `2160x3840`(4K)에 해당합니다. **LTX-2.5 (Fast)** 에서 10초를 초과하는 지속 시간은 720p 또는 1080p 해상도와 24/25 FPS가 필요합니다. **LTX-2 (Fast)** 에서 10초를 초과하는 지속 시간은 1080p 및 25 FPS로만 사용할 수 있습니다. ## Luma diff --git a/ko/tutorials/video/ltx/ltx-2-5.mdx b/ko/tutorials/video/ltx/ltx-2-5.mdx new file mode 100644 index 000000000..f7f4378f8 --- /dev/null +++ b/ko/tutorials/video/ltx/ltx-2-5.mdx @@ -0,0 +1,259 @@ +--- +title: "LTX-2.5: ComfyUI 워크플로 예시" +description: "ComfyUI에서 LTX-2.5를 사용하는 방법을 알아보세요. 텍스트 기반 비디오 생성, 이미지 기반 비디오 생성, 첫 번째-마지막 프레임 기반 비디오 생성용 네이티브 워크플로와 pixel-diffusion 품질, 네이티브 멀티샷 장면, 동기화된 오디오를 제공합니다." +sidebarTitle: "LTX-2.5" +translationSourceHash: 34280157 +translationFrom: tutorials/video/ltx/ltx-2-5.mdx +translationBlockHashes: + "_intro": b8e0c526 + "Key features": c81ccf83 + "Getting started": d9d37046 + "ComfyUI Native Workflows": 616316a6 + "Resources": f89b58b6 +--- + +import UpdateReminder from "/snippets/ko/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5)는 Lightricks의 오픈 가중치 비디오 및 월드 모델로, 이제 ComfyUI에서 기본 지원됩니다. 이 모델은 LTX-2.3을 기반으로 장면 복잡도에 따라 렌더링 컴퓨팅을 할당하는 Diffusion Fidelity Rendering 파이프라인을 도입했습니다. LTX-2.3의 동기화된 오디오 및 비디오 생성을 유지하면서 네이티브 멀티샷을 추가했습니다. 한 번의 생성으로 여러 개의 연결된 샷을 만들어 컷 전반에 걸쳐 캐릭터, 환경, 조명, 음성을 유지합니다. + +LTX-2.5는 복잡한 프롬프트에서 여러 주제, 작업, 조명, 카메라 방향을 유지하는 맞춤형 Gemma 4 12B 텍스트 인코더와, 짧은 프롬프트를 상세한 지시문으로 확장하는 프롬프트 강화기를 제공합니다. 자동 재생 시간은 디퓨전 시작 이전에 설명된 작업에서 클립 길이를 예측합니다. 증류 모델은 더 빠른 생성을 제공하며, 출력은 네이티브 4K HDR, 최대 50 FPS, 전문가용 후반 작업 파이프라인을 위한 RAW 워크플로를 지원합니다. + +ComfyUI는 LTX-2.5용 기본 워크플로 3개를 제공합니다: + +- **텍스트 기반 비디오 생성** (T2V): 텍스트 프롬프트에서 비디오 생성 +- **이미지 기반 비디오 생성** (I2V): 입력 이미지에서 비디오 생성 +- **FLF2V**: 첫 번째 프레임과 마지막 프레임 이미지 사이를 보간 + + + +## 주요 기능 + +- **Diffusion 충실도 렌더링**: 키프레임 우선 생성 방식으로 모든 장면을 고충실도 키프레임의 그리드 위에 구성하여 픽셀 단위 품질을 구현합니다. +- **Diffusion 비디오 디코더**: 빠른 모션에서도 더 선명한 얼굴, 읽기 쉬운 텍스트, 더 적은 번짐을 제공합니다. +- **네이티브 멀티샷**: 한 번의 생성으로 여러 개의 연결된 샷을 만들어내며, 컷 사이에서도 캐릭터, 환경, 조명, 음성, 스타일을 유지합니다. +- **Gemma 4 12B 텍스트 인코더**: 복잡한 프롬프트에서도 여러 대상, 작업, 조명, 디테일, 카메라 방향을 유지합니다. +- **프롬프트 강화기**: 짧은 프롬프트를 풍부한 시네마틱 지시문으로 확장하는 경량 모델입니다. +- **자동 재생 시간**: diffusion이 시작되기 이전에 설명된 작업으로부터 적절한 클립 길이를 예측합니다. +- **네이티브 4K HDR**: 전문 후반 작업용으로 제작된 고해상도 HDR 출력으로, 최대 50 FPS에서 동기화된 오디오-비디오를 지원합니다. + +## 시작하기 + +LTX-2.5는 ComfyUI에서 기본적으로 지원됩니다. 시작하려면: + +1. ComfyUI를 최신 버전으로 업데이트하세요 +2. **템플릿 라이브러리** > **비디오** > 원하는 LTX-2.5 워크플로를 선택하세요 +3. Hugging Face에서 모델 저장소에 대한 액세스를 요청한 다음(아래 참조), 팝업에 따라 모델을 다운로드하고 워크플로를 실행하세요 + + + LTX-2.5 모델 파일은 Hugging Face의 액세스 제한(gated) 저장소인 [Lightricks/LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5)에 호스팅되어 있습니다. 저장소 페이지를 열고 모델 라이선스에 동의한 후 액세스 요청이 승인될 때까지 기다렸다가 모델을 다운로드하세요. 액세스 권한이 없으면 모델 다운로드가 실패합니다. + + +## ComfyUI 네이티브 워크플로 + +### LTX-2.5 텍스트 기반 비디오 생성 (T2V) + +텍스트 프롬프트로 비디오를 생성하고, 선택적으로 공간 업스케일링을 적용하여 해상도를 향상시킬 수 있습니다. + + + + + + Comfy Cloud에서 열기 + + + JSON을 다운로드하거나 템플릿 라이브러리에서 "LTX-2.5 T2V"를 검색하세요. + + + +#### 모델 다운로드 + + + + ComfyUI/models/diffusion_models/ 폴더에 배치하세요. + + + ComfyUI/models/text_encoders/ 폴더에 배치하세요. + + + ComfyUI/models/text_encoders/ 폴더에 배치하세요. + + + ComfyUI/models/vae/ 폴더에 배치하세요. + + + ComfyUI/models/vae/ 폴더에 배치하세요. + + + ComfyUI/models/latent_upscale_models/ 폴더에 배치하세요. + + + +#### 모델 저장 위치 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### 프롬프트 작성 팁 + +1. **전체 장면 설명**: 샷 유형, 장면, 동작, 캐릭터, 카메라 움직임을 하나의 자연스러운 문단에 포함하세요. +2. **오디오**: 장면에 필요한 소리와 대사를 설명하세요. 모델이 비디오와 동기화된 오디오를 함께 생성합니다. +3. **프롬프트 강화기 사용**: 이 워크플로는 짧은 프롬프트를 상세한 시네마틱 연출 지침으로 자동 확장합니다. 간단한 장면 아이디어를 작성하고 나머지는 강화기가 처리하도록 하세요. + +--- + +### LTX-2.5 이미지 기반 비디오 생성 (I2V) + +입력 이미지에서 비디오를 생성하며, 모델이 제공된 첫 프레임으로부터 장면을 애니메이션화합니다. + + + + + + Comfy Cloud에서 열기 + + + JSON 다운로드 또는 템플릿 라이브러리에서 "LTX-2.5 I2V" 검색 + + + + + + 기본 입력 이미지를 다운로드하거나 자신의 이미지를 사용하세요. + + + +#### 모델 다운로드 + +I2V 워크플로는 텍스트 기반 비디오 생성과 동일한 모델 세트를 사용합니다. + + + + ComfyUI/models/diffusion_models/에 배치하세요. + + + ComfyUI/models/text_encoders/에 배치하세요. + + + ComfyUI/models/text_encoders/에 배치하세요. + + + ComfyUI/models/vae/에 배치하세요. + + + ComfyUI/models/vae/에 배치하세요. + + + ComfyUI/models/latent_upscale_models/에 배치하세요. + + + +#### 모델 저장 위치 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### 프롬프트 작성 팁 + +1. **다음에 일어날 일을 설명하세요**: 입력 이미지 이후의 움직임, 카메라 이동, 소리를 작성하세요. 이미 보이는 내용을 다시 설명하지 마세요. +2. **첫 프레임을 고정하세요**: 이어지는 내용을 작성할 때 "제공된 시작 이미지를 첫 프레임으로 사용하세요"와 같은 표현을 사용하세요. +3. **오디오**: 장면의 대화와 소리를 설명하세요. 모델이 동기화된 오디오를 생성합니다. + +--- + +### LTX-2.5 FLF2V + +첫 번째 프레임과 마지막 프레임 이미지 사이를 연결하는 비디오를 생성하여, 두 프레임을 하나의 애니메이션 시퀀스로 혼합합니다. + + + + + + Comfy Cloud에서 열기 + + + JSON을 다운로드하거나 템플릿 라이브러리에서 "LTX-2.5 FLF2V"를 검색하세요. + + + + + + 비디오의 첫 번째 프레임입니다. + + + 비디오의 마지막 프레임입니다. + + + +#### 모델 다운로드 + +FLF2V 워크플로는 공간 업스케일러 없이 텍스트 기반 비디오 생성과 동일한 모델 세트를 사용합니다. + + + + ComfyUI/models/diffusion_models/에 넣으세요. + + + ComfyUI/models/text_encoders/에 넣으세요. + + + ComfyUI/models/text_encoders/에 넣으세요. + + + ComfyUI/models/vae/에 넣으세요. + + + ComfyUI/models/vae/에 넣으세요. + + + +#### 모델 저장 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ └── 📂 vae/ +│ ├── ltx-2.5-video-vae-bf16.safetensors +│ └── ltx-2.5-audio-vae-bf16.safetensors +``` + +#### 프롬프트 팁 + +1. **전환 설명**: 두 프레임 사이에 일어나는 일을 카메라 움직임과 오디오를 포함해 작성하세요. +2. **프레임 정렬 유지**: 부드러운 보간을 위해 동일한 화면 비율의 이미지를 사용하세요. +3. **오디오**: 장면의 대화와 소리를 설명하세요. 모델이 동기화된 오디오를 생성합니다. + +## 리소스 + +- [Hugging Face 모델](https://huggingface.co/Lightricks/LTX-2.5) +- [LTX 모델 페이지](https://ltx.io/model/ltx-2-5) +- [GitHub 저장소](https://github.com/Lightricks/LTX-Video) +- [ComfyUI-LTXVideo](https://github.com/Lightricks/ComfyUI-LTXVideo/) diff --git a/tutorials/partner-nodes/lightricks/ltx-2-5.mdx b/tutorials/partner-nodes/lightricks/ltx-2-5.mdx new file mode 100644 index 000000000..cef7bcc74 --- /dev/null +++ b/tutorials/partner-nodes/lightricks/ltx-2-5.mdx @@ -0,0 +1,106 @@ +--- +title: "LTX-2.5 API Video Generation in ComfyUI" +description: "Learn how to use the LTX-2.5 API nodes in ComfyUI with cloud workflows for text-to-video, image-to-video, and first-last-frame-to-video, with Fast and Pro model tiers." +sidebarTitle: "LTX-2.5 API" +--- + +import ReqHint from "/snippets/tutorials/partner-nodes/req-hint.mdx"; +import UpdateReminder from "/snippets/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://ltx.io/model/ltx-2-5) is Lightricks' open-weights video and world model, available in ComfyUI through the LTX-2.5 API nodes. The API workflows run generation on Lightricks' servers, so no model downloads or local GPU is required, and each second of generated video is billed to your Comfy API account. + +The LTX-2.5 API offers two model tiers: **LTX-2.5 (Pro)** for maximum quality, and **LTX-2.5 (Fast)** for faster generation with the distilled model. Both tiers support synchronized audio and video generation with native multishot scenes, and output up to 4K resolution on the Fast tier. + + + + +## LTX-2.5 API workflows + +The template library ships with three API example workflows: + +- **Text to Video**: Generate videos from text prompts +- **Image to Video**: Generate videos from a single start image +- **FLF2V**: Generate a video between a first and last frame image + +### LTX-2.5 Text to Video + +Generate videos from text prompts through the LTX-2.5 API, with synchronized audio. + + + + + + Open in Comfy Cloud + + + Download JSON or search "LTX-2.5 T2V" in Template Library + + + +#### Prompting tips + +1. **Model**: Choose `LTX-2.5 (Pro)` for the highest quality, or `LTX-2.5 (Fast)` for faster generation +2. **Duration**: Pro supports 2-10 seconds; Fast supports 2-20 seconds. Durations over 10 seconds require a 720p or 1080p resolution and 24/25 FPS +3. **Resolution and ratio**: Choose a resolution preset such as `1280x720` (landscape) or `720x1280` (portrait); Fast adds `2560x1440` and `3840x2160` options +4. **FPS**: Choose between 24, 25, 48, or 50 FPS on Fast, and 24, 25, or 50 FPS on Pro +5. **Audio**: Audio generation is on by default; disable it in the node's advanced settings if you only need a silent video +6. **Prompt**: Describe the whole scene in one flowing paragraph, including the shot, action, characters, camera movement, and the audio you want + +### LTX-2.5 Image to Video + +Generate a video from a single start image through the LTX-2.5 API, with synchronized audio. + + + + + + Open in Comfy Cloud + + + Download JSON or search "LTX-2.5 I2V" in Template Library + + + + + + Start frame for the workflow, or use your own image. + + + +#### Prompting tips + +1. **Model**: This workflow ships with `LTX-2.5 (Fast)` selected for faster iteration; switch to `LTX-2.5 (Pro)` for maximum quality +2. **Single image only**: The image input accepts exactly one start image +3. **Prompt**: Describe what happens next: the motion, camera movement, and sounds that follow from the start image. Phrasing like "Use the provided start image as the first frame" helps anchor the generation +4. **Duration and resolution**: The Fast tier supports up to 20 seconds and resolutions up to 4K; durations over 10 seconds require 720p/1080p at 24/25 FPS + +### LTX-2.5 FLF2V + +Generate a video between a first and last frame image through the LTX-2.5 API, blending both frames into a single animated sequence. + + + + + + Open in Comfy Cloud + + + Download JSON or search "LTX-2.5 FLF2V" in Template Library + + + + + + First frame of the video. + + + Last frame of the video, connected to the `last_frame` input of the node. + + + +#### Prompting tips + +1. **Frames**: The workflow connects the first image to the `image` input and the second image to the `last_frame` input; the model generates the motion between them +2. **Keep frames aligned**: Use images with the same aspect ratio for a smooth interpolation +3. **Model**: This workflow ships with `LTX-2.5 (Pro)` selected; switch to `LTX-2.5 (Fast)` for faster generation +4. **Prompt**: Describe the transition between the frames, and the audio (dialogue, SFX, music) you want diff --git a/tutorials/partner-nodes/pricing.mdx b/tutorials/partner-nodes/pricing.mdx index 69c52e1c9..899af261a 100644 --- a/tutorials/partner-nodes/pricing.mdx +++ b/tutorials/partner-nodes/pricing.mdx @@ -589,16 +589,22 @@ Total image credits = **(credits / image) × `n`** on **Kling 3.0 Image**; multi Total credits = **(credits / sec) × `duration`**. -| Model | resolution | Credits | -| :----------- | :--------- | :----------- | -| LTX-2 (Pro) | 1080p | 12.66 / sec | -| LTX-2 (Pro) | 1440p | 25.32 / sec | -| LTX-2 (Pro) | 4K | 50.64 / sec | -| LTX-2 (Fast) | 1080p | 8.44 / sec | -| LTX-2 (Fast) | 1440p | 16.88 / sec | -| LTX-2 (Fast) | 4K | 33.76 / sec | - -Resolution maps to `1920x1080`, `2560x1440`, and `3840x2160`. Durations over 10s are only available on **LTX-2 (Fast)** at 1080p and 25 FPS. +| Model | resolution | Credits | +| :------------ | :--------- | :---------- | +| LTX-2.5 (Pro) | 720p | 36.21 / sec | +| LTX-2.5 (Pro) | 1080p | 51.29 / sec | +| LTX-2.5 (Fast)| 720p | 27.16 / sec | +| LTX-2.5 (Fast)| 1080p | 39.22 / sec | +| LTX-2.5 (Fast)| 1440p | 57.33 / sec | +| LTX-2.5 (Fast)| 4K | 90.52 / sec | +| LTX-2 (Pro) | 1080p | 12.66 / sec | +| LTX-2 (Pro) | 1440p | 25.32 / sec | +| LTX-2 (Pro) | 4K | 50.64 / sec | +| LTX-2 (Fast) | 1080p | 8.44 / sec | +| LTX-2 (Fast) | 1440p | 16.88 / sec | +| LTX-2 (Fast) | 4K | 33.76 / sec | + +Resolution maps to `1280x720` / `720x1280` (720p), `1920x1080` / `1080x1920` (1080p), `2560x1440` / `1440x2560` (1440p), and `3840x2160` / `2160x3840` (4K). Durations over 10s on **LTX-2.5 (Fast)** require a 720p or 1080p resolution and 24/25 FPS; on **LTX-2 (Fast)** they require 1080p and 25 FPS. ## Luma diff --git a/tutorials/video/ltx/ltx-2-5.mdx b/tutorials/video/ltx/ltx-2-5.mdx new file mode 100644 index 000000000..b4a730e15 --- /dev/null +++ b/tutorials/video/ltx/ltx-2-5.mdx @@ -0,0 +1,251 @@ +--- +title: "LTX-2.5: ComfyUI Workflow Examples" +description: "Learn how to use LTX-2.5 in ComfyUI with native workflows for text-to-video, image-to-video, and first-last-frame-to-video, with pixel-diffusion quality, native multishot scenes, and synchronized audio." +sidebarTitle: "LTX-2.5" +--- + +import UpdateReminder from "/snippets/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5) is Lightricks' open-weights video and world model, natively supported in ComfyUI. It builds on LTX-2.3 with a Diffusion Fidelity Rendering pipeline that allocates rendering compute by scene complexity. The model keeps the synchronized audio and video generation of LTX-2.3 and adds native multishot: one generation produces multiple connected shots that hold character, environment, lighting, and voice across cuts. + +LTX-2.5 ships with a custom Gemma 4 12B text encoder that retains multiple subjects, actions, lighting, and camera direction across complex prompts, plus a prompt enhancer that expands short prompts into detailed instructions. Auto duration predicts the clip length from the described action before diffusion begins. The distilled model provides faster generation, and output supports native 4K HDR at up to 50 FPS with a RAW workflow for professional finishing pipelines. + +ComfyUI ships with three native workflows for LTX-2.5: + +- **Text-to-Video** (T2V): Generate videos from text prompts +- **Image-to-Video** (I2V): Generate videos from an input image +- **FLF2V**: Interpolate between a first and last frame image + + + +## Key features + +- **Diffusion Fidelity Rendering**: keyframes-first generation that builds every scene on a grid of high-fidelity keyframes +- **Diffusion Video Decoder**: sharper faces, legible text, and fewer smears in fast motion +- **Native multishot**: one generation yields multiple connected shots holding character, environment, lighting, voice, and style across cuts +- **Gemma 4 12B text encoder**: retains multiple subjects, actions, lighting, detail, and camera direction across complex prompts +- **Prompt enhancer**: a lightweight model that expands short prompts into detailed instructions +- **Auto duration**: predicts the right clip length from the described action before diffusion begins +- **Native 4K HDR**: high-resolution HDR output built for professional finishing, with synchronized audio-video at up to 50 FPS + +## Getting started + +LTX-2.5 is natively supported in ComfyUI. To get started: + +1. Update ComfyUI to the latest version +2. Go to **Template Library** > **Video** > choose any LTX-2.5 workflow +3. Request access to the model repository on Hugging Face (see below), then follow the pop-up to download models and run the workflow + + + The LTX-2.5 model files are hosted in the gated [Lightricks/LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5) repository on Hugging Face. Open the repository page, accept the model license, and wait for your access request to be approved before downloading the models. Model downloads will fail without access. + + +## ComfyUI Native Workflows + +### LTX-2.5 Text to Video (T2V) + +Generate videos from text prompts with optional spatial upscaling for enhanced resolution. + + + + + + Open in Comfy Cloud + + + Download JSON or search "LTX-2.5 T2V" in Template Library + + + +#### Model downloads + + + + Place in ComfyUI/models/diffusion_models/ + + + Place in ComfyUI/models/text_encoders/ + + + Place in ComfyUI/models/text_encoders/ + + + Place in ComfyUI/models/vae/ + + + Place in ComfyUI/models/vae/ + + + Place in ComfyUI/models/latent_upscale_models/ + + + +#### Model storage + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### Prompting tips + +1. **Describe the whole scene**: Include the shot type, scene, action, characters, and camera movement in one flowing paragraph +2. **Audio**: Describe sounds and dialogue needed for the scene; the model generates synchronized audio together with the video +3. **Use the prompt enhancer**: The workflow expands short prompts into detailed cinematic instructions automatically; write a simple scene idea and let the enhancer handle the rest + +--- + +### LTX-2.5 Image to Video (I2V) + +Generate videos from an input image, with the model animating the scene from the provided first frame. + + + + + + Open in Comfy Cloud + + + Download JSON or search "LTX-2.5 I2V" in Template Library + + + + + + Download the default input image, or use your own image. + + + +#### Model downloads + +The I2V workflow uses the same model set as Text-to-Video. + + + + Place in ComfyUI/models/diffusion_models/ + + + Place in ComfyUI/models/text_encoders/ + + + Place in ComfyUI/models/text_encoders/ + + + Place in ComfyUI/models/vae/ + + + Place in ComfyUI/models/vae/ + + + Place in ComfyUI/models/latent_upscale_models/ + + + +#### Model storage + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### Prompting tips + +1. **Describe what happens next**: Write the motion, camera movement, and sounds that follow from the input image; do not re-describe what is already visible +2. **Anchor the first frame**: Use phrasing like "Use the provided start image as the first frame" when writing a continuation +3. **Audio**: Describe dialogue and sounds for the scene; the model generates synchronized audio + +--- + +### LTX-2.5 FLF2V + +Generate a video between a first and last frame image, blending both frames into a single animated sequence. + + + + + + Open in Comfy Cloud + + + Download JSON or search "LTX-2.5 FLF2V" in Template Library + + + + + + First frame of the video. + + + Last frame of the video. + + + +#### Model downloads + +The FLF2V workflow uses the same model set as Text-to-Video, without the spatial upscaler. + + + + Place in ComfyUI/models/diffusion_models/ + + + Place in ComfyUI/models/text_encoders/ + + + Place in ComfyUI/models/text_encoders/ + + + Place in ComfyUI/models/vae/ + + + Place in ComfyUI/models/vae/ + + + +#### Model storage + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ └── 📂 vae/ +│ ├── ltx-2.5-video-vae-bf16.safetensors +│ └── ltx-2.5-audio-vae-bf16.safetensors +``` + +#### Prompting tips + +1. **Describe the transition**: Write what happens between the two frames, including camera movement and audio +2. **Keep frames aligned**: Use images with the same aspect ratio for a smooth interpolation +3. **Audio**: Describe dialogue and sounds for the scene; the model generates synchronized audio + +## Resources + +- [Hugging Face Model](https://huggingface.co/Lightricks/LTX-2.5) +- [LTX Model Page](https://ltx.io/model/ltx-2-5) +- [GitHub Repository](https://github.com/Lightricks/LTX-Video) +- [ComfyUI-LTXVideo](https://github.com/Lightricks/ComfyUI-LTXVideo/) diff --git a/zh/tutorials/partner-nodes/lightricks/ltx-2-5.mdx b/zh/tutorials/partner-nodes/lightricks/ltx-2-5.mdx new file mode 100644 index 000000000..5744082c1 --- /dev/null +++ b/zh/tutorials/partner-nodes/lightricks/ltx-2-5.mdx @@ -0,0 +1,108 @@ +--- +title: "在 ComfyUI 中使用 LTX-2.5 API 生成视频" +description: "了解如何结合云端工作流,在 ComfyUI 中使用 LTX-2.5 API 节点进行文生视频、图生视频和首尾帧生视频,并提供 Fast 与 Pro 两种模型档位。" +sidebarTitle: "LTX-2.5 API" +translationSourceHash: 131a7a07 +translationFrom: tutorials/partner-nodes/lightricks/ltx-2-5.mdx +--- + +import ReqHint from "/snippets/zh/tutorials/partner-nodes/req-hint.mdx"; +import UpdateReminder from "/snippets/zh/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://ltx.io/model/ltx-2-5) 是 Lightricks 的开源权重视频与世界模型,可通过 LTX-2.5 API 节点在 ComfyUI 中使用。API 工作流在 Lightricks 的服务器上运行生成任务,因此无需下载模型,也无需本地 GPU,生成的视频将按秒计入您的 Comfy API 账户。 + +LTX-2.5 API 提供两种模型档位:**LTX-2.5 (Pro)** 面向最高质量,**LTX-2.5 (Fast)** 采用蒸馏模型,生成速度更快。两种档位均支持音视频同步生成,并原生支持多镜头场景,其中 Fast 档位最高可输出 4K 分辨率。 + + + + +## LTX-2.5 API 工作流 + +模板库内置了三个 API 示例工作流: + +- **文生视频**:根据文本提示词生成视频 +- **图生视频**:根据单张起始图像生成视频 +- **FLF2V(首尾帧生视频)**:在首帧和尾帧图像之间生成视频 + +### LTX-2.5 文生视频 + +通过 LTX-2.5 API 根据文本提示词生成视频,并同步生成音频。 + + + + + + 在 Comfy Cloud 中打开 + + + 在模板库中下载 JSON 或搜索"LTX-2.5 T2V" + + + +#### 提示技巧 + +1. **模型**:选择 `LTX-2.5 (Pro)` 以获得最高质量,或选择 `LTX-2.5 (Fast)` 以获得更快的生成速度 +2. **时长**:Pro 支持 2-10 秒,Fast 支持 2-20 秒。时长超过 10 秒时,需使用 720p 或 1080p 分辨率以及 24/25 FPS +3. **分辨率与比例**:选择分辨率预设,如 `1280x720`(横屏)或 `720x1280`(竖屏);Fast 额外提供 `2560x1440` 和 `3840x2160` 选项 +4. **FPS**:Fast 可选 24、25、48 或 50 FPS,Pro 可选 24、25 或 50 FPS +5. **音频**:音频生成默认开启;如果只需要无声视频,可在节点的高级设置中关闭 +6. **提示词**:用一段连贯的文字描述整个场景,包括镜头、动作、角色、相机运动以及您想要的音频 + +### LTX-2.5 图生视频 + +通过 LTX-2.5 API 根据单张起始图像生成视频,并同步生成音频。 + + + + + + 在 Comfy Cloud 中打开 + + + 在模板库中下载 JSON 或搜索"LTX-2.5 I2V" + + + + + + 作为工作流的起始帧,或使用您自己的图像。 + + + +#### 提示技巧 + +1. **模型**:此工作流默认选用 `LTX-2.5 (Fast)`,便于快速迭代;如需最高质量,可切换为 `LTX-2.5 (Pro)` +2. **仅限单个图像**:图像输入仅接受一张起始图像 +3. **提示词**:描述接下来发生的内容:从起始图像延续的运动、相机运动和声音。使用类似"将提供的起始图像作为第一帧"的措辞,有助于锚定生成方向 +4. **时长与分辨率**:Fast 档位支持最长 20 秒、最高 4K 分辨率;时长超过 10 秒时,需使用 720p/1080p 并设为 24/25 FPS + +### LTX-2.5 FLF2V(首尾帧生视频) + +通过 LTX-2.5 API 在首帧和尾帧图像之间生成视频,将两帧画面融合为一段完整的动画序列。 + + + + + + 在 Comfy Cloud 中打开 + + + 在模板库中下载 JSON 或搜索"LTX-2.5 FLF2V" + + + + + + 视频的首帧。 + + + 视频的尾帧,连接到节点的 `last_frame` 输入。 + + + +#### 提示技巧 + +1. **帧**:工作流将第一张图像连接到 `image` 输入,将第二张图像连接到 `last_frame` 输入;模型负责生成两帧之间的运动 +2. **保持帧对齐**:使用相同宽高比的图像,以获得平滑的插值效果 +3. **模型**:此工作流默认选用 `LTX-2.5 (Pro)`;切换为 `LTX-2.5 (Fast)` 可加快生成速度 +4. **提示词**:描述两帧之间的过渡,以及您想要的音频(对白、音效、音乐) \ No newline at end of file diff --git a/zh/tutorials/partner-nodes/pricing.mdx b/zh/tutorials/partner-nodes/pricing.mdx index cd229d8e6..df2effd18 100644 --- a/zh/tutorials/partner-nodes/pricing.mdx +++ b/zh/tutorials/partner-nodes/pricing.mdx @@ -3,7 +3,7 @@ title: "定价" description: "本文列出了当前合作伙伴节点的定价。" sidebarTitle: "定价" mode: "wide" -translationSourceHash: 6ae911be +translationSourceHash: 320cca0f translationFrom: tutorials/partner-nodes/pricing.mdx translationBlockHashes: "_intro": 471c1619 @@ -21,7 +21,7 @@ translationBlockHashes: "Ideogram": a6ecfb20 "Krea": fb20c257 "Kling": d61f855f - "Lightricks": 4870964f + "Lightricks": 838c1ed7 "Luma": 655207a5 "Meshy": 6cd06881 "Minimax": bf30472e @@ -637,15 +637,21 @@ Seedream 5.0 Pro 使用基于分辨率的计费方式:**1K**(≤2.36 百 总积分 = **(积分/秒) × `时长`**。 | 模型 | 分辨率 | 积分 | -| :----------- | :--------- | :----------- | -| LTX-2 (Pro) | 1080p | 12.66 / 秒 | -| LTX-2 (Pro) | 1440p | 25.32 / 秒 | -| LTX-2 (Pro) | 4K | 50.64 / 秒 | -| LTX-2 (Fast) | 1080p | 8.44 / 秒 | -| LTX-2 (Fast) | 1440p | 16.88 / 秒 | -| LTX-2 (Fast) | 4K | 33.76 / 秒 | - -分辨率映射到 `1920x1080`、`2560x1440` 和 `3840x2160`。超过 10 秒的时长仅 **LTX-2 (Fast)** 在 1080p 和 25 FPS 下可用。 +| :------------ | :--------- | :---------- | +| LTX-2.5 (Pro) | 720p | 36.21 / 秒 | +| LTX-2.5 (Pro) | 1080p | 51.29 / 秒 | +| LTX-2.5 (Fast)| 720p | 27.16 / 秒 | +| LTX-2.5 (Fast)| 1080p | 39.22 / 秒 | +| LTX-2.5 (Fast)| 1440p | 57.33 / 秒 | +| LTX-2.5 (Fast)| 4K | 90.52 / 秒 | +| LTX-2 (Pro) | 1080p | 12.66 / 秒 | +| LTX-2 (Pro) | 1440p | 25.32 / 秒 | +| LTX-2 (Pro) | 4K | 50.64 / 秒 | +| LTX-2 (Fast) | 1080p | 8.44 / 秒 | +| LTX-2 (Fast) | 1440p | 16.88 / 秒 | +| LTX-2 (Fast) | 4K | 33.76 / 秒 | + +分辨率映射到 `1280x720` / `720x1280`(720p)、`1920x1080` / `1080x1920`(1080p)、`2560x1440` / `1440x2560`(1440p)和 `3840x2160` / `2160x3840`(4K)。**LTX-2.5 (Fast)** 超过 10 秒的时长需要 720p 或 1080p 分辨率和 24/25 FPS;**LTX-2 (Fast)** 超过 10 秒的时长仅 1080p 和 25 FPS 下可用。 ## Luma diff --git a/zh/tutorials/video/ltx/ltx-2-5.mdx b/zh/tutorials/video/ltx/ltx-2-5.mdx new file mode 100644 index 000000000..15cb6f7bc --- /dev/null +++ b/zh/tutorials/video/ltx/ltx-2-5.mdx @@ -0,0 +1,258 @@ +--- +title: "LTX-2.5:ComfyUI 工作流示例" +description: "了解如何通过原生工作流在 ComfyUI 中使用 LTX-2.5,实现文生视频、图生视频和首尾帧生视频,具备像素级 diffusion 质量、原生多镜头场景和同步音频。" +sidebarTitle: "LTX-2.5" +translationSourceHash: 34280157 +translationFrom: tutorials/video/ltx/ltx-2-5.mdx +translationBlockHashes: + "_intro": b8e0c526 + "Key features": c81ccf83 + "Getting started": d9d37046 + "ComfyUI Native Workflows": 616316a6 + "Resources": f89b58b6 +--- +import UpdateReminder from "/snippets/zh/tutorials/update-reminder.mdx"; + +[LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5) 是 Lightricks 推出的一款开放权重视频与世界模型,现已获得 ComfyUI 原生支持。它以 LTX-2.3 为基础,采用 Diffusion Fidelity Rendering pipeline,根据场景复杂度分配渲染算力。该模型保留了 LTX-2.3 的音频与视频同步生成能力,并新增了原生多镜头(multishot)功能:一次生成即可产出多个连接式镜头,在剪辑切换之间保持角色、环境、光影与声音的一致性。 + +LTX-2.5 内置定制的 Gemma 4 12B 文本编码器,面对复杂提示词时仍能保留多个主体、动作、光影与相机方向;同时还配备了提示词增强器,可将简短提示词扩展为详细指令。自动时长功能会在扩散开始之前,根据所描述的动作预测片段长度。蒸馏模型提供更快的生成速度,输出支持原生 4K HDR、最高 50 FPS,并为专业后期精修 pipeline 提供 RAW 工作流。 + +ComfyUI 为 LTX-2.5 内置了三个原生工作流: + +- **文生视频**(T2V):根据文本提示词生成视频 +- **图生视频**(I2V):根据输入图像生成视频 +- **FLF2V**:在首帧图像与末帧图像之间进行插值 + + + +## 主要功能 + +- **扩散保真渲染**:关键帧优先生成,在高质量关键帧网格上构建每个场景,实现像素级画质 +- **扩散视频解码器**:快速运动下呈现更清晰的面部、可读的文本和更少的拖影 +- **原生多镜头**:一次生成即可产出多个连接式镜头,在镜头切换中保持角色、环境、光照、声音和风格一致 +- **Gemma 4 12B 文本编码器**:在复杂提示词中保留多个主体、动作、光照、细节和相机方向 +- **提示词增强器**:一个轻量级模型,可将简短提示词扩展为丰富的电影级指令 +- **自动时长**:在扩散开始之前,根据所描述的动作预测正确的片段长度 +- **原生 4K HDR**:专为专业后期制作打造的高分辨率 HDR 输出,支持最高 50 FPS 的音频与视频同步 + +## 入门 + +ComfyUI 原生支持 LTX-2.5。开始使用: + +1. 将 ComfyUI 更新到最新版本 +2. 前往 **模板库** > **视频** > 选择任意 LTX-2.5 工作流 +3. 在 Hugging Face 上申请模型仓库的访问权限(见下文),然后按照弹出窗口提示下载模型并运行工作流 + + + LTX-2.5 模型文件托管在 Hugging Face 上受控访问(gated)的 [Lightricks/LTX-2.5](https://huggingface.co/Lightricks/LTX-2.5) 仓库中。打开仓库页面,接受模型许可协议,并等待访问申请获批后再下载模型。未获得访问权限时,模型下载会失败。 + + +## ComfyUI 原生工作流 + +### LTX-2.5 文生视频 (T2V) + +根据文本提示生成视频,可选择使用空间放大以获得更高分辨率。 + + + + + + 在 Comfy Cloud 中打开 + + + 下载 JSON,或在模板库中搜索“LTX-2.5 T2V” + + + +#### 模型下载 + + + + 放入 ComfyUI/models/diffusion_models/ + + + 放入 ComfyUI/models/text_encoders/ + + + 放入 ComfyUI/models/text_encoders/ + + + 放入 ComfyUI/models/vae/ + + + 放入 ComfyUI/models/vae/ + + + 放入 ComfyUI/models/latent_upscale_models/ + + + +#### 模型存储 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### 提示词技巧 + +1. **描述完整场景**:用一段连贯的文字描述镜头类型、场景、动作、角色和相机移动 +2. **音频**:描述场景所需的声音和对话;模型会与视频一同生成同步音频 +3. **使用提示词增强器**:工作流会自动将简短提示词扩展为详细的电影级指令;写下简单的场景构思,剩下的交给增强器处理 + +--- + +### LTX-2.5 图生视频 (I2V) + +根据输入图像生成视频,模型以提供的第一帧为起点对场景进行动画处理。 + + + + + + 在 Comfy Cloud 中打开 + + + 下载 JSON,或在模板库中搜索“LTX-2.5 I2V” + + + + + + 下载默认输入图像,或使用你自己的图像。 + + + +#### 模型下载 + +I2V 工作流使用的模型与文生视频相同。 + + + + 放入 ComfyUI/models/diffusion_models/ + + + 放入 ComfyUI/models/text_encoders/ + + + 放入 ComfyUI/models/text_encoders/ + + + 放入 ComfyUI/models/vae/ + + + 放入 ComfyUI/models/vae/ + + + 放入 ComfyUI/models/latent_upscale_models/ + + + +#### 模型存储 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ ├── 📂 vae/ +│ │ ├── ltx-2.5-video-vae-bf16.safetensors +│ │ └── ltx-2.5-audio-vae-bf16.safetensors +│ └── 📂 latent_upscale_models/ +│ └── ltx-2.5-latent-spatial-upscaler-x2-bf16-1.0.safetensors +``` + +#### 提示技巧 + +1. **描述接下来发生的事情**:写出输入图像之后发生的运动、相机移动和声音;不要重复描述画面中已可见的内容 +2. **锚定第一帧**:续写时使用类似“将提供的起始图像作为第一帧”的措辞 +3. **音频**:描述场景中的对话和声音;模型会生成同步音频 + +--- + +### LTX-2.5 FLF2V + +在首帧图像和末帧图像之间生成视频,将两帧融合为单个动画序列。 + + + + + + 在 Comfy Cloud 中打开 + + + 下载 JSON 或在模板库中搜索 "LTX-2.5 FLF2V" + + + + + + 视频的首帧。 + + + 视频的末帧。 + + + +#### 模型下载 + +FLF2V 工作流使用与文生视频相同的模型集,但不包含空间放大模型。 + + + + 放到 ComfyUI/models/diffusion_models/ + + + 放到 ComfyUI/models/text_encoders/ + + + 放到 ComfyUI/models/text_encoders/ + + + 放到 ComfyUI/models/vae/ + + + 放到 ComfyUI/models/vae/ + + + +#### 模型存储 + +``` +ComfyUI/ +├── 📂 models/ +│ ├── 📂 diffusion_models/ +│ │ └── ltx-2.5-22b-distilled-transformer-comfy-int8-convrot.safetensors +│ ├── 📂 text_encoders/ +│ │ ├── gemma4-12b-with-proj-ltx-2.5-comfy-int8-convrot.safetensors +│ │ └── gemma4_e2b_it_bf16.safetensors +│ └── 📂 vae/ +│ ├── ltx-2.5-video-vae-bf16.safetensors +│ └── ltx-2.5-audio-vae-bf16.safetensors +``` + +#### 提示词技巧 + +1. **描述过渡**:写出两帧之间发生的内容,包括相机移动和音频 +2. **保持帧对齐**:使用相同宽高比的图像以获得平滑的插值效果 +3. **音频**:描述场景中的对话和声音;模型会生成同步的音频 + +## 资源 + +- [Hugging Face 模型](https://huggingface.co/Lightricks/LTX-2.5) +- [LTX 模型页面](https://ltx.io/model/ltx-2-5) +- [GitHub 仓库](https://github.com/Lightricks/LTX-Video) +- [ComfyUI-LTXVideo](https://github.com/Lightricks/ComfyUI-LTXVideo/)