@nodaro/prompts 1.1.0 → 1.2.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/dist/index.d.cts CHANGED
@@ -3305,10 +3305,18 @@ declare function applyTemplate(template: string, vars: Record<string, string>):
3305
3305
  * 4. compact recipes in MCP tool descriptions point here via get_node_skill
3306
3306
  *
3307
3307
  * Sources, in precedence order (conflicts resolve top-down):
3308
+ * Seedance 2.0:
3308
3309
  * - Official BytePlus ModelArk "Dreamina Seedance 2.0 series prompt guide"
3309
3310
  * https://docs.byteplus.com/en/docs/ModelArk/2222480
3310
3311
  * - Official launch post https://seed.bytedance.com/en/blog/official-launch-of-seedance-2-0
3311
3312
  * - KIE API docs https://docs.kie.ai/market/bytedance/seedance-2
3313
+ * Kling:
3314
+ * - Official "Kling Video 2.6 Audio User Guide"
3315
+ * https://kling.ai/quickstart/klingai-video-26-audio-user-guide
3316
+ * - fal.ai "Kling 3.0 Prompting Guide" https://blog.fal.ai/kling-3-0-prompting-guide/
3317
+ * - KIE API docs https://docs.kie.ai/market/kling/kling-3-0
3318
+ * - Live KIE-path dialogue probe 2026-07-16 (scripted lines spoken verbatim,
3319
+ * lip-synced, on both kling-2.6 and kling-3.0 with sound=true)
3312
3320
  */
3313
3321
  interface ProviderPromptDoctrine {
3314
3322
  /** MODEL_CATALOG ids this doctrine covers. */
package/dist/index.d.ts CHANGED
@@ -3305,10 +3305,18 @@ declare function applyTemplate(template: string, vars: Record<string, string>):
3305
3305
  * 4. compact recipes in MCP tool descriptions point here via get_node_skill
3306
3306
  *
3307
3307
  * Sources, in precedence order (conflicts resolve top-down):
3308
+ * Seedance 2.0:
3308
3309
  * - Official BytePlus ModelArk "Dreamina Seedance 2.0 series prompt guide"
3309
3310
  * https://docs.byteplus.com/en/docs/ModelArk/2222480
3310
3311
  * - Official launch post https://seed.bytedance.com/en/blog/official-launch-of-seedance-2-0
3311
3312
  * - KIE API docs https://docs.kie.ai/market/bytedance/seedance-2
3313
+ * Kling:
3314
+ * - Official "Kling Video 2.6 Audio User Guide"
3315
+ * https://kling.ai/quickstart/klingai-video-26-audio-user-guide
3316
+ * - fal.ai "Kling 3.0 Prompting Guide" https://blog.fal.ai/kling-3-0-prompting-guide/
3317
+ * - KIE API docs https://docs.kie.ai/market/kling/kling-3-0
3318
+ * - Live KIE-path dialogue probe 2026-07-16 (scripted lines spoken verbatim,
3319
+ * lip-synced, on both kling-2.6 and kling-3.0 with sound=true)
3312
3320
  */
3313
3321
  interface ProviderPromptDoctrine {
3314
3322
  /** MODEL_CATALOG ids this doctrine covers. */
package/dist/index.js CHANGED
@@ -26056,8 +26056,45 @@ precise subject \u2192 action details \u2192 scene/environment \u2192 lighting &
26056
26056
  - More than 4 referenced people gets unstable: group people into composite images of \u22644 first (image generation), then reference those composites.
26057
26057
  - Repeated extension degrades quality: prefer high-definition reference assets and avoid stacking many continuations.`
26058
26058
  };
26059
+ var KLING_AUDIO_DOCTRINE = {
26060
+ providers: ["kling", "kling-3.0", "kling-3-omni"],
26061
+ heading: "Kling 2.6 / 3.0 / 3 Omni (kling, kling-3.0, kling-3-omni)",
26062
+ tips: [
26063
+ 'Kling speaks scripted dialogue natively with lip sync \u2014 quote the line and enable sound: [Anna: warm calm voice]: "We made it." On kling/kling-3.0 audio raises the credit cost; kling-3-omni includes it.',
26064
+ "Structure prompts as Scene \u2192 character/element \u2192 Motion \u2192 Audio \u2192 style. Put ALL sound in one 'Audio:' block: dialogue in quotes, then SFX and ambience described plainly ('rain tapping on glass, no music').",
26065
+ "Give each speaker a stable label + voice description and reuse it exactly \u2014 [Detective: low raspy voice, tired] \u2014 pronouns or renamed speakers break voice binding in multi-character scenes.",
26066
+ "Tone words in the bracket steer delivery (whispering, crying, fast urgent voice); pace with 'Immediately' / 'after a pause'. 2.6 voices are English/Chinese only; 3.0 adds dialects and code-switching.",
26067
+ "kling-3.0 wired references become @element_name mentions (handled automatically by the editor); multi-shot mode forces sound ON. Kling 2.6 prompts cap at 1000 chars \u2014 keep the Audio block tight.",
26068
+ "Say what should NOT sound: 'no background music, no other sounds' \u2014 otherwise Kling invents a music bed under dialogue."
26069
+ ],
26070
+ doctrine: `Prompt structure: Scene (setting, light) \u2192 Character/Element (who, appearance) \u2192 Motion (action, camera) \u2192 Audio (dialogue / SFX / ambience / music) \u2192 Others (style, emotion).
26071
+
26072
+ **Dialogue (native speech + lip sync \u2014 verified on the KIE path 2026-07-16)**
26073
+ - Quote the spoken line and enable the sound toggle; the model bakes the voice AND matching lip movement: the woman says "The quick brown fox jumps over the lazy dog."
26074
+ - Prefer labeled dialogue with a voice description: [Character label: voice/tone description]: "line". Example: [Exhausted Partner: trembling frustrated voice]: "You never listen to me."
26075
+ - Keep character labels unique and reuse them verbatim \u2014 never switch to pronouns mid-prompt; the label is what binds a voice to a speaker across lines. Kling 2.6 additionally supports [Character@VoiceName] platform-voice binding.
26076
+ - Tone words inside the bracket steer delivery: whispering, crying voice, controlled serious voice, fast urgent voice. Sequence speech with temporal markers ("Immediately", "after a pause") when two lines must not overlap.
26077
+ - Languages: Kling 2.6 outputs English/Chinese voices only (other languages are auto-translated to English). Kling 3.0 supports multiple languages, dialects, accents, and code-switching within one scene \u2014 mark the language explicitly ("says in Japanese \u2026").
26078
+
26079
+ **SFX / ambience / music**
26080
+ - Put them in the same Audio block, described plainly: "Rain tapping softly on the window, distant thunder, no music."
26081
+ - State exclusions explicitly \u2014 "no background music, no other sounds" \u2014 or the model tends to add a bed under dialogue.
26082
+
26083
+ **Toggle + cost**
26084
+ - The audio lever is the node's sound toggle (KIE \`sound\` param). On kling (2.6) and kling-3.0 enabling audio raises the credit cost (the \`:audio\` composite); kling-3.0 generates audio by DEFAULT \u2014 pass sound: false for the cheaper silent tier. kling-3-omni (Replicate) includes audio in its flat per-duration rate.
26085
+ - Multi-shot kling-3.0 (\`multi_shots\`) forces sound ON \u2014 budget for the audio rate.
26086
+
26087
+ **References & elements (kling-3.0 / omni)**
26088
+ - Wired references are injected as \`kling_elements\` and MUST be mentioned as @element_name in the prompt \u2014 the editor's {image:N} tokens and the server prefixer handle this automatically; when hand-writing prompts, mention every element or it is silently ignored.
26089
+ - kling-3-omni is image-to-video only (start frame required) and accepts up to 7 reference images; element voice references (element_input_audio_urls, 5-30s clips) bind a voice to an element.
26090
+
26091
+ **Limits**
26092
+ - Kling 2.6 prompts cap at 1000 characters \u2014 front-load scene + dialogue and trim style tails first. kling-3.0 accepts long prompts.
26093
+ - Durations: 2.6 = 5/10s; 3.0/omni = 3-15s. A spoken line needs roughly 1s per 2-3 words \u2014 don't script more dialogue than the clip can hold.`
26094
+ };
26059
26095
  var PROVIDER_PROMPT_DOCTRINES = [
26060
- SEEDANCE_2_DOCTRINE
26096
+ SEEDANCE_2_DOCTRINE,
26097
+ KLING_AUDIO_DOCTRINE
26061
26098
  ];
26062
26099
  var DOCTRINE_BY_PROVIDER = new Map(
26063
26100
  PROVIDER_PROMPT_DOCTRINES.flatMap((d) => d.providers.map((p) => [p, d]))
@@ -26157,6 +26194,7 @@ var PROVIDER_CAPABILITIES = {
26157
26194
  "nano-banana": "Fast generation, style flexibility, reference image support",
26158
26195
  "nano-banana-pro": "Higher quality Nano Banana with better detail",
26159
26196
  "nano-banana-2": "Latest Nano Banana with resolution options (1K/2K/4K)",
26197
+ "nano-banana-2-lite": "Fast, low-cost 1K generation for drafts and iteration",
26160
26198
  "gpt-image": "Creative concepts, illustration, variable quality tiers",
26161
26199
  "gpt-image-2": "Latest GPT Image \u2014 sharp text, photorealism, 1K/2K/4K resolution",
26162
26200
  "grok": "General purpose, good text understanding",
@@ -26167,6 +26205,7 @@ var PROVIDER_CAPABILITIES = {
26167
26205
  "qwen": "Versatile, good prompt adherence",
26168
26206
  "seedream": "Artistic, painterly styles, creative interpretation",
26169
26207
  "seedream-5-lite": "Lighter Seedream, faster artistic generation",
26208
+ "seedream-5-pro": "Flagship Seedream \u2014 strongest instruction following, 2K via high quality",
26170
26209
  "z-image": "Experimental, novel generation approaches",
26171
26210
  "wan-2.7": "Wan 2.7 T2I \u2014 1K/2K/4K, up to 9 ref images",
26172
26211
  "wan-2.7-pro": "Wan 2.7 Pro T2I \u2014 higher quality, 1K/2K/4K",
@@ -26189,6 +26228,7 @@ var PROVIDER_CAPABILITIES = {
26189
26228
  "qwen-edit": "Instruction-based editing",
26190
26229
  "seedream-edit": "Artistic style editing",
26191
26230
  "seedream-5-lite-i2i": "Light artistic transformation",
26231
+ "seedream-5-pro-i2i": "Pro-grade instruction-based edits, multi-reference",
26192
26232
  "flux-kontext": "Character-consistent edits with reference awareness",
26193
26233
  "flux-kontext-max": "Premium character-consistent editing",
26194
26234
  "kontext-multi": "Multi-image Kontext via Replicate \u2014 up to 4 refs, no safety filter",
@@ -26211,6 +26251,7 @@ var PROVIDER_CAPABILITIES = {
26211
26251
  "qwen-edit": "Instruction-based editing",
26212
26252
  "seedream-edit": "Artistic style editing",
26213
26253
  "seedream-5-lite-i2i": "Light artistic transformation",
26254
+ "seedream-5-pro-i2i": "Pro-grade instruction-based edits, multi-reference",
26214
26255
  "flux-kontext": "Character-consistent edits with reference awareness",
26215
26256
  "flux-kontext-max": "Premium character-consistent editing",
26216
26257
  "kontext-multi": "Multi-image Kontext via Replicate \u2014 up to 4 refs, no safety filter",
@@ -26240,7 +26281,7 @@ var PROVIDER_CAPABILITIES = {
26240
26281
  "bytedance-pro": "Higher quality ByteDance",
26241
26282
  "runway-kie": "Runway via KIE, strong cinematic quality",
26242
26283
  "wan-2.7-t2v": "Wan 2.7 T2V \u2014 2\u201315s, 720p/1080p",
26243
- "happyhorse": "HappyHorse T2V \u2014 3\u201315s, 720p/1080p",
26284
+ "happyhorse": "HappyHorse 1.1 T2V \u2014 3\u201315s, 720p/1080p, 9 aspect ratios incl. 21:9/9:21",
26244
26285
  "ltx-2.3-pro": "Lightricks LTX 2.3 Pro \u2014 text/image/audio\u2192video, 6\u201310s, up to 4K",
26245
26286
  "ltx-2.3-fast": "Lightricks LTX 2.3 Fast \u2014 text/image\u2192video, 6\u201320s, up to 4K",
26246
26287
  "gemini-omni-video": "Google Gemini Omni \u2014 multimodal video with native audio, 4\u201310s, up to 4K.",
@@ -26270,8 +26311,8 @@ var PROVIDER_CAPABILITIES = {
26270
26311
  "grok-i2v": "General purpose animation",
26271
26312
  "runway-kie": "Cinematic image animation",
26272
26313
  "wan-2.7-i2v": "Wan 2.7 I2V \u2014 2\u201315s, 720p/1080p, start+end frame",
26273
- "happyhorse-i2v": "HappyHorse I2V \u2014 3\u201315s, 720p/1080p",
26274
- "happyhorse-ref2v": "HappyHorse Ref2V \u2014 multi-ref image to video, 3\u201315s",
26314
+ "happyhorse-i2v": "HappyHorse 1.1 I2V \u2014 3\u201315s, 720p/1080p",
26315
+ "happyhorse-ref2v": "HappyHorse 1.1 Ref2V \u2014 multi-ref image to video (1\u20139 refs), 3\u201315s",
26275
26316
  "ltx-2.3-pro": "Lightricks LTX 2.3 Pro \u2014 start/end frame i2v + audio\u2192video, 6\u201310s, up to 4K",
26276
26317
  "ltx-2.3-fast": "Lightricks LTX 2.3 Fast \u2014 start/end frame i2v, 6\u201320s, up to 4K",
26277
26318
  "gemini-omni-video": "Google Gemini Omni \u2014 multimodal video with native audio, 4\u201310s, up to 4K.",