@nodaro/prompts 1.1.1 → 1.2.1

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/dist/index.cjs CHANGED
@@ -8392,6 +8392,8 @@ function getParameterPromptHint(node, ctx) {
8392
8392
  return withCustomText(data, buildPostProcessHints(data.postProcess).join(", "));
8393
8393
  case "tone":
8394
8394
  return asStr(data.tone).trim();
8395
+ case "style-guide":
8396
+ return asStr(data.text).trim();
8395
8397
  case "text-prompt":
8396
8398
  return asStr(data.text).trim();
8397
8399
  default:
@@ -26058,8 +26060,45 @@ precise subject \u2192 action details \u2192 scene/environment \u2192 lighting &
26058
26060
  - More than 4 referenced people gets unstable: group people into composite images of \u22644 first (image generation), then reference those composites.
26059
26061
  - Repeated extension degrades quality: prefer high-definition reference assets and avoid stacking many continuations.`
26060
26062
  };
26063
+ var KLING_AUDIO_DOCTRINE = {
26064
+ providers: ["kling", "kling-3.0", "kling-3-omni"],
26065
+ heading: "Kling 2.6 / 3.0 / 3 Omni (kling, kling-3.0, kling-3-omni)",
26066
+ tips: [
26067
+ 'Kling speaks scripted dialogue natively with lip sync \u2014 quote the line and enable sound: [Anna: warm calm voice]: "We made it." On kling/kling-3.0 audio raises the credit cost; kling-3-omni includes it.',
26068
+ "Structure prompts as Scene \u2192 character/element \u2192 Motion \u2192 Audio \u2192 style. Put ALL sound in one 'Audio:' block: dialogue in quotes, then SFX and ambience described plainly ('rain tapping on glass, no music').",
26069
+ "Give each speaker a stable label + voice description and reuse it exactly \u2014 [Detective: low raspy voice, tired] \u2014 pronouns or renamed speakers break voice binding in multi-character scenes.",
26070
+ "Tone words in the bracket steer delivery (whispering, crying, fast urgent voice); pace with 'Immediately' / 'after a pause'. 2.6 voices are English/Chinese only; 3.0 adds dialects and code-switching.",
26071
+ "kling-3.0 wired references become @element_name mentions (handled automatically by the editor); multi-shot mode forces sound ON. Kling 2.6 prompts cap at 1000 chars \u2014 keep the Audio block tight.",
26072
+ "Say what should NOT sound: 'no background music, no other sounds' \u2014 otherwise Kling invents a music bed under dialogue."
26073
+ ],
26074
+ doctrine: `Prompt structure: Scene (setting, light) \u2192 Character/Element (who, appearance) \u2192 Motion (action, camera) \u2192 Audio (dialogue / SFX / ambience / music) \u2192 Others (style, emotion).
26075
+
26076
+ **Dialogue (native speech + lip sync \u2014 verified on the KIE path 2026-07-16)**
26077
+ - Quote the spoken line and enable the sound toggle; the model bakes the voice AND matching lip movement: the woman says "The quick brown fox jumps over the lazy dog."
26078
+ - Prefer labeled dialogue with a voice description: [Character label: voice/tone description]: "line". Example: [Exhausted Partner: trembling frustrated voice]: "You never listen to me."
26079
+ - Keep character labels unique and reuse them verbatim \u2014 never switch to pronouns mid-prompt; the label is what binds a voice to a speaker across lines. Kling 2.6 additionally supports [Character@VoiceName] platform-voice binding.
26080
+ - Tone words inside the bracket steer delivery: whispering, crying voice, controlled serious voice, fast urgent voice. Sequence speech with temporal markers ("Immediately", "after a pause") when two lines must not overlap.
26081
+ - Languages: Kling 2.6 outputs English/Chinese voices only (other languages are auto-translated to English). Kling 3.0 supports multiple languages, dialects, accents, and code-switching within one scene \u2014 mark the language explicitly ("says in Japanese \u2026").
26082
+
26083
+ **SFX / ambience / music**
26084
+ - Put them in the same Audio block, described plainly: "Rain tapping softly on the window, distant thunder, no music."
26085
+ - State exclusions explicitly \u2014 "no background music, no other sounds" \u2014 or the model tends to add a bed under dialogue.
26086
+
26087
+ **Toggle + cost**
26088
+ - The audio lever is the node's sound toggle (KIE \`sound\` param). On kling (2.6) and kling-3.0 enabling audio raises the credit cost (the \`:audio\` composite); kling-3.0 generates audio by DEFAULT \u2014 pass sound: false for the cheaper silent tier. kling-3-omni (Replicate) includes audio in its flat per-duration rate.
26089
+ - Multi-shot kling-3.0 (\`multi_shots\`) forces sound ON \u2014 budget for the audio rate.
26090
+
26091
+ **References & elements (kling-3.0 / omni)**
26092
+ - Wired references are injected as \`kling_elements\` and MUST be mentioned as @element_name in the prompt \u2014 the editor's {image:N} tokens and the server prefixer handle this automatically; when hand-writing prompts, mention every element or it is silently ignored.
26093
+ - kling-3-omni is image-to-video only (start frame required) and accepts up to 7 reference images; element voice references (element_input_audio_urls, 5-30s clips) bind a voice to an element.
26094
+
26095
+ **Limits**
26096
+ - Kling 2.6 prompts cap at 1000 characters \u2014 front-load scene + dialogue and trim style tails first. kling-3.0 accepts long prompts.
26097
+ - Durations: 2.6 = 5/10s; 3.0/omni = 3-15s. A spoken line needs roughly 1s per 2-3 words \u2014 don't script more dialogue than the clip can hold.`
26098
+ };
26061
26099
  var PROVIDER_PROMPT_DOCTRINES = [
26062
- SEEDANCE_2_DOCTRINE
26100
+ SEEDANCE_2_DOCTRINE,
26101
+ KLING_AUDIO_DOCTRINE
26063
26102
  ];
26064
26103
  var DOCTRINE_BY_PROVIDER = new Map(
26065
26104
  PROVIDER_PROMPT_DOCTRINES.flatMap((d) => d.providers.map((p) => [p, d]))
@@ -26159,6 +26198,7 @@ var PROVIDER_CAPABILITIES = {
26159
26198
  "nano-banana": "Fast generation, style flexibility, reference image support",
26160
26199
  "nano-banana-pro": "Higher quality Nano Banana with better detail",
26161
26200
  "nano-banana-2": "Latest Nano Banana with resolution options (1K/2K/4K)",
26201
+ "nano-banana-2-lite": "Fast, low-cost 1K generation for drafts and iteration",
26162
26202
  "gpt-image": "Creative concepts, illustration, variable quality tiers",
26163
26203
  "gpt-image-2": "Latest GPT Image \u2014 sharp text, photorealism, 1K/2K/4K resolution",
26164
26204
  "grok": "General purpose, good text understanding",
@@ -26169,6 +26209,7 @@ var PROVIDER_CAPABILITIES = {
26169
26209
  "qwen": "Versatile, good prompt adherence",
26170
26210
  "seedream": "Artistic, painterly styles, creative interpretation",
26171
26211
  "seedream-5-lite": "Lighter Seedream, faster artistic generation",
26212
+ "seedream-5-pro": "Flagship Seedream \u2014 strongest instruction following, 2K via high quality",
26172
26213
  "z-image": "Experimental, novel generation approaches",
26173
26214
  "wan-2.7": "Wan 2.7 T2I \u2014 1K/2K/4K, up to 9 ref images",
26174
26215
  "wan-2.7-pro": "Wan 2.7 Pro T2I \u2014 higher quality, 1K/2K/4K",
@@ -26191,6 +26232,7 @@ var PROVIDER_CAPABILITIES = {
26191
26232
  "qwen-edit": "Instruction-based editing",
26192
26233
  "seedream-edit": "Artistic style editing",
26193
26234
  "seedream-5-lite-i2i": "Light artistic transformation",
26235
+ "seedream-5-pro-i2i": "Pro-grade instruction-based edits, multi-reference",
26194
26236
  "flux-kontext": "Character-consistent edits with reference awareness",
26195
26237
  "flux-kontext-max": "Premium character-consistent editing",
26196
26238
  "kontext-multi": "Multi-image Kontext via Replicate \u2014 up to 4 refs, no safety filter",
@@ -26213,6 +26255,7 @@ var PROVIDER_CAPABILITIES = {
26213
26255
  "qwen-edit": "Instruction-based editing",
26214
26256
  "seedream-edit": "Artistic style editing",
26215
26257
  "seedream-5-lite-i2i": "Light artistic transformation",
26258
+ "seedream-5-pro-i2i": "Pro-grade instruction-based edits, multi-reference",
26216
26259
  "flux-kontext": "Character-consistent edits with reference awareness",
26217
26260
  "flux-kontext-max": "Premium character-consistent editing",
26218
26261
  "kontext-multi": "Multi-image Kontext via Replicate \u2014 up to 4 refs, no safety filter",