@nodaro/prompts 1.1.1 → 1.2.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/index.cjs +42 -1
- package/dist/index.cjs.map +1 -1
- package/dist/index.d.cts +8 -0
- package/dist/index.d.ts +8 -0
- package/dist/index.js +42 -1
- package/dist/index.js.map +1 -1
- package/package.json +1 -1
- package/src/__tests__/seedance-extend.test.ts +2 -1
- package/src/prompt-wizard-categories.ts +4 -0
- package/src/provider-prompt-doctrine.ts +46 -0
package/dist/index.cjs
CHANGED
|
@@ -26058,8 +26058,45 @@ precise subject \u2192 action details \u2192 scene/environment \u2192 lighting &
|
|
|
26058
26058
|
- More than 4 referenced people gets unstable: group people into composite images of \u22644 first (image generation), then reference those composites.
|
|
26059
26059
|
- Repeated extension degrades quality: prefer high-definition reference assets and avoid stacking many continuations.`
|
|
26060
26060
|
};
|
|
26061
|
+
var KLING_AUDIO_DOCTRINE = {
|
|
26062
|
+
providers: ["kling", "kling-3.0", "kling-3-omni"],
|
|
26063
|
+
heading: "Kling 2.6 / 3.0 / 3 Omni (kling, kling-3.0, kling-3-omni)",
|
|
26064
|
+
tips: [
|
|
26065
|
+
'Kling speaks scripted dialogue natively with lip sync \u2014 quote the line and enable sound: [Anna: warm calm voice]: "We made it." On kling/kling-3.0 audio raises the credit cost; kling-3-omni includes it.',
|
|
26066
|
+
"Structure prompts as Scene \u2192 character/element \u2192 Motion \u2192 Audio \u2192 style. Put ALL sound in one 'Audio:' block: dialogue in quotes, then SFX and ambience described plainly ('rain tapping on glass, no music').",
|
|
26067
|
+
"Give each speaker a stable label + voice description and reuse it exactly \u2014 [Detective: low raspy voice, tired] \u2014 pronouns or renamed speakers break voice binding in multi-character scenes.",
|
|
26068
|
+
"Tone words in the bracket steer delivery (whispering, crying, fast urgent voice); pace with 'Immediately' / 'after a pause'. 2.6 voices are English/Chinese only; 3.0 adds dialects and code-switching.",
|
|
26069
|
+
"kling-3.0 wired references become @element_name mentions (handled automatically by the editor); multi-shot mode forces sound ON. Kling 2.6 prompts cap at 1000 chars \u2014 keep the Audio block tight.",
|
|
26070
|
+
"Say what should NOT sound: 'no background music, no other sounds' \u2014 otherwise Kling invents a music bed under dialogue."
|
|
26071
|
+
],
|
|
26072
|
+
doctrine: `Prompt structure: Scene (setting, light) \u2192 Character/Element (who, appearance) \u2192 Motion (action, camera) \u2192 Audio (dialogue / SFX / ambience / music) \u2192 Others (style, emotion).
|
|
26073
|
+
|
|
26074
|
+
**Dialogue (native speech + lip sync \u2014 verified on the KIE path 2026-07-16)**
|
|
26075
|
+
- Quote the spoken line and enable the sound toggle; the model bakes the voice AND matching lip movement: the woman says "The quick brown fox jumps over the lazy dog."
|
|
26076
|
+
- Prefer labeled dialogue with a voice description: [Character label: voice/tone description]: "line". Example: [Exhausted Partner: trembling frustrated voice]: "You never listen to me."
|
|
26077
|
+
- Keep character labels unique and reuse them verbatim \u2014 never switch to pronouns mid-prompt; the label is what binds a voice to a speaker across lines. Kling 2.6 additionally supports [Character@VoiceName] platform-voice binding.
|
|
26078
|
+
- Tone words inside the bracket steer delivery: whispering, crying voice, controlled serious voice, fast urgent voice. Sequence speech with temporal markers ("Immediately", "after a pause") when two lines must not overlap.
|
|
26079
|
+
- Languages: Kling 2.6 outputs English/Chinese voices only (other languages are auto-translated to English). Kling 3.0 supports multiple languages, dialects, accents, and code-switching within one scene \u2014 mark the language explicitly ("says in Japanese \u2026").
|
|
26080
|
+
|
|
26081
|
+
**SFX / ambience / music**
|
|
26082
|
+
- Put them in the same Audio block, described plainly: "Rain tapping softly on the window, distant thunder, no music."
|
|
26083
|
+
- State exclusions explicitly \u2014 "no background music, no other sounds" \u2014 or the model tends to add a bed under dialogue.
|
|
26084
|
+
|
|
26085
|
+
**Toggle + cost**
|
|
26086
|
+
- The audio lever is the node's sound toggle (KIE \`sound\` param). On kling (2.6) and kling-3.0 enabling audio raises the credit cost (the \`:audio\` composite); kling-3.0 generates audio by DEFAULT \u2014 pass sound: false for the cheaper silent tier. kling-3-omni (Replicate) includes audio in its flat per-duration rate.
|
|
26087
|
+
- Multi-shot kling-3.0 (\`multi_shots\`) forces sound ON \u2014 budget for the audio rate.
|
|
26088
|
+
|
|
26089
|
+
**References & elements (kling-3.0 / omni)**
|
|
26090
|
+
- Wired references are injected as \`kling_elements\` and MUST be mentioned as @element_name in the prompt \u2014 the editor's {image:N} tokens and the server prefixer handle this automatically; when hand-writing prompts, mention every element or it is silently ignored.
|
|
26091
|
+
- kling-3-omni is image-to-video only (start frame required) and accepts up to 7 reference images; element voice references (element_input_audio_urls, 5-30s clips) bind a voice to an element.
|
|
26092
|
+
|
|
26093
|
+
**Limits**
|
|
26094
|
+
- Kling 2.6 prompts cap at 1000 characters \u2014 front-load scene + dialogue and trim style tails first. kling-3.0 accepts long prompts.
|
|
26095
|
+
- Durations: 2.6 = 5/10s; 3.0/omni = 3-15s. A spoken line needs roughly 1s per 2-3 words \u2014 don't script more dialogue than the clip can hold.`
|
|
26096
|
+
};
|
|
26061
26097
|
var PROVIDER_PROMPT_DOCTRINES = [
|
|
26062
|
-
SEEDANCE_2_DOCTRINE
|
|
26098
|
+
SEEDANCE_2_DOCTRINE,
|
|
26099
|
+
KLING_AUDIO_DOCTRINE
|
|
26063
26100
|
];
|
|
26064
26101
|
var DOCTRINE_BY_PROVIDER = new Map(
|
|
26065
26102
|
PROVIDER_PROMPT_DOCTRINES.flatMap((d) => d.providers.map((p) => [p, d]))
|
|
@@ -26159,6 +26196,7 @@ var PROVIDER_CAPABILITIES = {
|
|
|
26159
26196
|
"nano-banana": "Fast generation, style flexibility, reference image support",
|
|
26160
26197
|
"nano-banana-pro": "Higher quality Nano Banana with better detail",
|
|
26161
26198
|
"nano-banana-2": "Latest Nano Banana with resolution options (1K/2K/4K)",
|
|
26199
|
+
"nano-banana-2-lite": "Fast, low-cost 1K generation for drafts and iteration",
|
|
26162
26200
|
"gpt-image": "Creative concepts, illustration, variable quality tiers",
|
|
26163
26201
|
"gpt-image-2": "Latest GPT Image \u2014 sharp text, photorealism, 1K/2K/4K resolution",
|
|
26164
26202
|
"grok": "General purpose, good text understanding",
|
|
@@ -26169,6 +26207,7 @@ var PROVIDER_CAPABILITIES = {
|
|
|
26169
26207
|
"qwen": "Versatile, good prompt adherence",
|
|
26170
26208
|
"seedream": "Artistic, painterly styles, creative interpretation",
|
|
26171
26209
|
"seedream-5-lite": "Lighter Seedream, faster artistic generation",
|
|
26210
|
+
"seedream-5-pro": "Flagship Seedream \u2014 strongest instruction following, 2K via high quality",
|
|
26172
26211
|
"z-image": "Experimental, novel generation approaches",
|
|
26173
26212
|
"wan-2.7": "Wan 2.7 T2I \u2014 1K/2K/4K, up to 9 ref images",
|
|
26174
26213
|
"wan-2.7-pro": "Wan 2.7 Pro T2I \u2014 higher quality, 1K/2K/4K",
|
|
@@ -26191,6 +26230,7 @@ var PROVIDER_CAPABILITIES = {
|
|
|
26191
26230
|
"qwen-edit": "Instruction-based editing",
|
|
26192
26231
|
"seedream-edit": "Artistic style editing",
|
|
26193
26232
|
"seedream-5-lite-i2i": "Light artistic transformation",
|
|
26233
|
+
"seedream-5-pro-i2i": "Pro-grade instruction-based edits, multi-reference",
|
|
26194
26234
|
"flux-kontext": "Character-consistent edits with reference awareness",
|
|
26195
26235
|
"flux-kontext-max": "Premium character-consistent editing",
|
|
26196
26236
|
"kontext-multi": "Multi-image Kontext via Replicate \u2014 up to 4 refs, no safety filter",
|
|
@@ -26213,6 +26253,7 @@ var PROVIDER_CAPABILITIES = {
|
|
|
26213
26253
|
"qwen-edit": "Instruction-based editing",
|
|
26214
26254
|
"seedream-edit": "Artistic style editing",
|
|
26215
26255
|
"seedream-5-lite-i2i": "Light artistic transformation",
|
|
26256
|
+
"seedream-5-pro-i2i": "Pro-grade instruction-based edits, multi-reference",
|
|
26216
26257
|
"flux-kontext": "Character-consistent edits with reference awareness",
|
|
26217
26258
|
"flux-kontext-max": "Premium character-consistent editing",
|
|
26218
26259
|
"kontext-multi": "Multi-image Kontext via Replicate \u2014 up to 4 refs, no safety filter",
|