@slatesvideo/shared 0.5.0 → 0.5.2
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/clients/cloud.d.ts +7 -7
- package/dist/operations/index.d.ts +10 -2
- package/dist/operations/index.js +217 -94
- package/dist/prompts/model-facts.js +18 -2
- package/dist/skills/content.js +13 -12
- package/package.json +1 -1
- package/skills/slates-content-policy.md +14 -1
- package/skills/slates-cost-discipline.md +15 -11
- package/skills/slates-direct-response-ad.md +2 -2
- package/skills/slates-edit-and-iterate.md +1 -1
- package/skills/slates-model-selection.md +11 -7
- package/skills/slates-one-prompt-film.md +1 -1
- package/skills/slates-prompting-lip-sync.md +12 -12
- package/skills/slates-prompting-motion-transfer.md +11 -11
- package/skills/slates-prompting-omni-flash.md +44 -0
- package/skills/slates-prompting-seedance.md +1 -1
- package/skills/slates-prompting-veo-3.md +1 -1
- package/skills/slates-storyboard-from-script.md +1 -1
- package/skills/slates-vision-feedback-loop.md +1 -1
|
@@ -58,7 +58,7 @@ export const MODEL_FACTS = [
|
|
|
58
58
|
kind: 'video',
|
|
59
59
|
maxRefImages: null,
|
|
60
60
|
maxIngredients: 9, // ingredient images per video gen
|
|
61
|
-
notes: 'PREMIUM video tier — route here the moment physics, effects, destruction, or scale matter, and for hero shots. VIDEO-ONLY: cannot generate standalone images (use NB2/FLUX.2/Seedream for those). Up to 9 ingredient images. Strong I2V / own-footage restyle. Native 4K. Also the PREMIUM engine inside the Motion Transfer and Lip Sync tools (single-pass: driving video / dialogue are native conditioning signals — better motion fidelity, natural speech, voice cloned from a video source; video references bill input+output seconds).',
|
|
61
|
+
notes: 'PREMIUM video tier — route here the moment physics, effects, destruction, or scale matter, and for hero shots. VIDEO-ONLY: cannot generate standalone images (use NB2/FLUX.2/Seedream for those). Up to 9 ingredient images. Strong I2V / own-footage restyle. Native 4K, but 4K VIDEO is a Pro-only tier gate (base maxes at 1080p; server returns PRO_REQUIRED) — default 1080p unless the user is on Pro. Also the PREMIUM engine inside the Motion Transfer and Lip Sync tools (single-pass: driving video / dialogue are native conditioning signals — better motion fidelity, natural speech, voice cloned from a video source; video references bill input+output seconds).',
|
|
62
62
|
},
|
|
63
63
|
{
|
|
64
64
|
id: 'kling-v3',
|
|
@@ -74,7 +74,7 @@ export const MODEL_FACTS = [
|
|
|
74
74
|
kind: 'video',
|
|
75
75
|
maxRefImages: null,
|
|
76
76
|
maxIngredients: 4, // combined subject elements + style refs per edit
|
|
77
|
-
notes: 'VIDEO-TO-VIDEO EDIT
|
|
77
|
+
notes: 'VIDEO-TO-VIDEO EDIT — the REF-DRIVEN edit tool: takes an EXISTING 3–15s clip and changes what the prompt names, with element/style reference images (@ElementN = frontal + angles) locking subject identity; max 4 combined refs. keep_audio preserves the ORIGINAL audio verbatim (spoken words cannot drift) — but video lips can drift slightly against it, and multi-beat instructions get under-executed (7/09 receipt: missed a second action beat Omni Flash edit landed) — ONE beat per pass. Route here when an edit NEEDS reference images or bit-exact audio; for prompt-only footage-synced VFX, omni-flash-edit won the 7/09 fidelity head-to-head. Billed per second of output (≈ clip length, rounded up). Seedance edit/relocate is the alternative for style-transfer-heavy jobs.',
|
|
78
78
|
},
|
|
79
79
|
{
|
|
80
80
|
id: 'veo-3.1',
|
|
@@ -84,6 +84,22 @@ export const MODEL_FACTS = [
|
|
|
84
84
|
maxIngredients: 3,
|
|
85
85
|
notes: 'NICHE, never the default — pick only when native synchronized audio must generate WITH the video in one gen. 16:9 only, 4/6/8s only. Otherwise Kling (default) or Seedance (physics/premium) win.',
|
|
86
86
|
},
|
|
87
|
+
{
|
|
88
|
+
id: 'omni-flash',
|
|
89
|
+
label: 'Gemini Omni Flash',
|
|
90
|
+
kind: 'video',
|
|
91
|
+
maxRefImages: null,
|
|
92
|
+
maxIngredients: 7, // ref2v image_urls; 7 mirrors Google's own reference limit
|
|
93
|
+
notes: 'CHEAP 720p tier with native synced audio included — t2v, single-start-frame i2v, or reference-to-video with up to 7 reference images. 3-10s, 16:9/9:16 only. No last frame, no video/audio references. VIDEO-ONLY. New seat: quality vs Kling/Seedance unproven pending comparison gens — do not route hero shots here; use it for cheap drafts, audio-in-one-gen at low cost, ref2v character consistency trials, and its edit variant.',
|
|
94
|
+
},
|
|
95
|
+
{
|
|
96
|
+
id: 'omni-flash-edit',
|
|
97
|
+
label: 'Omni Flash Edit',
|
|
98
|
+
kind: 'video',
|
|
99
|
+
maxRefImages: null,
|
|
100
|
+
maxIngredients: 0, // prompt + source clip ONLY — no element/style refs on this endpoint
|
|
101
|
+
notes: 'VIDEO-TO-VIDEO EDIT, prompt-only — THE EDIT-FIDELITY WINNER (7/09 head-to-head vs Kling edit on real talking footage: lips held perfectly, audio near-identical, both action beats landed). Takes an EXISTING 3-10s clip and changes what the prompt names, footage-synced (prop/effect/environment/lighting swaps). Fidelity is EARNED by prompt discipline: ONE short instruction + "Keep everything else the same." — long descriptive prompts DESTROY it (Google-documented + 7/09 receipt). Never name objects as metaphors ("candle-like" → literal candle). Quirk: occasional tail jitter/doubled last speech beat — trim the tail. NO reference images (identity swaps needing refs → Kling edit); bit-exact audio needs → Kling keep_audio or segment-splice. 720p output, cheapest edit seat (~2/3 of Kling edit Std).',
|
|
102
|
+
},
|
|
87
103
|
];
|
|
88
104
|
const FACT_BY_ID = new Map(MODEL_FACTS.map((m) => [m.id, m]));
|
|
89
105
|
export function getModelFact(id) {
|