@mandujs/core 0.47.0 → 0.48.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@mandujs/core",
3
- "version": "0.47.0",
3
+ "version": "0.48.0",
4
4
  "description": "Mandu Framework Core - Spec, Generator, Guard, Runtime",
5
5
  "type": "module",
6
6
  "main": "./src/index.ts",
@@ -52,6 +52,11 @@ export {
52
52
  type InferenceResult,
53
53
  } from "./inference/heuristic";
54
54
 
55
+ export {
56
+ inferDeployIntentWithBrain,
57
+ type BrainInferenceOptions,
58
+ } from "./inference/brain";
59
+
55
60
  export {
56
61
  planDeploy,
57
62
  planHasChanges,
@@ -0,0 +1,268 @@
1
+ /**
2
+ * Brain-based deploy intent inferer.
3
+ *
4
+ * Issue #250 — Phase 1 M4.
5
+ *
6
+ * The brain inferer is **not** a replacement for the heuristic — it
7
+ * runs AFTER the heuristic and decides whether to confirm or refine
8
+ * the first-pass result. Three reasons for the wrap-not-replace shape:
9
+ *
10
+ * 1. **Cost cap.** The heuristic answers ~80% of routes correctly
11
+ * and pays nothing. Brain only weighs in on the rest.
12
+ * 2. **Determinism floor.** When the brain is offline, fails to
13
+ * authenticate, returns malformed JSON, or its output fails Zod
14
+ * validation, the heuristic result stands. The pipeline never
15
+ * blocks on the brain.
16
+ * 3. **Auditability.** Every intent the cache stores carries a
17
+ * `rationale` string. Wrapping lets the brain rationale read
18
+ * "agreed with heuristic: ..." or "overridden because ..." —
19
+ * reviewers can see the reasoning in the diff.
20
+ *
21
+ * Failure modes (all fall back to heuristic):
22
+ *
23
+ * - LLM call throws (network, auth, rate limit) → heuristic.
24
+ * - LLM returns empty / non-JSON content → heuristic.
25
+ * - LLM returns JSON that doesn't satisfy `DeployIntent` Zod → heuristic.
26
+ * - LLM returns valid intent but `runtime: "static"` on a route
27
+ * the manifest knows can't prerender → heuristic (caught by
28
+ * `isStaticIntentValidFor`).
29
+ *
30
+ * The fallback is silent on `failOnError: false` (default — agents
31
+ * should still get a useful plan even with brain hiccups). Tests can
32
+ * pass `failOnError: true` to surface every brain miss.
33
+ */
34
+
35
+ import { z } from "zod";
36
+ import {
37
+ DeployIntent,
38
+ isStaticIntentValidFor,
39
+ type DeployIntent as DeployIntentType,
40
+ } from "../intent";
41
+ import type { LLMAdapter } from "../../brain/adapters/base";
42
+ import type { DeployInferenceContext } from "./context";
43
+ import { inferDeployIntentHeuristic, type InferenceResult } from "./heuristic";
44
+
45
+ // ─── Public surface ───────────────────────────────────────────────────
46
+
47
+ export interface BrainInferenceOptions {
48
+ /**
49
+ * The LLM adapter to call. Get one from
50
+ * `resolveBrainAdapter({ adapter: "auto" })`.
51
+ */
52
+ adapter: LLMAdapter;
53
+ /**
54
+ * Hard-fail on any brain error instead of falling back to the
55
+ * heuristic. Tests use this; production leaves it false.
56
+ */
57
+ failOnError?: boolean;
58
+ /**
59
+ * Timeout in ms for the LLM call. Defaults to 15s — deploy:plan is
60
+ * usually run interactively, agents prefer quick failure over long
61
+ * waits when the network is bad.
62
+ */
63
+ timeoutMs?: number;
64
+ /**
65
+ * Maximum source bytes sent to the brain. Anything beyond this is
66
+ * truncated with a `[…truncated]` marker. Default 4096 — enough for
67
+ * a route handler, well under any cloud token cap.
68
+ */
69
+ maxSourceBytes?: number;
70
+ }
71
+
72
+ /**
73
+ * Build the brain-based inferer. The returned function has the same
74
+ * signature `planDeploy({ infer })` expects, so it plugs in without
75
+ * any changes to the plan engine.
76
+ */
77
+ export function inferDeployIntentWithBrain(
78
+ options: BrainInferenceOptions,
79
+ ): (ctx: DeployInferenceContext) => Promise<InferenceResult> {
80
+ return async (ctx) => {
81
+ const heuristic = inferDeployIntentHeuristic(ctx);
82
+ try {
83
+ const refined = await callBrain(ctx, heuristic, options);
84
+ if (!refined) return heuristic;
85
+
86
+ // Re-validate the brain's intent against the route shape — the
87
+ // brain is good at semantics but may miss the "static requires
88
+ // generateStaticParams" rule.
89
+ const validation = isStaticIntentValidFor(refined.intent, {
90
+ isDynamic: ctx.isDynamic,
91
+ hasGenerateStaticParams: ctx.hasGenerateStaticParams,
92
+ kind: ctx.kind,
93
+ });
94
+ if (!validation.ok) {
95
+ return {
96
+ intent: heuristic.intent,
97
+ rationale: `${heuristic.rationale} (brain proposed ${refined.intent.runtime} but it conflicts with route shape: ${validation.reason})`,
98
+ };
99
+ }
100
+
101
+ return refined;
102
+ } catch (err) {
103
+ if (options.failOnError) throw err;
104
+ return {
105
+ intent: heuristic.intent,
106
+ rationale: `${heuristic.rationale} (brain unavailable: ${err instanceof Error ? err.message : String(err)})`,
107
+ };
108
+ }
109
+ };
110
+ }
111
+
112
+ // ─── Internals ────────────────────────────────────────────────────────
113
+
114
+ /** Schema for the brain's JSON response — slightly looser than DeployIntent
115
+ * so we can validate field-by-field with helpful errors. */
116
+ const BrainInferenceResponse = z
117
+ .object({
118
+ runtime: z.enum(["static", "edge", "node", "bun"]),
119
+ cache: z.union([
120
+ z.literal("no-store"),
121
+ z.literal("public"),
122
+ z.object({
123
+ maxAge: z.number().int().nonnegative().optional(),
124
+ sMaxAge: z.number().int().nonnegative().optional(),
125
+ swr: z.number().int().nonnegative().optional(),
126
+ }),
127
+ ]).optional(),
128
+ regions: z.array(z.string().min(1)).optional(),
129
+ timeout: z.number().int().positive().optional(),
130
+ visibility: z.enum(["public", "private"]).optional(),
131
+ rationale: z.string().min(1),
132
+ /** Brain's confidence — used only for telemetry; not persisted. */
133
+ confidence: z.number().min(0).max(1).optional(),
134
+ })
135
+ .strict();
136
+
137
+ async function callBrain(
138
+ ctx: DeployInferenceContext,
139
+ heuristic: InferenceResult,
140
+ options: BrainInferenceOptions,
141
+ ): Promise<InferenceResult | null> {
142
+ const timeout = options.timeoutMs ?? 15_000;
143
+ // maxBytes reserved for future source-passing; for now the prompt
144
+ // works off context metadata (imports, deps) without raw source.
145
+ void options.maxSourceBytes;
146
+
147
+ const prompt = buildPrompt(ctx, heuristic);
148
+
149
+ // Race the LLM call against a timeout.
150
+ const controller = new AbortController();
151
+ const tHandle = setTimeout(() => controller.abort(), timeout);
152
+ let raw: string;
153
+ try {
154
+ raw = await options.adapter.generate(prompt, {
155
+ // CompletionOptions doesn't formally include `signal` today; we
156
+ // pass it via the loose extras object so adapters that honour it
157
+ // get the cancel signal, while older ones gracefully ignore.
158
+ ...({ signal: controller.signal } as object),
159
+ temperature: 0.0,
160
+ maxTokens: 512,
161
+ });
162
+ } finally {
163
+ clearTimeout(tHandle);
164
+ }
165
+
166
+ const parsed = parseBrainResponse(raw);
167
+ if (!parsed) return null;
168
+
169
+ // Merge against heuristic defaults — brain doesn't have to repeat
170
+ // every field, only the ones it overrides.
171
+ const merged: DeployIntentType = DeployIntent.parse({
172
+ runtime: parsed.runtime,
173
+ cache: parsed.cache ?? heuristic.intent.cache,
174
+ regions: parsed.regions ?? heuristic.intent.regions,
175
+ timeout: parsed.timeout ?? heuristic.intent.timeout,
176
+ visibility: parsed.visibility ?? heuristic.intent.visibility,
177
+ });
178
+
179
+ const rationale =
180
+ parsed.runtime === heuristic.intent.runtime &&
181
+ sameCache(parsed.cache ?? heuristic.intent.cache, heuristic.intent.cache)
182
+ ? `agreed with heuristic: ${parsed.rationale}`
183
+ : `brain refined: ${parsed.rationale}`;
184
+
185
+ return { intent: merged, rationale };
186
+ }
187
+
188
+ function buildPrompt(
189
+ ctx: DeployInferenceContext,
190
+ heuristic: InferenceResult,
191
+ ): string {
192
+ const truncatedImports = ctx.imports.slice(0, 30).join(", ");
193
+ const dependencyClasses = [...ctx.dependencyClasses].join(", ");
194
+
195
+ return [
196
+ "You are an expert deploy engineer for the Mandu framework.",
197
+ "",
198
+ "TASK: Decide the deploy runtime + cache directive for one route.",
199
+ "Output strict JSON ONLY (no markdown fences, no commentary).",
200
+ "",
201
+ "SCHEMA:",
202
+ `{
203
+ "runtime": "static" | "edge" | "node" | "bun",
204
+ "cache": "no-store" | "public" | { "maxAge"?: number, "sMaxAge"?: number, "swr"?: number },
205
+ "regions": string[],
206
+ "timeout": number_in_ms,
207
+ "visibility": "public" | "private",
208
+ "rationale": "1-sentence explanation",
209
+ "confidence": 0..1
210
+ }`,
211
+ "",
212
+ "RULES:",
213
+ "- runtime=static is only valid for prerenderable pages (no dynamic segments OR has generateStaticParams).",
214
+ "- API routes (kind=api) MUST use edge / node / bun, never static.",
215
+ "- Prefer edge for low-latency stateless work; node/bun for DB / native deps / long-latency calls.",
216
+ "- bun is required when the route imports `bun:*` modules.",
217
+ "- Default cache for API: no-store. Default for static pages: { sMaxAge: 31536000, swr: 86400 }.",
218
+ "- Override the heuristic only when there is CLEAR evidence in the source. Otherwise echo the heuristic.",
219
+ "",
220
+ `ROUTE: ${ctx.routeId} (${ctx.pattern})`,
221
+ `KIND: ${ctx.kind}`,
222
+ `DYNAMIC: ${ctx.isDynamic}`,
223
+ `HAS_GENERATE_STATIC_PARAMS: ${ctx.hasGenerateStaticParams}`,
224
+ `IMPORTS: ${truncatedImports || "(none)"}`,
225
+ `DEPENDENCY_CLASSES: ${dependencyClasses || "(none)"}`,
226
+ "",
227
+ "HEURISTIC FIRST-PASS:",
228
+ `- runtime: ${heuristic.intent.runtime}`,
229
+ `- cache: ${formatCacheForPrompt(heuristic.intent.cache)}`,
230
+ `- visibility: ${heuristic.intent.visibility}`,
231
+ `- rationale: ${heuristic.rationale}`,
232
+ "",
233
+ "RESPOND WITH JSON ONLY:",
234
+ ].join("\n");
235
+ }
236
+
237
+ function parseBrainResponse(raw: string): z.infer<typeof BrainInferenceResponse> | null {
238
+ if (!raw || raw.trim() === "") return null;
239
+ // The brain might wrap JSON in fences despite the rules — strip the
240
+ // most common wrappers before parsing.
241
+ const stripped = stripCodeFences(raw).trim();
242
+ let json: unknown;
243
+ try {
244
+ json = JSON.parse(stripped);
245
+ } catch {
246
+ return null;
247
+ }
248
+ const result = BrainInferenceResponse.safeParse(json);
249
+ if (!result.success) return null;
250
+ return result.data;
251
+ }
252
+
253
+ function stripCodeFences(s: string): string {
254
+ // ```json ... ``` or ``` ... ``` — capture the inner block.
255
+ const fenced = /^```(?:json)?\s*\n([\s\S]*?)\n```\s*$/m.exec(s.trim());
256
+ if (fenced) return fenced[1]!;
257
+ return s;
258
+ }
259
+
260
+ function formatCacheForPrompt(cache: DeployIntentType["cache"]): string {
261
+ if (cache === "no-store" || cache === "public") return cache;
262
+ return JSON.stringify(cache);
263
+ }
264
+
265
+ function sameCache(a: DeployIntentType["cache"], b: DeployIntentType["cache"]): boolean {
266
+ if (typeof a === "string" || typeof b === "string") return a === b;
267
+ return JSON.stringify(a) === JSON.stringify(b);
268
+ }