@tangle-network/agent-eval 0.115.2 → 0.116.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/CHANGELOG.md +34 -0
- package/dist/analyst/index.d.ts +8 -10
- package/dist/analyst/index.js +28 -23
- package/dist/analyst/index.js.map +1 -1
- package/dist/{analyst-C8HHvfJp.d.ts → analyst-CFBc14Wc.d.ts} +1 -1
- package/dist/{analyze-runs-BYHg6Irm.d.ts → analyze-runs-0rz_m29H.d.ts} +3 -3
- package/dist/belief-state/index.d.ts +3 -3
- package/dist/benchmarks/index.d.ts +7 -3
- package/dist/benchmarks/index.js +5 -5
- package/dist/campaign/index.d.ts +212 -23
- package/dist/campaign/index.js +20 -5
- package/dist/{chunk-N6MTC3GK.js → chunk-3274WNK7.js} +428 -94
- package/dist/chunk-3274WNK7.js.map +1 -0
- package/dist/{chunk-DRPIZQIT.js → chunk-4D5RVB3W.js} +2 -2
- package/dist/{chunk-LVTGFSHF.js → chunk-7GKEAIAD.js} +2 -2
- package/dist/{chunk-DWLIGZBX.js → chunk-CIUOICJT.js} +748 -3
- package/dist/chunk-CIUOICJT.js.map +1 -0
- package/dist/{chunk-5NVBGKPH.js → chunk-GSW3OBHK.js} +1284 -182
- package/dist/chunk-GSW3OBHK.js.map +1 -0
- package/dist/{chunk-FUCQVFMU.js → chunk-GY4SYVPJ.js} +12 -3
- package/dist/chunk-GY4SYVPJ.js.map +1 -0
- package/dist/chunk-MPHTT5HE.js +74 -0
- package/dist/chunk-MPHTT5HE.js.map +1 -0
- package/dist/{chunk-I2HNIE6N.js → chunk-NBSS5NDZ.js} +4 -4
- package/dist/{chunk-QG5F6463.js → chunk-ONM6PEAE.js} +2 -2
- package/dist/cli.js +2 -2
- package/dist/{code-agent-session-D-g04tcy.d.ts → code-agent-session-CdxteG0y.d.ts} +1 -1
- package/dist/contract/index.d.ts +19 -19
- package/dist/contract/index.js +6 -4
- package/dist/contract/index.js.map +1 -1
- package/dist/{control-CcBiAEnn.d.ts → control-DbcDxouY.d.ts} +1 -1
- package/dist/control.d.ts +2 -2
- package/dist/{default-registry-DltpYR5u.d.ts → default-registry-DDfv22MQ.d.ts} +2 -1
- package/dist/{gepa-dne9JDPL.d.ts → gepa-CQelRtuC.d.ts} +10 -8
- package/dist/hosted/index.d.ts +8 -4
- package/dist/{index-BTEpx9He.d.ts → index-DbCXJfZ1.d.ts} +2 -2
- package/dist/index.d.ts +27 -30
- package/dist/index.js +28 -22
- package/dist/index.js.map +1 -1
- package/dist/{insight-report-IwwvqZZv.d.ts → insight-report-oMVxDTxl.d.ts} +1 -1
- package/dist/{integrity-qemeBAyx.d.ts → integrity-C6PZ73iC.d.ts} +1 -1
- package/dist/kind-factory-DWOvXjR_.d.ts +171 -0
- package/dist/meta-eval/index.d.ts +2 -2
- package/dist/multishot/index.d.ts +6 -2
- package/dist/openapi.json +1 -1
- package/dist/policy-edit-Clb2v6Oa.d.ts +708 -0
- package/dist/{pre-registration-D8h7ZxNL.d.ts → pre-registration--vU0mMtD.d.ts} +4 -4
- package/dist/{provenance-Bibyg1U9.d.ts → provenance-BbVagC68.d.ts} +26 -14
- package/dist/{release-report-CCtzajxP.d.ts → release-report-CamNDe90.d.ts} +2 -2
- package/dist/reporting.d.ts +4 -4
- package/dist/{researcher-Dq-EtpbE.d.ts → researcher-Dwbo_Fxx.d.ts} +5 -5
- package/dist/rl.d.ts +11 -9
- package/dist/rl.js +2 -2
- package/dist/{rubric-predictive-validity-DYTLjGWu.d.ts → rubric-predictive-validity-BIdf9h4R.d.ts} +1 -1
- package/dist/{run-record-B7RTi_ix.d.ts → run-record-CZmcpWPo.d.ts} +1 -1
- package/dist/{runtime-trajectory-Dws7Kpgi.d.ts → runtime-trajectory-CC0jx9ql.d.ts} +1 -1
- package/dist/{semantic-concept-judge-DxJmRkyJ.d.ts → semantic-concept-judge-CKjePUMh.d.ts} +3 -3
- package/dist/{store-C1YxJDEK.d.ts → store-9cAScOcb.d.ts} +132 -1
- package/dist/{summary-report-BJ5aNwZ1.d.ts → summary-report-DTNgQycC.d.ts} +1 -1
- package/dist/traces.d.ts +6 -8
- package/dist/{types-C5gJrOVT.d.ts → types-Ca_63YSD.d.ts} +59 -2
- package/dist/wire/index.js +2 -2
- package/docs/design/loop-taxonomy.md +1 -2
- package/package.json +1 -1
- package/dist/chunk-5NVBGKPH.js.map +0 -1
- package/dist/chunk-AN5UYSVD.js +0 -761
- package/dist/chunk-AN5UYSVD.js.map +0 -1
- package/dist/chunk-DWLIGZBX.js.map +0 -1
- package/dist/chunk-FUCQVFMU.js.map +0 -1
- package/dist/chunk-N6MTC3GK.js.map +0 -1
- package/dist/kind-factory-DcNg13sZ.d.ts +0 -508
- package/dist/llm-client-DyqEH4jH.d.ts +0 -265
- package/dist/policy-edit-RLn8GWof.d.ts +0 -103
- package/dist/raw-provider-sink-C46HDghv.d.ts +0 -132
- /package/dist/{chunk-DRPIZQIT.js.map → chunk-4D5RVB3W.js.map} +0 -0
- /package/dist/{chunk-LVTGFSHF.js.map → chunk-7GKEAIAD.js.map} +0 -0
- /package/dist/{chunk-I2HNIE6N.js.map → chunk-NBSS5NDZ.js.map} +0 -0
- /package/dist/{chunk-QG5F6463.js.map → chunk-ONM6PEAE.js.map} +0 -0
package/dist/campaign/index.d.ts
CHANGED
|
@@ -1,33 +1,32 @@
|
|
|
1
|
-
import {
|
|
2
|
-
export {
|
|
3
|
-
import {
|
|
4
|
-
|
|
5
|
-
|
|
6
|
-
import {
|
|
7
|
-
export {
|
|
1
|
+
import { L as LlmClientOptions, w as PolicyEditAdmissionOptions, v as PolicyEditAdmission, u as PolicyEdit, d as AnalystFinding, F as FindingToPolicyEditOptions } from '../policy-edit-Clb2v6Oa.js';
|
|
2
|
+
export { s as POLICY_EDIT_CANDIDATE_RECORD_SCHEMA, P as PolicyEditCandidateRecord, a4 as validatePolicyEditCandidateRecord } from '../policy-edit-Clb2v6Oa.js';
|
|
3
|
+
import { P as PairedArmsComparison, S as SignedManifest, B as BackendIntegrityReport, C as CompletionRequirement, R as RuntimeEventLike, a as CompletionVerdict, b as ProducedState, c as CorrectnessChecker } from '../pre-registration--vU0mMtD.js';
|
|
4
|
+
export { L as LlmJudgeDimension, d as LlmJudgeOptions, l as llmJudge } from '../pre-registration--vU0mMtD.js';
|
|
5
|
+
import { A as AnalyzeTracesOptions, a as AnalyzeTracesInput, b as AnalyzeTracesResult } from '../analyst-CFBc14Wc.js';
|
|
6
|
+
import { S as Scenario, M as MutableSurface, D as DispatchContext, b as JudgeConfig, g as Gate, e as GenerationRecord, J as JudgeScore, L as LabeledScenarioStore, s as LabeledScenarioWrite, t as LabeledScenarioSampleArgs, u as LabeledScenarioRecord, v as LabelTrust, f as SurfaceProposer, w as ProposedCandidate, x as ProposeContext, m as CodeSurface, y as LabeledScenarioSource, C as CampaignResult } from '../types-Ca_63YSD.js';
|
|
7
|
+
export { i as CampaignAggregates, j as CampaignArtifactWriter, k as CampaignCellResult, l as CampaignCostMeter, z as CampaignTokenUsage, d as CampaignTraceWriter, c as DispatchFn, n as GateContext, h as GateDecision, G as GateResult, o as GenerationCandidate, A as JudgeAggregate, a as JudgeDimension, p as Mutator, O as OptimizationProposer, q as OptimizerConfig, P as ParetoParent, R as RedactionStatus, B as ScenarioAggregate, E as ScoredSurfaceOutcome, r as SessionScript, T as TraceSpan, F as isProposedCandidate, H as labelTrustRank } from '../types-Ca_63YSD.js';
|
|
8
|
+
import { C as CampaignRunPlan, P as PlanCampaignRunOptions, b as RunCampaignOptions, c as RunImprovementLoopOptions } from '../gepa-CQelRtuC.js';
|
|
9
|
+
export { f as CampaignRunPlanCell, h as GepaProposerConstraints, G as GepaProposerOptions, O as OpenAutoPrOptions, i as OpenAutoPrResult, a as RunImprovementLoopResult, R as RunOptimizationOptions, j as RunOptimizationResult, k as countSentenceEdits, l as defaultRenderDiff, m as extractH2Sections, g as gepaProposer, o as openAutoPr, p as planCampaignRun, r as runCampaign, d as runImprovementLoop, n as runOptimization } from '../gepa-CQelRtuC.js';
|
|
8
10
|
import { a as PairedBootstrapResult, E as EProcessState } from '../statistics-oUbOJe-S.js';
|
|
9
11
|
import { C as CampaignStorage } from '../storage-Dw_f7WMt.js';
|
|
10
12
|
export { f as fsCampaignStorage, i as inMemoryCampaignStorage } from '../storage-Dw_f7WMt.js';
|
|
11
|
-
export { A as AxisEvidence, a as AxisVerdict, B as BuildEvidenceVectorOptions, l as BuildLoopProvenanceArgs, D as DefaultProductionGateOptions, m as EmitLoopProvenanceArgs, n as EmitLoopProvenanceResult, E as EvidenceVector, b as EvolutionaryProposerOptions, H as HeldOutGateOptions, o as LoopProvenanceBackend, q as LoopProvenanceCandidate, L as LoopProvenanceRecord, O as ObjectiveSource, c as ParetoSignificanceGateOptions, P as PowerPreflight, s as PowerPreflightOptions, d as PromotionObjective, e as PromotionPolicy, R as RunEvalOptions, f as buildEvidenceVector, t as buildLoopProvenanceRecord, g as composeGate, h as defaultProductionGate, u as emitLoopProvenance, i as evolutionaryProposer, j as heldOutGate, v as loopProvenanceSpans, p as paretoPolicy, k as paretoSignificanceGate, w as powerPreflight, x as provenanceRecordPath, y as provenanceSpansPath, r as runEval } from '../provenance-
|
|
12
|
-
import { L as LlmClientOptions } from '../llm-client-DyqEH4jH.js';
|
|
13
|
+
export { A as AxisEvidence, a as AxisVerdict, B as BuildEvidenceVectorOptions, l as BuildLoopProvenanceArgs, D as DefaultProductionGateOptions, m as EmitLoopProvenanceArgs, n as EmitLoopProvenanceResult, E as EvidenceVector, b as EvolutionaryProposerOptions, H as HeldOutGateOptions, o as LoopProvenanceBackend, q as LoopProvenanceCandidate, L as LoopProvenanceRecord, O as ObjectiveSource, c as ParetoSignificanceGateOptions, P as PowerPreflight, s as PowerPreflightOptions, d as PromotionObjective, e as PromotionPolicy, R as RunEvalOptions, f as buildEvidenceVector, t as buildLoopProvenanceRecord, g as composeGate, h as defaultProductionGate, u as emitLoopProvenance, i as evolutionaryProposer, j as heldOutGate, v as loopProvenanceSpans, p as paretoPolicy, k as paretoSignificanceGate, w as powerPreflight, x as provenanceRecordPath, y as provenanceSpansPath, r as runEval } from '../provenance-BbVagC68.js';
|
|
13
14
|
import { AgentProfile } from '@tangle-network/agent-interface';
|
|
14
15
|
import { A as AgentEvalError, V as ValidationError } from '../errors-oeQrLqXC.js';
|
|
15
|
-
import {
|
|
16
|
-
import {
|
|
17
|
-
import
|
|
16
|
+
import { a as RunSplitTag, R as RunRecord } from '../run-record-CZmcpWPo.js';
|
|
17
|
+
import { T as TraceAnalystKindSpec } from '../kind-factory-DWOvXjR_.js';
|
|
18
|
+
import '../store-9cAScOcb.js';
|
|
19
|
+
import '../types-C7DGg5ex.js';
|
|
18
20
|
import '@tangle-network/tcloud';
|
|
19
|
-
import '../raw-provider-sink-C46HDghv.js';
|
|
20
21
|
import '../verdict-C9MlYujm.js';
|
|
21
22
|
import '@ax-llm/ax';
|
|
22
|
-
import '../store-C1YxJDEK.js';
|
|
23
23
|
import '../dataset-NENEzRgk.js';
|
|
24
24
|
import '../store-BsVi7ncX.js';
|
|
25
25
|
import '../schema-SGWcK9wa.js';
|
|
26
26
|
import '../judge-calibration-7C-IDmKr.js';
|
|
27
|
-
import '../types-C7DGg5ex.js';
|
|
28
27
|
import '../hosted/index.js';
|
|
29
|
-
import '../insight-report-
|
|
30
|
-
import '../summary-report-
|
|
28
|
+
import '../insight-report-oMVxDTxl.js';
|
|
29
|
+
import '../summary-report-DTNgQycC.js';
|
|
31
30
|
import '../failure-cluster-C48PiReX.js';
|
|
32
31
|
import 'zod';
|
|
33
32
|
|
|
@@ -1989,8 +1988,8 @@ declare function parseSkillPatchResponse(raw: string, maxPatches: number, editBu
|
|
|
1989
1988
|
|
|
1990
1989
|
/**
|
|
1991
1990
|
* `runSkillOpt` — the SkillOpt epoch hill-climb (Microsoft, arXiv:2605.23904).
|
|
1992
|
-
* Unlike `runOptimization`'s population
|
|
1993
|
-
* sequential, held-out-gated hill-climb on ONE skill document:
|
|
1991
|
+
* Unlike `runOptimization`'s population search around one global incumbent,
|
|
1992
|
+
* SkillOpt is a sequential, held-out-gated hill-climb on ONE skill document:
|
|
1994
1993
|
*
|
|
1995
1994
|
* each epoch:
|
|
1996
1995
|
* 1. reflect on the CURRENT surface's weakest TRAIN scenarios/dimensions
|
|
@@ -2208,6 +2207,167 @@ interface HaloProposerOptions {
|
|
|
2208
2207
|
/** Wrap the real halo-engine CLI as a SurfaceProposer (prompt-tier). */
|
|
2209
2208
|
declare function haloProposer(opts: HaloProposerOptions): SurfaceProposer;
|
|
2210
2209
|
|
|
2210
|
+
declare const JSON_POLICY_EDIT_TARGET_SURFACES: readonly ["prompt", "tool-contract", "runtime-config", "memory", "agent-profile"];
|
|
2211
|
+
type JsonPolicyEditTargetSurface = (typeof JSON_POLICY_EDIT_TARGET_SURFACES)[number];
|
|
2212
|
+
declare const DEFAULT_POLICY_EDIT_HISTORY_LIMITS: Readonly<{
|
|
2213
|
+
generations: 4;
|
|
2214
|
+
candidatesPerGeneration: 16;
|
|
2215
|
+
scenariosPerCandidate: 12;
|
|
2216
|
+
findings: 32;
|
|
2217
|
+
authorContextChars: 200000;
|
|
2218
|
+
}>;
|
|
2219
|
+
interface PolicyEditObjective {
|
|
2220
|
+
/** Stable objective key cited by forecasts, for example `search.composite`. */
|
|
2221
|
+
key: string;
|
|
2222
|
+
/** Steering objectives are search-only; fresh-task results never enter author context. */
|
|
2223
|
+
split: 'search';
|
|
2224
|
+
/** Search promotes only larger composite scores. */
|
|
2225
|
+
direction: 'increase';
|
|
2226
|
+
scale: {
|
|
2227
|
+
min: number;
|
|
2228
|
+
max: number;
|
|
2229
|
+
};
|
|
2230
|
+
/** Forecast amounts are absolute deltas on the declared score scale. */
|
|
2231
|
+
unit: 'score';
|
|
2232
|
+
}
|
|
2233
|
+
type PolicyEditFindingSource = {
|
|
2234
|
+
kind: 'surface';
|
|
2235
|
+
surfaceHash: string;
|
|
2236
|
+
generation: number;
|
|
2237
|
+
} | {
|
|
2238
|
+
kind: 'global';
|
|
2239
|
+
label: string;
|
|
2240
|
+
};
|
|
2241
|
+
/** Trace-derived findings must name the measured profile that produced them.
|
|
2242
|
+
* Cross-run doctrine can be explicitly global; an unwrapped finding is rejected. */
|
|
2243
|
+
interface PolicyEditFindingInput {
|
|
2244
|
+
finding: AnalystFinding;
|
|
2245
|
+
source: PolicyEditFindingSource;
|
|
2246
|
+
}
|
|
2247
|
+
interface PolicyEditHistoryProjectionOptions {
|
|
2248
|
+
/** Number of most recent generations retained. Default: 4. */
|
|
2249
|
+
maxGenerations?: number;
|
|
2250
|
+
/** Number of candidates retained per generation. Default: 16. */
|
|
2251
|
+
maxCandidatesPerGeneration?: number;
|
|
2252
|
+
/** Scored tasks retained per candidate/outcome after deterministic extreme selection. */
|
|
2253
|
+
maxScenariosPerCandidate?: number;
|
|
2254
|
+
/** Optional pseudonymizer applied before scenario IDs enter author text. */
|
|
2255
|
+
scenarioIdTransform?: (scenarioId: string) => string;
|
|
2256
|
+
/** Objectives used to compute forecast residuals from measured composite deltas. */
|
|
2257
|
+
objectives?: readonly PolicyEditObjective[];
|
|
2258
|
+
}
|
|
2259
|
+
interface PolicyEditCandidateSummary {
|
|
2260
|
+
editId: string;
|
|
2261
|
+
axis: PolicyEdit['axis'];
|
|
2262
|
+
target: PolicyEdit['target'];
|
|
2263
|
+
change: PolicyEdit['change'];
|
|
2264
|
+
claim: string;
|
|
2265
|
+
expectedGain: PolicyEdit['expectedGain'];
|
|
2266
|
+
confidence: number;
|
|
2267
|
+
risk: PolicyEdit['risk'];
|
|
2268
|
+
sourceFindingIds: string[];
|
|
2269
|
+
rationale: string | null;
|
|
2270
|
+
validationPlan: string | null;
|
|
2271
|
+
}
|
|
2272
|
+
interface PolicyEditHistoryCandidateContext {
|
|
2273
|
+
surfaceHash: string;
|
|
2274
|
+
parentSurfaceHash: string | null;
|
|
2275
|
+
parentComposite: number | null;
|
|
2276
|
+
label: string | null;
|
|
2277
|
+
rationale: string | null;
|
|
2278
|
+
composite: number;
|
|
2279
|
+
observedDeltaFromParent: number | null;
|
|
2280
|
+
eligibleForPromotion: boolean | null;
|
|
2281
|
+
coverage: {
|
|
2282
|
+
expectedCells: number;
|
|
2283
|
+
scorableCells: number;
|
|
2284
|
+
unscorableCells: Array<{
|
|
2285
|
+
reason: string;
|
|
2286
|
+
}>;
|
|
2287
|
+
} | null;
|
|
2288
|
+
dimensions: Record<string, number>;
|
|
2289
|
+
scenarios: Array<{
|
|
2290
|
+
scenarioId: string;
|
|
2291
|
+
composite: number;
|
|
2292
|
+
notes: string | null;
|
|
2293
|
+
}>;
|
|
2294
|
+
candidateEdit: PolicyEditCandidateSummary | null;
|
|
2295
|
+
forecastCalibration: {
|
|
2296
|
+
objectiveKey: string;
|
|
2297
|
+
predictedDelta: number;
|
|
2298
|
+
observedDelta: number;
|
|
2299
|
+
residual: number;
|
|
2300
|
+
} | null;
|
|
2301
|
+
}
|
|
2302
|
+
interface PolicyEditOutcomeContext {
|
|
2303
|
+
split: 'search';
|
|
2304
|
+
generation: number;
|
|
2305
|
+
surfaceHash: string;
|
|
2306
|
+
composite: number;
|
|
2307
|
+
dimensions: Record<string, number>;
|
|
2308
|
+
scenarios: Array<{
|
|
2309
|
+
scenarioId: string;
|
|
2310
|
+
composite: number;
|
|
2311
|
+
notes: string | null;
|
|
2312
|
+
}>;
|
|
2313
|
+
coverage: {
|
|
2314
|
+
expectedCells: number;
|
|
2315
|
+
scorableCells: number;
|
|
2316
|
+
};
|
|
2317
|
+
}
|
|
2318
|
+
interface PolicyEditHistoryGenerationContext {
|
|
2319
|
+
generationIndex: number;
|
|
2320
|
+
promoted: string[];
|
|
2321
|
+
candidates: PolicyEditHistoryCandidateContext[];
|
|
2322
|
+
}
|
|
2323
|
+
interface LlmPolicyEditProposerOptions {
|
|
2324
|
+
llm: LlmClientOptions;
|
|
2325
|
+
model: string;
|
|
2326
|
+
/** Plain-language description of the JSON surface being improved. */
|
|
2327
|
+
target: string;
|
|
2328
|
+
/** PolicyEdit target surface every authored edit must retain. */
|
|
2329
|
+
targetSurface: JsonPolicyEditTargetSurface;
|
|
2330
|
+
/** Exact JSON paths the author may change. Prefix or fuzzy matches are not accepted. */
|
|
2331
|
+
allowedJsonPaths: readonly string[];
|
|
2332
|
+
/** Exact search objectives forecasts may name. Unknown keys or mismatched directions fail. */
|
|
2333
|
+
objectives: readonly PolicyEditObjective[];
|
|
2334
|
+
/** Default: evidence-only, so uncertain edits are measured rather than
|
|
2335
|
+
* suppressed by their own model-authored predictions. */
|
|
2336
|
+
admissionMode?: 'evidence-only' | 'strict';
|
|
2337
|
+
/** Readiness thresholds used only when admissionMode is explicitly strict. */
|
|
2338
|
+
admission?: PolicyEditAdmissionOptions;
|
|
2339
|
+
maxCandidates?: number;
|
|
2340
|
+
temperature?: number;
|
|
2341
|
+
maxTokens?: number;
|
|
2342
|
+
timeoutMs?: number;
|
|
2343
|
+
/** Number of most recent scored generations sent to the author. Default: 4. */
|
|
2344
|
+
maxHistoryGenerations?: number;
|
|
2345
|
+
/** Candidates retained per admitted generation. Default: 16. */
|
|
2346
|
+
maxHistoryCandidatesPerGeneration?: number;
|
|
2347
|
+
/** Scored tasks retained per candidate/outcome after deterministic extreme selection. */
|
|
2348
|
+
maxScenariosPerCandidate?: number;
|
|
2349
|
+
/** Evidence-bearing findings retained after deterministic severity/confidence ordering. */
|
|
2350
|
+
maxFindings?: number;
|
|
2351
|
+
/** Hard character limit over system + schema + serialized author context. */
|
|
2352
|
+
maxAuthorContextChars?: number;
|
|
2353
|
+
/** Optional one-to-one pseudonymizer applied to every author-visible evidence field. */
|
|
2354
|
+
scenarioIdTransform?: (scenarioId: string) => string;
|
|
2355
|
+
onAdmission?: (admission: PolicyEditAdmission) => void;
|
|
2356
|
+
}
|
|
2357
|
+
/**
|
|
2358
|
+
* LLM-backed PolicyEdit author. It reads only the current JSON surface,
|
|
2359
|
+
* evidence-bearing analyst findings, and scored generation history. Model output
|
|
2360
|
+
* is validated and rebound to exact finding evidence before the deterministic
|
|
2361
|
+
* policyEditProposer applies, admits, and deduplicates candidates.
|
|
2362
|
+
*/
|
|
2363
|
+
declare function llmPolicyEditProposer(opts: LlmPolicyEditProposerOptions): SurfaceProposer<PolicyEditFindingInput>;
|
|
2364
|
+
/**
|
|
2365
|
+
* Projects scored history into the only fields a policy author may consume.
|
|
2366
|
+
* It admits recent generations and a bounded candidate count, while retaining
|
|
2367
|
+
* every dimension and scenario score for each admitted candidate.
|
|
2368
|
+
*/
|
|
2369
|
+
declare function projectPolicyEditHistory(history: readonly GenerationRecord[], options?: PolicyEditHistoryProjectionOptions): PolicyEditHistoryGenerationContext[];
|
|
2370
|
+
|
|
2211
2371
|
/**
|
|
2212
2372
|
* `memoryCurationProposer` — a CURATOR `SurfaceProposer`, the complement to the
|
|
2213
2373
|
* OPTIMIZER proposers (`gepaProposer` rewrites the prompt; this one BUILDS a
|
|
@@ -2281,6 +2441,33 @@ interface PolicyEditProposerOptions {
|
|
|
2281
2441
|
*/
|
|
2282
2442
|
declare function policyEditProposer(opts?: PolicyEditProposerOptions): SurfaceProposer;
|
|
2283
2443
|
|
|
2444
|
+
/** One measured scenario row eligible for PolicyEdit author context. */
|
|
2445
|
+
interface PolicyEditAuthorScenarioRow {
|
|
2446
|
+
scenarioId: string;
|
|
2447
|
+
composite: number;
|
|
2448
|
+
}
|
|
2449
|
+
interface SelectPolicyEditAuthorRowsOptions {
|
|
2450
|
+
/** Maximum returned rows. Must be a positive safe integer. */
|
|
2451
|
+
limit: number;
|
|
2452
|
+
/** Optional score to compare against, keyed by scenario ID. */
|
|
2453
|
+
referenceByScenario?: ReadonlyMap<string, number>;
|
|
2454
|
+
}
|
|
2455
|
+
interface SerializedJsonBudget {
|
|
2456
|
+
json: string;
|
|
2457
|
+
actualChars: number;
|
|
2458
|
+
maxChars: number;
|
|
2459
|
+
}
|
|
2460
|
+
/**
|
|
2461
|
+
* Select a bounded, deterministic evidence slice for a PolicyEdit author.
|
|
2462
|
+
*
|
|
2463
|
+
* Rows are deduplicated by scenario ID, keeping the first measured row. The
|
|
2464
|
+
* result then interleaves three ranked views: hardest score, largest regression,
|
|
2465
|
+
* and largest improvement. A row selected by multiple views appears once.
|
|
2466
|
+
*/
|
|
2467
|
+
declare function selectPolicyEditAuthorRows<T extends PolicyEditAuthorScenarioRow>(rows: readonly T[], options: SelectPolicyEditAuthorRowsOptions): T[];
|
|
2468
|
+
/** Serialize once and fail before dispatch when author context exceeds its budget. */
|
|
2469
|
+
declare function assertPolicyEditAuthorContextBudget(value: unknown, maxChars: number): SerializedJsonBudget;
|
|
2470
|
+
|
|
2284
2471
|
/**
|
|
2285
2472
|
* `traceAnalystProposer` — wraps agent-eval's OWN trace-analyst engine
|
|
2286
2473
|
* (`AnalystRegistry` over the agentic OTLP reader) as a `SurfaceProposer`.
|
|
@@ -2399,9 +2586,11 @@ declare function selectDiscriminative(signals: ScenarioSignal[], k: number, opts
|
|
|
2399
2586
|
* the optimizers cannot drift on how a surface's score is computed.
|
|
2400
2587
|
*/
|
|
2401
2588
|
|
|
2402
|
-
/** Mean composite across a campaign: per cell, the mean of its
|
|
2403
|
-
* composites; then the mean across cells.
|
|
2404
|
-
*
|
|
2589
|
+
/** Mean composite across a campaign: per cell, the mean of its finite,
|
|
2590
|
+
* successful judge composites; then the mean across cells. Invalid scores
|
|
2591
|
+
* remain visible on raw cells and coverage receipts but never poison the
|
|
2592
|
+
* descriptive aggregate with NaN. Cells with no valid scores are skipped.
|
|
2593
|
+
* Empty ⇒ 0. */
|
|
2405
2594
|
declare function campaignMeanComposite<TArtifact, TScenario extends Scenario>(campaign: CampaignResult<TArtifact, TScenario>): number;
|
|
2406
2595
|
interface CampaignBreakdown {
|
|
2407
2596
|
/** Mean score per judge dimension across all cells. */
|
|
@@ -2899,4 +3088,4 @@ declare function verifyCodeSurface(surface: CodeSurface, worktreeDir?: string):
|
|
|
2899
3088
|
* identity against the checkout at `worktreeRef`. */
|
|
2900
3089
|
declare function resolveWorktreePath(surface: CodeSurface, worktreeDir?: string): string;
|
|
2901
3090
|
|
|
2902
|
-
export { type AcceptedEdit, type AceProposerOptions, type AnalystArtifact, type AnalystScenario, type AnalyzeCrossSurfaceInteractionsInput, type ApplySkillPatchResult, type BuildAnalystSurfaceDispatchOptions, type CampaignBreakdown, CampaignResult, CampaignRunPlan, CampaignStorage, CodeSurface, type CodeSurfaceVerification, type CompareProposersOptions, type CompositeProposerOptions, type CrossSurfaceAdditionDecision, type CrossSurfaceAdditionRejectionReason, type CrossSurfaceAttemptCompleteness, type CrossSurfaceBestSingleSelection, type CrossSurfaceBootstrapPolicy, type CrossSurfaceCandidate, type CrossSurfaceCandidateComparison, type CrossSurfaceCandidateEvidence, type CrossSurfaceCandidateOutcome, type CrossSurfaceCandidateSummary, type CrossSurfaceComponent, type CrossSurfaceComponentEvidence, type CrossSurfaceCompositionStep, type CrossSurfaceDistribution, type CrossSurfaceEligibility, type CrossSurfaceEvidenceBreakdown, type CrossSurfaceIneligibilityReason, type CrossSurfaceInteractionAwareSelection, type CrossSurfaceInteractionEffect, type CrossSurfaceInteractionPath, type CrossSurfaceInteractionReport, type CrossSurfaceInteractionTask, type CrossSurfaceNaiveStackSelection, type CrossSurfacePairCompatibility, type CrossSurfacePairEvidence, type CrossSurfacePairIncompatibilityReason, type CrossSurfacePairwiseEntry, type CrossSurfaceRankedSingle, type CrossSurfaceRelativeCost, type CrossSurfaceSelectionPolicy, type CrossSurfaceSelections, type CrossSurfaceTaskRow, type DimensionRegression, type DiscriminationScore, DispatchContext, type EvalFixture, type EvalFixtureFile, type EvalFixtureLoadOptions, type EvalFixtureRunPlan, type EvalFixtureScenario, type EvalFixtureValidationMode, type FailureModeRecallJudgeOptions, type FapoAttributionSignals, type FapoEntryConfig, type FapoFailureCluster, type FapoOptimizationLevel, type FapoProposerOptions, type FapoReviewInput, type FapoReviewIssue, type FapoReviewResult, type FapoScopeContract, FileSearchLedger, FsLabeledScenarioStore, type FsLabeledScenarioStoreOptions, Gate, GenerationRecord, type GitWorktreeAdapterOptions, type Governor, type GovernorContext, type GovernorOp, type HaloProposerOptions, type HeldoutSignificance, type HeldoutSignificanceOptions, type HeuristicGovernorOptions, type JsonPrimitive, type JsonValue, JudgeConfig, JudgeScore, LabelTrust, LabeledScenarioRecord, LabeledScenarioSampleArgs, LabeledScenarioSource, LabeledScenarioStore, LabeledScenarioStoreError, LabeledScenarioWrite, Lineage, type LineageEdge, type LineageGraph, type LineageNode, type LineageNodeInput, type LineageStore, type LoadEvalFixtureScenariosOptions, type MemoryCurationProposerOptions, MutableSurface, type NeutralizationGateOptions, type OpenSearchLedgerOptions, type OptimizerEntryConfig, type PairedHoldout, type ParameterCandidate, type ParameterChange, type ParameterSweepProposerOptions, PlanCampaignRunOptions, type PlanEvalFixtureRunOptions, type PlaybackContext, type PlaybackDriver, type PlaybackStep, type PolicyEditProposerOptions, type ProfileDispatchFn, ProfileMatrixError, type ProfileSummary, ProposeContext, type ProposePatchesArgs, ProposedCandidate, type ProposerComparison, type ProposerEntry, type ProposerPairwise, type ProposerScore, type RejectedEdit, type RolloutArgumentDiff, type RolloutArgumentDiffOptions, type RolloutCall, RunCampaignOptions, RunImprovementLoopOptions, type RunLineageLoopOptions, type RunLineageLoopResult, type RunLineageLoopSeed, type RunLineageOptions, type RunLineageResult, type RunLineageSeed, type RunLineageStepResult, type RunProfileMatrixOptions, type RunProfileMatrixResult, type RunSkillOptOptions, type RunSkillOptResult, SEARCH_LEDGER_SCHEMA, Scenario, type ScenarioRollup, type ScenarioSignal, type ScoreboardRenderOptions, type ScoreboardRow, type ScoreboardSummary, type ScoredRollout, type SearchAccountingAudit, type SearchArtifactRef, type SearchAttemptAccounting, type SearchCandidateDecidedEvent, type SearchCandidateLineage, type SearchCandidateRegisteredEvent, type SearchCandidateSlot, type SearchCandidateSlotClosedEvent, type SearchCandidateSurface, type SearchCompletedEvent, type SearchCostAccounting, type SearchFailureReason, type SearchLedger, type SearchLedgerAppendResult, SearchLedgerConflictError, type SearchLedgerEntry, SearchLedgerError, type SearchLedgerEvent, type SearchLedgerHash, SearchLedgerIntegrityError, type SearchLedgerReplay, type SearchModelIdentity, type SearchOperationKind, type SearchOperationRecordedEvent, type SearchPlan, type SearchPlannedEvent, type SearchPlannedOperation, type SearchPlannedTask, type SearchSourceRef, type SearchSurfaceEffect, type SearchSurfaceEvidence, type SearchSurfaceKind, type SearchTaskAttemptedEvent, type SearchTaskOutcome, type SearchTokenAccounting, type SequentialDecideFn, type SequentialDecideOptions, type SequentialDecision, type SequentialObservation, type SequentialPairedGate, type SequentialPairedGateOptions, type SingleRunLock, type SingleRunLockOptions, type SkillOptEpochRecord, type SkillOptEvidence, type SkillOptProposer, type SkillOptProposerOptions, type SkillPatch, type SkillPatchOp, SkillPatchParseError, type SkillPatchRejection, SurfaceProposer, type SurfaceScore, type TraceAnalystProposerOptions, type TransientFailureOptions, type UngroundedLiteralReport, type UserStory, type UserStoryVerdict, type Worktree, type WorktreeAdapter, WorktreeAdapterError, aceProposer, acquireSingleRunLock, analyzeCrossSurfaceInteractions, applySkillPatch, assertCodeSurfaceIdentity, buildAnalystSurfaceDispatch, callbackGovernor, campaignBreakdown, campaignMeanComposite, classifyUngroundedLiterals, codeSurfaceIdentityMaterial, compareProposers, compositeProposer, detectScale, dimensionRegressions, discoverEvalFixtures, extractFapoAttributionSignals, failureModeRecallJudge, fapoEscalationEntry, fapoProposer, fsLineageStore, gepaParetoEntry, gepaReflectionEntry, gitWorktreeAdapter, haloProposer, heldoutSignificance, heuristicGovernor, isTransientTransportFailure, lineageNodeId, loadEvalFixture, loadEvalFixtureScenarios, makePlaybackDispatch, memLineageStore, memoryCurationProposer, neutralizationGate, neutralizeText, openSearchLedger, pairHoldout, parameterSweepProposer, parseSkillPatchResponse, patchEditCount, planEvalFixtureRun, policyEditProposer, renderScoreboardMarkdown, resolveRunDir, resolveWorktreePath, rolloutArgumentDiff, runLineage, runLineageLoop, runProfileMatrix, runSkillOpt, scoreDiscrimination, scoreUserStory, scoreboardSummary, selectDiscriminative, sequentialDecide, sequentialPairedGate, skillOptEntry, skillOptProposer, surfaceContentHash, surfaceHash, tangleTracesRoot, traceAnalystProposer, userStoryScoreboard, validateSearchLedgerEvent, verifyCodeSurface };
|
|
3091
|
+
export { type AcceptedEdit, type AceProposerOptions, type AnalystArtifact, type AnalystScenario, type AnalyzeCrossSurfaceInteractionsInput, type ApplySkillPatchResult, type BuildAnalystSurfaceDispatchOptions, type CampaignBreakdown, CampaignResult, CampaignRunPlan, CampaignStorage, CodeSurface, type CodeSurfaceVerification, type CompareProposersOptions, type CompositeProposerOptions, type CrossSurfaceAdditionDecision, type CrossSurfaceAdditionRejectionReason, type CrossSurfaceAttemptCompleteness, type CrossSurfaceBestSingleSelection, type CrossSurfaceBootstrapPolicy, type CrossSurfaceCandidate, type CrossSurfaceCandidateComparison, type CrossSurfaceCandidateEvidence, type CrossSurfaceCandidateOutcome, type CrossSurfaceCandidateSummary, type CrossSurfaceComponent, type CrossSurfaceComponentEvidence, type CrossSurfaceCompositionStep, type CrossSurfaceDistribution, type CrossSurfaceEligibility, type CrossSurfaceEvidenceBreakdown, type CrossSurfaceIneligibilityReason, type CrossSurfaceInteractionAwareSelection, type CrossSurfaceInteractionEffect, type CrossSurfaceInteractionPath, type CrossSurfaceInteractionReport, type CrossSurfaceInteractionTask, type CrossSurfaceNaiveStackSelection, type CrossSurfacePairCompatibility, type CrossSurfacePairEvidence, type CrossSurfacePairIncompatibilityReason, type CrossSurfacePairwiseEntry, type CrossSurfaceRankedSingle, type CrossSurfaceRelativeCost, type CrossSurfaceSelectionPolicy, type CrossSurfaceSelections, type CrossSurfaceTaskRow, DEFAULT_POLICY_EDIT_HISTORY_LIMITS, type DimensionRegression, type DiscriminationScore, DispatchContext, type EvalFixture, type EvalFixtureFile, type EvalFixtureLoadOptions, type EvalFixtureRunPlan, type EvalFixtureScenario, type EvalFixtureValidationMode, type FailureModeRecallJudgeOptions, type FapoAttributionSignals, type FapoEntryConfig, type FapoFailureCluster, type FapoOptimizationLevel, type FapoProposerOptions, type FapoReviewInput, type FapoReviewIssue, type FapoReviewResult, type FapoScopeContract, FileSearchLedger, FsLabeledScenarioStore, type FsLabeledScenarioStoreOptions, Gate, GenerationRecord, type GitWorktreeAdapterOptions, type Governor, type GovernorContext, type GovernorOp, type HaloProposerOptions, type HeldoutSignificance, type HeldoutSignificanceOptions, type HeuristicGovernorOptions, type JsonPolicyEditTargetSurface, type JsonPrimitive, type JsonValue, JudgeConfig, JudgeScore, LabelTrust, LabeledScenarioRecord, LabeledScenarioSampleArgs, LabeledScenarioSource, LabeledScenarioStore, LabeledScenarioStoreError, LabeledScenarioWrite, Lineage, type LineageEdge, type LineageGraph, type LineageNode, type LineageNodeInput, type LineageStore, type LlmPolicyEditProposerOptions, type LoadEvalFixtureScenariosOptions, type MemoryCurationProposerOptions, MutableSurface, type NeutralizationGateOptions, type OpenSearchLedgerOptions, type OptimizerEntryConfig, type PairedHoldout, type ParameterCandidate, type ParameterChange, type ParameterSweepProposerOptions, PlanCampaignRunOptions, type PlanEvalFixtureRunOptions, type PlaybackContext, type PlaybackDriver, type PlaybackStep, type PolicyEditAuthorScenarioRow, type PolicyEditCandidateSummary, type PolicyEditFindingInput, type PolicyEditFindingSource, type PolicyEditHistoryCandidateContext, type PolicyEditHistoryGenerationContext, type PolicyEditHistoryProjectionOptions, type PolicyEditObjective, type PolicyEditOutcomeContext, type PolicyEditProposerOptions, type ProfileDispatchFn, ProfileMatrixError, type ProfileSummary, ProposeContext, type ProposePatchesArgs, ProposedCandidate, type ProposerComparison, type ProposerEntry, type ProposerPairwise, type ProposerScore, type RejectedEdit, type RolloutArgumentDiff, type RolloutArgumentDiffOptions, type RolloutCall, RunCampaignOptions, RunImprovementLoopOptions, type RunLineageLoopOptions, type RunLineageLoopResult, type RunLineageLoopSeed, type RunLineageOptions, type RunLineageResult, type RunLineageSeed, type RunLineageStepResult, type RunProfileMatrixOptions, type RunProfileMatrixResult, type RunSkillOptOptions, type RunSkillOptResult, SEARCH_LEDGER_SCHEMA, Scenario, type ScenarioRollup, type ScenarioSignal, type ScoreboardRenderOptions, type ScoreboardRow, type ScoreboardSummary, type ScoredRollout, type SearchAccountingAudit, type SearchArtifactRef, type SearchAttemptAccounting, type SearchCandidateDecidedEvent, type SearchCandidateLineage, type SearchCandidateRegisteredEvent, type SearchCandidateSlot, type SearchCandidateSlotClosedEvent, type SearchCandidateSurface, type SearchCompletedEvent, type SearchCostAccounting, type SearchFailureReason, type SearchLedger, type SearchLedgerAppendResult, SearchLedgerConflictError, type SearchLedgerEntry, SearchLedgerError, type SearchLedgerEvent, type SearchLedgerHash, SearchLedgerIntegrityError, type SearchLedgerReplay, type SearchModelIdentity, type SearchOperationKind, type SearchOperationRecordedEvent, type SearchPlan, type SearchPlannedEvent, type SearchPlannedOperation, type SearchPlannedTask, type SearchSourceRef, type SearchSurfaceEffect, type SearchSurfaceEvidence, type SearchSurfaceKind, type SearchTaskAttemptedEvent, type SearchTaskOutcome, type SearchTokenAccounting, type SelectPolicyEditAuthorRowsOptions, type SequentialDecideFn, type SequentialDecideOptions, type SequentialDecision, type SequentialObservation, type SequentialPairedGate, type SequentialPairedGateOptions, type SerializedJsonBudget, type SingleRunLock, type SingleRunLockOptions, type SkillOptEpochRecord, type SkillOptEvidence, type SkillOptProposer, type SkillOptProposerOptions, type SkillPatch, type SkillPatchOp, SkillPatchParseError, type SkillPatchRejection, SurfaceProposer, type SurfaceScore, type TraceAnalystProposerOptions, type TransientFailureOptions, type UngroundedLiteralReport, type UserStory, type UserStoryVerdict, type Worktree, type WorktreeAdapter, WorktreeAdapterError, aceProposer, acquireSingleRunLock, analyzeCrossSurfaceInteractions, applySkillPatch, assertCodeSurfaceIdentity, assertPolicyEditAuthorContextBudget, buildAnalystSurfaceDispatch, callbackGovernor, campaignBreakdown, campaignMeanComposite, classifyUngroundedLiterals, codeSurfaceIdentityMaterial, compareProposers, compositeProposer, detectScale, dimensionRegressions, discoverEvalFixtures, extractFapoAttributionSignals, failureModeRecallJudge, fapoEscalationEntry, fapoProposer, fsLineageStore, gepaParetoEntry, gepaReflectionEntry, gitWorktreeAdapter, haloProposer, heldoutSignificance, heuristicGovernor, isTransientTransportFailure, lineageNodeId, llmPolicyEditProposer, loadEvalFixture, loadEvalFixtureScenarios, makePlaybackDispatch, memLineageStore, memoryCurationProposer, neutralizationGate, neutralizeText, openSearchLedger, pairHoldout, parameterSweepProposer, parseSkillPatchResponse, patchEditCount, planEvalFixtureRun, policyEditProposer, projectPolicyEditHistory, renderScoreboardMarkdown, resolveRunDir, resolveWorktreePath, rolloutArgumentDiff, runLineage, runLineageLoop, runProfileMatrix, runSkillOpt, scoreDiscrimination, scoreUserStory, scoreboardSummary, selectDiscriminative, selectPolicyEditAuthorRows, sequentialDecide, sequentialPairedGate, skillOptEntry, skillOptProposer, surfaceContentHash, surfaceHash, tangleTracesRoot, traceAnalystProposer, userStoryScoreboard, validateSearchLedgerEvent, verifyCodeSurface };
|
package/dist/campaign/index.js
CHANGED
|
@@ -1,4 +1,5 @@
|
|
|
1
1
|
import {
|
|
2
|
+
DEFAULT_POLICY_EDIT_HISTORY_LIMITS,
|
|
2
3
|
FileSearchLedger,
|
|
3
4
|
FsLabeledScenarioStore,
|
|
4
5
|
LabeledScenarioStoreError,
|
|
@@ -14,6 +15,7 @@ import {
|
|
|
14
15
|
acquireSingleRunLock,
|
|
15
16
|
analyzeCrossSurfaceInteractions,
|
|
16
17
|
applySkillPatch,
|
|
18
|
+
assertPolicyEditAuthorContextBudget,
|
|
17
19
|
buildAnalystSurfaceDispatch,
|
|
18
20
|
callbackGovernor,
|
|
19
21
|
classifyUngroundedLiterals,
|
|
@@ -33,6 +35,7 @@ import {
|
|
|
33
35
|
isTransientTransportFailure,
|
|
34
36
|
lineageNodeId,
|
|
35
37
|
llmJudge,
|
|
38
|
+
llmPolicyEditProposer,
|
|
36
39
|
loadEvalFixture,
|
|
37
40
|
loadEvalFixtureScenarios,
|
|
38
41
|
makePlaybackDispatch,
|
|
@@ -46,6 +49,7 @@ import {
|
|
|
46
49
|
patchEditCount,
|
|
47
50
|
planEvalFixtureRun,
|
|
48
51
|
policyEditProposer,
|
|
52
|
+
projectPolicyEditHistory,
|
|
49
53
|
renderScoreboardMarkdown,
|
|
50
54
|
resolveWorktreePath,
|
|
51
55
|
rolloutArgumentDiff,
|
|
@@ -57,6 +61,7 @@ import {
|
|
|
57
61
|
scoreUserStory,
|
|
58
62
|
scoreboardSummary,
|
|
59
63
|
selectDiscriminative,
|
|
64
|
+
selectPolicyEditAuthorRows,
|
|
60
65
|
sequentialDecide,
|
|
61
66
|
sequentialPairedGate,
|
|
62
67
|
skillOptEntry,
|
|
@@ -65,7 +70,7 @@ import {
|
|
|
65
70
|
userStoryScoreboard,
|
|
66
71
|
validateSearchLedgerEvent,
|
|
67
72
|
verifyCodeSurface
|
|
68
|
-
} from "../chunk-
|
|
73
|
+
} from "../chunk-GSW3OBHK.js";
|
|
69
74
|
import {
|
|
70
75
|
assertCodeSurfaceIdentity,
|
|
71
76
|
buildEvidenceVector,
|
|
@@ -100,7 +105,7 @@ import {
|
|
|
100
105
|
runOptimization,
|
|
101
106
|
surfaceContentHash,
|
|
102
107
|
surfaceHash
|
|
103
|
-
} from "../chunk-
|
|
108
|
+
} from "../chunk-3274WNK7.js";
|
|
104
109
|
import "../chunk-VI2UW6B6.js";
|
|
105
110
|
import {
|
|
106
111
|
fsCampaignStorage,
|
|
@@ -110,8 +115,11 @@ import {
|
|
|
110
115
|
runCampaign,
|
|
111
116
|
tangleTracesRoot
|
|
112
117
|
} from "../chunk-FAOEFFRT.js";
|
|
113
|
-
import "../chunk-
|
|
114
|
-
import
|
|
118
|
+
import "../chunk-MPHTT5HE.js";
|
|
119
|
+
import {
|
|
120
|
+
POLICY_EDIT_CANDIDATE_RECORD_SCHEMA,
|
|
121
|
+
validatePolicyEditCandidateRecord
|
|
122
|
+
} from "../chunk-CIUOICJT.js";
|
|
115
123
|
import "../chunk-ARU2PZFM.js";
|
|
116
124
|
import "../chunk-PJQFMIOX.js";
|
|
117
125
|
import "../chunk-RPDDVKI7.js";
|
|
@@ -120,15 +128,17 @@ import "../chunk-LNQEP766.js";
|
|
|
120
128
|
import "../chunk-5UF54T55.js";
|
|
121
129
|
import "../chunk-XJYR7XFV.js";
|
|
122
130
|
import "../chunk-VSMTAMNK.js";
|
|
123
|
-
import "../chunk-
|
|
131
|
+
import "../chunk-GY4SYVPJ.js";
|
|
124
132
|
import "../chunk-PC4UYEBM.js";
|
|
125
133
|
import "../chunk-ONWEPEDO.js";
|
|
126
134
|
import "../chunk-PZ5AY32C.js";
|
|
127
135
|
export {
|
|
136
|
+
DEFAULT_POLICY_EDIT_HISTORY_LIMITS,
|
|
128
137
|
FileSearchLedger,
|
|
129
138
|
FsLabeledScenarioStore,
|
|
130
139
|
LabeledScenarioStoreError,
|
|
131
140
|
Lineage,
|
|
141
|
+
POLICY_EDIT_CANDIDATE_RECORD_SCHEMA,
|
|
132
142
|
ProfileMatrixError,
|
|
133
143
|
SEARCH_LEDGER_SCHEMA,
|
|
134
144
|
SearchLedgerConflictError,
|
|
@@ -141,6 +151,7 @@ export {
|
|
|
141
151
|
analyzeCrossSurfaceInteractions,
|
|
142
152
|
applySkillPatch,
|
|
143
153
|
assertCodeSurfaceIdentity,
|
|
154
|
+
assertPolicyEditAuthorContextBudget,
|
|
144
155
|
buildAnalystSurfaceDispatch,
|
|
145
156
|
buildEvidenceVector,
|
|
146
157
|
buildLoopProvenanceRecord,
|
|
@@ -181,6 +192,7 @@ export {
|
|
|
181
192
|
labelTrustRank,
|
|
182
193
|
lineageNodeId,
|
|
183
194
|
llmJudge,
|
|
195
|
+
llmPolicyEditProposer,
|
|
184
196
|
loadEvalFixture,
|
|
185
197
|
loadEvalFixtureScenarios,
|
|
186
198
|
loopProvenanceSpans,
|
|
@@ -201,6 +213,7 @@ export {
|
|
|
201
213
|
planEvalFixtureRun,
|
|
202
214
|
policyEditProposer,
|
|
203
215
|
powerPreflight,
|
|
216
|
+
projectPolicyEditHistory,
|
|
204
217
|
provenanceRecordPath,
|
|
205
218
|
provenanceSpansPath,
|
|
206
219
|
renderScoreboardMarkdown,
|
|
@@ -219,6 +232,7 @@ export {
|
|
|
219
232
|
scoreUserStory,
|
|
220
233
|
scoreboardSummary,
|
|
221
234
|
selectDiscriminative,
|
|
235
|
+
selectPolicyEditAuthorRows,
|
|
222
236
|
sequentialDecide,
|
|
223
237
|
sequentialPairedGate,
|
|
224
238
|
skillOptEntry,
|
|
@@ -228,6 +242,7 @@ export {
|
|
|
228
242
|
tangleTracesRoot,
|
|
229
243
|
traceAnalystProposer,
|
|
230
244
|
userStoryScoreboard,
|
|
245
|
+
validatePolicyEditCandidateRecord,
|
|
231
246
|
validateSearchLedgerEvent,
|
|
232
247
|
verifyCodeSurface
|
|
233
248
|
};
|