@tangle-network/agent-eval 0.95.0 → 0.96.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (968) hide show
  1. package/CHANGELOG.md +17 -0
  2. package/dist/adapters/http.d.ts +17 -10
  3. package/dist/adapters/langchain.d.ts +14 -7
  4. package/dist/adapters/otel.d.ts +25 -13
  5. package/dist/{rl/adversarial.d.ts → adversarial-DIVcDoI_.d.ts} +7 -6
  6. package/dist/analyst/index.d.ts +236 -28
  7. package/dist/analyst/index.js +1 -1
  8. package/dist/{trace-analyst/analyst.d.ts → analyst-C8HHvfJp.d.ts} +14 -12
  9. package/dist/{contract/analyze-runs.d.ts → analyze-runs-DtT6F_6T.d.ts} +10 -9
  10. package/dist/authenticity/index.d.ts +16 -15
  11. package/dist/{baseline.d.ts → baseline-Bbid3WoO.d.ts} +44 -8
  12. package/dist/belief-state/index.d.ts +605 -14
  13. package/dist/benchmarks/index.d.ts +5 -23
  14. package/dist/builder-eval/index.d.ts +250 -5
  15. package/dist/calibration-BPmzuVPk.d.ts +101 -0
  16. package/dist/campaign/index.d.ts +1452 -38
  17. package/dist/{chunk-AQ5WQAIV.js → chunk-3NHEO6ZC.js} +2 -2
  18. package/dist/chunk-3NHEO6ZC.js.map +1 -0
  19. package/dist/cli.d.ts +0 -2
  20. package/dist/{contract/intake/code-agent-session.d.ts → code-agent-session-CPHRCb4-.d.ts} +17 -15
  21. package/dist/contract/index.d.ts +824 -107
  22. package/dist/contract/index.js +145 -1
  23. package/dist/contract/index.js.map +1 -1
  24. package/dist/control-Doncu-B_.d.ts +259 -0
  25. package/dist/{control-runtime.d.ts → control-runtime-Acf9CGhw.d.ts} +26 -23
  26. package/dist/control.d.ts +10 -11
  27. package/dist/corpus-D4YW9UoJ.d.ts +560 -0
  28. package/dist/{cost-ledger.d.ts → cost-ledger-DuSqlw5B.d.ts} +11 -10
  29. package/dist/{counterfactual.d.ts → counterfactual-DlOz8PBx.d.ts} +13 -12
  30. package/dist/{dataset.d.ts → dataset-BbGkaN2I.d.ts} +14 -11
  31. package/dist/{analyst/registry.d.ts → default-registry-GyE8X5SP.d.ts} +37 -8
  32. package/dist/diagnose.d.ts +252 -1
  33. package/dist/{trace/emitter.d.ts → emitter-C2rqGH_l.d.ts} +12 -9
  34. package/dist/{errors.d.ts → errors-CzMUYo7b.d.ts} +11 -10
  35. package/dist/failure-cluster-DH9Flgcf.d.ts +76 -0
  36. package/dist/{feedback-trajectory.d.ts → feedback-trajectory-BxY0cKfs.d.ts} +38 -36
  37. package/dist/fuzz.d.ts +547 -1
  38. package/dist/gepa-Dprxvz8r.d.ts +414 -0
  39. package/dist/governance/index.d.ts +135 -5
  40. package/dist/harness-optimizer-mOl9XX_O.d.ts +106 -0
  41. package/dist/hosted/index.d.ts +239 -10
  42. package/dist/index-_Y4oNOOb.d.ts +159 -0
  43. package/dist/index.d.ts +5660 -277
  44. package/dist/index.js +1 -1
  45. package/dist/{contract/insight-report.d.ts → insight-report-BnRjTibG.d.ts} +19 -16
  46. package/dist/{trace/integrity.d.ts → integrity-D2t12mMw.d.ts} +14 -11
  47. package/dist/{judge-calibration.d.ts → judge-calibration-0p2QcWNE.d.ts} +17 -16
  48. package/dist/{analyst/kind-factory.d.ts → kind-factory-X3eDYbKn.d.ts} +61 -10
  49. package/dist/knowledge/index.d.ts +103 -3
  50. package/dist/{llm-client.d.ts → llm-client-Bj7g0rqu.d.ts} +23 -21
  51. package/dist/matrix/index.d.ts +30 -12
  52. package/dist/meta-eval/index.d.ts +182 -6
  53. package/dist/{multi-layer-verifier.d.ts → multi-layer-verifier-DUZXrPDA.d.ts} +15 -12
  54. package/dist/multishot/index.d.ts +290 -7
  55. package/dist/{rl/off-policy.d.ts → off-policy-DiwuKKg7.d.ts} +9 -8
  56. package/dist/openapi.json +1 -1
  57. package/dist/{meta-eval/outcome-store.d.ts → outcome-store-rnXLEqSn.d.ts} +8 -7
  58. package/dist/{pareto.d.ts → pareto-E-pembql.d.ts} +10 -9
  59. package/dist/perf/index.d.ts +119 -13
  60. package/dist/pipelines/index.d.ts +173 -8
  61. package/dist/pre-registration-nfUdc9EQ.d.ts +483 -0
  62. package/dist/prm/index.d.ts +104 -5
  63. package/dist/provenance-CIxfBnkl.d.ts +426 -0
  64. package/dist/query-B7GGjRox.d.ts +32 -0
  65. package/dist/{trace/raw-provider-sink.d.ts → raw-provider-sink-C46HDghv.d.ts} +14 -13
  66. package/dist/{red-team.d.ts → red-team-BWdoyleI.d.ts} +16 -14
  67. package/dist/{trace/redact.d.ts → redact-B40YG2M_.d.ts} +8 -7
  68. package/dist/release-report-pidWUMZ2.d.ts +233 -0
  69. package/dist/reporting.d.ts +16 -15
  70. package/dist/{eval-campaign.d.ts → researcher-Jr8ME1dZ.d.ts} +156 -34
  71. package/dist/rl.d.ts +1193 -1
  72. package/dist/{prm/rubric.d.ts → rubric-Cc6UHvUb.d.ts} +13 -10
  73. package/dist/{meta-eval/rubric-predictive-validity.d.ts → rubric-predictive-validity-C2hDKM8Z.d.ts} +11 -8
  74. package/dist/run-critic-CmMf05uV.d.ts +56 -0
  75. package/dist/{run-record.d.ts → run-record-CP2ObebC.d.ts} +117 -15
  76. package/dist/runtime-trajectory-BOUUjI0y.d.ts +49 -0
  77. package/dist/{trace/schema.d.ts → schema-m0gsnbt3.d.ts} +30 -29
  78. package/dist/semantic-concept-judge-DSBB2Cfp.d.ts +624 -0
  79. package/dist/{sequential.d.ts → sequential-5iSVfzl2.d.ts} +10 -9
  80. package/dist/{series-convergence.d.ts → series-convergence-D5OWMBg6.d.ts} +5 -4
  81. package/dist/sink-fetch-B1Yg4Til.d.ts +101 -0
  82. package/dist/{statistics.d.ts → statistics-CCJpTGOS.d.ts} +48 -46
  83. package/dist/{trace/store.d.ts → store-BcFXE6LG.d.ts} +12 -21
  84. package/dist/{trace-analyst/types.d.ts → store-C1YxJDEK.d.ts} +74 -18
  85. package/dist/storyboard/index.d.ts +81 -16
  86. package/dist/{summary-report.d.ts → summary-report-CInXwsza.d.ts} +160 -23
  87. package/dist/telemetry/{sink-file.d.ts → file.d.ts} +7 -5
  88. package/dist/telemetry/index.d.ts +35 -17
  89. package/dist/{sandbox-harness.d.ts → test-graded-scenario-DeODGLra.d.ts} +60 -15
  90. package/dist/{locked-jsonl-appender.d.ts → testing-C21CHsq2.d.ts} +4 -3
  91. package/dist/testing.d.ts +1 -5
  92. package/dist/traces.d.ts +976 -4
  93. package/dist/{trajectory.d.ts → trajectory-2TkpSEVh.d.ts} +9 -8
  94. package/dist/{analyst/types.d.ts → types-B5x54y6n.d.ts} +112 -19
  95. package/dist/{campaign/types.d.ts → types-BMahhhio.d.ts} +47 -44
  96. package/dist/{matrix/types.d.ts → types-BUxNaJ8c.d.ts} +11 -9
  97. package/dist/{types.d.ts → types-C7DGg5ex.d.ts} +33 -31
  98. package/dist/{verdict.d.ts → verdict-C9MlYujm.d.ts} +3 -2
  99. package/dist/wire/index.d.ts +570 -13
  100. package/dist/workflow/index.d.ts +496 -22
  101. package/package.json +2 -2
  102. package/dist/action-policy.d.ts +0 -24
  103. package/dist/action-policy.d.ts.map +0 -1
  104. package/dist/action-policy.test.d.ts +0 -2
  105. package/dist/action-policy.test.d.ts.map +0 -1
  106. package/dist/active-learning.d.ts +0 -41
  107. package/dist/active-learning.d.ts.map +0 -1
  108. package/dist/adapters/http.d.ts.map +0 -1
  109. package/dist/adapters/langchain.d.ts.map +0 -1
  110. package/dist/adapters/otel.d.ts.map +0 -1
  111. package/dist/agent-profile-cell.d.ts +0 -101
  112. package/dist/agent-profile-cell.d.ts.map +0 -1
  113. package/dist/agent-profile.d.ts +0 -27
  114. package/dist/agent-profile.d.ts.map +0 -1
  115. package/dist/agent-profile.test.d.ts +0 -2
  116. package/dist/agent-profile.test.d.ts.map +0 -1
  117. package/dist/analyst/adapters.d.ts +0 -62
  118. package/dist/analyst/adapters.d.ts.map +0 -1
  119. package/dist/analyst/analyst.test.d.ts +0 -2
  120. package/dist/analyst/analyst.test.d.ts.map +0 -1
  121. package/dist/analyst/ax-service.d.ts +0 -27
  122. package/dist/analyst/ax-service.d.ts.map +0 -1
  123. package/dist/analyst/behavioral-analyst.d.ts +0 -28
  124. package/dist/analyst/behavioral-analyst.d.ts.map +0 -1
  125. package/dist/analyst/chat-client.d.ts +0 -91
  126. package/dist/analyst/chat-client.d.ts.map +0 -1
  127. package/dist/analyst/default-registry.d.ts +0 -27
  128. package/dist/analyst/default-registry.d.ts.map +0 -1
  129. package/dist/analyst/default-registry.test.d.ts +0 -2
  130. package/dist/analyst/default-registry.test.d.ts.map +0 -1
  131. package/dist/analyst/finding-signature.d.ts +0 -48
  132. package/dist/analyst/finding-signature.d.ts.map +0 -1
  133. package/dist/analyst/finding-subject.d.ts +0 -146
  134. package/dist/analyst/finding-subject.d.ts.map +0 -1
  135. package/dist/analyst/finding-subject.test.d.ts +0 -2
  136. package/dist/analyst/finding-subject.test.d.ts.map +0 -1
  137. package/dist/analyst/findings-store.d.ts +0 -75
  138. package/dist/analyst/findings-store.d.ts.map +0 -1
  139. package/dist/analyst/index.d.ts.map +0 -1
  140. package/dist/analyst/kind-factory.d.ts.map +0 -1
  141. package/dist/analyst/kinds/failure-mode.d.ts +0 -19
  142. package/dist/analyst/kinds/failure-mode.d.ts.map +0 -1
  143. package/dist/analyst/kinds/improvement.d.ts +0 -23
  144. package/dist/analyst/kinds/improvement.d.ts.map +0 -1
  145. package/dist/analyst/kinds/index.d.ts +0 -22
  146. package/dist/analyst/kinds/index.d.ts.map +0 -1
  147. package/dist/analyst/kinds/kinds.test.d.ts +0 -2
  148. package/dist/analyst/kinds/kinds.test.d.ts.map +0 -1
  149. package/dist/analyst/kinds/knowledge-gap.d.ts +0 -28
  150. package/dist/analyst/kinds/knowledge-gap.d.ts.map +0 -1
  151. package/dist/analyst/kinds/knowledge-poisoning.d.ts +0 -22
  152. package/dist/analyst/kinds/knowledge-poisoning.d.ts.map +0 -1
  153. package/dist/analyst/kinds/skill-usage.d.ts +0 -84
  154. package/dist/analyst/kinds/skill-usage.d.ts.map +0 -1
  155. package/dist/analyst/kinds/skill-usage.test.d.ts +0 -2
  156. package/dist/analyst/kinds/skill-usage.test.d.ts.map +0 -1
  157. package/dist/analyst/parse-tolerant.d.ts +0 -26
  158. package/dist/analyst/parse-tolerant.d.ts.map +0 -1
  159. package/dist/analyst/parse-tolerant.test.d.ts +0 -2
  160. package/dist/analyst/parse-tolerant.test.d.ts.map +0 -1
  161. package/dist/analyst/registry.budget.test.d.ts +0 -2
  162. package/dist/analyst/registry.budget.test.d.ts.map +0 -1
  163. package/dist/analyst/registry.d.ts.map +0 -1
  164. package/dist/analyst/steer-firewall.d.ts +0 -35
  165. package/dist/analyst/steer-firewall.d.ts.map +0 -1
  166. package/dist/analyst/steer-firewall.test.d.ts +0 -2
  167. package/dist/analyst/steer-firewall.test.d.ts.map +0 -1
  168. package/dist/analyst/structure-findings.d.ts +0 -37
  169. package/dist/analyst/structure-findings.d.ts.map +0 -1
  170. package/dist/analyst/structure-findings.test.d.ts +0 -2
  171. package/dist/analyst/structure-findings.test.d.ts.map +0 -1
  172. package/dist/analyst/tool-groups.d.ts +0 -34
  173. package/dist/analyst/tool-groups.d.ts.map +0 -1
  174. package/dist/analyst/types.d.ts.map +0 -1
  175. package/dist/anti-slop.d.ts +0 -59
  176. package/dist/anti-slop.d.ts.map +0 -1
  177. package/dist/artifact-validator.d.ts +0 -74
  178. package/dist/artifact-validator.d.ts.map +0 -1
  179. package/dist/attestation.d.ts +0 -63
  180. package/dist/attestation.d.ts.map +0 -1
  181. package/dist/attestation.test.d.ts +0 -2
  182. package/dist/attestation.test.d.ts.map +0 -1
  183. package/dist/authenticity/index.d.ts.map +0 -1
  184. package/dist/authenticity/index.test.d.ts +0 -2
  185. package/dist/authenticity/index.test.d.ts.map +0 -1
  186. package/dist/auto-pr.d.ts +0 -120
  187. package/dist/auto-pr.d.ts.map +0 -1
  188. package/dist/baseline.d.ts.map +0 -1
  189. package/dist/behavior-dsl.d.ts +0 -73
  190. package/dist/behavior-dsl.d.ts.map +0 -1
  191. package/dist/belief-state/calibration.d.ts +0 -10
  192. package/dist/belief-state/calibration.d.ts.map +0 -1
  193. package/dist/belief-state/calibration.test.d.ts +0 -2
  194. package/dist/belief-state/calibration.test.d.ts.map +0 -1
  195. package/dist/belief-state/code-agent-corpus.d.ts +0 -66
  196. package/dist/belief-state/code-agent-corpus.d.ts.map +0 -1
  197. package/dist/belief-state/code-agent-corpus.test.d.ts +0 -2
  198. package/dist/belief-state/code-agent-corpus.test.d.ts.map +0 -1
  199. package/dist/belief-state/code-agent-evidence.d.ts +0 -22
  200. package/dist/belief-state/code-agent-evidence.d.ts.map +0 -1
  201. package/dist/belief-state/code-agent-evidence.test.d.ts +0 -2
  202. package/dist/belief-state/code-agent-evidence.test.d.ts.map +0 -1
  203. package/dist/belief-state/extract.d.ts +0 -7
  204. package/dist/belief-state/extract.d.ts.map +0 -1
  205. package/dist/belief-state/extract.test.d.ts +0 -2
  206. package/dist/belief-state/extract.test.d.ts.map +0 -1
  207. package/dist/belief-state/index.d.ts.map +0 -1
  208. package/dist/belief-state/ope.d.ts +0 -17
  209. package/dist/belief-state/ope.d.ts.map +0 -1
  210. package/dist/belief-state/ope.test.d.ts +0 -2
  211. package/dist/belief-state/ope.test.d.ts.map +0 -1
  212. package/dist/belief-state/phase0-measurement.d.ts +0 -55
  213. package/dist/belief-state/phase0-measurement.d.ts.map +0 -1
  214. package/dist/belief-state/report.d.ts +0 -17
  215. package/dist/belief-state/report.d.ts.map +0 -1
  216. package/dist/belief-state/report.test.d.ts +0 -2
  217. package/dist/belief-state/report.test.d.ts.map +0 -1
  218. package/dist/belief-state/research-evidence.d.ts +0 -23
  219. package/dist/belief-state/research-evidence.d.ts.map +0 -1
  220. package/dist/belief-state/research-evidence.test.d.ts +0 -2
  221. package/dist/belief-state/research-evidence.test.d.ts.map +0 -1
  222. package/dist/belief-state/runtime-benchmark-corpus.d.ts +0 -32
  223. package/dist/belief-state/runtime-benchmark-corpus.d.ts.map +0 -1
  224. package/dist/belief-state/runtime-hooks.d.ts +0 -87
  225. package/dist/belief-state/runtime-hooks.d.ts.map +0 -1
  226. package/dist/belief-state/runtime-hooks.test.d.ts +0 -2
  227. package/dist/belief-state/runtime-hooks.test.d.ts.map +0 -1
  228. package/dist/belief-state/selective.d.ts +0 -15
  229. package/dist/belief-state/selective.d.ts.map +0 -1
  230. package/dist/belief-state/selective.test.d.ts +0 -2
  231. package/dist/belief-state/selective.test.d.ts.map +0 -1
  232. package/dist/belief-state/shadow-probe.d.ts +0 -80
  233. package/dist/belief-state/shadow-probe.d.ts.map +0 -1
  234. package/dist/belief-state/shadow-probe.test.d.ts +0 -2
  235. package/dist/belief-state/shadow-probe.test.d.ts.map +0 -1
  236. package/dist/belief-state/types.d.ts +0 -195
  237. package/dist/belief-state/types.d.ts.map +0 -1
  238. package/dist/belief-state/types.test.d.ts +0 -2
  239. package/dist/belief-state/types.test.d.ts.map +0 -1
  240. package/dist/benchmark.d.ts +0 -14
  241. package/dist/benchmark.d.ts.map +0 -1
  242. package/dist/benchmarks/index.d.ts.map +0 -1
  243. package/dist/benchmarks/routing/dataset.d.ts +0 -34
  244. package/dist/benchmarks/routing/dataset.d.ts.map +0 -1
  245. package/dist/benchmarks/routing/index.d.ts +0 -34
  246. package/dist/benchmarks/routing/index.d.ts.map +0 -1
  247. package/dist/benchmarks/types.d.ts +0 -49
  248. package/dist/benchmarks/types.d.ts.map +0 -1
  249. package/dist/bisector.d.ts +0 -81
  250. package/dist/bisector.d.ts.map +0 -1
  251. package/dist/budget-guard.d.ts +0 -31
  252. package/dist/budget-guard.d.ts.map +0 -1
  253. package/dist/builder-eval/builder-session.d.ts +0 -111
  254. package/dist/builder-eval/builder-session.d.ts.map +0 -1
  255. package/dist/builder-eval/correlation.d.ts +0 -32
  256. package/dist/builder-eval/correlation.d.ts.map +0 -1
  257. package/dist/builder-eval/index.d.ts.map +0 -1
  258. package/dist/builder-eval/project-registry.d.ts +0 -51
  259. package/dist/builder-eval/project-registry.d.ts.map +0 -1
  260. package/dist/builder-eval/three-layer-eval.d.ts +0 -55
  261. package/dist/builder-eval/three-layer-eval.d.ts.map +0 -1
  262. package/dist/campaign/analyst-surface.d.ts +0 -108
  263. package/dist/campaign/analyst-surface.d.ts.map +0 -1
  264. package/dist/campaign/analyst-surface.test.d.ts +0 -2
  265. package/dist/campaign/analyst-surface.test.d.ts.map +0 -1
  266. package/dist/campaign/auto-pr.d.ts +0 -46
  267. package/dist/campaign/auto-pr.d.ts.map +0 -1
  268. package/dist/campaign/distillation/agreement-judge.d.ts +0 -69
  269. package/dist/campaign/distillation/agreement-judge.d.ts.map +0 -1
  270. package/dist/campaign/distillation/cli.d.ts +0 -35
  271. package/dist/campaign/distillation/cli.d.ts.map +0 -1
  272. package/dist/campaign/distillation/distillation.test.d.ts +0 -2
  273. package/dist/campaign/distillation/distillation.test.d.ts.map +0 -1
  274. package/dist/campaign/distillation/gold-scenarios.d.ts +0 -54
  275. package/dist/campaign/distillation/gold-scenarios.d.ts.map +0 -1
  276. package/dist/campaign/distillation/run-distillation.d.ts +0 -119
  277. package/dist/campaign/distillation/run-distillation.d.ts.map +0 -1
  278. package/dist/campaign/gates/compose.d.ts +0 -12
  279. package/dist/campaign/gates/compose.d.ts.map +0 -1
  280. package/dist/campaign/gates/default-production-gate.d.ts +0 -58
  281. package/dist/campaign/gates/default-production-gate.d.ts.map +0 -1
  282. package/dist/campaign/gates/heldout-gate.d.ts +0 -12
  283. package/dist/campaign/gates/heldout-gate.d.ts.map +0 -1
  284. package/dist/campaign/gates/promotion-policy.d.ts +0 -125
  285. package/dist/campaign/gates/promotion-policy.d.ts.map +0 -1
  286. package/dist/campaign/gates/promotion-policy.test.d.ts +0 -2
  287. package/dist/campaign/gates/promotion-policy.test.d.ts.map +0 -1
  288. package/dist/campaign/gates/sequential.d.ts +0 -146
  289. package/dist/campaign/gates/sequential.d.ts.map +0 -1
  290. package/dist/campaign/gates/sequential.test.d.ts +0 -2
  291. package/dist/campaign/gates/sequential.test.d.ts.map +0 -1
  292. package/dist/campaign/gates/statistical-heldout.d.ts +0 -99
  293. package/dist/campaign/gates/statistical-heldout.d.ts.map +0 -1
  294. package/dist/campaign/gates/statistical-heldout.test.d.ts +0 -2
  295. package/dist/campaign/gates/statistical-heldout.test.d.ts.map +0 -1
  296. package/dist/campaign/index.d.ts.map +0 -1
  297. package/dist/campaign/labeled-store/fs-adapter.d.ts +0 -59
  298. package/dist/campaign/labeled-store/fs-adapter.d.ts.map +0 -1
  299. package/dist/campaign/presets/compare-proposers.d.ts +0 -146
  300. package/dist/campaign/presets/compare-proposers.d.ts.map +0 -1
  301. package/dist/campaign/presets/playback.d.ts +0 -120
  302. package/dist/campaign/presets/playback.d.ts.map +0 -1
  303. package/dist/campaign/presets/playback.test.d.ts +0 -2
  304. package/dist/campaign/presets/playback.test.d.ts.map +0 -1
  305. package/dist/campaign/presets/run-eval.d.ts +0 -14
  306. package/dist/campaign/presets/run-eval.d.ts.map +0 -1
  307. package/dist/campaign/presets/run-improvement-loop.d.ts +0 -63
  308. package/dist/campaign/presets/run-improvement-loop.d.ts.map +0 -1
  309. package/dist/campaign/presets/run-improvement-loop.test.d.ts +0 -2
  310. package/dist/campaign/presets/run-improvement-loop.test.d.ts.map +0 -1
  311. package/dist/campaign/presets/run-optimization.d.ts +0 -92
  312. package/dist/campaign/presets/run-optimization.d.ts.map +0 -1
  313. package/dist/campaign/presets/run-profile-matrix.d.ts +0 -151
  314. package/dist/campaign/presets/run-profile-matrix.d.ts.map +0 -1
  315. package/dist/campaign/presets/run-skill-opt.d.ts +0 -96
  316. package/dist/campaign/presets/run-skill-opt.d.ts.map +0 -1
  317. package/dist/campaign/proposers/_findings-text.d.ts +0 -22
  318. package/dist/campaign/proposers/_findings-text.d.ts.map +0 -1
  319. package/dist/campaign/proposers/ace.d.ts +0 -33
  320. package/dist/campaign/proposers/ace.d.ts.map +0 -1
  321. package/dist/campaign/proposers/ace.test.d.ts +0 -2
  322. package/dist/campaign/proposers/ace.test.d.ts.map +0 -1
  323. package/dist/campaign/proposers/analysis-edit.d.ts +0 -32
  324. package/dist/campaign/proposers/analysis-edit.d.ts.map +0 -1
  325. package/dist/campaign/proposers/evolutionary.d.ts +0 -20
  326. package/dist/campaign/proposers/evolutionary.d.ts.map +0 -1
  327. package/dist/campaign/proposers/fapo.d.ts +0 -120
  328. package/dist/campaign/proposers/fapo.d.ts.map +0 -1
  329. package/dist/campaign/proposers/gepa.d.ts +0 -86
  330. package/dist/campaign/proposers/gepa.d.ts.map +0 -1
  331. package/dist/campaign/proposers/halo.d.ts +0 -44
  332. package/dist/campaign/proposers/halo.d.ts.map +0 -1
  333. package/dist/campaign/proposers/halo.test.d.ts +0 -2
  334. package/dist/campaign/proposers/halo.test.d.ts.map +0 -1
  335. package/dist/campaign/proposers/memory.d.ts +0 -47
  336. package/dist/campaign/proposers/memory.d.ts.map +0 -1
  337. package/dist/campaign/proposers/memory.test.d.ts +0 -2
  338. package/dist/campaign/proposers/memory.test.d.ts.map +0 -1
  339. package/dist/campaign/proposers/skill-opt.d.ts +0 -88
  340. package/dist/campaign/proposers/skill-opt.d.ts.map +0 -1
  341. package/dist/campaign/proposers/trace-analyst.d.ts +0 -48
  342. package/dist/campaign/proposers/trace-analyst.d.ts.map +0 -1
  343. package/dist/campaign/proposers/trace-analyst.test.d.ts +0 -2
  344. package/dist/campaign/proposers/trace-analyst.test.d.ts.map +0 -1
  345. package/dist/campaign/provenance.d.ts +0 -185
  346. package/dist/campaign/provenance.d.ts.map +0 -1
  347. package/dist/campaign/run-campaign.d.ts +0 -90
  348. package/dist/campaign/run-campaign.d.ts.map +0 -1
  349. package/dist/campaign/score-utils.d.ts +0 -26
  350. package/dist/campaign/score-utils.d.ts.map +0 -1
  351. package/dist/campaign/skill-patch.d.ts +0 -62
  352. package/dist/campaign/skill-patch.d.ts.map +0 -1
  353. package/dist/campaign/storage.d.ts +0 -38
  354. package/dist/campaign/storage.d.ts.map +0 -1
  355. package/dist/campaign/types.d.ts.map +0 -1
  356. package/dist/campaign/worktree/index.d.ts +0 -53
  357. package/dist/campaign/worktree/index.d.ts.map +0 -1
  358. package/dist/canary.d.ts +0 -101
  359. package/dist/canary.d.ts.map +0 -1
  360. package/dist/causal-attribution.d.ts +0 -45
  361. package/dist/causal-attribution.d.ts.map +0 -1
  362. package/dist/chunk-AQ5WQAIV.js.map +0 -1
  363. package/dist/ci-gate.d.ts +0 -44
  364. package/dist/ci-gate.d.ts.map +0 -1
  365. package/dist/cli.d.ts.map +0 -1
  366. package/dist/client.d.ts +0 -77
  367. package/dist/client.d.ts.map +0 -1
  368. package/dist/client.test.d.ts +0 -2
  369. package/dist/client.test.d.ts.map +0 -1
  370. package/dist/command-runner.d.ts +0 -74
  371. package/dist/command-runner.d.ts.map +0 -1
  372. package/dist/command-runner.test.d.ts +0 -2
  373. package/dist/command-runner.test.d.ts.map +0 -1
  374. package/dist/completion-verifier.d.ts +0 -147
  375. package/dist/completion-verifier.d.ts.map +0 -1
  376. package/dist/completion-verifier.test.d.ts +0 -9
  377. package/dist/completion-verifier.test.d.ts.map +0 -1
  378. package/dist/concurrency.d.ts +0 -23
  379. package/dist/concurrency.d.ts.map +0 -1
  380. package/dist/contamination-guard.d.ts +0 -81
  381. package/dist/contamination-guard.d.ts.map +0 -1
  382. package/dist/contract/analyze-runs.d.ts.map +0 -1
  383. package/dist/contract/define-agent-eval.d.ts +0 -52
  384. package/dist/contract/define-agent-eval.d.ts.map +0 -1
  385. package/dist/contract/diff.d.ts +0 -114
  386. package/dist/contract/diff.d.ts.map +0 -1
  387. package/dist/contract/index.d.ts.map +0 -1
  388. package/dist/contract/insight-report.d.ts.map +0 -1
  389. package/dist/contract/insight-types-fwd.d.ts +0 -7
  390. package/dist/contract/insight-types-fwd.d.ts.map +0 -1
  391. package/dist/contract/intake/agent-trace.d.ts +0 -97
  392. package/dist/contract/intake/agent-trace.d.ts.map +0 -1
  393. package/dist/contract/intake/code-agent-session.d.ts.map +0 -1
  394. package/dist/contract/intake/feedback-table.d.ts +0 -87
  395. package/dist/contract/intake/feedback-table.d.ts.map +0 -1
  396. package/dist/contract/intake/index.d.ts +0 -22
  397. package/dist/contract/intake/index.d.ts.map +0 -1
  398. package/dist/contract/intake/otel-spans.d.ts +0 -33
  399. package/dist/contract/intake/otel-spans.d.ts.map +0 -1
  400. package/dist/contract/self-improve.d.ts +0 -284
  401. package/dist/contract/self-improve.d.ts.map +0 -1
  402. package/dist/control-runtime.d.ts.map +0 -1
  403. package/dist/control-runtime.test.d.ts +0 -2
  404. package/dist/control-runtime.test.d.ts.map +0 -1
  405. package/dist/control.d.ts.map +0 -1
  406. package/dist/convergence.d.ts +0 -29
  407. package/dist/convergence.d.ts.map +0 -1
  408. package/dist/cost-ledger.d.ts.map +0 -1
  409. package/dist/cost-ledger.test.d.ts +0 -2
  410. package/dist/cost-ledger.test.d.ts.map +0 -1
  411. package/dist/cost-report.d.ts +0 -43
  412. package/dist/cost-report.d.ts.map +0 -1
  413. package/dist/cost-report.test.d.ts +0 -2
  414. package/dist/cost-report.test.d.ts.map +0 -1
  415. package/dist/cost-tracker.d.ts +0 -76
  416. package/dist/cost-tracker.d.ts.map +0 -1
  417. package/dist/counterfactual.d.ts.map +0 -1
  418. package/dist/cross-trace-diff.d.ts +0 -56
  419. package/dist/cross-trace-diff.d.ts.map +0 -1
  420. package/dist/dataset.d.ts.map +0 -1
  421. package/dist/deploy-gate-layer.d.ts +0 -125
  422. package/dist/deploy-gate-layer.d.ts.map +0 -1
  423. package/dist/deploy-gate-layer.test.d.ts +0 -2
  424. package/dist/deploy-gate-layer.test.d.ts.map +0 -1
  425. package/dist/description-length-gate.d.ts +0 -119
  426. package/dist/description-length-gate.d.ts.map +0 -1
  427. package/dist/detectors/edge.test.d.ts +0 -2
  428. package/dist/detectors/edge.test.d.ts.map +0 -1
  429. package/dist/detectors/index.d.ts +0 -81
  430. package/dist/detectors/index.d.ts.map +0 -1
  431. package/dist/detectors/index.test.d.ts +0 -2
  432. package/dist/detectors/index.test.d.ts.map +0 -1
  433. package/dist/diagnose/causal-sweep.d.ts +0 -100
  434. package/dist/diagnose/causal-sweep.d.ts.map +0 -1
  435. package/dist/diagnose/index.d.ts +0 -36
  436. package/dist/diagnose/index.d.ts.map +0 -1
  437. package/dist/diagnose/remediation.d.ts +0 -68
  438. package/dist/diagnose/remediation.d.ts.map +0 -1
  439. package/dist/diagnose/repair.d.ts +0 -77
  440. package/dist/diagnose/repair.d.ts.map +0 -1
  441. package/dist/discover-personas.d.ts +0 -35
  442. package/dist/discover-personas.d.ts.map +0 -1
  443. package/dist/driver.d.ts +0 -95
  444. package/dist/driver.d.ts.map +0 -1
  445. package/dist/driver.test.d.ts +0 -8
  446. package/dist/driver.test.d.ts.map +0 -1
  447. package/dist/dual-agent-bench.d.ts +0 -81
  448. package/dist/dual-agent-bench.d.ts.map +0 -1
  449. package/dist/error-count-extractor.d.ts +0 -47
  450. package/dist/error-count-extractor.d.ts.map +0 -1
  451. package/dist/error-count-extractor.test.d.ts +0 -2
  452. package/dist/error-count-extractor.test.d.ts.map +0 -1
  453. package/dist/errors.d.ts.map +0 -1
  454. package/dist/eval-campaign.d.ts.map +0 -1
  455. package/dist/eval-campaign.test.d.ts +0 -2
  456. package/dist/eval-campaign.test.d.ts.map +0 -1
  457. package/dist/eval-tools.d.ts +0 -55
  458. package/dist/eval-tools.d.ts.map +0 -1
  459. package/dist/eval-trace-store.d.ts +0 -107
  460. package/dist/eval-trace-store.d.ts.map +0 -1
  461. package/dist/eval-trace-store.test.d.ts +0 -2
  462. package/dist/eval-trace-store.test.d.ts.map +0 -1
  463. package/dist/executor.d.ts +0 -38
  464. package/dist/executor.d.ts.map +0 -1
  465. package/dist/executor.test.d.ts +0 -10
  466. package/dist/executor.test.d.ts.map +0 -1
  467. package/dist/experiment-tracker.d.ts +0 -178
  468. package/dist/experiment-tracker.d.ts.map +0 -1
  469. package/dist/experiment-tracker.test.d.ts +0 -2
  470. package/dist/experiment-tracker.test.d.ts.map +0 -1
  471. package/dist/failure-taxonomy.d.ts +0 -38
  472. package/dist/failure-taxonomy.d.ts.map +0 -1
  473. package/dist/feedback-trajectory.d.ts.map +0 -1
  474. package/dist/feedback-trajectory.test.d.ts +0 -2
  475. package/dist/feedback-trajectory.test.d.ts.map +0 -1
  476. package/dist/flow-layer.d.ts +0 -90
  477. package/dist/flow-layer.d.ts.map +0 -1
  478. package/dist/flow-layer.test.d.ts +0 -2
  479. package/dist/flow-layer.test.d.ts.map +0 -1
  480. package/dist/fuzz/capsule.d.ts +0 -46
  481. package/dist/fuzz/capsule.d.ts.map +0 -1
  482. package/dist/fuzz/cube.d.ts +0 -36
  483. package/dist/fuzz/cube.d.ts.map +0 -1
  484. package/dist/fuzz/explorer-cost.test.d.ts +0 -2
  485. package/dist/fuzz/explorer-cost.test.d.ts.map +0 -1
  486. package/dist/fuzz/explorer.d.ts +0 -64
  487. package/dist/fuzz/explorer.d.ts.map +0 -1
  488. package/dist/fuzz/fuzz-agent.d.ts +0 -16
  489. package/dist/fuzz/fuzz-agent.d.ts.map +0 -1
  490. package/dist/fuzz/fuzz-agent.test.d.ts +0 -2
  491. package/dist/fuzz/fuzz-agent.test.d.ts.map +0 -1
  492. package/dist/fuzz/gates.d.ts +0 -33
  493. package/dist/fuzz/gates.d.ts.map +0 -1
  494. package/dist/fuzz/index.d.ts +0 -26
  495. package/dist/fuzz/index.d.ts.map +0 -1
  496. package/dist/fuzz/policies.d.ts +0 -28
  497. package/dist/fuzz/policies.d.ts.map +0 -1
  498. package/dist/fuzz/tools.d.ts +0 -20
  499. package/dist/fuzz/tools.d.ts.map +0 -1
  500. package/dist/fuzz/types.d.ts +0 -307
  501. package/dist/fuzz/types.d.ts.map +0 -1
  502. package/dist/golden-matcher.d.ts +0 -71
  503. package/dist/golden-matcher.d.ts.map +0 -1
  504. package/dist/governance/eu-ai-act.d.ts +0 -37
  505. package/dist/governance/eu-ai-act.d.ts.map +0 -1
  506. package/dist/governance/index.d.ts.map +0 -1
  507. package/dist/governance/nist-ai-rmf.d.ts +0 -15
  508. package/dist/governance/nist-ai-rmf.d.ts.map +0 -1
  509. package/dist/governance/soc2.d.ts +0 -12
  510. package/dist/governance/soc2.d.ts.map +0 -1
  511. package/dist/governance/types.d.ts +0 -66
  512. package/dist/governance/types.d.ts.map +0 -1
  513. package/dist/harness-optimizer.d.ts +0 -82
  514. package/dist/harness-optimizer.d.ts.map +0 -1
  515. package/dist/held-out-gate.d.ts +0 -135
  516. package/dist/held-out-gate.d.ts.map +0 -1
  517. package/dist/hosted/client.d.ts +0 -73
  518. package/dist/hosted/client.d.ts.map +0 -1
  519. package/dist/hosted/from-env.test.d.ts +0 -8
  520. package/dist/hosted/from-env.test.d.ts.map +0 -1
  521. package/dist/hosted/index.d.ts.map +0 -1
  522. package/dist/hosted/types.d.ts +0 -159
  523. package/dist/hosted/types.d.ts.map +0 -1
  524. package/dist/index.d.ts.map +0 -1
  525. package/dist/integrity/backend-integrity.d.ts +0 -71
  526. package/dist/integrity/backend-integrity.d.ts.map +0 -1
  527. package/dist/integrity/preflight.d.ts +0 -72
  528. package/dist/integrity/preflight.d.ts.map +0 -1
  529. package/dist/integrity/preflight.test.d.ts +0 -2
  530. package/dist/integrity/preflight.test.d.ts.map +0 -1
  531. package/dist/integrity/single-backend.d.ts +0 -67
  532. package/dist/integrity/single-backend.d.ts.map +0 -1
  533. package/dist/intent-match-judge.d.ts +0 -69
  534. package/dist/intent-match-judge.d.ts.map +0 -1
  535. package/dist/intent-match-judge.test.d.ts +0 -2
  536. package/dist/intent-match-judge.test.d.ts.map +0 -1
  537. package/dist/judge-calibration.d.ts.map +0 -1
  538. package/dist/judge-ensemble.d.ts +0 -66
  539. package/dist/judge-ensemble.d.ts.map +0 -1
  540. package/dist/judge-ensemble.test.d.ts +0 -8
  541. package/dist/judge-ensemble.test.d.ts.map +0 -1
  542. package/dist/judge-families.d.ts +0 -38
  543. package/dist/judge-families.d.ts.map +0 -1
  544. package/dist/judge-panel.d.ts +0 -65
  545. package/dist/judge-panel.d.ts.map +0 -1
  546. package/dist/judge-retry.d.ts +0 -70
  547. package/dist/judge-retry.d.ts.map +0 -1
  548. package/dist/judge-runner.d.ts +0 -36
  549. package/dist/judge-runner.d.ts.map +0 -1
  550. package/dist/judge-runner.test.d.ts +0 -2
  551. package/dist/judge-runner.test.d.ts.map +0 -1
  552. package/dist/judges.d.ts +0 -74
  553. package/dist/judges.d.ts.map +0 -1
  554. package/dist/keyword-coverage-judge.d.ts +0 -89
  555. package/dist/keyword-coverage-judge.d.ts.map +0 -1
  556. package/dist/keyword-coverage-judge.test.d.ts +0 -2
  557. package/dist/keyword-coverage-judge.test.d.ts.map +0 -1
  558. package/dist/knowledge/index.d.ts.map +0 -1
  559. package/dist/knowledge/readiness.d.ts +0 -26
  560. package/dist/knowledge/readiness.d.ts.map +0 -1
  561. package/dist/knowledge/types.d.ts +0 -75
  562. package/dist/knowledge/types.d.ts.map +0 -1
  563. package/dist/live-proof.d.ts +0 -62
  564. package/dist/live-proof.d.ts.map +0 -1
  565. package/dist/llm-client.d.ts.map +0 -1
  566. package/dist/llm-client.test.d.ts +0 -2
  567. package/dist/llm-client.test.d.ts.map +0 -1
  568. package/dist/locked-jsonl-appender.d.ts.map +0 -1
  569. package/dist/matrix/aggregation.d.ts +0 -16
  570. package/dist/matrix/aggregation.d.ts.map +0 -1
  571. package/dist/matrix/index.d.ts.map +0 -1
  572. package/dist/matrix/runner.d.ts +0 -15
  573. package/dist/matrix/runner.d.ts.map +0 -1
  574. package/dist/matrix/types.d.ts.map +0 -1
  575. package/dist/meta-eval/calibration.d.ts +0 -47
  576. package/dist/meta-eval/calibration.d.ts.map +0 -1
  577. package/dist/meta-eval/correlation-study.d.ts +0 -53
  578. package/dist/meta-eval/correlation-study.d.ts.map +0 -1
  579. package/dist/meta-eval/index.d.ts.map +0 -1
  580. package/dist/meta-eval/outcome-store.d.ts.map +0 -1
  581. package/dist/meta-eval/rubric-predictive-validity.d.ts.map +0 -1
  582. package/dist/meta-eval/sentinel.d.ts +0 -169
  583. package/dist/meta-eval/sentinel.d.ts.map +0 -1
  584. package/dist/metrics.d.ts +0 -63
  585. package/dist/metrics.d.ts.map +0 -1
  586. package/dist/model-seats.d.ts +0 -71
  587. package/dist/model-seats.d.ts.map +0 -1
  588. package/dist/model-seats.test.d.ts +0 -2
  589. package/dist/model-seats.test.d.ts.map +0 -1
  590. package/dist/muffled-gate-scanner.d.ts +0 -102
  591. package/dist/muffled-gate-scanner.d.ts.map +0 -1
  592. package/dist/multi-layer-verifier.d.ts.map +0 -1
  593. package/dist/multi-layer-verifier.test.d.ts +0 -2
  594. package/dist/multi-layer-verifier.test.d.ts.map +0 -1
  595. package/dist/multi-toolchain-layer.d.ts +0 -80
  596. package/dist/multi-toolchain-layer.d.ts.map +0 -1
  597. package/dist/multi-toolchain-layer.test.d.ts +0 -2
  598. package/dist/multi-toolchain-layer.test.d.ts.map +0 -1
  599. package/dist/multishot/default-tools.d.ts +0 -34
  600. package/dist/multishot/default-tools.d.ts.map +0 -1
  601. package/dist/multishot/index.d.ts.map +0 -1
  602. package/dist/multishot/judges.d.ts +0 -32
  603. package/dist/multishot/judges.d.ts.map +0 -1
  604. package/dist/multishot/matrix.d.ts +0 -107
  605. package/dist/multishot/matrix.d.ts.map +0 -1
  606. package/dist/multishot/multishot.d.ts +0 -23
  607. package/dist/multishot/multishot.d.ts.map +0 -1
  608. package/dist/multishot/router.d.ts +0 -37
  609. package/dist/multishot/router.d.ts.map +0 -1
  610. package/dist/multishot/types.d.ts +0 -60
  611. package/dist/multishot/types.d.ts.map +0 -1
  612. package/dist/observability.d.ts +0 -71
  613. package/dist/observability.d.ts.map +0 -1
  614. package/dist/oracle.d.ts +0 -55
  615. package/dist/oracle.d.ts.map +0 -1
  616. package/dist/orthogonality.d.ts +0 -35
  617. package/dist/orthogonality.d.ts.map +0 -1
  618. package/dist/otel-pipeline.d.ts +0 -31
  619. package/dist/otel-pipeline.d.ts.map +0 -1
  620. package/dist/paraphrase.d.ts +0 -107
  621. package/dist/paraphrase.d.ts.map +0 -1
  622. package/dist/pareto.d.ts.map +0 -1
  623. package/dist/partition-held-out.d.ts +0 -70
  624. package/dist/partition-held-out.d.ts.map +0 -1
  625. package/dist/partition-held-out.test.d.ts +0 -2
  626. package/dist/partition-held-out.test.d.ts.map +0 -1
  627. package/dist/perf/index.d.ts.map +0 -1
  628. package/dist/perf/integrity.d.ts +0 -30
  629. package/dist/perf/integrity.d.ts.map +0 -1
  630. package/dist/perf/journey.d.ts +0 -45
  631. package/dist/perf/journey.d.ts.map +0 -1
  632. package/dist/perf/ratchet.d.ts +0 -47
  633. package/dist/perf/ratchet.d.ts.map +0 -1
  634. package/dist/pipelines/budget-breach.d.ts +0 -31
  635. package/dist/pipelines/budget-breach.d.ts.map +0 -1
  636. package/dist/pipelines/budget-breach.test.d.ts +0 -2
  637. package/dist/pipelines/budget-breach.test.d.ts.map +0 -1
  638. package/dist/pipelines/failure-cluster.d.ts +0 -38
  639. package/dist/pipelines/failure-cluster.d.ts.map +0 -1
  640. package/dist/pipelines/failure-cluster.test.d.ts +0 -2
  641. package/dist/pipelines/failure-cluster.test.d.ts.map +0 -1
  642. package/dist/pipelines/first-divergence.d.ts +0 -26
  643. package/dist/pipelines/first-divergence.d.ts.map +0 -1
  644. package/dist/pipelines/first-divergence.test.d.ts +0 -2
  645. package/dist/pipelines/first-divergence.test.d.ts.map +0 -1
  646. package/dist/pipelines/index.d.ts.map +0 -1
  647. package/dist/pipelines/judge-agreement.d.ts +0 -26
  648. package/dist/pipelines/judge-agreement.d.ts.map +0 -1
  649. package/dist/pipelines/judge-agreement.test.d.ts +0 -2
  650. package/dist/pipelines/judge-agreement.test.d.ts.map +0 -1
  651. package/dist/pipelines/regression.d.ts +0 -23
  652. package/dist/pipelines/regression.d.ts.map +0 -1
  653. package/dist/pipelines/regression.test.d.ts +0 -2
  654. package/dist/pipelines/regression.test.d.ts.map +0 -1
  655. package/dist/pipelines/stuck-loop.d.ts +0 -32
  656. package/dist/pipelines/stuck-loop.d.ts.map +0 -1
  657. package/dist/pipelines/stuck-loop.test.d.ts +0 -2
  658. package/dist/pipelines/stuck-loop.test.d.ts.map +0 -1
  659. package/dist/pipelines/tool-waste.d.ts +0 -34
  660. package/dist/pipelines/tool-waste.d.ts.map +0 -1
  661. package/dist/pipelines/tool-waste.test.d.ts +0 -2
  662. package/dist/pipelines/tool-waste.test.d.ts.map +0 -1
  663. package/dist/playbook.d.ts +0 -16
  664. package/dist/playbook.d.ts.map +0 -1
  665. package/dist/pr-review-benchmark.d.ts +0 -88
  666. package/dist/pr-review-benchmark.d.ts.map +0 -1
  667. package/dist/pr-review-benchmark.test.d.ts +0 -2
  668. package/dist/pr-review-benchmark.test.d.ts.map +0 -1
  669. package/dist/pre-registration.d.ts +0 -125
  670. package/dist/pre-registration.d.ts.map +0 -1
  671. package/dist/prm/builtin-rubrics.d.ts +0 -33
  672. package/dist/prm/builtin-rubrics.d.ts.map +0 -1
  673. package/dist/prm/index.d.ts.map +0 -1
  674. package/dist/prm/inference.d.ts +0 -29
  675. package/dist/prm/inference.d.ts.map +0 -1
  676. package/dist/prm/inference.test.d.ts +0 -2
  677. package/dist/prm/inference.test.d.ts.map +0 -1
  678. package/dist/prm/rubric.d.ts.map +0 -1
  679. package/dist/prm/training-export.d.ts +0 -38
  680. package/dist/prm/training-export.d.ts.map +0 -1
  681. package/dist/produced-state.d.ts +0 -63
  682. package/dist/produced-state.d.ts.map +0 -1
  683. package/dist/produced-state.test.d.ts +0 -8
  684. package/dist/produced-state.test.d.ts.map +0 -1
  685. package/dist/profile/baselines.d.ts +0 -37
  686. package/dist/profile/baselines.d.ts.map +0 -1
  687. package/dist/profile/index.d.ts +0 -105
  688. package/dist/profile/index.d.ts.map +0 -1
  689. package/dist/promotion-gate.d.ts +0 -93
  690. package/dist/promotion-gate.d.ts.map +0 -1
  691. package/dist/prompt-registry.d.ts +0 -41
  692. package/dist/prompt-registry.d.ts.map +0 -1
  693. package/dist/propose-review-control.d.ts +0 -49
  694. package/dist/propose-review-control.d.ts.map +0 -1
  695. package/dist/propose-review-control.test.d.ts +0 -2
  696. package/dist/propose-review-control.test.d.ts.map +0 -1
  697. package/dist/propose-review.d.ts +0 -155
  698. package/dist/propose-review.d.ts.map +0 -1
  699. package/dist/red-team.d.ts.map +0 -1
  700. package/dist/reference-replay-steering.d.ts +0 -11
  701. package/dist/reference-replay-steering.d.ts.map +0 -1
  702. package/dist/reference-replay.d.ts +0 -176
  703. package/dist/reference-replay.d.ts.map +0 -1
  704. package/dist/reflective-mutation.d.ts +0 -79
  705. package/dist/reflective-mutation.d.ts.map +0 -1
  706. package/dist/registry.d.ts +0 -31
  707. package/dist/registry.d.ts.map +0 -1
  708. package/dist/release-confidence.d.ts +0 -128
  709. package/dist/release-confidence.d.ts.map +0 -1
  710. package/dist/release-report.d.ts +0 -11
  711. package/dist/release-report.d.ts.map +0 -1
  712. package/dist/replay.d.ts +0 -120
  713. package/dist/replay.d.ts.map +0 -1
  714. package/dist/reporter.d.ts +0 -14
  715. package/dist/reporter.d.ts.map +0 -1
  716. package/dist/reporting.d.ts.map +0 -1
  717. package/dist/researcher.d.ts +0 -140
  718. package/dist/researcher.d.ts.map +0 -1
  719. package/dist/reviewer.d.ts +0 -118
  720. package/dist/reviewer.d.ts.map +0 -1
  721. package/dist/reviewer.test.d.ts +0 -2
  722. package/dist/reviewer.test.d.ts.map +0 -1
  723. package/dist/reward-model-export.d.ts +0 -60
  724. package/dist/reward-model-export.d.ts.map +0 -1
  725. package/dist/rl/active-curriculum.d.ts +0 -110
  726. package/dist/rl/active-curriculum.d.ts.map +0 -1
  727. package/dist/rl/adaptation-eval.d.ts +0 -109
  728. package/dist/rl/adaptation-eval.d.ts.map +0 -1
  729. package/dist/rl/adversarial.d.ts.map +0 -1
  730. package/dist/rl/compute-curves.d.ts +0 -127
  731. package/dist/rl/compute-curves.d.ts.map +0 -1
  732. package/dist/rl/contamination.d.ts +0 -117
  733. package/dist/rl/contamination.d.ts.map +0 -1
  734. package/dist/rl/corpus.d.ts +0 -55
  735. package/dist/rl/corpus.d.ts.map +0 -1
  736. package/dist/rl/corpus.test.d.ts +0 -2
  737. package/dist/rl/corpus.test.d.ts.map +0 -1
  738. package/dist/rl/dataset.d.ts +0 -102
  739. package/dist/rl/dataset.d.ts.map +0 -1
  740. package/dist/rl/dataset.test.d.ts +0 -2
  741. package/dist/rl/dataset.test.d.ts.map +0 -1
  742. package/dist/rl/exporters.d.ts +0 -141
  743. package/dist/rl/exporters.d.ts.map +0 -1
  744. package/dist/rl/index.d.ts +0 -49
  745. package/dist/rl/index.d.ts.map +0 -1
  746. package/dist/rl/off-policy.d.ts.map +0 -1
  747. package/dist/rl/predictive-validity-researcher.d.ts +0 -69
  748. package/dist/rl/predictive-validity-researcher.d.ts.map +0 -1
  749. package/dist/rl/preferences.d.ts +0 -141
  750. package/dist/rl/preferences.d.ts.map +0 -1
  751. package/dist/rl/process-reward.d.ts +0 -122
  752. package/dist/rl/process-reward.d.ts.map +0 -1
  753. package/dist/rl/reward-hacking.d.ts +0 -104
  754. package/dist/rl/reward-hacking.d.ts.map +0 -1
  755. package/dist/rl/rl-campaign.d.ts +0 -85
  756. package/dist/rl/rl-campaign.d.ts.map +0 -1
  757. package/dist/rl/run-record-adapters.d.ts +0 -56
  758. package/dist/rl/run-record-adapters.d.ts.map +0 -1
  759. package/dist/rl/sim-fidelity.d.ts +0 -166
  760. package/dist/rl/sim-fidelity.d.ts.map +0 -1
  761. package/dist/rl/sim-fidelity.test.d.ts +0 -2
  762. package/dist/rl/sim-fidelity.test.d.ts.map +0 -1
  763. package/dist/rl/tournament.d.ts +0 -115
  764. package/dist/rl/tournament.d.ts.map +0 -1
  765. package/dist/rl/verifiable-reward.d.ts +0 -124
  766. package/dist/rl/verifiable-reward.d.ts.map +0 -1
  767. package/dist/run-critic.d.ts +0 -23
  768. package/dist/run-critic.d.ts.map +0 -1
  769. package/dist/run-evidence.d.ts +0 -32
  770. package/dist/run-evidence.d.ts.map +0 -1
  771. package/dist/run-record.d.ts.map +0 -1
  772. package/dist/run-record.test.d.ts +0 -2
  773. package/dist/run-record.test.d.ts.map +0 -1
  774. package/dist/run-score.d.ts +0 -31
  775. package/dist/run-score.d.ts.map +0 -1
  776. package/dist/runtime-trajectory.d.ts +0 -47
  777. package/dist/runtime-trajectory.d.ts.map +0 -1
  778. package/dist/sandbox-harness.d.ts.map +0 -1
  779. package/dist/sandbox-harness.test.d.ts +0 -2
  780. package/dist/sandbox-harness.test.d.ts.map +0 -1
  781. package/dist/sandbox-pool.d.ts +0 -74
  782. package/dist/sandbox-pool.d.ts.map +0 -1
  783. package/dist/sandbox-pool.test.d.ts +0 -2
  784. package/dist/sandbox-pool.test.d.ts.map +0 -1
  785. package/dist/scorecard.d.ts +0 -133
  786. package/dist/scorecard.d.ts.map +0 -1
  787. package/dist/scorecard.test.d.ts +0 -2
  788. package/dist/scorecard.test.d.ts.map +0 -1
  789. package/dist/self-play.d.ts +0 -69
  790. package/dist/self-play.d.ts.map +0 -1
  791. package/dist/semantic-concept-judge.d.ts +0 -135
  792. package/dist/semantic-concept-judge.d.ts.map +0 -1
  793. package/dist/semantic-concept-judge.test.d.ts +0 -2
  794. package/dist/semantic-concept-judge.test.d.ts.map +0 -1
  795. package/dist/sequential.d.ts.map +0 -1
  796. package/dist/series-convergence.d.ts.map +0 -1
  797. package/dist/slo.d.ts +0 -48
  798. package/dist/slo.d.ts.map +0 -1
  799. package/dist/state-continuity.d.ts +0 -47
  800. package/dist/state-continuity.d.ts.map +0 -1
  801. package/dist/statistics.d.ts.map +0 -1
  802. package/dist/statistics.test.d.ts +0 -2
  803. package/dist/statistics.test.d.ts.map +0 -1
  804. package/dist/steering-optimizer.d.ts +0 -58
  805. package/dist/steering-optimizer.d.ts.map +0 -1
  806. package/dist/steering.d.ts +0 -24
  807. package/dist/steering.d.ts.map +0 -1
  808. package/dist/storyboard/code-edit.d.ts +0 -64
  809. package/dist/storyboard/code-edit.d.ts.map +0 -1
  810. package/dist/storyboard/code-edit.test.d.ts +0 -2
  811. package/dist/storyboard/code-edit.test.d.ts.map +0 -1
  812. package/dist/storyboard/index.d.ts.map +0 -1
  813. package/dist/storyboard/index.test.d.ts +0 -2
  814. package/dist/storyboard/index.test.d.ts.map +0 -1
  815. package/dist/summary-report.d.ts.map +0 -1
  816. package/dist/telemetry/client.d.ts +0 -35
  817. package/dist/telemetry/client.d.ts.map +0 -1
  818. package/dist/telemetry/index.d.ts.map +0 -1
  819. package/dist/telemetry/schema.d.ts +0 -61
  820. package/dist/telemetry/schema.d.ts.map +0 -1
  821. package/dist/telemetry/sink-fetch.d.ts +0 -39
  822. package/dist/telemetry/sink-fetch.d.ts.map +0 -1
  823. package/dist/telemetry/sink-file.d.ts.map +0 -1
  824. package/dist/test-graded-scenario.d.ts +0 -42
  825. package/dist/test-graded-scenario.d.ts.map +0 -1
  826. package/dist/testing.d.ts.map +0 -1
  827. package/dist/tool-use-metrics.d.ts +0 -35
  828. package/dist/tool-use-metrics.d.ts.map +0 -1
  829. package/dist/trace/capture-fetch.d.ts +0 -48
  830. package/dist/trace/capture-fetch.d.ts.map +0 -1
  831. package/dist/trace/capture-fetch.test.d.ts +0 -2
  832. package/dist/trace/capture-fetch.test.d.ts.map +0 -1
  833. package/dist/trace/emitter.d.ts.map +0 -1
  834. package/dist/trace/extract-usage.d.ts +0 -46
  835. package/dist/trace/extract-usage.d.ts.map +0 -1
  836. package/dist/trace/extract-usage.test.d.ts +0 -2
  837. package/dist/trace/extract-usage.test.d.ts.map +0 -1
  838. package/dist/trace/index.d.ts +0 -15
  839. package/dist/trace/index.d.ts.map +0 -1
  840. package/dist/trace/integrity.d.ts.map +0 -1
  841. package/dist/trace/otel-bridge.d.ts +0 -29
  842. package/dist/trace/otel-bridge.d.ts.map +0 -1
  843. package/dist/trace/otel-export.d.ts +0 -52
  844. package/dist/trace/otel-export.d.ts.map +0 -1
  845. package/dist/trace/otel.d.ts +0 -57
  846. package/dist/trace/otel.d.ts.map +0 -1
  847. package/dist/trace/otlp-attributes.d.ts +0 -17
  848. package/dist/trace/otlp-attributes.d.ts.map +0 -1
  849. package/dist/trace/query.d.ts +0 -29
  850. package/dist/trace/query.d.ts.map +0 -1
  851. package/dist/trace/query.test.d.ts +0 -2
  852. package/dist/trace/query.test.d.ts.map +0 -1
  853. package/dist/trace/raw-provider-sink.d.ts.map +0 -1
  854. package/dist/trace/redact.d.ts.map +0 -1
  855. package/dist/trace/schema.d.ts.map +0 -1
  856. package/dist/trace/store-to-otlp.d.ts +0 -72
  857. package/dist/trace/store-to-otlp.d.ts.map +0 -1
  858. package/dist/trace/store-to-otlp.test.d.ts +0 -2
  859. package/dist/trace/store-to-otlp.test.d.ts.map +0 -1
  860. package/dist/trace/store.d.ts.map +0 -1
  861. package/dist/trace/store.test.d.ts +0 -2
  862. package/dist/trace/store.test.d.ts.map +0 -1
  863. package/dist/trace-analyst/analyst.d.ts.map +0 -1
  864. package/dist/trace-analyst/analyst.test.d.ts +0 -2
  865. package/dist/trace-analyst/analyst.test.d.ts.map +0 -1
  866. package/dist/trace-analyst/behavioral-metrics.d.ts +0 -40
  867. package/dist/trace-analyst/behavioral-metrics.d.ts.map +0 -1
  868. package/dist/trace-analyst/behavioral-metrics.test.d.ts +0 -2
  869. package/dist/trace-analyst/behavioral-metrics.test.d.ts.map +0 -1
  870. package/dist/trace-analyst/hook.d.ts +0 -55
  871. package/dist/trace-analyst/hook.d.ts.map +0 -1
  872. package/dist/trace-analyst/index.d.ts +0 -18
  873. package/dist/trace-analyst/index.d.ts.map +0 -1
  874. package/dist/trace-analyst/insights.d.ts +0 -71
  875. package/dist/trace-analyst/insights.d.ts.map +0 -1
  876. package/dist/trace-analyst/insights.test.d.ts +0 -2
  877. package/dist/trace-analyst/insights.test.d.ts.map +0 -1
  878. package/dist/trace-analyst/otlp-flatten.d.ts +0 -42
  879. package/dist/trace-analyst/otlp-flatten.d.ts.map +0 -1
  880. package/dist/trace-analyst/otlp-span.d.ts +0 -85
  881. package/dist/trace-analyst/otlp-span.d.ts.map +0 -1
  882. package/dist/trace-analyst/otlp-span.test.d.ts +0 -2
  883. package/dist/trace-analyst/otlp-span.test.d.ts.map +0 -1
  884. package/dist/trace-analyst/otlp-to-run-records.d.ts +0 -115
  885. package/dist/trace-analyst/otlp-to-run-records.d.ts.map +0 -1
  886. package/dist/trace-analyst/otlp-to-run-records.test.d.ts +0 -2
  887. package/dist/trace-analyst/otlp-to-run-records.test.d.ts.map +0 -1
  888. package/dist/trace-analyst/otlp-to-run-records.timestamps.test.d.ts +0 -2
  889. package/dist/trace-analyst/otlp-to-run-records.timestamps.test.d.ts.map +0 -1
  890. package/dist/trace-analyst/prompts.d.ts +0 -6
  891. package/dist/trace-analyst/prompts.d.ts.map +0 -1
  892. package/dist/trace-analyst/store-otlp.d.ts +0 -126
  893. package/dist/trace-analyst/store-otlp.d.ts.map +0 -1
  894. package/dist/trace-analyst/store-otlp.test.d.ts +0 -8
  895. package/dist/trace-analyst/store-otlp.test.d.ts.map +0 -1
  896. package/dist/trace-analyst/store-otlp.timestamps.test.d.ts +0 -2
  897. package/dist/trace-analyst/store-otlp.timestamps.test.d.ts.map +0 -1
  898. package/dist/trace-analyst/store.d.ts +0 -63
  899. package/dist/trace-analyst/store.d.ts.map +0 -1
  900. package/dist/trace-analyst/tools.d.ts +0 -44
  901. package/dist/trace-analyst/tools.d.ts.map +0 -1
  902. package/dist/trace-analyst/tools.test.d.ts +0 -10
  903. package/dist/trace-analyst/tools.test.d.ts.map +0 -1
  904. package/dist/trace-analyst/types.d.ts.map +0 -1
  905. package/dist/trace-contracts.d.ts +0 -180
  906. package/dist/trace-contracts.d.ts.map +0 -1
  907. package/dist/traced-analyst.d.ts +0 -26
  908. package/dist/traced-analyst.d.ts.map +0 -1
  909. package/dist/traced-judges.d.ts +0 -27
  910. package/dist/traced-judges.d.ts.map +0 -1
  911. package/dist/traces.d.ts.map +0 -1
  912. package/dist/trajectory.d.ts.map +0 -1
  913. package/dist/types.d.ts.map +0 -1
  914. package/dist/ui-finding.d.ts +0 -104
  915. package/dist/ui-finding.d.ts.map +0 -1
  916. package/dist/verdict-cache.d.ts +0 -78
  917. package/dist/verdict-cache.d.ts.map +0 -1
  918. package/dist/verdict-cache.test.d.ts +0 -2
  919. package/dist/verdict-cache.test.d.ts.map +0 -1
  920. package/dist/verdict.d.ts.map +0 -1
  921. package/dist/visual-diff.d.ts +0 -32
  922. package/dist/visual-diff.d.ts.map +0 -1
  923. package/dist/wire/handlers.d.ts +0 -54
  924. package/dist/wire/handlers.d.ts.map +0 -1
  925. package/dist/wire/index.d.ts.map +0 -1
  926. package/dist/wire/openapi.d.ts +0 -3
  927. package/dist/wire/openapi.d.ts.map +0 -1
  928. package/dist/wire/rpc.d.ts +0 -21
  929. package/dist/wire/rpc.d.ts.map +0 -1
  930. package/dist/wire/rubrics.d.ts +0 -34
  931. package/dist/wire/rubrics.d.ts.map +0 -1
  932. package/dist/wire/schemas.d.ts +0 -410
  933. package/dist/wire/schemas.d.ts.map +0 -1
  934. package/dist/wire/server.d.ts +0 -60
  935. package/dist/wire/server.d.ts.map +0 -1
  936. package/dist/workflow/event-schema.d.ts +0 -5
  937. package/dist/workflow/event-schema.d.ts.map +0 -1
  938. package/dist/workflow/feedback-pack.d.ts +0 -99
  939. package/dist/workflow/feedback-pack.d.ts.map +0 -1
  940. package/dist/workflow/index.d.ts.map +0 -1
  941. package/dist/workflow/intelligence-export.d.ts +0 -62
  942. package/dist/workflow/intelligence-export.d.ts.map +0 -1
  943. package/dist/workflow/partner-report.d.ts +0 -49
  944. package/dist/workflow/partner-report.d.ts.map +0 -1
  945. package/dist/workflow/phase-graph.d.ts +0 -43
  946. package/dist/workflow/phase-graph.d.ts.map +0 -1
  947. package/dist/workflow/promotion-gate.d.ts +0 -61
  948. package/dist/workflow/promotion-gate.d.ts.map +0 -1
  949. package/dist/workflow/run-record.d.ts +0 -12
  950. package/dist/workflow/run-record.d.ts.map +0 -1
  951. package/dist/workflow/runtime-adapter.d.ts +0 -20
  952. package/dist/workflow/runtime-adapter.d.ts.map +0 -1
  953. package/dist/workflow/sanitize.d.ts +0 -21
  954. package/dist/workflow/sanitize.d.ts.map +0 -1
  955. package/dist/workflow/schema.d.ts +0 -5
  956. package/dist/workflow/schema.d.ts.map +0 -1
  957. package/dist/workflow/summary.d.ts +0 -43
  958. package/dist/workflow/summary.d.ts.map +0 -1
  959. package/dist/workflow/trace-event-fields.d.ts +0 -6
  960. package/dist/workflow/trace-event-fields.d.ts.map +0 -1
  961. package/dist/workflow/trajectory.d.ts +0 -15
  962. package/dist/workflow/trajectory.d.ts.map +0 -1
  963. package/dist/workflow/types.d.ts +0 -68
  964. package/dist/workflow/types.d.ts.map +0 -1
  965. package/dist/workspace-inspector.d.ts +0 -67
  966. package/dist/workspace-inspector.d.ts.map +0 -1
  967. package/dist/wrangler-deploy-runner.test.d.ts +0 -2
  968. package/dist/wrangler-deploy-runner.test.d.ts.map +0 -1
package/dist/index.js CHANGED
@@ -132,7 +132,7 @@ import {
132
132
  import {
133
133
  buildDefaultAnalystRegistry,
134
134
  computeTraceMetrics
135
- } from "./chunk-AQ5WQAIV.js";
135
+ } from "./chunk-3NHEO6ZC.js";
136
136
  import {
137
137
  LockedJsonlAppender,
138
138
  Mutex
@@ -1,3 +1,6 @@
1
+ import { G as GainDistributionBin, P as ParetoFigureSpec } from './summary-report-CInXwsza.js';
2
+ import { C as ContinuousAgreement } from './judge-calibration-0p2QcWNE.js';
3
+
1
4
  /**
2
5
  * # InsightReport — the rigorous decision packet for any set of agent runs.
3
6
  *
@@ -27,9 +30,8 @@
27
30
  * Consumers read the `recommendations` array first — that's the
28
31
  * actionable layer, ranked by priority. The numeric sections back it up.
29
32
  */
30
- import type { GainDistributionBin, ParetoFigureSpec } from '../summary-report';
31
- import type { ContinuousAgreement } from './insight-types-fwd';
32
- export interface InsightReport {
33
+
34
+ interface InsightReport {
33
35
  /** Number of runs analyzed. */
34
36
  n: number;
35
37
  /** Composite-score distribution across all runs. Always present. */
@@ -94,7 +96,7 @@ export interface InsightReport {
94
96
  recommendations: Recommendation[];
95
97
  }
96
98
  /** Distributional summary of a scalar-valued metric. */
97
- export interface ScalarDistribution {
99
+ interface ScalarDistribution {
98
100
  /** Sample count after dropping non-finite values. */
99
101
  n: number;
100
102
  mean: number;
@@ -114,7 +116,7 @@ export interface ScalarDistribution {
114
116
  score: number;
115
117
  }>;
116
118
  }
117
- export interface JudgeInsight {
119
+ interface JudgeInsight {
118
120
  /** Number of times this judge scored a run. */
119
121
  n: number;
120
122
  /** Mean composite over this judge's runs. */
@@ -132,7 +134,7 @@ export interface JudgeInsight {
132
134
  * quality? */
133
135
  verbosityBias?: number;
134
136
  }
135
- export interface InterRaterInsight {
137
+ interface InterRaterInsight {
136
138
  /** Number of raters whose scores were aggregated. */
137
139
  raters: number;
138
140
  /** Number of runs every rater scored. */
@@ -151,7 +153,7 @@ export interface InterRaterInsight {
151
153
  range: number;
152
154
  }>;
153
155
  }
154
- export interface LiftInsight {
156
+ interface LiftInsight {
155
157
  baselineMean: number;
156
158
  candidateMean: number;
157
159
  /** Candidate − baseline. */
@@ -169,7 +171,7 @@ export interface LiftInsight {
169
171
  /** Sample size needed to detect the observed delta at 80% power. */
170
172
  requiredN: number;
171
173
  }
172
- export interface FailureClusterInsight {
174
+ interface FailureClusterInsight {
173
175
  /** All clusters identified by the registry, ranked by share descending. */
174
176
  clusters: Array<{
175
177
  id: string;
@@ -188,7 +190,7 @@ export interface FailureClusterInsight {
188
190
  * is computed directly from the tags the harness already recorded — so a
189
191
  * customer ingesting one batch with no judge/analyst still learns which
190
192
  * named failure dominates. */
191
- export interface FailureModeTally {
193
+ interface FailureModeTally {
192
194
  /** The `failureMode` tag. */
193
195
  mode: string;
194
196
  /** Number of runs carrying this tag. */
@@ -196,7 +198,7 @@ export interface FailureModeTally {
196
198
  /** Share of the whole corpus, 0..1. */
197
199
  share: number;
198
200
  }
199
- export interface ContaminationInsight {
201
+ interface ContaminationInsight {
200
202
  /** Canary phrases that leaked into outputs. */
201
203
  leaks: number;
202
204
  /** Holdout audit verdict — did any holdout-tagged run end up in the
@@ -208,7 +210,7 @@ export interface ContaminationInsight {
208
210
  matched: string;
209
211
  }>;
210
212
  }
211
- export interface OutcomeCorrelationInsight {
213
+ interface OutcomeCorrelationInsight {
212
214
  /** What outcome the consumer is correlating against (e.g.
213
215
  * `'engagement_rate'`, `'approval_rate'`, `'downstream_pass'`). */
214
216
  metric: string;
@@ -225,7 +227,7 @@ export interface OutcomeCorrelationInsight {
225
227
  r2: number;
226
228
  };
227
229
  }
228
- export interface ReleaseSummary {
230
+ interface ReleaseSummary {
229
231
  /** Overall verdict across axes — fail if any axis fails, else warn if any
230
232
  * warns, else pass. */
231
233
  status: 'pass' | 'warn' | 'fail';
@@ -238,7 +240,7 @@ export interface ReleaseSummary {
238
240
  * consumers can post-process to populate. */
239
241
  issues: string[];
240
242
  }
241
- export interface MetricDelta {
243
+ interface MetricDelta {
242
244
  /** Current-period mean. */
243
245
  current: number;
244
246
  /** Baseline-period mean. */
@@ -262,7 +264,7 @@ export interface MetricDelta {
262
264
  * tiny from triggering recommendations. */
263
265
  significant: boolean;
264
266
  }
265
- export interface PriorPeriodComparison {
267
+ interface PriorPeriodComparison {
266
268
  /** Sample counts. */
267
269
  baselineN: number;
268
270
  currentN: number;
@@ -278,7 +280,7 @@ export interface PriorPeriodComparison {
278
280
  /** Metric names where current is significantly BETTER than baseline. */
279
281
  improvedMetrics: string[];
280
282
  }
281
- export interface Recommendation {
283
+ interface Recommendation {
282
284
  priority: 'critical' | 'high' | 'medium' | 'low';
283
285
  kind: 'ship' | 'hold' | 'investigate' | 'fix' | 'recalibrate' | 'expand-corpus';
284
286
  title: string;
@@ -286,4 +288,5 @@ export interface Recommendation {
286
288
  /** Optional pointer back into the report for the evidence. */
287
289
  evidencePath?: string;
288
290
  }
289
- //# sourceMappingURL=insight-report.d.ts.map
291
+
292
+ export type { FailureClusterInsight as F, InsightReport as I, JudgeInsight as J, LiftInsight as L, OutcomeCorrelationInsight as O, Recommendation as R, ScalarDistribution as S, InterRaterInsight as a, ReleaseSummary as b };
@@ -1,3 +1,7 @@
1
+ import { C as CaptureIntegrityError } from './errors-CzMUYo7b.js';
2
+ import { R as RawProviderSink } from './raw-provider-sink-C46HDghv.js';
3
+ import { T as TraceStore } from './store-BcFXE6LG.js';
4
+
1
5
  /**
2
6
  * Run-completion integrity check — at end of run, verify the expected event
3
7
  * types were actually captured. The point is the launch-review failure mode:
@@ -18,10 +22,8 @@
18
22
  * the caller chooses the failure mode (throw, mark run failed, log warning).
19
23
  * `throwIfRunIncomplete` is the convenient strict mode.
20
24
  */
21
- import { CaptureIntegrityError } from '../errors';
22
- import type { RawProviderSink } from './raw-provider-sink';
23
- import type { TraceStore } from './store';
24
- export interface RunIntegrityExpectations {
25
+
26
+ interface RunIntegrityExpectations {
25
27
  /** Minimum LLM span count. Default 0 (no requirement). */
26
28
  llmSpansMin?: number;
27
29
  /** Minimum judge span count. Default 0. */
@@ -44,13 +46,13 @@ export interface RunIntegrityExpectations {
44
46
  /** Run outcome must be set (not null/undefined). Default false. */
45
47
  requireOutcome?: boolean;
46
48
  }
47
- export type RunIntegrityIssueCode = 'no_run' | 'missing_llm_spans' | 'missing_judge_spans' | 'missing_tool_spans' | 'missing_raw_events' | 'no_raw_sink' | 'orphan_llm_span' | 'missing_outcome';
48
- export interface RunIntegrityIssue {
49
+ type RunIntegrityIssueCode = 'no_run' | 'missing_llm_spans' | 'missing_judge_spans' | 'missing_tool_spans' | 'missing_raw_events' | 'no_raw_sink' | 'orphan_llm_span' | 'missing_outcome';
50
+ interface RunIntegrityIssue {
49
51
  code: RunIntegrityIssueCode;
50
52
  message: string;
51
53
  detail?: Record<string, unknown>;
52
54
  }
53
- export interface RunIntegrityReport {
55
+ interface RunIntegrityReport {
54
56
  ok: boolean;
55
57
  runId: string;
56
58
  llmSpanCount: number;
@@ -68,11 +70,12 @@ export interface RunIntegrityReport {
68
70
  };
69
71
  issues: RunIntegrityIssue[];
70
72
  }
71
- export declare class RunIntegrityError extends CaptureIntegrityError {
73
+ declare class RunIntegrityError extends CaptureIntegrityError {
72
74
  readonly report: RunIntegrityReport;
73
75
  constructor(report: RunIntegrityReport);
74
76
  }
75
- export declare function assertRunCaptured(store: TraceStore, runId: string, expectations?: RunIntegrityExpectations): Promise<RunIntegrityReport>;
77
+ declare function assertRunCaptured(store: TraceStore, runId: string, expectations?: RunIntegrityExpectations): Promise<RunIntegrityReport>;
76
78
  /** Strict mode: throws `RunIntegrityError` when the report isn't ok. */
77
- export declare function throwIfRunIncomplete(report: RunIntegrityReport): void;
78
- //# sourceMappingURL=integrity.d.ts.map
79
+ declare function throwIfRunIncomplete(report: RunIntegrityReport): void;
80
+
81
+ export { type RunIntegrityExpectations as R, type RunIntegrityReport as a, RunIntegrityError as b, type RunIntegrityIssue as c, type RunIntegrityIssueCode as d, assertRunCaptured as e, throwIfRunIncomplete as t };
@@ -18,19 +18,19 @@
18
18
  * Returns actionable diagnostics, not a single number. Consumers then
19
19
  * decide whether to trust the judge, retrain it, or add a tie-breaker.
20
20
  */
21
- export interface GoldenItem {
21
+ interface GoldenItem {
22
22
  itemId: string;
23
23
  humanScore: number;
24
24
  /** Optional group used for per-group bias audits (e.g. model-of-output family). */
25
25
  group?: string;
26
26
  }
27
- export interface CandidateScore {
27
+ interface CandidateScore {
28
28
  itemId: string;
29
29
  score: number;
30
30
  /** Optional — enables positional-bias analysis (did order matter?). */
31
31
  positionOfAInput?: 'first' | 'second';
32
32
  }
33
- export interface CalibrationResult {
33
+ interface CalibrationResult {
34
34
  n: number;
35
35
  pearson: number;
36
36
  /** Cohen's κ with quadratic weights over integer-rounded scores. */
@@ -45,8 +45,8 @@ export interface CalibrationResult {
45
45
  delta: number;
46
46
  }>;
47
47
  }
48
- export declare function calibrateJudge(golden: GoldenItem[], candidate: CandidateScore[]): CalibrationResult;
49
- export interface PositionalBiasResult {
48
+ declare function calibrateJudge(golden: GoldenItem[], candidate: CandidateScore[]): CalibrationResult;
49
+ interface PositionalBiasResult {
50
50
  /**
51
51
  * Score delta (first-position - second-position) averaged across items
52
52
  * presented in both positions. Non-zero = positional bias.
@@ -58,17 +58,17 @@ export interface PositionalBiasResult {
58
58
  * Feed the same items to the judge twice with A/B swapped and pass all
59
59
  * results here. Items that don't appear in both positions are ignored.
60
60
  */
61
- export declare function positionalBias(scores: CandidateScore[]): PositionalBiasResult;
62
- export interface VerbosityBiasResult {
61
+ declare function positionalBias(scores: CandidateScore[]): PositionalBiasResult;
62
+ interface VerbosityBiasResult {
63
63
  /** Pearson correlation between output length and score. Strong positive = verbosity bias. */
64
64
  pearson: number;
65
65
  n: number;
66
66
  }
67
- export declare function verbosityBias(samples: Array<{
67
+ declare function verbosityBias(samples: Array<{
68
68
  outputLen: number;
69
69
  score: number;
70
70
  }>): VerbosityBiasResult;
71
- export interface SelfPreferenceResult {
71
+ interface SelfPreferenceResult {
72
72
  /** Mean judge score when judge's family matches output's family. */
73
73
  inFamilyMean: number;
74
74
  outOfFamilyMean: number;
@@ -80,11 +80,11 @@ export interface SelfPreferenceResult {
80
80
  * model X (in-family) and model Y (out-of-family). Non-zero delta
81
81
  * indicates self-preference.
82
82
  */
83
- export declare function selfPreference(samples: Array<{
83
+ declare function selfPreference(samples: Array<{
84
84
  score: number;
85
85
  inFamily: boolean;
86
86
  }>): SelfPreferenceResult;
87
- export interface ContinuousAgreement {
87
+ interface ContinuousAgreement {
88
88
  /** Cohen's κ_w with quadratic weights, computed on raw [0,1] scores. */
89
89
  weightedKappa: number;
90
90
  /** ICC(2,1): two-way random effects, absolute agreement, single rater. */
@@ -103,7 +103,7 @@ export interface ContinuousAgreement {
103
103
  /** Number of raters. */
104
104
  raters: number;
105
105
  }
106
- export interface ContinuousAgreementOptions {
106
+ interface ContinuousAgreementOptions {
107
107
  /** Bootstrap iterations. Default 1000. Set to 0 to skip CIs (CI = [NaN, NaN]). */
108
108
  bootstrap?: number;
109
109
  /** κ weighting scheme. Default 'quadratic'. */
@@ -120,8 +120,8 @@ export interface ContinuousAgreementOptions {
120
120
  * are dropped. Returns NaN metrics if fewer than 2 raters or 2 complete
121
121
  * items remain.
122
122
  */
123
- export declare function continuousAgreement(scores: number[][], opts?: ContinuousAgreementOptions): ContinuousAgreement;
124
- export interface ContinuousCalibrationResult extends CalibrationResult {
123
+ declare function continuousAgreement(scores: number[][], opts?: ContinuousAgreementOptions): ContinuousAgreement;
124
+ interface ContinuousCalibrationResult extends CalibrationResult {
125
125
  /** Cohen's κ_w computed on raw (un-rounded) scores. */
126
126
  weightedKappaContinuous: number;
127
127
  /** ICC(2,1) treating golden + candidate as two raters. */
@@ -137,5 +137,6 @@ export interface ContinuousCalibrationResult extends CalibrationResult {
137
137
  * agreement metrics. The old fields (n, pearson, kappa, mae, worstItems)
138
138
  * are preserved unchanged so existing callers continue to work.
139
139
  */
140
- export declare function calibrateJudgeContinuous(golden: GoldenItem[], candidate: CandidateScore[], opts?: ContinuousAgreementOptions): ContinuousCalibrationResult;
141
- //# sourceMappingURL=judge-calibration.d.ts.map
140
+ declare function calibrateJudgeContinuous(golden: GoldenItem[], candidate: CandidateScore[], opts?: ContinuousAgreementOptions): ContinuousCalibrationResult;
141
+
142
+ export { type ContinuousAgreement as C, type GoldenItem as G, type PositionalBiasResult as P, type SelfPreferenceResult as S, type VerbosityBiasResult as V, type CalibrationResult as a, type ContinuousCalibrationResult as b, type CandidateScore as c, type ContinuousAgreementOptions as d, calibrateJudge as e, calibrateJudgeContinuous as f, continuousAgreement as g, positionalBias as p, selfPreference as s, verbosityBias as v };
@@ -1,3 +1,56 @@
1
+ import { AxAIService, AxFunction } from '@ax-llm/ax';
2
+ import { T as TraceAnalysisStore } from './store-C1YxJDEK.js';
3
+ import { z } from 'zod';
4
+ import { g as AnalystCost, a as AnalystContext, A as Analyst } from './types-B5x54y6n.js';
5
+
6
+ /**
7
+ * Typed Ax output for analyst findings.
8
+ *
9
+ * Replaces the legacy `findings:string[]` pattern (where every bullet
10
+ * became a flat-severity `AnalystFinding`) with a structured object
11
+ * array. Ax binds the field as `findings:json[]` so the provider emits
12
+ * native structured output; at the kind-factory boundary we Zod-validate
13
+ * each emitted finding so malformed rows fail loud instead of being
14
+ * silently lifted with default severity.
15
+ *
16
+ * Why not `f.object().array()` directly in the signature? The Ax
17
+ * signature string `question:string -> findings:json[]` already lets
18
+ * the provider emit JSON arrays. A Zod boundary is required either
19
+ * way (the provider can return any JSON), and Zod gives us a single
20
+ * validation surface independent of which Ax version is installed.
21
+ */
22
+
23
+ declare const ANALYST_SEVERITIES: readonly ["critical", "high", "medium", "low", "info"];
24
+ declare const RawAnalystFindingSchema: z.ZodObject<{
25
+ severity: z.ZodEnum<{
26
+ info: "info";
27
+ critical: "critical";
28
+ medium: "medium";
29
+ low: "low";
30
+ high: "high";
31
+ }>;
32
+ claim: z.ZodString;
33
+ subject: z.ZodOptional<z.ZodString>;
34
+ evidence_uri: z.ZodString;
35
+ evidence_excerpt: z.ZodOptional<z.ZodString>;
36
+ confidence: z.ZodNumber;
37
+ rationale: z.ZodOptional<z.ZodString>;
38
+ recommended_action: z.ZodOptional<z.ZodString>;
39
+ }, z.core.$strict>;
40
+ type RawAnalystFinding = z.infer<typeof RawAnalystFindingSchema>;
41
+ /**
42
+ * Description embedded into the actor prompt so the LLM knows what
43
+ * shape to emit. Kept here so kinds share one source of truth rather
44
+ * than restating the schema in every prompt.
45
+ */
46
+ declare const RAW_FINDING_SCHEMA_PROMPT = "Each finding MUST be a JSON object with these fields:\n - severity: one of \"critical\" | \"high\" | \"medium\" | \"low\" | \"info\"\n - claim: one-sentence statement (max 2000 chars)\n - subject?: the routing locus this finding is about. It MUST be one of the exact subject forms listed in this kind's instructions above (e.g. `system-prompt:<section>`, `agent-knowledge:wiki:<slug>`, `tool-doc:<tool>`). A free phrase, a bare noun, or any form not in that list is REJECTED at parse time and the finding is discarded \u2014 omit subject entirely rather than guess a form.\n - evidence_uri: REQUIRED, never blank. Exactly one of \"span://<trace_id>/<span_id>\" (trace evidence), \"artifact://<relative-path>\" (files), \"metric://<name>\" (named scalars) \u2014 ALWAYS cite a real id surfaced by the tools. If you have no citable id, do not emit the finding.\n - evidence_excerpt?: short quote (<=2000 chars) from the cited span/artifact\n - confidence: number 0..1 \u2014 0.9+ when backed by exact quotes, 0.6-0.8 for inferred patterns, <0.5 for speculative\n - rationale?: one or two sentences explaining the reasoning\n - recommended_action?: concrete change phrased as an imperative (\"Add ...\", \"Replace ...\", \"Stop ...\") \u2014 omit when the finding is purely descriptive\n\nEmit an empty array when the question has no findings to report. Do not fabricate evidence.";
47
+ /**
48
+ * Validate one row emitted by the LLM. Returns the typed finding on
49
+ * success; returns `null` and logs the reason on failure so the kind
50
+ * factory can skip-and-count rather than abort the whole analyst run.
51
+ */
52
+ declare function parseRawFinding(row: unknown, log?: (msg: string, fields?: Record<string, unknown>) => void): RawAnalystFinding | null;
53
+
1
54
  /**
2
55
  * Analyst-kind factory — the typed way to define trace analysts.
3
56
  *
@@ -23,15 +76,12 @@
23
76
  * description programmatically. Stored on the kind, not the registry,
24
77
  * because the right metric is kind-specific.
25
78
  */
26
- import type { AxAIService, AxFunction } from '@ax-llm/ax';
27
- import type { TraceAnalysisStore } from '../trace-analyst/store';
28
- import { type RawAnalystFinding } from './finding-signature';
29
- import type { Analyst, AnalystContext, AnalystCost } from './types';
79
+
30
80
  /**
31
81
  * Per-kind specification. The factory turns this into a regular
32
82
  * `Analyst<TraceAnalysisStore>` ready for `AnalystRegistry.register()`.
33
83
  */
34
- export interface TraceAnalystKindSpec {
84
+ interface TraceAnalystKindSpec {
35
85
  /** Stable id. Appears in finding_id, telemetry, and registry exclusions. */
36
86
  id: string;
37
87
  /** One-sentence description shown in `registry.list()`. */
@@ -68,11 +118,11 @@ export interface TraceAnalystKindSpec {
68
118
  * is the ground-truth finding set a fitted prompt should produce on this
69
119
  * input. Metric: kind-specific (default: F1 on `finding_id` overlap).
70
120
  */
71
- export interface TraceAnalystGolden {
121
+ interface TraceAnalystGolden {
72
122
  question: string;
73
123
  expected: ReadonlyArray<Omit<RawAnalystFinding, 'confidence'>>;
74
124
  }
75
- export interface CreateTraceAnalystKindOpts {
125
+ interface CreateTraceAnalystKindOpts {
76
126
  /** AxAIService bound at registration time. */
77
127
  ai: AxAIService;
78
128
  /** Optional model override; falls back to the AI service's default. */
@@ -101,7 +151,7 @@ export interface CreateTraceAnalystKindOpts {
101
151
  * `analyze()` call (the agent carries chat-log + usage state we don't
102
152
  * want shared across analyst runs).
103
153
  */
104
- export declare function createTraceAnalystKind(spec: TraceAnalystKindSpec, opts: CreateTraceAnalystKindOpts): Analyst<TraceAnalysisStore>;
154
+ declare function createTraceAnalystKind(spec: TraceAnalystKindSpec, opts: CreateTraceAnalystKindOpts): Analyst<TraceAnalysisStore>;
105
155
  /**
106
156
  * Render a compact prior-findings block the actor reads alongside its
107
157
  * brief. Each row is one line so the actor can scan dozens cheaply.
@@ -116,5 +166,6 @@ export declare function createTraceAnalystKind(spec: TraceAnalystKindSpec, opts:
116
166
  * Exported for tests + for consumers that build their own actor
117
167
  * prompts (e.g. specialized analysts living outside the default kinds).
118
168
  */
119
- export declare function renderPriorFindings(prior: AnalystContext['priorFindings']): string;
120
- //# sourceMappingURL=kind-factory.d.ts.map
169
+ declare function renderPriorFindings(prior: AnalystContext['priorFindings']): string;
170
+
171
+ export { ANALYST_SEVERITIES as A, type CreateTraceAnalystKindOpts as C, RAW_FINDING_SCHEMA_PROMPT as R, type TraceAnalystKindSpec as T, type RawAnalystFinding as a, RawAnalystFindingSchema as b, type TraceAnalystGolden as c, createTraceAnalystKind as d, parseRawFinding as p, renderPriorFindings as r };
@@ -1,3 +1,103 @@
1
- export * from './readiness';
2
- export * from './types';
3
- //# sourceMappingURL=index.d.ts.map
1
+ import { j as ControlSeverity, C as ControlEvalResult } from '../control-runtime-Acf9CGhw.js';
2
+ import { T as TraceEmitter } from '../emitter-C2rqGH_l.js';
3
+ import '../schema-m0gsnbt3.js';
4
+ import '../store-BcFXE6LG.js';
5
+
6
+ type KnowledgeRequirementCategory = 'user_specific' | 'company_specific' | 'domain_specific' | 'codebase_specific' | 'market_specific' | 'regulatory' | 'tool_api' | 'credential_or_secret' | 'runtime_environment' | 'preference' | 'historical_context';
7
+ type KnowledgeAcquisitionMode = 'ask_user' | 'search_web' | 'query_connector' | 'inspect_repo' | 'run_command' | 'infer_low_confidence' | 'not_available';
8
+ type KnowledgeImportance = 'blocking' | 'high' | 'medium' | 'low';
9
+ type KnowledgeFreshness = 'static' | 'monthly' | 'weekly' | 'daily' | 'realtime';
10
+ type KnowledgeSensitivity = 'public' | 'private' | 'secret';
11
+ type KnowledgeFallbackPolicy = 'block' | 'ask' | 'continue_with_caveat' | 'use_default';
12
+ interface KnowledgeRequirement {
13
+ id: string;
14
+ description: string;
15
+ requiredFor: string[];
16
+ category: KnowledgeRequirementCategory;
17
+ acquisitionMode: KnowledgeAcquisitionMode;
18
+ importance: KnowledgeImportance;
19
+ freshness: KnowledgeFreshness;
20
+ sensitivity: KnowledgeSensitivity;
21
+ confidenceNeeded: number;
22
+ currentConfidence: number;
23
+ evidenceIds: string[];
24
+ fallbackPolicy: KnowledgeFallbackPolicy;
25
+ /**
26
+ * ISO timestamp after which this requirement must be treated as stale.
27
+ * Stale requirements score as missing even when they still have evidence.
28
+ */
29
+ validUntil?: string;
30
+ /** ISO timestamp for the last source-grounding or human verification pass. */
31
+ lastVerifiedAt?: string;
32
+ metadata?: Record<string, unknown>;
33
+ }
34
+ interface KnowledgeBundle {
35
+ taskId: string;
36
+ requirements: KnowledgeRequirement[];
37
+ evidenceIds: string[];
38
+ claimIds: string[];
39
+ wikiPageIds: string[];
40
+ userAnswers: Record<string, string>;
41
+ missing: KnowledgeRequirement[];
42
+ readinessScore: number;
43
+ metadata?: Record<string, unknown>;
44
+ }
45
+ type KnowledgeRecommendedAction = 'run_agent' | 'ask_user' | 'collect_web_data' | 'query_connectors' | 'inspect_repo' | 'build_domain_wiki' | 'continue_with_caveat' | 'abort_or_rescope';
46
+ interface KnowledgeReadinessReport {
47
+ taskId: string;
48
+ readinessScore: number;
49
+ blockingMissingRequirements: KnowledgeRequirement[];
50
+ nonBlockingGaps: KnowledgeRequirement[];
51
+ recommendedAction: KnowledgeRecommendedAction;
52
+ bundle: KnowledgeBundle;
53
+ severity: ControlSeverity;
54
+ reason: string;
55
+ }
56
+ interface UserQuestion {
57
+ id: string;
58
+ question: string;
59
+ reason: string;
60
+ requirementId: string;
61
+ importance: KnowledgeImportance;
62
+ answerType: 'free_text' | 'select_one' | 'multi_select' | 'file_upload' | 'credential' | 'url';
63
+ defaultIfSkipped?: string;
64
+ impactIfUnknown: string;
65
+ options?: string[];
66
+ metadata?: Record<string, unknown>;
67
+ }
68
+ interface DataAcquisitionPlan {
69
+ id: string;
70
+ requirementIds: string[];
71
+ mode: Exclude<KnowledgeAcquisitionMode, 'not_available' | 'infer_low_confidence'> | 'build_domain_wiki';
72
+ description: string;
73
+ priority: KnowledgeImportance;
74
+ expectedEvidenceIds?: string[];
75
+ questions?: UserQuestion[];
76
+ metadata?: Record<string, unknown>;
77
+ }
78
+ type KnowledgeResponsibleSurface = 'knowledge-requirements' | 'data-acquisition' | 'retrieval-policy' | 'user-question-policy';
79
+
80
+ interface ScoreKnowledgeReadinessOptions {
81
+ taskId: string;
82
+ requirements: KnowledgeRequirement[];
83
+ evidenceIds?: string[];
84
+ claimIds?: string[];
85
+ wikiPageIds?: string[];
86
+ userAnswers?: Record<string, string>;
87
+ metadata?: Record<string, unknown>;
88
+ now?: Date;
89
+ }
90
+ declare function scoreKnowledgeReadiness(options: ScoreKnowledgeReadinessOptions): KnowledgeReadinessReport;
91
+ declare function blockingKnowledgeEval(report: KnowledgeReadinessReport, options?: {
92
+ id?: string;
93
+ minimumScore?: number;
94
+ emitter?: TraceEmitter;
95
+ }): ControlEvalResult;
96
+ declare function knowledgeReadinessTracePayload(report: KnowledgeReadinessReport, options?: {
97
+ passed?: boolean;
98
+ minimumScore?: number;
99
+ }): Record<string, unknown>;
100
+ declare function userQuestionsForKnowledgeGaps(gaps: KnowledgeRequirement[]): UserQuestion[];
101
+ declare function acquisitionPlansForKnowledgeGaps(gaps: KnowledgeRequirement[]): DataAcquisitionPlan[];
102
+
103
+ export { type DataAcquisitionPlan, type KnowledgeAcquisitionMode, type KnowledgeBundle, type KnowledgeFallbackPolicy, type KnowledgeFreshness, type KnowledgeImportance, type KnowledgeReadinessReport, type KnowledgeRecommendedAction, type KnowledgeRequirement, type KnowledgeRequirementCategory, type KnowledgeResponsibleSurface, type KnowledgeSensitivity, type ScoreKnowledgeReadinessOptions, type UserQuestion, acquisitionPlansForKnowledgeGaps, blockingKnowledgeEval, knowledgeReadinessTracePayload, scoreKnowledgeReadiness, userQuestionsForKnowledgeGaps };
@@ -1,3 +1,6 @@
1
+ import { A as AgentEvalError, C as CaptureIntegrityError } from './errors-CzMUYo7b.js';
2
+ import { R as RawProviderSink, P as ProviderRedactor } from './raw-provider-sink-C46HDghv.js';
3
+
1
4
  /**
2
5
  * LLM client with graceful degrade.
3
6
  *
@@ -19,9 +22,8 @@
19
22
  * output (semantic concept judge, reviewer directives, critic scores). Primitives
20
23
  * that need free-form text use `callLlm` and parse output themselves.
21
24
  */
22
- import { AgentEvalError, CaptureIntegrityError } from './errors';
23
- import { type ProviderRedactor, type RawProviderSink } from './trace/raw-provider-sink';
24
- export interface LlmMessage {
25
+
26
+ interface LlmMessage {
25
27
  role: 'system' | 'user' | 'assistant';
26
28
  /**
27
29
  * Either a plain text content string OR a multimodal content array
@@ -38,7 +40,7 @@ export interface LlmMessage {
38
40
  };
39
41
  }>;
40
42
  }
41
- export interface LlmCallRequest {
43
+ interface LlmCallRequest {
42
44
  model: string;
43
45
  messages: LlmMessage[];
44
46
  /** Optional JSON-mode response format (response_format: json_object). */
@@ -53,14 +55,14 @@ export interface LlmCallRequest {
53
55
  /** Per-call timeout, default 300s. */
54
56
  timeoutMs?: number;
55
57
  }
56
- export interface LlmUsage {
58
+ interface LlmUsage {
57
59
  promptTokens: number;
58
60
  completionTokens: number;
59
61
  totalTokens: number;
60
62
  /** Proxies populate this when prompt caching is on. */
61
63
  cachedPromptTokens?: number;
62
64
  }
63
- export interface LlmCallResult {
65
+ interface LlmCallResult {
64
66
  /** The text content of the first choice. Empty string if none. */
65
67
  content: string;
66
68
  usage: LlmUsage;
@@ -92,13 +94,13 @@ export interface LlmCallResult {
92
94
  /** Raw response body. */
93
95
  raw: Record<string, unknown>;
94
96
  }
95
- export declare class LlmCallError extends AgentEvalError {
97
+ declare class LlmCallError extends AgentEvalError {
96
98
  readonly status: number;
97
99
  readonly body: string;
98
100
  readonly model: string;
99
101
  constructor(message: string, status: number, body: string, model: string);
100
102
  }
101
- export interface LlmClientOptions {
103
+ interface LlmClientOptions {
102
104
  /** Base URL (without trailing slash). Must end at the `/v1` prefix. */
103
105
  baseUrl?: string;
104
106
  /** Bearer token — either `apiKey` or `bearer` populates `Authorization: Bearer ...`. */
@@ -163,39 +165,38 @@ export interface LlmClientOptions {
163
165
  * treated identically whether it surfaces in the HTTP client or a
164
166
  * TCloud-backed judge.
165
167
  */
166
- export declare function isTransientLlmError(err: unknown): boolean;
168
+ declare function isTransientLlmError(err: unknown): boolean;
167
169
  /** Exponential backoff: 500ms, 1s, 2s, 4s, ... capped at 16s. Attempt is 0-indexed. */
168
- export declare function backoffMs(attempt: number): number;
170
+ declare function backoffMs(attempt: number): number;
169
171
  /**
170
172
  * Strip a ```json / ``` code fence if the model emitted one.
171
173
  * Idempotent for naked JSON. Some models (claude-code via router, certain
172
174
  * deepseek models) wrap output even under json_object.
173
175
  */
174
- export declare function stripFencedJson(raw: string): string;
175
- export declare function extractJsonPayload(raw: string): string;
176
+ declare function stripFencedJson(raw: string): string;
176
177
  /**
177
178
  * Low-level call. Returns raw content + usage + cost. Retries on transient
178
179
  * failures; does NOT degrade schema here — callers that want graceful
179
180
  * degrade use `callLlmJson`.
180
181
  */
181
- export declare function callLlm(req: LlmCallRequest, opts?: LlmClientOptions): Promise<LlmCallResult>;
182
+ declare function callLlm(req: LlmCallRequest, opts?: LlmClientOptions): Promise<LlmCallResult>;
182
183
  /**
183
184
  * Structured-output call. Returns parsed JSON plus the raw result envelope.
184
185
  * Degrades `jsonSchema` → `jsonMode` on a 400 that names the schema param —
185
186
  * critical for deepseek-v3/v4, kimi-k2.6, and other models that don't accept
186
187
  * the `response_format.json_schema` shape but DO accept `json_object`.
187
188
  */
188
- export declare function callLlmJson<T = unknown>(req: LlmCallRequest, opts?: LlmClientOptions): Promise<{
189
+ declare function callLlmJson<T = unknown>(req: LlmCallRequest, opts?: LlmClientOptions): Promise<{
189
190
  value: T;
190
191
  result: LlmCallResult;
191
192
  }>;
192
- export type LlmRouteAssertionReason = 'no_explicit_base_url' | 'base_url_blocked' | 'base_url_not_allowed' | 'no_auth' | 'wrong_provider';
193
- export declare class LlmRouteAssertionError extends CaptureIntegrityError {
193
+ type LlmRouteAssertionReason = 'no_explicit_base_url' | 'base_url_blocked' | 'base_url_not_allowed' | 'no_auth' | 'wrong_provider';
194
+ declare class LlmRouteAssertionError extends CaptureIntegrityError {
194
195
  readonly reason: LlmRouteAssertionReason;
195
196
  readonly baseUrl: string;
196
197
  constructor(message: string, reason: LlmRouteAssertionReason, baseUrl: string);
197
198
  }
198
- export interface LlmRouteRequirements {
199
+ interface LlmRouteRequirements {
199
200
  /**
200
201
  * Throw if `opts.baseUrl` is undefined, i.e. the call would fall back to
201
202
  * `DEFAULT_BASE_URL`. Set this for evaluation runs where silently using
@@ -227,7 +228,7 @@ export interface LlmRouteRequirements {
227
228
  * Throws `LlmRouteAssertionError`. Pure — no I/O — so it's safe to call
228
229
  * from constructors and CI gates.
229
230
  */
230
- export declare function assertLlmRoute(opts: LlmClientOptions, req?: LlmRouteRequirements): void;
231
+ declare function assertLlmRoute(opts: LlmClientOptions, req?: LlmRouteRequirements): void;
231
232
  /**
232
233
  * Probe whether a model is reachable. Returns latency + null error on
233
234
  * success; `ok=false` + error message on any failure (HTTP, timeout,
@@ -239,7 +240,7 @@ export declare function assertLlmRoute(opts: LlmClientOptions, req?: LlmRouteReq
239
240
  * for short prompts, so don't tighten this further. We don't validate
240
241
  * content; HTTP 200 means reachable.
241
242
  */
242
- export declare function probeLlm(model: string, opts?: LlmClientOptions & {
243
+ declare function probeLlm(model: string, opts?: LlmClientOptions & {
243
244
  timeoutMs?: number;
244
245
  }): Promise<{
245
246
  ok: boolean;
@@ -251,7 +252,7 @@ export declare function probeLlm(model: string, opts?: LlmClientOptions & {
251
252
  * Thin wrapper around the free functions; exists for callers that want
252
253
  * to inject a single configured instance into multiple primitives.
253
254
  */
254
- export declare class LlmClient {
255
+ declare class LlmClient {
255
256
  private readonly opts;
256
257
  constructor(opts?: LlmClientOptions);
257
258
  call(req: LlmCallRequest, per?: LlmClientOptions): Promise<LlmCallResult>;
@@ -260,4 +261,5 @@ export declare class LlmClient {
260
261
  result: LlmCallResult;
261
262
  }>;
262
263
  }
263
- //# sourceMappingURL=llm-client.d.ts.map
264
+
265
+ export { type LlmClientOptions as L, type LlmRouteRequirements as a, type LlmCallRequest as b, type LlmCallResult as c, LlmCallError as d, LlmClient as e, type LlmMessage as f, LlmRouteAssertionError as g, type LlmUsage as h, assertLlmRoute as i, backoffMs as j, callLlm as k, callLlmJson as l, isTransientLlmError as m, probeLlm as p, stripFencedJson as s };