@tangle-network/agent-eval 0.150.1 → 0.161.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (353) hide show
  1. package/CHANGELOG.md +178 -1
  2. package/README.md +7 -3
  3. package/dist/{active-curriculum-C4mk67HP.js → active-curriculum-CD5TU2yW.js} +3 -13
  4. package/dist/active-curriculum-CD5TU2yW.js.map +1 -0
  5. package/dist/{agent-profile-cell-BkcRDikH.d.ts → agent-profile-cell-CTOZJUuE.d.ts} +4 -2
  6. package/dist/agent-profile-cell-CTOZJUuE.d.ts.map +1 -0
  7. package/dist/analyst/index.d.ts +19 -36
  8. package/dist/analyst/index.d.ts.map +1 -1
  9. package/dist/analyst/index.js +8 -8
  10. package/dist/analyst/index.js.map +1 -1
  11. package/dist/{backend-integrity-DOCa_QrR.d.ts → backend-integrity-DxuQCu_A.d.ts} +4 -3
  12. package/dist/backend-integrity-DxuQCu_A.d.ts.map +1 -0
  13. package/dist/{benchmark-BtAWA8nT.d.ts → benchmark-CGPp-kDC.d.ts} +3 -3
  14. package/dist/{benchmark-BtAWA8nT.d.ts.map → benchmark-CGPp-kDC.d.ts.map} +1 -1
  15. package/dist/{benchmark-command-CAFwbH0L.js → benchmark-command-BVtaq_ve.js} +26 -31
  16. package/dist/benchmark-command-BVtaq_ve.js.map +1 -0
  17. package/dist/benchmarks/index.d.ts +6 -19
  18. package/dist/benchmarks/index.d.ts.map +1 -1
  19. package/dist/benchmarks/index.js +4 -4
  20. package/dist/benchmarks/index.js.map +1 -1
  21. package/dist/builder-eval/index.d.ts +3 -3
  22. package/dist/builder-eval/index.js +3 -3
  23. package/dist/campaign/index.d.ts +9 -9
  24. package/dist/campaign/index.js +7 -7
  25. package/dist/{campaign-CN_7xJdV.js → campaign-BSmOwskD.js} +77 -795
  26. package/dist/campaign-BSmOwskD.js.map +1 -0
  27. package/dist/{canonical-D-XsTQ6_.js → canonical-IL-Bu-14.js} +26 -2
  28. package/dist/canonical-IL-Bu-14.js.map +1 -0
  29. package/dist/{capture-fetch-BBVFzhkk.d.ts → capture-fetch-CqwsJkkG.d.ts} +3 -3
  30. package/dist/{capture-fetch-BBVFzhkk.d.ts.map → capture-fetch-CqwsJkkG.d.ts.map} +1 -1
  31. package/dist/{chat-client-2bVfrzhN.js → chat-client-DlMlAeYI.js} +5 -55
  32. package/dist/{chat-client-2bVfrzhN.js.map → chat-client-DlMlAeYI.js.map} +1 -1
  33. package/dist/chat-json-call-6g5sJobJ.js +53 -0
  34. package/dist/chat-json-call-6g5sJobJ.js.map +1 -0
  35. package/dist/cli.js +54 -19
  36. package/dist/cli.js.map +1 -1
  37. package/dist/{client-LIuo-KPv.js → client-CX7KqIdB.js} +3 -3
  38. package/dist/client-CX7KqIdB.js.map +1 -0
  39. package/dist/{client-kPQYT_56.d.ts → client-L9VVPkim.d.ts} +4 -4
  40. package/dist/{client-kPQYT_56.d.ts.map → client-L9VVPkim.d.ts.map} +1 -1
  41. package/dist/contract/index.d.ts +13 -27
  42. package/dist/contract/index.d.ts.map +1 -1
  43. package/dist/contract/index.js +14 -17
  44. package/dist/contract/index.js.map +1 -1
  45. package/dist/{counterfactual--bpysZF0.d.ts → counterfactual-BaFUWK3H.d.ts} +4 -4
  46. package/dist/{counterfactual--bpysZF0.d.ts.map → counterfactual-BaFUWK3H.d.ts.map} +1 -1
  47. package/dist/{counterfactual-lDfCx0Uz.js → counterfactual-D_VWavVm.js} +2 -2
  48. package/dist/{counterfactual-lDfCx0Uz.js.map → counterfactual-D_VWavVm.js.map} +1 -1
  49. package/dist/{dataset-CJjKqQfA.d.ts → dataset-DQqhOCPt.d.ts} +5 -4
  50. package/dist/{dataset-CJjKqQfA.d.ts.map → dataset-DQqhOCPt.d.ts.map} +1 -1
  51. package/dist/{default-registry-Cw0Ohdoj.d.ts → default-registry-G9CKMNkc.d.ts} +7 -8
  52. package/dist/{default-registry-Cw0Ohdoj.d.ts.map → default-registry-G9CKMNkc.d.ts.map} +1 -1
  53. package/dist/{define-agent-eval-CEQWL9Hy.d.ts → define-agent-eval-Dx1JnPEa.d.ts} +26 -6
  54. package/dist/define-agent-eval-Dx1JnPEa.d.ts.map +1 -0
  55. package/dist/{define-agent-eval-rqNyVhVV.js → define-agent-eval-h-s-sI-v.js} +17 -11
  56. package/dist/define-agent-eval-h-s-sI-v.js.map +1 -0
  57. package/dist/{descriptive-B5MwKfbf.js → descriptive-jDOuI6mz.js} +22 -2
  58. package/dist/descriptive-jDOuI6mz.js.map +1 -0
  59. package/dist/{dspy-rlm-engine-DHI0WrUU.js → dspy-rlm-engine-DptEII26.js} +95 -15
  60. package/dist/dspy-rlm-engine-DptEII26.js.map +1 -0
  61. package/dist/{emitter-CPBAhxum.js → emitter-BpYFQPj4.js} +2 -18
  62. package/dist/emitter-BpYFQPj4.js.map +1 -0
  63. package/dist/{emitter-DGQGoLyj.d.ts → emitter-D_jYSGRd.d.ts} +4 -20
  64. package/dist/{emitter-DGQGoLyj.d.ts.map → emitter-D_jYSGRd.d.ts.map} +1 -1
  65. package/dist/{engine-BLzhNzoY.d.ts → engine-Cu5qD5Fc.d.ts} +9 -11
  66. package/dist/{engine-BLzhNzoY.d.ts.map → engine-Cu5qD5Fc.d.ts.map} +1 -1
  67. package/dist/{eval-campaign-C4jmuM-b.js → eval-campaign-BsXWL2-2.js} +17 -28
  68. package/dist/eval-campaign-BsXWL2-2.js.map +1 -0
  69. package/dist/{exact-types-ccQAyut1.d.ts → exact-types-qnexxJ1Z.d.ts} +2 -2
  70. package/dist/{exact-types-ccQAyut1.d.ts.map → exact-types-qnexxJ1Z.d.ts.map} +1 -1
  71. package/dist/{exec-y-DCLqK7.js → exec-D9WpA2p-.js} +2 -2
  72. package/dist/exec-D9WpA2p-.js.map +1 -0
  73. package/dist/experiment/index.d.ts +45 -11
  74. package/dist/experiment/index.d.ts.map +1 -1
  75. package/dist/experiment/index.js +30 -11
  76. package/dist/experiment/index.js.map +1 -1
  77. package/dist/{experiment-tracker-0MhuPArU.d.ts → experiment-tracker-DCO6Cz4s.d.ts} +2 -2
  78. package/dist/{experiment-tracker-0MhuPArU.d.ts.map → experiment-tracker-DCO6Cz4s.d.ts.map} +1 -1
  79. package/dist/{exporters-q9iL-2Jf.js → exporters-Df7TgHFv.js} +3 -3
  80. package/dist/exporters-Df7TgHFv.js.map +1 -0
  81. package/dist/{external-optimizer-contracts-DbLsm4Po.d.ts → external-optimizer-contracts-szBJ_1vh.d.ts} +2 -2
  82. package/dist/{external-optimizer-contracts-DbLsm4Po.d.ts.map → external-optimizer-contracts-szBJ_1vh.d.ts.map} +1 -1
  83. package/dist/{external-optimizer-process-Bhmzngf-.js → external-optimizer-process-WosTBChy.js} +4 -4
  84. package/dist/{external-optimizer-process-Bhmzngf-.js.map → external-optimizer-process-WosTBChy.js.map} +1 -1
  85. package/dist/{external-optimizer-subprocess-Dn90UqN2.js → external-optimizer-subprocess-BIWbHpgD.js} +10 -6
  86. package/dist/external-optimizer-subprocess-BIWbHpgD.js.map +1 -0
  87. package/dist/{failure-cluster-BLURuWG4.d.ts → failure-cluster-CXL8NbEw.d.ts} +3 -3
  88. package/dist/{failure-cluster-BLURuWG4.d.ts.map → failure-cluster-CXL8NbEw.d.ts.map} +1 -1
  89. package/dist/{feedback-trajectory-DpTTjo0q.d.ts → feedback-trajectory-B3ZHaHV_.d.ts} +7 -7
  90. package/dist/{feedback-trajectory-DpTTjo0q.d.ts.map → feedback-trajectory-B3ZHaHV_.d.ts.map} +1 -1
  91. package/dist/fuzz.js +1 -1
  92. package/dist/{hf-dataset-XggBupCr.js → hf-dataset-D8_RNIis.js} +4 -4
  93. package/dist/{hf-dataset-XggBupCr.js.map → hf-dataset-D8_RNIis.js.map} +1 -1
  94. package/dist/hosted/index.d.ts +3 -14
  95. package/dist/hosted/index.d.ts.map +1 -1
  96. package/dist/hosted/index.js +2 -2
  97. package/dist/{index-CTKpu9ry.d.ts → index-CGtH1piv.d.ts} +48 -26
  98. package/dist/index-CGtH1piv.d.ts.map +1 -0
  99. package/dist/{index-BNPtkBPf.d.ts → index-D-IiQIBB.d.ts} +5 -10
  100. package/dist/index-D-IiQIBB.d.ts.map +1 -0
  101. package/dist/{index-IQccV3Ou.d.ts → index-D-V8gCs_.d.ts} +13 -90
  102. package/dist/index-D-V8gCs_.d.ts.map +1 -0
  103. package/dist/{index-B8Ui1mr1.d.ts → index-lfaSeKSD.d.ts} +18 -2
  104. package/dist/index-lfaSeKSD.d.ts.map +1 -0
  105. package/dist/index-vrJugRal.d.ts +1 -0
  106. package/dist/index.d.ts +68 -125
  107. package/dist/index.d.ts.map +1 -1
  108. package/dist/index.js +61 -112
  109. package/dist/index.js.map +1 -1
  110. package/dist/{insight-report-BeT8KCgI.d.ts → insight-report-DRe8LB6d.d.ts} +4 -4
  111. package/dist/{insight-report-BeT8KCgI.d.ts.map → insight-report-DRe8LB6d.d.ts.map} +1 -1
  112. package/dist/{integrity-DysDBWDu.js → integrity-CyWSSoQS.js} +17 -4
  113. package/dist/integrity-CyWSSoQS.js.map +1 -0
  114. package/dist/{integrity-B0dZ96EO.d.ts → integrity-DUNX9Fao.d.ts} +3 -3
  115. package/dist/{integrity-B0dZ96EO.d.ts.map → integrity-DUNX9Fao.d.ts.map} +1 -1
  116. package/dist/{judge-calibration-DZkWrm5H.js → judge-calibration-zZjLz8hr.js} +2 -2
  117. package/dist/{judge-calibration-DZkWrm5H.js.map → judge-calibration-zZjLz8hr.js.map} +1 -1
  118. package/dist/{kind-factory-DmAa0h3K.js → kind-factory-DY8FdoXf.js} +3 -71
  119. package/dist/kind-factory-DY8FdoXf.js.map +1 -0
  120. package/dist/ledger-core/index.d.ts +2 -2
  121. package/dist/ledger-core/index.js +3 -3
  122. package/dist/{ledger-core-DTae9rv_.js → ledger-core-BOzlRygb.js} +2 -2
  123. package/dist/{ledger-core-DTae9rv_.js.map → ledger-core-BOzlRygb.js.map} +1 -1
  124. package/dist/{llm-client-Bg32RW0j.js → llm-client-hgDieDNN.js} +53 -98
  125. package/dist/llm-client-hgDieDNN.js.map +1 -0
  126. package/dist/{llm-judge-CVq33oz1.js → llm-judge-BhasIPFT.js} +1136 -78
  127. package/dist/llm-judge-BhasIPFT.js.map +1 -0
  128. package/dist/{matrix-DrVnRp4G.d.ts → matrix-eXKRMHnL.d.ts} +74 -72
  129. package/dist/matrix-eXKRMHnL.d.ts.map +1 -0
  130. package/dist/meta-eval/index.d.ts +8 -6
  131. package/dist/meta-eval/index.d.ts.map +1 -1
  132. package/dist/meta-eval/index.js +9 -7
  133. package/dist/meta-eval/index.js.map +1 -1
  134. package/dist/{mint-BV6tLVWl.js → mint-DfODW1KW.js} +3 -3
  135. package/dist/{mint-BV6tLVWl.js.map → mint-DfODW1KW.js.map} +1 -1
  136. package/dist/multishot/golden/index.d.ts +2 -8
  137. package/dist/multishot/golden/index.d.ts.map +1 -1
  138. package/dist/multishot/golden/index.js +56 -86
  139. package/dist/multishot/golden/index.js.map +1 -1
  140. package/dist/multishot/index.d.ts +11 -46
  141. package/dist/multishot/index.d.ts.map +1 -1
  142. package/dist/multishot/index.js +30 -83
  143. package/dist/multishot/index.js.map +1 -1
  144. package/dist/openapi.json +1 -1
  145. package/dist/{opencode-sqlite-DJWAXLms.js → opencode-sqlite-eK6HW6dr.js} +2 -6
  146. package/dist/{opencode-sqlite-DJWAXLms.js.map → opencode-sqlite-eK6HW6dr.js.map} +1 -1
  147. package/dist/pipelines/index.d.ts +5 -5
  148. package/dist/pipelines/index.js +3 -3
  149. package/dist/{pre-registration-zFSLEiFU.d.ts → pre-registration-CzFCcwYk.d.ts} +55 -40
  150. package/dist/pre-registration-CzFCcwYk.d.ts.map +1 -0
  151. package/dist/pre-registration-KN9jkh58.js +110 -0
  152. package/dist/pre-registration-KN9jkh58.js.map +1 -0
  153. package/dist/{produced-state-jfk8Du3b.js → produced-state-DZ89riy5.js} +8 -8
  154. package/dist/produced-state-DZ89riy5.js.map +1 -0
  155. package/dist/profile-cell.d.ts +1 -1
  156. package/dist/profile-cell.js +31 -5
  157. package/dist/profile-cell.js.map +1 -1
  158. package/dist/{promotion-policy-DLOUkYhI.d.ts → promotion-policy-DtnOIZvk.d.ts} +2 -2
  159. package/dist/{promotion-policy-DLOUkYhI.d.ts.map → promotion-policy-DtnOIZvk.d.ts.map} +1 -1
  160. package/dist/{query-Di7eEQ79.js → query-CHmMP42p.js} +20 -11
  161. package/dist/query-CHmMP42p.js.map +1 -0
  162. package/dist/{query-CJ_DX8vl.d.ts → query-DxPYqpmT.d.ts} +10 -4
  163. package/dist/query-DxPYqpmT.d.ts.map +1 -0
  164. package/dist/raw-provider-sink-BU29Sh8h.d.ts +134 -0
  165. package/dist/raw-provider-sink-BU29Sh8h.d.ts.map +1 -0
  166. package/dist/{registry-BQwrSYpC.d.ts → registry-8You7OK1.d.ts} +5 -7
  167. package/dist/{registry-BQwrSYpC.d.ts.map → registry-8You7OK1.d.ts.map} +1 -1
  168. package/dist/{release-confidence-BknrpBnO.js → release-confidence-DKfD2RYU.js} +28 -14
  169. package/dist/release-confidence-DKfD2RYU.js.map +1 -0
  170. package/dist/{release-confidence-4XrqlpFD.d.ts → release-confidence-Dqt0NFep.d.ts} +7 -6
  171. package/dist/release-confidence-Dqt0NFep.d.ts.map +1 -0
  172. package/dist/reporting.d.ts +3 -3
  173. package/dist/reporting.js +3 -3
  174. package/dist/{researcher-DJnoUE8c.d.ts → researcher-Cz565b7D.d.ts} +34 -21
  175. package/dist/researcher-Cz565b7D.d.ts.map +1 -0
  176. package/dist/{reward-hacking-62tojkQd.d.ts → reward-hacking-MBf7qpSB.d.ts} +2 -2
  177. package/dist/{reward-hacking-62tojkQd.d.ts.map → reward-hacking-MBf7qpSB.d.ts.map} +1 -1
  178. package/dist/{reward-hacking-DKI9T52l.js → reward-hacking-t4lB1yt8.js} +3 -3
  179. package/dist/{reward-hacking-DKI9T52l.js.map → reward-hacking-t4lB1yt8.js.map} +1 -1
  180. package/dist/rl.d.ts +11 -42
  181. package/dist/rl.d.ts.map +1 -1
  182. package/dist/rl.js +41 -24
  183. package/dist/rl.js.map +1 -1
  184. package/dist/rollout/index.d.ts +3 -3
  185. package/dist/rollout/index.js +7 -7
  186. package/dist/{rollout-ytVQ7WT8.js → rollout-Dm2tSdiQ.js} +6 -6
  187. package/dist/{rollout-ytVQ7WT8.js.map → rollout-Dm2tSdiQ.js.map} +1 -1
  188. package/dist/{rubric-predictive-validity-Cwwyd7ah.js → rubric-predictive-validity-CK8SCOg-.js} +6 -16
  189. package/dist/rubric-predictive-validity-CK8SCOg-.js.map +1 -0
  190. package/dist/{rubric-predictive-validity-C7LnNvF2.d.ts → rubric-predictive-validity-CxycqzX5.d.ts} +4 -3
  191. package/dist/rubric-predictive-validity-CxycqzX5.d.ts.map +1 -0
  192. package/dist/{run-record-D2lDdSAz.js → run-record-BC0ebuRP.js} +2 -2
  193. package/dist/{run-record-D2lDdSAz.js.map → run-record-BC0ebuRP.js.map} +1 -1
  194. package/dist/{run-record-DVV82Gwh.d.ts → run-record-VVy4T9OW.d.ts} +3 -3
  195. package/dist/{run-record-DVV82Gwh.d.ts.map → run-record-VVy4T9OW.d.ts.map} +1 -1
  196. package/dist/{schema-BtVldJ3T.d.ts → schema-Bjgdsn73.d.ts} +2 -4
  197. package/dist/{schema-BtVldJ3T.d.ts.map → schema-Bjgdsn73.d.ts.map} +1 -1
  198. package/dist/{schema-Cef2cFmb.d.ts → schema-BzWDXhOR.d.ts} +2 -5
  199. package/dist/schema-BzWDXhOR.d.ts.map +1 -0
  200. package/dist/{schema-C6DW4ZHR.js → schema-C1aaAxTf.js} +2 -2
  201. package/dist/schema-C1aaAxTf.js.map +1 -0
  202. package/dist/{schema-CRhEY1SO.js → schema-k6ZBftVv.js} +2 -8
  203. package/dist/{schema-CRhEY1SO.js.map → schema-k6ZBftVv.js.map} +1 -1
  204. package/dist/{semantic-concept-judge-laMCnTLn.js → semantic-concept-judge-BSkKKHeq.js} +14 -38
  205. package/dist/semantic-concept-judge-BSkKKHeq.js.map +1 -0
  206. package/dist/{sequential-eprocess-CbUt2htw.js → sequential-eprocess-D1jKoihe.js} +49 -2
  207. package/dist/sequential-eprocess-D1jKoihe.js.map +1 -0
  208. package/dist/{sequential-C458DXNf.js → sequential-rYW-Ophm.js} +41 -16
  209. package/dist/sequential-rYW-Ophm.js.map +1 -0
  210. package/dist/{series-convergence-BxKEgBwA.d.ts → series-convergence-D9WgpXGi.d.ts} +2 -2
  211. package/dist/{series-convergence-BxKEgBwA.d.ts.map → series-convergence-D9WgpXGi.d.ts.map} +1 -1
  212. package/dist/{server-dIWwF3j_.js → server-BtFd4uzB.js} +19 -42
  213. package/dist/server-BtFd4uzB.js.map +1 -0
  214. package/dist/{skillopt-optimization-method-DLeUcK-K.js → skillopt-optimization-method-DbaekMcn.js} +794 -8
  215. package/dist/skillopt-optimization-method-DbaekMcn.js.map +1 -0
  216. package/dist/{skillopt-optimization-method-BDD_o1xE.d.ts → skillopt-optimization-method-x7TTF23P.d.ts} +20 -7
  217. package/dist/{skillopt-optimization-method-BDD_o1xE.d.ts.map → skillopt-optimization-method-x7TTF23P.d.ts.map} +1 -1
  218. package/dist/{statistical-heldout-_woZ9q9j.d.ts → statistical-heldout-Cy3EhjlC.d.ts} +21 -8
  219. package/dist/statistical-heldout-Cy3EhjlC.d.ts.map +1 -0
  220. package/dist/{steps-AmkT-GIM.d.ts → steps-CiNVJry_.d.ts} +2 -17
  221. package/dist/steps-CiNVJry_.d.ts.map +1 -0
  222. package/dist/{store-CT9YIIve.d.ts → store-B06JdC56.d.ts} +2 -2
  223. package/dist/{store-CT9YIIve.d.ts.map → store-B06JdC56.d.ts.map} +1 -1
  224. package/dist/{store-otlp-CDYWW_8N.js → store-otlp-C_Rq5I4D.js} +2 -2
  225. package/dist/{store-otlp-CDYWW_8N.js.map → store-otlp-C_Rq5I4D.js.map} +1 -1
  226. package/dist/{store-tool-spans-Br2_IUhm.d.ts → store-tool-spans-DPUG7UUY.d.ts} +6 -6
  227. package/dist/{store-tool-spans-Br2_IUhm.d.ts.map → store-tool-spans-DPUG7UUY.d.ts.map} +1 -1
  228. package/dist/{store-tool-spans-CykkbOlv.js → store-tool-spans-Dlh9vkFK.js} +3 -3
  229. package/dist/{store-tool-spans-CykkbOlv.js.map → store-tool-spans-Dlh9vkFK.js.map} +1 -1
  230. package/dist/storyboard/index.d.ts +1 -1
  231. package/dist/{summary-report-Blysd6Z2.js → summary-report-BI5hUtvK.js} +7 -17
  232. package/dist/summary-report-BI5hUtvK.js.map +1 -0
  233. package/dist/{summary-report-B__Y5ub3.d.ts → summary-report-CC07PhEL.d.ts} +6 -5
  234. package/dist/summary-report-CC07PhEL.d.ts.map +1 -0
  235. package/dist/supervisor-run/index.d.ts +27 -18
  236. package/dist/supervisor-run/index.d.ts.map +1 -1
  237. package/dist/supervisor-run/index.js +104 -26
  238. package/dist/supervisor-run/index.js.map +1 -1
  239. package/dist/{task-failure-attributes--ZTP3tYO.js → task-failure-attributes-DTl-7-Kw.js} +3 -3
  240. package/dist/{task-failure-attributes--ZTP3tYO.js.map → task-failure-attributes-DTl-7-Kw.js.map} +1 -1
  241. package/dist/{tool-groups-B4tqh8jB.d.ts → tool-groups-Ci8i9ErB.d.ts} +3 -3
  242. package/dist/tool-groups-Ci8i9ErB.d.ts.map +1 -0
  243. package/dist/{tool-waste-BDdBZG1F.js → tool-waste-BqzmVdJk.js} +4 -4
  244. package/dist/{tool-waste-BDdBZG1F.js.map → tool-waste-BqzmVdJk.js.map} +1 -1
  245. package/dist/{tool-waste-DjRDEsuI.d.ts → tool-waste-Dro0gJi3.d.ts} +4 -4
  246. package/dist/{tool-waste-DjRDEsuI.d.ts.map → tool-waste-Dro0gJi3.d.ts.map} +1 -1
  247. package/dist/trace-repair/index.d.ts +4 -77
  248. package/dist/trace-repair/index.d.ts.map +1 -1
  249. package/dist/trace-repair/index.js +5 -15
  250. package/dist/trace-repair/index.js.map +1 -1
  251. package/dist/traces.d.ts +13 -23
  252. package/dist/traces.d.ts.map +1 -1
  253. package/dist/traces.js +9 -20
  254. package/dist/traces.js.map +1 -1
  255. package/dist/{trajectory-YC15QDYQ.d.ts → trajectory-Bi157Gun.d.ts} +3 -3
  256. package/dist/{trajectory-YC15QDYQ.d.ts.map → trajectory-Bi157Gun.d.ts.map} +1 -1
  257. package/dist/trajectory-replay/index.d.ts +5 -5
  258. package/dist/trajectory-replay/index.js +5 -5
  259. package/dist/{provenance-oA4-zUqm.d.ts → transient-failure-DKF5Mofa.d.ts} +468 -13
  260. package/dist/transient-failure-DKF5Mofa.d.ts.map +1 -0
  261. package/dist/types-B3jzCp0p.js.map +1 -1
  262. package/dist/{types-CLAwnY-L.d.ts → types-BPb2Kf_C2.d.ts} +3 -3
  263. package/dist/types-BPb2Kf_C2.d.ts.map +1 -0
  264. package/dist/types-Bfk0uxRj.d.ts +443 -0
  265. package/dist/types-Bfk0uxRj.d.ts.map +1 -0
  266. package/dist/{types-DdFNuyxQ.d.ts → types-D4s7Z6nq.d.ts} +30 -6
  267. package/dist/types-D4s7Z6nq.d.ts.map +1 -0
  268. package/dist/{types-B2NsbrNy.d.ts → types-D9ssmxKL.d.ts} +3 -3
  269. package/dist/{types-B2NsbrNy.d.ts.map → types-D9ssmxKL.d.ts.map} +1 -1
  270. package/dist/{types-yLK8gXE9.d.ts → types-DeIUdzNd.d.ts} +160 -10
  271. package/dist/types-DeIUdzNd.d.ts.map +1 -0
  272. package/dist/{verdict-BndeTAh_.js → verdict-B0xltqu6.js} +2 -2
  273. package/dist/{verdict-BndeTAh_.js.map → verdict-B0xltqu6.js.map} +1 -1
  274. package/dist/verdict-cache-CdVVTVmn.js +88 -0
  275. package/dist/verdict-cache-CdVVTVmn.js.map +1 -0
  276. package/dist/wire/index.d.ts +21 -111
  277. package/dist/wire/index.d.ts.map +1 -1
  278. package/dist/wire/index.js +2 -2
  279. package/docs/adapters-observability.md +9 -23
  280. package/docs/building-doctrine.md +3 -3
  281. package/docs/campaign-proposers.md +54 -15
  282. package/docs/concepts.md +3 -4
  283. package/docs/design/statistics-decisions.md +89 -1
  284. package/docs/eval-surface-map.md +14 -0
  285. package/docs/experiment.md +19 -2
  286. package/docs/feedback-trajectories.md +1 -1
  287. package/docs/multishot-golden-records.md +4 -4
  288. package/docs/public-api.md +1616 -0
  289. package/docs/research-report-methodology.md +1 -1
  290. package/docs/search-history-receipts.md +39 -1
  291. package/docs/trace-analysis.md +1 -1
  292. package/docs/trace-repair-admission.md +1 -1
  293. package/docs/trace-repair-continuation.md +1 -1
  294. package/docs/verdicts.md +24 -0
  295. package/docs/wire-protocol.md +1 -1
  296. package/package.json +6 -2
  297. package/dist/active-curriculum-C4mk67HP.js.map +0 -1
  298. package/dist/agent-profile-cell-BkcRDikH.d.ts.map +0 -1
  299. package/dist/backend-integrity-DOCa_QrR.d.ts.map +0 -1
  300. package/dist/benchmark-command-CAFwbH0L.js.map +0 -1
  301. package/dist/campaign-CN_7xJdV.js.map +0 -1
  302. package/dist/canonical-D-XsTQ6_.js.map +0 -1
  303. package/dist/client-LIuo-KPv.js.map +0 -1
  304. package/dist/define-agent-eval-CEQWL9Hy.d.ts.map +0 -1
  305. package/dist/define-agent-eval-rqNyVhVV.js.map +0 -1
  306. package/dist/descriptive-B5MwKfbf.js.map +0 -1
  307. package/dist/dspy-rlm-engine-DHI0WrUU.js.map +0 -1
  308. package/dist/emitter-CPBAhxum.js.map +0 -1
  309. package/dist/eval-campaign-C4jmuM-b.js.map +0 -1
  310. package/dist/exec-y-DCLqK7.js.map +0 -1
  311. package/dist/exporters-q9iL-2Jf.js.map +0 -1
  312. package/dist/external-optimizer-subprocess-Dn90UqN2.js.map +0 -1
  313. package/dist/index-B8Ui1mr1.d.ts.map +0 -1
  314. package/dist/index-BNPtkBPf.d.ts.map +0 -1
  315. package/dist/index-C1ravkGA.d.ts +0 -1
  316. package/dist/index-CTKpu9ry.d.ts.map +0 -1
  317. package/dist/index-IQccV3Ou.d.ts.map +0 -1
  318. package/dist/integrity-DysDBWDu.js.map +0 -1
  319. package/dist/kind-factory-DmAa0h3K.js.map +0 -1
  320. package/dist/llm-client-Bg32RW0j.js.map +0 -1
  321. package/dist/llm-judge-CVq33oz1.js.map +0 -1
  322. package/dist/matrix-DrVnRp4G.d.ts.map +0 -1
  323. package/dist/pre-registration-DakwTRXk.js +0 -96
  324. package/dist/pre-registration-DakwTRXk.js.map +0 -1
  325. package/dist/pre-registration-zFSLEiFU.d.ts.map +0 -1
  326. package/dist/produced-state-jfk8Du3b.js.map +0 -1
  327. package/dist/provenance-oA4-zUqm.d.ts.map +0 -1
  328. package/dist/query-CJ_DX8vl.d.ts.map +0 -1
  329. package/dist/query-Di7eEQ79.js.map +0 -1
  330. package/dist/release-confidence-4XrqlpFD.d.ts.map +0 -1
  331. package/dist/release-confidence-BknrpBnO.js.map +0 -1
  332. package/dist/researcher-DJnoUE8c.d.ts.map +0 -1
  333. package/dist/rubric-predictive-validity-C7LnNvF2.d.ts.map +0 -1
  334. package/dist/rubric-predictive-validity-Cwwyd7ah.js.map +0 -1
  335. package/dist/schema-C6DW4ZHR.js.map +0 -1
  336. package/dist/schema-Cef2cFmb.d.ts.map +0 -1
  337. package/dist/semantic-concept-judge-laMCnTLn.js.map +0 -1
  338. package/dist/sequential-C458DXNf.js.map +0 -1
  339. package/dist/sequential-eprocess-CbUt2htw.js.map +0 -1
  340. package/dist/server-dIWwF3j_.js.map +0 -1
  341. package/dist/skillopt-optimization-method-DLeUcK-K.js.map +0 -1
  342. package/dist/statistical-heldout-_woZ9q9j.d.ts.map +0 -1
  343. package/dist/steps-AmkT-GIM.d.ts.map +0 -1
  344. package/dist/summary-report-B__Y5ub3.d.ts.map +0 -1
  345. package/dist/summary-report-Blysd6Z2.js.map +0 -1
  346. package/dist/tool-groups-B4tqh8jB.d.ts.map +0 -1
  347. package/dist/types-CLAwnY-L.d.ts.map +0 -1
  348. package/dist/types-DdFNuyxQ.d.ts.map +0 -1
  349. package/dist/types-jUBXJ7Iz.d.ts +0 -884
  350. package/dist/types-jUBXJ7Iz.d.ts.map +0 -1
  351. package/dist/types-yLK8gXE9.d.ts.map +0 -1
  352. package/dist/verdict-cache-mZf5FEiY.js +0 -107
  353. package/dist/verdict-cache-mZf5FEiY.js.map +0 -1
@@ -0,0 +1,1616 @@
1
+ # Public API census
2
+
3
+ Generated by `pnpm api:census` on demand — this is a dated reading, not a gate. Every named value export of every `package.json#exports` subpath, with the consumer that justifies publishing it. Regenerate it when the surface changes; do not edit it by hand.
4
+
5
+ ## Totals
6
+
7
+ | measure | count |
8
+ | --- | --- |
9
+ | export subpaths | 27 |
10
+ | published value exports (subpath x symbol) | 1352 |
11
+ | distinct symbols | 1183 |
12
+ | production | 903 |
13
+ | planned | 211 |
14
+ | none | 238 |
15
+
16
+ Type-only exports are not listed: a type binds no runtime surface, and removing one cannot break a caller at run time.
17
+
18
+ ## Classification
19
+
20
+ - **production** — a consumer repository imports it in production code, this package's own CLI or wire server imports it, a consumer binds it in a type position, or another production module of this package imports it.
21
+ - **planned** — no production caller, but a consumer's tests, a runnable example, a Markdown front door, or a consumer repository that names the symbol where the import graph cannot see the bind.
22
+ - **none** — no evidence in any channel above. This is the delete list.
23
+
24
+ ## Consumer sweep
25
+
26
+ Swept 2026-08-21 across 31 repositories, each on its default branch:
27
+
28
+ | repository | ref | commit |
29
+ | --- | --- | --- |
30
+ | agent-app | main | `da061b0` |
31
+ | agent-builder | main | `8350c93` |
32
+ | agent-dev-container | unknown | `unknown` |
33
+ | agent-knowledge | HEAD | `02c9abe` |
34
+ | agent-lab | unknown | `unknown` |
35
+ | agent-runtime | HEAD | `6d11021e` |
36
+ | agent-sdk | main | `5241bf6` |
37
+ | ai-trading-blueprint | main | `c0070e8` |
38
+ | blueprint-agent | develop | `38f6f93` |
39
+ | braid | HEAD | `3d7b79f` |
40
+ | browser-agent-driver | main | `7fa06d5` |
41
+ | creative-agent | master | `895ff9c` |
42
+ | discovery-lab | master | `d7b5703` |
43
+ | gtm-agent | master | `ca1eb45` |
44
+ | insurance-agent | main | `caf9d7f` |
45
+ | legal-agent | main | `3962c63` |
46
+ | loops | unknown | `unknown` |
47
+ | phony | unknown | `unknown` |
48
+ | physim | main | `826c877` |
49
+ | playproof | unknown | `unknown` |
50
+ | redteam | main | `5c0411c` |
51
+ | relationships-agent | main | `3997f03` |
52
+ | run-capsule | main | `5d756f8` |
53
+ | skeletal-os | unknown | `unknown` |
54
+ | starter-foundry | main | `94225e3` |
55
+ | supervisor-lab | main | `c95d732` |
56
+ | tangle-router | main | `2b770d4` |
57
+ | tax-agent | main | `e65b319` |
58
+ | traces | unknown | `unknown` |
59
+ | tuner-agent | main | `090cdd1` |
60
+ | workcomp-agent | master | `20a68e7` |
61
+
62
+ ## What this census cannot see
63
+
64
+ Every gap below makes a `none` less certain, so each one is answered by keeping the symbol, never by deleting it.
65
+
66
+ 1. **Dynamic imports and namespace binds.** `const { x } = await import('@tangle-network/agent-eval')` and `import * as evaluation from ...` name no symbol the import graph can resolve. The sweep therefore also records every census name mentioned anywhere in a consumer repository; a symbol with only that evidence is `planned`, never `none`.
67
+ 2. **Repositories outside the sweep.** The list above is every repository in the `tangle-network` organisation whose `package.json` names this package, plus the local checkouts. A private fork, a consumer outside the organisation, or an npm consumer nobody told us about returns no hits and reads exactly like an unused symbol.
68
+ 3. **Pinned versions.** Each consumer is read at its default branch HEAD, which may resolve an older published version whose surface differs from this one.
69
+ 4. **The wire and RPC surfaces.** A JSON-RPC method name or an OpenAPI schema reached over the wire binds no TypeScript symbol. `docs/wire-protocol.md` and `clients/python/` own that contract.
70
+ 5. **String-keyed dispatch.** A symbol reached through a registry keyed by string is invisible to both channels.
71
+
72
+ ## Why a `none` can still be published
73
+
74
+ A `none` row is a deletion candidate, not a deletion order. A symbol stays when removing it would lose something the census cannot weigh: a documented historical constant, a value another module of this package still needs, or a name whose only reference is a contract test over an on-disk artifact. Those are kept deliberately and stay listed here as `none`, so the next reader sees the same evidence and can decide again.
75
+
76
+ ## Exports
77
+
78
+ ### `.`
79
+
80
+ 318 value exports — 267 production, 32 planned, 19 none.
81
+
82
+ | symbol | consumer | evidence |
83
+ | --- | --- | --- |
84
+ | `acquisitionPlansForKnowledgeGaps` | production | agent-builder:src/lib/.server/eval/loops/data-acquisition-engine.ts |
85
+ | `AGENT_PROFILE_KINDS` | production | blueprint-agent:scripts/experiments/lib/agent-profile-cell.ts |
86
+ | `AgentEvalError` | production | agent-runtime:src/errors.ts |
87
+ | `agentProfileCellHashMaterial` | planned | consumer tests: blueprint-agent:scripts/experiments/lib/__tests__/agent-profile-cell.test.ts |
88
+ | `agentProfileCellKey` | planned | consumer tests: blueprint-agent:scripts/experiments/lib/__tests__/agent-profile-cell.test.ts |
89
+ | `agentProfileHash` | production | ai-trading-blueprint:evals/src/trading/persona-agent-eval.ts |
90
+ | `agentProfileId` | production | agent-dev-container:products/intelligence/api/src/lib/project-profile-change-preflight.ts |
91
+ | `aggregateJudgeVerdicts` | production | agent-app:src/eval-campaign/index.ts |
92
+ | `aggregateRunScore` | production | agent-dev-container:products/sandbox/evals/src/auto-optimization.ts |
93
+ | `analystFindingDigest` | production | this package: src/feedback-trajectory.ts:4 |
94
+ | `AnalystRegistry` | production | blueprint-agent:scripts/experiments/vb-multi-analyst.ts |
95
+ | `analystRunDigest` | production | this package: src/feedback-trajectory.ts:4 |
96
+ | `analystRunToFeedbackTrajectory` | planned | doc: docs/feedback-trajectories.md |
97
+ | `analystRunToReviewRequests` | planned | doc: docs/feedback-trajectories.md |
98
+ | `analyzeAntiSlop` | production | physim:apps/server/src/lib/eval.ts |
99
+ | `analyzeRuns` | production | agent-builder:eval/scripts/insight-report.ts |
100
+ | `analyzeSeries` | production | agent-dev-container:products/sandbox/evals/run.ts |
101
+ | `analyzeTraces` | production | agent-builder:src/lib/.server/eval/analysts/canonical-trace-analyst.ts |
102
+ | `argHash` | production | agent-runtime:src/runtime/supervise/detector-monitor.ts |
103
+ | `assertCapabilityHeadroom` | production | blueprint-agent:scripts/experiments/lib/validity-gates.ts |
104
+ | `assertCrossFamily` | production | agent-builder:frontier/judges/artifact-head-to-head.ts |
105
+ | `assertCrossFamilyServed` | planned | doc: docs/building-doctrine.md |
106
+ | `assertNoHiddenLeak` | planned | doc: docs/design/statistics-decisions.md |
107
+ | `assertProductBenchmarkRun` | production | creative-agent:eval/research-package.ts |
108
+ | `assertRealBackend` | production | blueprint-agent:scripts/experiments/lib/canonical/assert-real-backend.ts |
109
+ | `assertRunCaptured` | production | creative-agent:eval/canonical-runner.ts |
110
+ | `assertServedModel` | production | creative-agent:eval/lib/persona-driver.ts |
111
+ | `assertServedModels` | planned | doc: docs/building-doctrine.md |
112
+ | `assertSingleBackend` | planned | consumer tests: gtm-agent:tests/agent-eval.smoke.test.ts |
113
+ | `assignFeedbackSplit` | production | agent-builder:src/lib/.server/eval/feedback/feedback-capture.ts |
114
+ | `BackendIntegrityError` | production | creative-agent:eval/leaderboard.ts |
115
+ | `benjaminiHochberg` | production | agent-runtime:bench/src/corpus-report.mts |
116
+ | `blendHeldout` | production | blueprint-agent:scripts/experiments/lib/qa-grader.ts |
117
+ | `blockingKnowledgeEval` | production | agent-dev-container:products/intelligence/api/src/routes/project-engines.ts |
118
+ | `bonferroni` | planned | doc: docs/design/statistics-decisions.md |
119
+ | `BOOTSTRAP_GATE_MIN_N` | production | discovery-lab:tools/confirmation-activation.mjs |
120
+ | `bootstrapCi` | production | agent-dev-container:products/intelligence/api/src/lib/analysis-worker.ts |
121
+ | `BudgetBreachError` | production | creative-agent:scripts/evals/run-driver-soak.ts |
122
+ | `budgetBreachView` | production | insurance-agent:scripts/analyze-agent-eval-evidence.ts |
123
+ | `BudgetGuard` | production | creative-agent:scripts/evals/run-driver-soak.ts |
124
+ | `buildAgentProfileCell` | production | ai-trading-blueprint:evals/src/trading/agent-profile-cell.ts |
125
+ | `buildDefaultAnalystRegistry` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
126
+ | `buildEquivalenceRecord` | none | only this package's tests: src/verification-strategy.test.ts:11 |
127
+ | `buildReflectionPrompt` | production | blueprint-agent:scripts/experiments/lib/gepa-reflective-proposer.ts |
128
+ | `buildTraceInsightContext` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/report-html-sections.ts |
129
+ | `buildTraceInsightPrompt` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/brief.ts |
130
+ | `buildTrajectory` | production | agent-runtime:src/runtime/supervise/trajectory-recorder.ts |
131
+ | `calibrateJudge` | production | creative-agent:eval/calibrate-judges.ts |
132
+ | `calibrateJudgeContinuous` | planned | consumer tests: tax-agent:tests/eval/lib/judge-sentinel.ts |
133
+ | `canonicalJson` | production | agent-knowledge:src/benchmarks/memory-recovery.ts |
134
+ | `capabilityHeadroom` | production | blueprint-agent:scripts/experiments/lib/validity-gates.ts |
135
+ | `captureFetchToRawSink` | production | creative-agent:eval/canonical-runner.ts |
136
+ | `certificationEvidenceDigest` | production | this package: src/completion-verifier.ts:38 |
137
+ | `checkCanaries` | production | blueprint-agent:scripts/experiments/lib/contamination-preflight.ts |
138
+ | `checkServedModel` | production | this package: src/integrity/preflight.ts:33 |
139
+ | `checkTraceContracts` | production | creative-agent:eval/lib/trace-contracts.ts |
140
+ | `clamp01` | production | agent-dev-container:products/intelligence/api/src/lib/consultant/executor.ts |
141
+ | `classifyFailure` | production | phony:products/builder/api/src/eval/autoresearch/trace-analyst.ts |
142
+ | `cliffsDelta` | production | agent-builder:src/lib/.server/eval/loops/differential-eval.ts |
143
+ | `CODING_HARNESSES` | production | agent-runtime:src/runtime/define-leaderboard.ts |
144
+ | `cohensD` | production | starter-foundry:examples/recruiter-eval-workspace/eval/src/eval/regression/gate.ts |
145
+ | `comparePairedArms` | production | agent-knowledge:src/memory/experiment/learning-metrics.ts |
146
+ | `completionVerdict` | production | blueprint-agent:scripts/experiments/eval/launch-scoreboard/scoreboard-core.ts |
147
+ | `computeExperimentStats` | production | blueprint-agent:scripts/experiments/lib/vb-improve-harness.ts |
148
+ | `computeFindingId` | production | agent-runtime:src/runtime/index.ts |
149
+ | `computeToolUseMetrics` | production | traces:src/pipelines.ts |
150
+ | `confidenceInterval` | production | agent-runtime:src/runtime/benchmark-report.ts |
151
+ | `ConfigError` | production | agent-runtime:src/errors.ts |
152
+ | `contentHash` | production | agent-knowledge:src/file-transaction.ts |
153
+ | `continuousAgreement` | production | agent-builder:frontier/judges/eval-validity.ts |
154
+ | `controlRunToFeedbackTrajectory` | production | creative-agent:eval/control/creative-onboarding.ts |
155
+ | `corpusInterRaterAgreement` | planned | doc: docs/design/statistics-decisions.md |
156
+ | `corpusInterRaterAgreementFromJudgeScores` | planned | consumer tests: tax-agent:tests/eval/lib/judge-irr.ts |
157
+ | `CostAccountingIncompleteError` | production | this package: src/analyst/benchmark-command.ts:6 |
158
+ | `CostCallConflictError` | production | this package: src/analyst/benchmark-public-model.ts:8 |
159
+ | `CostCeilingReachedError` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
160
+ | `costForTokenPricing` | production | agent-runtime:src/runtime/profile-chat-client.ts |
161
+ | `costForUsage` | production | agent-builder:src/lib/.server/llm-cost.ts |
162
+ | `CostLedger` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
163
+ | `CostLedgerPersistenceError` | production | this package: src/analyst/benchmark-public-model.ts:8 |
164
+ | `CostReceiptCaptureError` | production | this package: src/analyst/benchmark-public-model.ts:8 |
165
+ | `costReceiptFromLlm` | production | agent-dev-container:products/intelligence/api/src/lib/router-model-owner.ts |
166
+ | `costReceiptFromLlmError` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
167
+ | `CostReservationExceededError` | production | this package: src/analyst/benchmark-public-errors.ts:2 |
168
+ | `CostTracker` | production | creative-agent:eval/canonical-runner.ts |
169
+ | `createAntiSlopJudge` | production | phony:products/builder/api/src/eval/judges.ts |
170
+ | `createBoundedTraceAnalysisStore` | production | braid:src/adapters/analysis/trace-store.ts |
171
+ | `createChatClient` | production | agent-dev-container:products/intelligence/api/src/lib/intent-audit-analyst.ts |
172
+ | `createDspyRlmTraceEngine` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
173
+ | `createFeedbackTrajectory` | production | agent-builder:src/lib/.server/eval/feedback/feedback-capture.ts |
174
+ | `createLlmCorrectnessChecker` | production | agent-app:src/eval/index.ts |
175
+ | `createLlmReviewer` | production | agent-builder:src/lib/.server/eval/judges/reviewer-llm.ts |
176
+ | `createTraceAnalyst` | production | blueprint-agent:scripts/experiments/lib/sandbox-driver/analyst-steer.ts |
177
+ | `CrossFamilyError` | production | creative-agent:eval/lib/judge-ensemble.ts |
178
+ | `decidePairedPromotion` | production | discovery-lab:tools/decide.mjs |
179
+ | `DECISION_PAIRED_DELTA_STATISTIC` | production | this package: src/campaign/gates/promotion-policy.ts:33 |
180
+ | `DEFAULT_PERMUTATIONS` | none | — |
181
+ | `DEFAULT_RED_TEAM_CORPUS` | production | starter-foundry:registry/layers/agent-eval/redteam/files/src/eval/redteam/runner.ts |
182
+ | `DEFAULT_REDACTION_RULES` | production | agent-dev-container:products/intelligence/api/src/lib/redact.ts |
183
+ | `DEFAULT_TRACE_ANALYST_BUDGETS` | production | agent-builder:src/lib/.server/eval/stores/d1-trace-analysis-store-adapter.ts |
184
+ | `DEFAULT_TRACE_ANALYST_KINDS` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
185
+ | `defaultBlendWeights` | production | blueprint-agent:scripts/experiments/lib/qa-grader.ts |
186
+ | `defineAgentEval` | planned | example: examples/evaluate-a-change/index.ts:10 |
187
+ | `defineEquivalenceCheck` | planned | example: examples/verify-without-an-answer-key/index.ts:13 |
188
+ | `deployGateLayer` | production | blueprint-agent:scripts/experiments/lib/deploy-runner.ts |
189
+ | `describeTraceInsightScope` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/report-html.ts |
190
+ | `diffFindings` | production | agent-runtime:src/analyst-loop/run-analyst-loop.ts |
191
+ | `diffScorecard` | production | ai-trading-blueprint:evals/src/trading/scorecard-integration.ts |
192
+ | `discoverPersonas` | production | insurance-agent:scripts/harvest-candidate-outputs.ts |
193
+ | `domainEvidencePattern` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/session-loader.ts |
194
+ | `dominates` | production | blueprint-agent:scripts/experiments/lib/competition-analytics.ts |
195
+ | `ensembleJudge` | production | blueprint-agent:scripts/experiments/lib/qa-grader.ts |
196
+ | `eProcess` | production | this package: src/campaign/gates/sequential.ts:35 |
197
+ | `EquivalenceProtocolError` | planned | doc: docs/verification-strategies.md |
198
+ | `equivalenceVerdict` | planned | doc: docs/verdicts.md |
199
+ | `ERROR_COUNT_PATTERNS` | production | blueprint-agent:scripts/experiments/lib/error-count-extractor.ts |
200
+ | `errorStreakDetector` | production | agent-runtime:src/runtime/supervise/detector-monitor.ts |
201
+ | `estimateCost` | production | agent-builder:src/lib/.server/llm-cost.ts |
202
+ | `estimateTokens` | production | physim:apps/server/src/lib/runner.ts |
203
+ | `evaluateActionPolicy` | planned | doc: docs/control-runtime.md |
204
+ | `evaluateInterimReleaseConfidence` | production | agent-builder:scripts/eval.ts |
205
+ | `evaluateOracles` | planned | doc: docs/verdicts.md |
206
+ | `evaluateReleaseConfidence` | production | agent-knowledge:src/release.ts |
207
+ | `expandProfileAxes` | production | agent-runtime:src/runtime/define-leaderboard.ts |
208
+ | `exportProductBenchmark` | production | creative-agent:eval/research-package.ts |
209
+ | `exportProductBenchmarkRuns` | production | creative-agent:eval/research-package.ts |
210
+ | `exportRunAsOtlp` | production | agent-builder:src/lib/.server/eval/analysts/canonical-trace-analyst.ts |
211
+ | `extractErrorCount` | production | blueprint-agent:scripts/experiments/lib/error-count-extractor.ts |
212
+ | `extractProducedState` | production | agent-app:src/eval/index.ts |
213
+ | `extractUsage` | production | this package: src/contract/intake/code-agent-session.ts:12 |
214
+ | `extractUsageFromSse` | production | this package: src/trace/capture-fetch.ts:8 |
215
+ | `FAILURE_CLASSES` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
216
+ | `FAILURE_MODE_KIND_SPEC` | planned | consumer tests: legal-agent:tests/eval/analysts/completion-coverage.ts |
217
+ | `failureClusterView` | production | insurance-agent:scripts/analyze-agent-eval-evidence.ts |
218
+ | `feedbackTrajectoriesToDatasetScenarios` | production | creative-agent:eval/control/creative-onboarding.ts |
219
+ | `feedbackTrajectoriesToOptimizerRows` | production | agent-dev-container:products/intelligence/api/src/lib/feedback-memory.ts |
220
+ | `feedbackTrajectoryToOptimizerRow` | planned | doc: docs/feedback-trajectories.md |
221
+ | `FileSystemFeedbackTrajectoryStore` | production | ai-trading-blueprint:evals/src/product/autoresearch-loop-runner.ts |
222
+ | `FileSystemRawProviderSink` | production | agent-builder:scripts/eval.ts |
223
+ | `FileSystemTraceStore` | production | agent-builder:scripts/eval.ts |
224
+ | `fileVerdictCache` | production | blueprint-agent:scripts/experiments/calibrate-instrument.ts |
225
+ | `FindingsStore` | production | agent-builder:scripts/run-canonical-analyst-loop.ts |
226
+ | `formatScorecardDiff` | production | ai-trading-blueprint:evals/src/trading/scorecard-integration.ts |
227
+ | `gainHistogram` | production | blueprint-agent:scripts/experiments/vb-runrecord-analysis.ts |
228
+ | `gateTreatmentApplied` | production | blueprint-agent:scripts/experiments/lib/validity-gates.ts |
229
+ | `gradeOnHidden` | none | only this package's tests: src/hidden-criteria-grading.test.ts:4 |
230
+ | `gradeSemanticStatus` | production | blueprint-agent:scripts/experiments/lib/verification-harness/primitives.ts |
231
+ | `groupRunsByAgentProfileCell` | planned | consumer tests: legal-agent:tests/eval/matrix-multi.ts |
232
+ | `HARNESS_NATIVE_MODEL` | production | agent-runtime:src/runtime/supervise/model-policy.ts |
233
+ | `harnessAxisOf` | production | gtm-agent:eval/matrix/harnesses.ts |
234
+ | `hashContent` | production | agent-runtime:bench/src/corpus.ts |
235
+ | `hashJson` | production | agent-dev-container:products/intelligence/api/src/lib/ingest-optimization-mapping.ts |
236
+ | `HeldOutGate` | production | agent-builder:src/lib/.server/eval/gates/promotion.ts |
237
+ | `hiddenGrade` | none | only this package's tests: src/hidden-criteria-grading.test.ts:4 |
238
+ | `holm` | planned | doc: docs/design/statistics-decisions.md |
239
+ | `improvementVerdict` | production | blueprint-agent:scripts/experiments/lib/vb-improve-harness.ts |
240
+ | `inferDomainKeywords` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/suite-projection.ts |
241
+ | `InMemoryFeedbackTrajectoryStore` | production | creative-agent:eval/control/creative-onboarding.ts |
242
+ | `InMemoryRawProviderSink` | production | agent-builder:src/lib/.server/eval/loops/canonical-campaign.ts |
243
+ | `inMemoryReviewStore` | production | agent-builder:scripts/propose-review-smoke.ts |
244
+ | `InMemoryTraceStore` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
245
+ | `interpretCliffs` | production | agent-builder:src/lib/.server/eval/loops/differential-eval.ts |
246
+ | `interRaterReliability` | production | agent-app:src/eval-campaign/trust-gate.ts |
247
+ | `iqr` | production | this package: src/experiment-tracker.ts:27 |
248
+ | `isBinaryOutcomeVector` | none | only this package's tests: src/statistics/paired-binary.test.ts:2 |
249
+ | `isJudgeSpan` | production | this package: src/trace/query.ts:12 |
250
+ | `isLlmSpan` | production | agent-runtime:src/candidate-execution/finalize.ts |
251
+ | `isModelPriced` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
252
+ | `isRunRecord` | production | this package: src/eval-trace-store.ts:21 |
253
+ | `isToolSpan` | production | agent-runtime:src/runtime/supervise/trace-evidence.ts |
254
+ | `isTransientLlmError` | production | supervisor-lab:bench/comms/judges.ts |
255
+ | `jsonlReviewStore` | production | starter-foundry:scripts/enrich-family.ts |
256
+ | `jsonlRunRecordBackend` | production | workcomp-agent:eval/benchmark/store.ts |
257
+ | `jsonShape` | planned | named in creative-agent (bind not in the import graph) |
258
+ | `judgeAgreementView` | production | insurance-agent:scripts/analyze-agent-eval-evidence.ts |
259
+ | `JudgeError` | production | agent-runtime:src/errors.ts |
260
+ | `judgeFamily` | production | agent-builder:frontier/judges/llm.ts |
261
+ | `judgeSpans` | production | this package: src/builder-eval/three-layer-eval.ts:25 |
262
+ | `knowledgeReadinessTracePayload` | production | agent-dev-container:products/intelligence/api/src/routes/project-engines.ts |
263
+ | `leaderboard` | production | blueprint-agent:scripts/experiments/lib/cross-harness-leaderboard.ts |
264
+ | `LlmCallError` | production | agent-dev-container:products/intelligence/api/src/lib/router-model-owner.ts |
265
+ | `llmJudge` | production | agent-runtime:examples/agentic-data-creation/offline-fixtures.ts |
266
+ | `LlmResponseError` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
267
+ | `loadScorecard` | production | ai-trading-blueprint:evals/src/trading/scorecard-integration.ts |
268
+ | `localCommandRunner` | production | blueprint-agent:scripts/experiments/lib/command-runner.ts |
269
+ | `makeFinding` | production | agent-dev-container:products/intelligence/api/src/lib/consultant/candidates.ts |
270
+ | `makeProposalFinding` | production | agent-runtime:bench/src/swe-arena/outer-loop.mts |
271
+ | `manifestContentDigest` | production | this package: src/campaign/gates/sequential.ts:34 |
272
+ | `MANN_WHITNEY_EXACT_MAX_STATES` | none | — |
273
+ | `MANN_WHITNEY_EXACT_MAX_WORK` | none | — |
274
+ | `mannWhitneyU` | production | blueprint-agent:scripts/experiments/lib/competition-analytics.ts |
275
+ | `maximumChargeForLlmRequest` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
276
+ | `mcnemar` | production | agent-runtime:bench/src/swe-arena/analyze.ts |
277
+ | `mcnemarPower` | production | supervisor-lab:bench/deepswe/headroom.ts |
278
+ | `mcnemarRequiredN` | production | blueprint-agent:scripts/experiments/analyze-convergence.ts |
279
+ | `minimumPairsForPairedDeltaTest` | production | agent-dev-container:products/intelligence/api/src/routes/optimizations.ts |
280
+ | `mintRolloutRows` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
281
+ | `MODEL_PRICING` | production | supervisor-lab:bench/profile-arena.ts |
282
+ | `modelHasSnapshot` | production | agent-dev-container:products/intelligence/api/src/lib/project-profile-change-preflight.ts |
283
+ | `modelPriceKey` | none | only this package's tests: src/cost-ledger.test.ts:7 |
284
+ | `ModelSubstitutionError` | production | this package: src/campaign/external-optimizer-model-proxy.ts:11 |
285
+ | `mulberry32` | production | agent-knowledge:src/memory/holdout.ts |
286
+ | `MultiLayerVerifier` | production | agent-builder:src/lib/.server/eval/loops/forge-refinement.ts |
287
+ | `NoopRawProviderSink` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
288
+ | `notBlocked` | none | only this package's tests: src/oracle.test.ts:2 |
289
+ | `NotFoundError` | production | agent-runtime:src/errors.ts |
290
+ | `objectiveEval` | production | agent-knowledge:src/research-loop.ts |
291
+ | `observeAll` | production | agent-runtime:src/runtime/supervise/detector-monitor.ts |
292
+ | `OtlpFileTraceStore` | production | agent-builder:scripts/run-canonical-analyst-loop.ts |
293
+ | `otlpTextToTraceAnalysisStore` | production | braid:src/adapters/analysis/trace-store.ts |
294
+ | `OUTPUT_VALUE` | production | agent-runtime:src/runtime/supervise-surface.ts |
295
+ | `pairArms` | production | agent-knowledge:src/memory/experiment/learning-pairs.ts |
296
+ | `pairedBinaryScale` | production | this package: src/paired-promotion-decision.ts:51 |
297
+ | `pairedBootstrap` | production | agent-builder:frontier/adapters/fleet-eval-runner.ts |
298
+ | `pairedCohensDz` | production | this package: src/contract/analyze-runs.ts:37 |
299
+ | `pairedDeltaTest` | production | gtm-agent:eval/matrix/report.ts |
300
+ | `pairedDeltaTieFraction` | production | this package: src/paired-promotion-decision.ts:51 |
301
+ | `pairedEvalueSequence` | production | creative-agent:src/lib/experiments/ab-design.ts |
302
+ | `pairedMde` | production | discovery-lab:tools/design-gate.mjs |
303
+ | `pairedRiskDifference` | production | blueprint-agent:scripts/experiments/lib/validity-gates.ts |
304
+ | `pairedRiskDifferenceExact` | production | this package: src/paired-promotion-decision.ts:51 |
305
+ | `pairedRiskDifferenceScore` | production | this package: src/paired-promotion-decision.ts:51 |
306
+ | `pairedSignTest` | production | agent-dev-container:products/intelligence/api/src/lib/paired-stats.ts |
307
+ | `pairedTTest` | production | this package: src/contract/analyze-runs.ts:37 |
308
+ | `pairRunRecords` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
309
+ | `PairwiseSteeringOptimizer` | production | browser-agent-driver:bench/research/webvoyager-agent-eval-loop.mjs |
310
+ | `paretoChart` | production | workcomp-agent:eval/benchmark/select.ts |
311
+ | `paretoFrontier` | production | agent-dev-container:products/sandbox/evals/run.ts |
312
+ | `parseReflectionResponse` | production | blueprint-agent:scripts/experiments/lib/gepa-reflective-proposer.ts |
313
+ | `parseRunRecordSafe` | production | gtm-agent:scripts/canonical-completeness/records.ts |
314
+ | `partialCredit` | planned | doc: docs/design/statistics-decisions.md |
315
+ | `partitionHeldOut` | production | blueprint-agent:scripts/experiments/calibrate-instrument.ts |
316
+ | `passAtK` | planned | doc: docs/design/statistics-decisions.md |
317
+ | `pearsonR` | production | this package: src/builder-eval/correlation.ts:18 |
318
+ | `preflightModels` | production | loops:src/preflight.ts |
319
+ | `productBenchmarkRepoIdentity` | production | creative-agent:eval/research-package.ts |
320
+ | `ProductClient` | production | creative-agent:eval/e2e/creative-product-harness.ts |
321
+ | `profile` | production | agent-app:src/profile/index.ts |
322
+ | `projectRuntimeTrajectoryEvidence` | production | agent-dev-container:products/intelligence/api/src/lib/workflow-trace-preflight.ts |
323
+ | `PromptRegistry` | production | blueprint-agent:scripts/experiments/lib/surface-prompt-registry.ts |
324
+ | `proposeSynthesisTargets` | production | blueprint-agent:scripts/experiments/researcher/inspect.ts |
325
+ | `ranks` | planned | doc: docs/design/statistics-decisions.md |
326
+ | `readProductBenchmarkManifest` | production | creative-agent:eval/research-package.ts |
327
+ | `recordRuns` | planned | consumer tests: tax-agent:tests/eval/benchmarks/taxcalc/analyze.mjs |
328
+ | `recordRunsToScorecard` | production | ai-trading-blueprint:evals/src/trading/persona-agent-eval.ts |
329
+ | `REDACTION_VERSION` | production | agent-dev-container:products/intelligence/api/src/lib/redact.ts |
330
+ | `redactString` | production | agent-dev-container:products/intelligence/api/src/lib/redact.ts |
331
+ | `redTeamDataset` | production | creative-agent:eval/red-team.ts |
332
+ | `redTeamReport` | production | creative-agent:eval/red-team.ts |
333
+ | `regexMatches` | none | only this package's tests: src/oracle.test.ts:2 |
334
+ | `renderPreferenceMemoryMarkdown` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
335
+ | `repeatedActionDetector` | production | agent-runtime:src/runtime/supervise/detector-monitor.ts |
336
+ | `requiredPairedSampleSize` | production | discovery-lab:tools/design-gate.mjs |
337
+ | `requiredSampleSize` | planned | doc: docs/design/statistics-decisions.md |
338
+ | `resolveModelPricing` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-primitives.ts |
339
+ | `resolveSeat` | production | creative-agent:eval/agent.config.ts |
340
+ | `roundTripRunRecord` | none | only this package's tests: tests/run-record.test.ts:3 |
341
+ | `routeFields` | none | only this package's tests: src/hidden-criteria-grading.test.ts:4 |
342
+ | `runAgentControlLoop` | production | agent-runtime:src/run.ts |
343
+ | `runCampaign` | production | agent-app:src/eval-campaign/index.ts |
344
+ | `runCanaries` | production | agent-builder:src/lib/.server/eval/loops/canary-cron.ts |
345
+ | `runCounterfactual` | production | traces:src/replay-batch.ts |
346
+ | `runEquivalenceCheck` | planned | example: examples/verify-without-an-answer-key/index.ts:13 |
347
+ | `runEvalCampaign` | production | agent-builder:scripts/eval.ts |
348
+ | `RunIntegrityError` | production | creative-agent:eval/canonical-runner.ts |
349
+ | `runIntentMatchJudge` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
350
+ | `runKeywordCoverageJudge` | production | blueprint-agent:scripts/experiments/calibration/scorers.ts |
351
+ | `runKeywordCoverageJudgeUrl` | production | blueprint-agent:scripts/experiments/lib/completeness-audit.ts |
352
+ | `runProposeReview` | production | agent-builder:scripts/propose-review-smoke.ts |
353
+ | `runProposeReviewAsControlLoop` | none | — |
354
+ | `RunRecordValidationError` | production | starter-foundry:src/lib/index.ts |
355
+ | `runSemanticConceptJudge` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
356
+ | `runsForScenario` | production | phony:products/builder/api/src/eval/gate.ts |
357
+ | `runTaskScore` | production | discovery-lab:tools/strict-screen.mjs |
358
+ | `scoreKnowledgeReadiness` | production | agent-builder:src/lib/.server/runtime/agent-task-readiness.ts |
359
+ | `scoreRedTeamOutput` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-gate.ts |
360
+ | `scoreTraceInsightReadiness` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/report-html-sections.ts |
361
+ | `seatPresets` | production | supervisor-lab:bench/comms/judges.ts |
362
+ | `selfImprove` | production | agent-app:src/eval-campaign/index.ts |
363
+ | `SEMANTIC_CONCEPT_JUDGE_VERSION` | production | blueprint-agent:scripts/experiments/lib/semantic-audit.ts |
364
+ | `ServedCrossFamilyError` | none | only this package's tests: src/integrity/served-model.test.ts:2 |
365
+ | `servedModelAcceptable` | production | this package: src/integrity/preflight.ts:33 |
366
+ | `spearmanR` | production | blueprint-agent:scripts/experiments/lib/benchmark-validity-report.mjs |
367
+ | `stripFencedJson` | production | blueprint-agent:scripts/experiments/compare-analysis-reports.ts |
368
+ | `subjectiveEval` | production | creative-agent:eval/control/creative-onboarding.ts |
369
+ | `summarizeBackendIntegrity` | production | agent-dev-container:products/intelligence/api/src/lib/run-provenance.ts |
370
+ | `summarizeNumberSeries` | production | this package: src/supervisor-run/analyze.ts:8 |
371
+ | `summarizePreferenceMemory` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
372
+ | `summaryTable` | production | blueprint-agent:scripts/experiments/vb-runrecord-analysis.ts |
373
+ | `textInSnapshot` | none | only this package's tests: src/oracle.test.ts:2 |
374
+ | `toAgentProfileJson` | production | blueprint-agent:scripts/experiments/lib/agent-profile-cell.ts |
375
+ | `tokenizeDomainWords` | planned | consumer tests: blueprint-agent:scripts/experiments/__tests__/analyze-vb-run.test.ts |
376
+ | `toolSpansToTraceAnalysisStore` | production | agent-runtime:src/runtime/supervise/trace-evidence.ts |
377
+ | `toolWasteView` | production | agent-runtime:src/runtime/supervise/trajectory-recorder.ts |
378
+ | `TRACE_ANALYST_TRUNCATION_MARKER_PREFIX` | production | agent-runtime:src/analyst-loop/iterations-to-trace-store.ts |
379
+ | `traceContract` | production | creative-agent:eval/lib/trace-contracts.ts |
380
+ | `TraceEmitter` | production | agent-builder:src/lib/.server/runtime/trace-runtime.ts |
381
+ | `transientDispatchFailure` | planned | doc: docs/eval-surface-map.md |
382
+ | `UNKNOWN_MODEL` | production | this package: src/campaign/presets/run-profile-matrix.ts:54 |
383
+ | `urlContains` | none | only this package's tests: src/oracle.test.ts:2 |
384
+ | `userQuestionsForKnowledgeGaps` | production | agent-dev-container:products/intelligence/api/src/routes/project-engines.ts |
385
+ | `validateRunRecord` | production | agent-builder:src/lib/.server/eval/stores/run-record-store.ts |
386
+ | `ValidationError` | production | agent-runtime:src/errors.ts |
387
+ | `verbosityBias` | production | blueprint-agent:scripts/experiments/calibration/judge-agreement.ts |
388
+ | `VERIFICATION_STRATEGIES` | planned | example: examples/verify-without-an-answer-key/index.ts:13 |
389
+ | `VERIFICATION_STRATEGY_SOURCES` | none | only this package's tests: src/verification-strategy.test.ts:11 |
390
+ | `verifyAgentProfileCell` | production | this package: src/eval-campaign.ts:39 |
391
+ | `verifyCompletion` | production | agent-app:src/eval/index.ts |
392
+ | `viteDeployRunner` | production | blueprint-agent:scripts/experiments/lib/deploy-runner.ts |
393
+ | `weightedComposite` | production | agent-app:src/eval/index.ts |
394
+ | `weightedMean` | production | blueprint-agent:scripts/experiments/lib/synthesis-quality-judge.ts |
395
+ | `WILCOXON_EXACT_MAX_N` | none | — |
396
+ | `wilcoxonSignedRank` | production | agent-builder:src/lib/.server/eval/loops/differential-eval.ts |
397
+ | `wilson` | production | agent-dev-container:products/intelligence/api/src/lib/project-analysis-coverage-recommendations.ts |
398
+ | `withAssignedFeedbackSplit` | production | phony:products/builder/api/src/feedback/production-trace.ts |
399
+ | `withHeldoutBlend` | none | only this package's tests: src/hidden-criteria-grading.test.ts:4 |
400
+ | `withJudgeRetry` | production | creative-agent:eval/lib/creative-judges.ts |
401
+ | `wranglerDeployRunner` | production | blueprint-agent:scripts/experiments/lib/deploy-runner.ts |
402
+
403
+ ### `./analyst`
404
+
405
+ 131 value exports — 75 production, 16 planned, 40 none.
406
+
407
+ | symbol | consumer | evidence |
408
+ | --- | --- | --- |
409
+ | `adaptPublicBenchmarkFindings` | production | this package: src/analyst/benchmark-public-rlm.ts:14 |
410
+ | `AGENT_RX_UPSTREAM_REVISION` | production | this package: src/analyst/benchmark-command-result.ts:2 |
411
+ | `agentRxBenchmarkCase` | production | this package: src/analyst/benchmark-public-data.ts:11 |
412
+ | `agentRxPredictionsToFindings` | production | this package: src/analyst/benchmark-public-adapters.ts:3 |
413
+ | `ANALYST_BENCHMARK_COST_LEDGER_FILE` | production | this package: src/analyst/benchmark-command-persistence.ts:9 |
414
+ | `ANALYST_BENCHMARK_DEPENDENCY_LOCK_DIGEST_ALGORITHM` | none | only this package's tests: src/analyst/benchmark-implementation.test.ts:7 |
415
+ | `ANALYST_BENCHMARK_DEPENDENCY_LOCK_FILES` | none | only this package's tests: src/analyst/benchmark-implementation.test.ts:7 |
416
+ | `ANALYST_BENCHMARK_DEPENDENCY_LOCK_SHA256` | production | this package: src/analyst/benchmark-command-persistence.ts:27 |
417
+ | `ANALYST_BENCHMARK_EVIDENCE_DEPENDENCY_LOCK_SHA256` | none | only this package's tests: src/analyst/benchmark-reference-result.test.ts:6 |
418
+ | `ANALYST_BENCHMARK_EVIDENCE_IMPLEMENTATION_SHA256` | none | only this package's tests: src/analyst/benchmark-reference-result.test.ts:6 |
419
+ | `ANALYST_BENCHMARK_IMPLEMENTATION_DIGEST_ALGORITHM` | none | only this package's tests: src/analyst/benchmark-implementation.test.ts:7 |
420
+ | `ANALYST_BENCHMARK_IMPLEMENTATION_FILES` | none | only this package's tests: src/analyst/benchmark-implementation.test.ts:7 |
421
+ | `ANALYST_BENCHMARK_IMPLEMENTATION_SHA256` | production | this package: src/analyst/benchmark-command-persistence.ts:27 |
422
+ | `ANALYST_BENCHMARK_LOCAL_RECEIPT_FILE` | production | this package: src/analyst/benchmark-command-persistence.ts:9 |
423
+ | `ANALYST_BENCHMARK_MANIFEST_FILE` | production | this package: src/analyst/benchmark-command-persistence.ts:9 |
424
+ | `ANALYST_BENCHMARK_OBSERVATIONS_FILE` | production | this package: src/analyst/benchmark-command-persistence.ts:9 |
425
+ | `analystBenchmarkDependencyLockDigest` | none | only this package's tests: src/analyst/benchmark-implementation.test.ts:7 |
426
+ | `analystBenchmarkImplementationDigest` | none | only this package's tests: src/analyst/benchmark-implementation.test.ts:7 |
427
+ | `analystDefinitionAsymmetries` | planned | doc: docs/trace-analysis.md |
428
+ | `analystDefinitionProtocolSha256` | planned | doc: docs/trace-analysis.md |
429
+ | `AnalystExpressivenessError` | production | this package: src/analyst/benchmark-public-model.ts:56 |
430
+ | `analystInstructionsOverrideFromText` | none | only this package's tests: src/analyst/benchmark-instructions-override.test.ts:5 |
431
+ | `AnalystRegistry` | production | blueprint-agent:scripts/experiments/vb-multi-analyst.ts |
432
+ | `analystUsageReceiptFromPrimeUsage` | production | this package: src/analyst/benchmark-runner-prime.ts:19 |
433
+ | `appendVerificationArtifactsToOtlp` | production | this package: src/analyst/benchmark-public-data.ts:31 |
434
+ | `assertProposalFindings` | production | agent-runtime:src/improvement/code-execution.ts |
435
+ | `behavioralAnalyst` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
436
+ | `bindAnalyst` | planned | doc: docs/trace-analysis.md |
437
+ | `buildDefaultAnalystRegistry` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
438
+ | `buildPrimePrompt` | production | this package: src/analyst/benchmark-runner-prime.ts:19 |
439
+ | `buildPrimeRepairPrompt` | planned | doc: docs/prime-analyst.md |
440
+ | `buildSkillUsageReport` | none | only this package's tests: src/analyst/kinds/skill-usage.test.ts:2 |
441
+ | `buildTraceToolsForGroup` | production | this package: src/analyst/kind-factory.ts:15 |
442
+ | `CODE_TRACE_BENCH_ANALYST_PROMPT` | production | this package: src/analyst/benchmark-runner-prime.ts:11 |
443
+ | `codeTraceBenchCase` | production | this package: src/analyst/benchmark-public-data.ts:11 |
444
+ | `codeTracerPredictionsToFindings` | none | only this package's tests: src/analyst/benchmark-datasets.test.ts:3 |
445
+ | `coerceJson` | production | this package: src/analyst/finding-signature.ts:10 |
446
+ | `compareAnalystRunners` | production | this package: src/analyst/benchmark-command-result.ts:15 |
447
+ | `computeFindingId` | production | agent-runtime:src/runtime/index.ts |
448
+ | `CONTROL_INTEGRITY_ANALYST` | none | only this package's tests: src/analyst/kinds/control-integrity.test.ts:5 |
449
+ | `ControlIntegrityAnalyst` | none | — |
450
+ | `createChatClient` | production | agent-dev-container:products/intelligence/api/src/lib/intent-audit-analyst.ts |
451
+ | `createDspyRlmTraceEngine` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
452
+ | `createJudgeAdapter` | none | — |
453
+ | `createPrimeBenchmarkRunner` | production | this package: src/analyst/benchmark-command.ts:62 |
454
+ | `createPublicBenchmarkDirectRunner` | production | this package: src/analyst/benchmark-command.ts:60 |
455
+ | `createPublicBenchmarkRlmRunner` | production | this package: src/analyst/benchmark-command.ts:61 |
456
+ | `createRunCriticAdapter` | none | — |
457
+ | `createSemanticConceptJudgeAdapter` | none | only this package's tests: src/analyst/adapters.test.ts:2 |
458
+ | `createTraceAnalyst` | production | blueprint-agent:scripts/experiments/lib/sandbox-driver/analyst-steer.ts |
459
+ | `createVerifierAdapter` | none | — |
460
+ | `decodeReplyRows` | production | this package: src/analyst/benchmark-public-model.ts:57 |
461
+ | `DEFAULT_MAX_VERIFICATION_ARTIFACT_BYTES` | production | this package: src/analyst/benchmark-command.ts:73 |
462
+ | `DEFAULT_TRACE_ANALYST_KINDS` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
463
+ | `defaultIsMaterial` | none | only this package's tests: src/analyst/analyst.test.ts:8 |
464
+ | `defineCustomAnalyst` | planned | doc: docs/trace-analysis.md |
465
+ | `defineTraceAnalyst` | planned | doc: docs/charter.md |
466
+ | `deriveEfficiencyFindings` | none | only this package's tests: src/trace-analyst/behavioral-metrics.test.ts:2 |
467
+ | `diffFindings` | production | agent-runtime:src/analyst-loop/run-analyst-loop.ts |
468
+ | `effectiveAnalystProtocolSha256` | production | this package: src/analyst/benchmark-command-persistence.ts:31 |
469
+ | `emitControlIntegrityFindings` | none | only this package's tests: src/analyst/kinds/control-integrity.test.ts:5 |
470
+ | `emitSkillUsageFindings` | none | only this package's tests: src/analyst/kinds/skill-usage.test.ts:2 |
471
+ | `emptyPrimeRawUsage` | none | only this package's tests: src/analyst/prime-protocol.test.ts:4 |
472
+ | `emptyPublicBenchmarkRunner` | production | this package: src/analyst/benchmark-command.ts:62 |
473
+ | `evidenceRefsFromRawFinding` | production | this package: src/analyst/benchmark-public-rlm.ts:36 |
474
+ | `ExactAnalystRunExecutionError` | planned | doc: docs/trace-analysis.md |
475
+ | `expandCodeTraceFailureBlocks` | production | this package: src/analyst/benchmark-public-model.ts:29 |
476
+ | `extractPrimeJsonObject` | none | only this package's tests: src/analyst/prime-protocol.test.ts:4 |
477
+ | `FAILURE_MODE_KIND_SPEC` | planned | consumer tests: legal-agent:tests/eval/analysts/completion-coverage.ts |
478
+ | `FINDING_SUBJECT_KINDS` | planned | consumer tests: tax-agent:tests/eval/benchmarks/taxcalc/compile-and-prove/carrier.test.ts |
479
+ | `FINDING_SUBJECT_SYNTAX` | none | only this package's tests: src/analyst/finding-subject.test.ts:2 |
480
+ | `FindingsStore` | production | agent-builder:scripts/run-canonical-analyst-loop.ts |
481
+ | `findingSubjectGrammarPromptFor` | production | this package: src/analyst/kinds/failure-mode.ts:38 |
482
+ | `IMPROVEMENT_KIND_SPEC` | none | only this package's tests: src/analyst/finding-subject.test.ts:11 |
483
+ | `isProposalFinding` | none | only this package's tests: src/analyst/proposal-findings.test.ts:2 |
484
+ | `KIND_EXPECTED_SUBJECTS` | production | this package: src/analyst/kind-factory.ts:14 |
485
+ | `KNOWLEDGE_GAP_KIND_SPEC` | none | only this package's tests: src/analyst/finding-subject.test.ts:11 |
486
+ | `KNOWLEDGE_POISONING_KIND_SPEC` | none | only this package's tests: src/analyst/finding-subject.test.ts:11 |
487
+ | `loadCodeTraceVerificationArtifacts` | production | this package: src/analyst/benchmark-public-data.ts:31 |
488
+ | `loadPublicBenchmarkRows` | none | only this package's tests: src/analyst/benchmark-command-public-data.test.ts:13 |
489
+ | `makeFinding` | production | agent-dev-container:products/intelligence/api/src/lib/consultant/candidates.ts |
490
+ | `makeProposalFinding` | production | agent-runtime:bench/src/swe-arena/outer-loop.mts |
491
+ | `MAX_INCORRECT_BLOCK_STEPS` | production | this package: src/analyst/benchmark-public-adapters.ts:9 |
492
+ | `MAX_INCORRECT_BLOCKS` | production | this package: src/analyst/benchmark-public-adapters.ts:9 |
493
+ | `mergePrimeRawUsage` | planned | doc: docs/prime-analyst.md |
494
+ | `nodeHttpPrimeBridgeTransport` | production | this package: src/analyst/benchmark-runner-prime.ts:18 |
495
+ | `normalizeAgentRxCategory` | none | only this package's tests: src/analyst/benchmark-datasets.test.ts:3 |
496
+ | `normalizeBenchmarkLabel` | production | this package: src/analyst/benchmark-dataset-agentrx.ts:10 |
497
+ | `normalizePrimeUsage` | planned | doc: docs/prime-analyst.md |
498
+ | `parseFindingSubject` | production | agent-dev-container:products/intelligence/api/src/lib/improvement-loop.ts |
499
+ | `parseRawFinding` | production | this package: src/analyst/kind-factory.ts:7 |
500
+ | `parseVerificationOutcome` | production | discovery-lab:tools/check-technical-campaign.mjs |
501
+ | `preparePublicAnalystBenchmark` | production | this package: src/analyst/benchmark-command.ts:62 |
502
+ | `primeAnalystProtocolSha256` | planned | doc: docs/prime-analyst.md |
503
+ | `primeCodeTraceAnalystDefinition` | planned | doc: docs/prime-analyst.md |
504
+ | `primeProtocolSha256` | production | this package: src/analyst/benchmark-runner-prime.ts:19 |
505
+ | `primeReplyDefect` | none | only this package's tests: src/analyst/prime-protocol.test.ts:4 |
506
+ | `projectPrimeTrajectory` | production | this package: src/analyst/benchmark-runner-prime.ts:19 |
507
+ | `publicBenchmarkDistributions` | none | — |
508
+ | `publicBenchmarkProtocolSha256` | production | this package: src/analyst/benchmark-instructions-override.ts:2 |
509
+ | `publicBenchmarkRlmInstructions` | production | this package: src/analyst/benchmark-public-rlm.ts:24 |
510
+ | `publicBenchmarkSelectionReport` | none | only this package's tests: src/analyst/benchmark-command-public-data.test.ts:13 |
511
+ | `publicBenchmarkSystemPrompt` | none | only this package's tests: src/analyst/benchmark-command-public-data.test.ts:13 |
512
+ | `publicDirectAnalystDefinition` | planned | doc: docs/trace-analysis.md |
513
+ | `publicRlmAnalystDefinition` | planned | doc: docs/trace-analysis.md |
514
+ | `RAW_FINDING_SCHEMA_PROMPT` | production | this package: src/analyst/benchmark-public-rlm.ts:36 |
515
+ | `RawAnalystFindingSchema` | production | this package: src/analyst/benchmark-public-rlm.ts:36 |
516
+ | `readAnalystBenchmarkArtifact` | production | this package: src/analyst/benchmark-command.ts:42 |
517
+ | `readAnalystInstructionsOverride` | production | this package: src/analyst/benchmark-command.ts:52 |
518
+ | `registryBenchmarkRunner` | none | only this package's tests: src/analyst/benchmark.test.ts:3 |
519
+ | `renderAgentRxCalibrationMarkdown` | production | this package: src/analyst/benchmark-command.ts:19 |
520
+ | `renderAnalystBenchmarkMarkdown` | production | this package: src/analyst/benchmark-command.ts:72 |
521
+ | `renderCodeTraceCalibrationMarkdown` | production | this package: src/analyst/benchmark-command.ts:56 |
522
+ | `renderFindingSubject` | production | creative-agent:eval/analyst-loop.ts |
523
+ | `renderPriorFindings` | none | — |
524
+ | `resolveTraceAnalystLimits` | production | this package: src/analyst/kind-factory.ts:5 |
525
+ | `rlmEngineLimits` | none | only this package's tests: src/analyst/definition-parity.test.ts:26 |
526
+ | `roundAgentRxStep` | production | this package: src/analyst/benchmark-agentrx-calibration.ts:2 |
527
+ | `runAnalystBenchmark` | production | this package: src/analyst/benchmark-command.ts:12 |
528
+ | `runAnalystBenchmarkCommand` | production | this package: src/cli.ts:16 |
529
+ | `runPrimeExchange` | production | this package: src/analyst/benchmark-runner-prime.ts:19 |
530
+ | `runTraceAnalyst` | production | this package: src/analyst/benchmark-public-rlm.ts:42 |
531
+ | `scoreAnalystFindings` | production | this package: src/analyst/benchmark.ts:3 |
532
+ | `selectPublicBenchmarkRows` | none | only this package's tests: src/analyst/benchmark-command-public-data.test.ts:13 |
533
+ | `SKILL_USAGE_ANALYST` | none | only this package's tests: src/analyst/kinds/skill-usage.test.ts:2 |
534
+ | `stripCodeFences` | planned | named in agent-builder (bind not in the import graph) |
535
+ | `summarizeAgentRxCalibration` | production | this package: src/analyst/benchmark-command-result.ts:2 |
536
+ | `summarizeAnalystBenchmarkRunner` | production | this package: src/analyst/benchmark-command-result.ts:18 |
537
+ | `summarizeCodeTraceCalibration` | production | this package: src/analyst/benchmark-command-result.ts:16 |
538
+ | `traceStoreEvidenceResolver` | production | this package: src/analyst/benchmark-command.ts:12 |
539
+
540
+ ### `./authenticity`
541
+
542
+ 5 value exports — 3 production, 0 planned, 2 none.
543
+
544
+ | symbol | consumer | evidence |
545
+ | --- | --- | --- |
546
+ | `gateRealness` | production | blueprint-agent:scripts/experiments/lib/enrich/realness.ts |
547
+ | `judgeRealnessLlm` | none | — |
548
+ | `scoreAuthenticity` | production | blueprint-agent:scripts/experiments/lib/enrich/realness.ts |
549
+ | `scoreAuthenticityNuance` | none | only this package's tests: src/authenticity/index.test.ts:3 |
550
+ | `scoreRealnessBlended` | production | physim:apps/server/src/scripts/agent-eval/investigators/realness.ts |
551
+
552
+ ### `./benchmarks`
553
+
554
+ 12 value exports — 2 production, 9 planned, 1 none.
555
+
556
+ | symbol | consumer | evidence |
557
+ | --- | --- | --- |
558
+ | `BENCHMARK_SPLIT_SEED` | production | tuner-agent:src/lib/eval-benchmarks.ts |
559
+ | `buildStandardRetrievalItems` | planned | doc: examples/benchmarks/README.md |
560
+ | `calibrateBenchmarkMetric` | planned | doc: examples/benchmarks/README.md |
561
+ | `createRetrievalIdBenchmarkAdapter` | planned | doc: examples/benchmarks/README.md |
562
+ | `deterministicSplit` | production | tuner-agent:src/lib/eval-benchmarks.ts |
563
+ | `evaluateStandardRetrieval` | planned | doc: examples/benchmarks/README.md |
564
+ | `parseBeirCorpusJsonl` | planned | doc: examples/benchmarks/README.md |
565
+ | `parseBeirQueriesJsonl` | planned | doc: examples/benchmarks/README.md |
566
+ | `parseJsonlRows` | none | only this package's tests: tests/benchmarks.test.ts:8 |
567
+ | `parseQrels` | planned | doc: examples/benchmarks/README.md |
568
+ | `routing` | planned | doc: docs/building-doctrine.md |
569
+ | `runBenchmarkAdapter` | planned | doc: examples/benchmarks/README.md |
570
+
571
+ ### `./builder-eval`
572
+
573
+ 6 value exports — 6 production, 0 planned, 0 none.
574
+
575
+ | symbol | consumer | evidence |
576
+ | --- | --- | --- |
577
+ | `BuilderSession` | production | agent-builder:src/lib/.server/eval/stores/session.ts |
578
+ | `correlateLayers` | production | agent-builder:src/lib/.server/eval/stores/session.ts |
579
+ | `ProjectRegistry` | production | agent-builder:src/lib/.server/eval/stores/session.ts |
580
+ | `resumeBuilderSession` | production | agent-builder:src/lib/.server/eval/stores/session.ts |
581
+ | `scoreAllProjects` | production | agent-builder:src/lib/.server/eval/stores/session.ts |
582
+ | `scoreProject` | production | agent-builder:src/lib/.server/eval/stores/session.ts |
583
+
584
+ ### `./campaign`
585
+
586
+ 117 value exports — 75 production, 15 planned, 27 none.
587
+
588
+ | symbol | consumer | evidence |
589
+ | --- | --- | --- |
590
+ | `acquireSingleRunLock` | production | agent-knowledge:src/memory/run-control.ts |
591
+ | `analyzeCrossSurfaceInteractions` | none | only this package's tests: src/campaign/cross-surface-interaction.test.ts:2 |
592
+ | `assertCampaignDesign` | production | this package: src/campaign/plan-campaign-run.ts:10 |
593
+ | `assertCampaignSplitIdentity` | production | agent-knowledge:src/memory/experiment/learning-pairs.ts |
594
+ | `assertCodeSurfaceIdentity` | none | — |
595
+ | `assertCompleteSearchHistory` | production | this package: src/campaign/presets/compare-optimization-methods.ts:31 |
596
+ | `assertSearchHistoryMatchesReplay` | planned | doc: docs/search-history-receipts.md |
597
+ | `autoevalsScorerJudge` | none | only this package's tests: src/campaign/upstream-evaluators.test.ts:6 |
598
+ | `buildEvidenceVector` | planned | doc: docs/experiment.md |
599
+ | `buildLoopProvenanceRecord` | none | only this package's tests: src/campaign/provenance-integrity.test.ts:5 |
600
+ | `buildTraceAnalystSurfaceDispatch` | none | only this package's tests: src/campaign/analyst-surface.test.ts:4 |
601
+ | `campaignBreakdown` | production | this package: src/campaign/external-text-evaluation.ts:11 |
602
+ | `campaignMeanComposite` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
603
+ | `campaignMeasurementDigest` | none | only this package's tests: src/campaign/provenance-integrity.test.ts:5 |
604
+ | `campaignScenarioIdentity` | production | agent-runtime:bench/src/swe-arena/premeasured-from-cells.mts |
605
+ | `campaignSplitDigest` | production | agent-runtime:bench/src/swe-arena/premeasured-from-cells.mts |
606
+ | `campaignSplitDigestFromIdentities` | production | agent-runtime:src/improvement/method-identity.ts |
607
+ | `canonicalDigest` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
608
+ | `cellCachePath` | production | discovery-lab:tools/strict-screen.mjs |
609
+ | `classifyUngroundedLiterals` | none | only this package's tests: src/campaign/grounded-reflection.test.ts:2 |
610
+ | `codeSurfaceIdentityMaterial` | none | — |
611
+ | `combineComparisonCosts` | production | this package: src/campaign/external-text-optimization.ts:41 |
612
+ | `compareOptimizationMethods` | production | agent-app:src/eval-campaign/index.ts |
613
+ | `compareRankKeys` | production | this package: src/campaign/presets/run-optimization.ts:29 |
614
+ | `componentSurfaceIdentityMaterial` | none | — |
615
+ | `composeGate` | production | discovery-lab:tools/confirmation-gate.mjs |
616
+ | `costFromLedgerSummary` | production | supervisor-lab:bench/drain/seat.ts |
617
+ | `createProfileMatrixPlan` | production | discovery-lab:tools/run-profile-confirmation.mjs |
618
+ | `createReferenceEquivalenceJudge` | none | only this package's tests: src/reference-equivalence-judge.test.ts:13 |
619
+ | `createRunCostLedger` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
620
+ | `createSearchHistoryReceipt` | production | this package: src/campaign/search-ledger-recording.ts:21 |
621
+ | `crowdedFrontierParent` | planned | doc: docs/campaign-proposers.md |
622
+ | `decodeExternalTextCandidate` | production | agent-runtime:src/improvement/method-execution.ts |
623
+ | `DEFAULT_EXTERNAL_OPTIMIZER_CALLBACK_LIMITS` | none | — |
624
+ | `DEFAULT_EXTERNAL_OPTIMIZER_PROCESS_LIMITS` | none | — |
625
+ | `defaultProductionGate` | production | agent-app:src/eval-campaign/index.ts |
626
+ | `detectScale` | production | this package: src/campaign/gates/promotion-policy.ts:39 |
627
+ | `dimensionRegressions` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-engine.ts |
628
+ | `discoverEvalFixtures` | planned | doc: docs/eval-fixtures.md |
629
+ | `emitLoopProvenance` | production | this package: src/contract/self-improve.ts:27 |
630
+ | `externalTextOptimizationMethod` | production | agent-app:src/eval-campaign/index.ts |
631
+ | `FileSearchLedger` | none | — |
632
+ | `finalizeProfileMatrix` | production | discovery-lab:tools/run-profile-confirmation.mjs |
633
+ | `fsCampaignStorage` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
634
+ | `FsLabeledScenarioStore` | production | agent-runtime:bench/src/swe-arena/outer-loop.mts |
635
+ | `gepaOptimizationMethod` | production | agent-app:src/eval-campaign/index.ts |
636
+ | `gitWorktreeAdapter` | production | agent-runtime:src/improvement/code-execution.ts |
637
+ | `heldOutGate` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-gate.ts |
638
+ | `heldoutSignificance` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
639
+ | `inMemoryCampaignStorage` | production | agent-dev-container:products/sandbox/evals/src/auto-optimization.ts |
640
+ | `isProposedCandidate` | production | supervisor-lab:bench/sequential-arena.ts |
641
+ | `isTransientTransportFailure` | production | gtm-agent:eval/suites/persona.ts |
642
+ | `LabeledScenarioStoreError` | none | only this package's tests: tests/campaign/run-campaign.test.ts:6 |
643
+ | `labelTrustRank` | production | this package: src/campaign/labeled-store/fs-adapter.ts:39 |
644
+ | `llmJudge` | production | agent-runtime:examples/agentic-data-creation/offline-fixtures.ts |
645
+ | `loadEvalFixture` | planned | doc: docs/eval-fixtures.md |
646
+ | `loadEvalFixtureScenarios` | planned | example: examples/eval-fixtures-quickstart/index.ts:11 |
647
+ | `loopProvenanceArgsFromResult` | production | this package: src/contract/self-improve.ts:27 |
648
+ | `loopProvenanceSpans` | none | only this package's tests: src/campaign/provenance-integrity.test.ts:5 |
649
+ | `makePlaybackDispatch` | none | only this package's tests: src/campaign/presets/playback.test.ts:6 |
650
+ | `makeProposalFinding` | production | agent-runtime:bench/src/swe-arena/outer-loop.mts |
651
+ | `neutralizationGate` | planned | consumer tests: tax-agent:tests/eval/benchmarks/taxcalc/evolve.ts |
652
+ | `neutralizeText` | planned | consumer tests: tax-agent:tests/eval/benchmarks/taxcalc/compile-and-prove/compile-diff.ts |
653
+ | `openAutoPr` | production | this package: src/campaign/presets/run-improvement-loop.ts:7 |
654
+ | `openSearchLedger` | production | this package: src/campaign/gepa-optimization-method.ts:66 |
655
+ | `optimizationTokenUsageFromSummary` | production | this package: src/campaign/external-text-optimization.ts:41 |
656
+ | `pairHoldout` | production | agent-dev-container:products/intelligence/api/src/lib/run-provenance.ts |
657
+ | `paretoPolicy` | none | only this package's tests: src/campaign/gates/promotion-policy.test.ts:3 |
658
+ | `paretoSignificanceGate` | production | agent-app:src/eval-campaign/index.ts |
659
+ | `phoenixEvaluatorJudge` | none | only this package's tests: src/campaign/upstream-evaluators.test.ts:6 |
660
+ | `planCampaignRun` | production | this package: src/campaign/fixtures.ts:5 |
661
+ | `planEvalFixtureRun` | planned | example: examples/eval-fixtures-quickstart/index.ts:11 |
662
+ | `powerPreflight` | production | discovery-lab:tools/run-sequential-profile-improvement.mjs |
663
+ | `ProfileMatrixError` | production | this package: src/campaign/presets/segmented-profile-matrix.ts:22 |
664
+ | `provenanceRecordPath` | planned | consumer tests: agent-dev-container:products/intelligence/api/tests/optimization-engine.test.ts |
665
+ | `provenanceSpansPath` | none | only this package's tests: tests/contract-self-improve.test.ts:25 |
666
+ | `readCachedCell` | production | discovery-lab:tools/strict-screen.mjs |
667
+ | `readExternalOptimizerObservationArtifact` | production | agent-runtime:src/improvement/method-execution.ts |
668
+ | `readGepaCandidatePopulationArtifact` | production | agent-runtime:src/improvement/method-execution.ts |
669
+ | `recordCandidatePopulationSearch` | production | this package: src/campaign/gepa-optimization-method.ts:67 |
670
+ | `renderScoreboardMarkdown` | production | blueprint-agent:scripts/experiments/eval/launch-scoreboard/run.ts |
671
+ | `renderSurfaceDiff` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-evidence-packet.ts |
672
+ | `resolveExternalOptimizerCallbackLimits` | production | this package: src/analyst/dspy-rlm-engine.ts:1 |
673
+ | `resolveExternalOptimizerProcessLimits` | production | this package: src/analyst/benchmark-command-persistence.ts:6 |
674
+ | `resolveRunDir` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
675
+ | `resolveWorktreePath` | none | only this package's tests: tests/campaign/worktree.test.ts:16 |
676
+ | `rolloutArgumentDiff` | none | only this package's tests: src/campaign/grounded-reflection.test.ts:2 |
677
+ | `runCampaign` | production | agent-app:src/eval-campaign/index.ts |
678
+ | `runEval` | production | ai-trading-blueprint:evals/src/sim/multishot-user-sim.ts |
679
+ | `runImprovementLoop` | production | blueprint-agent:scripts/experiments/vb-improve-trajectory.ts |
680
+ | `runOptimization` | production | agent-dev-container:products/sandbox/evals/src/auto-optimization.ts |
681
+ | `runProfileMatrix` | production | agent-runtime:examples/product-eval/product-eval.ts |
682
+ | `runProfileMatrixSegment` | production | discovery-lab:tools/run-profile-confirmation.mjs |
683
+ | `scoreboardSummary` | production | blueprint-agent:scripts/experiments/eval/launch-scoreboard/run.ts |
684
+ | `scoreDiscrimination` | production | supervisor-lab:bench/comms/seat-discrimination.ts |
685
+ | `scoreUserStory` | production | blueprint-agent:scripts/experiments/eval/launch-scoreboard/scoreboard-core.ts |
686
+ | `searchHistoryCoverageRow` | production | this package: src/campaign/presets/compare-optimization-methods.ts:31 |
687
+ | `SearchHistoryRequiredError` | none | only this package's tests: src/campaign/presets/compare-optimization-methods-history.test.ts:7 |
688
+ | `SearchLedgerConflictError` | production | this package: src/campaign/search-ledger.ts:33 |
689
+ | `SearchLedgerError` | production | this package: src/campaign/search-ledger.ts:33 |
690
+ | `SearchLedgerIntegrityError` | production | this package: src/campaign/search-ledger-file.ts:10 |
691
+ | `SearchRecorder` | production | this package: src/campaign/presets/run-optimization.ts:37 |
692
+ | `selectDiscriminative` | none | only this package's tests: src/campaign/scenario-selection.test.ts:2 |
693
+ | `sequentialDecide` | planned | doc: docs/experiment.md |
694
+ | `sequentialPairedGate` | planned | consumer tests: legal-agent:tests/eval/self-improve.ts |
695
+ | `skillOptOptimizationMethod` | production | agent-app:src/eval-campaign/index.ts |
696
+ | `surfaceContentHash` | production | discovery-lab:tools/run-profile-confirmation.mjs |
697
+ | `surfaceHash` | production | agent-builder:src/lib/.server/eval/loops/evolution-runner.ts |
698
+ | `tangleTracesRoot` | none | only this package's tests: src/campaign/run-dir.test.ts:4 |
699
+ | `traceAnalystQualityJudge` | none | only this package's tests: src/campaign/analyst-surface.test.ts:4 |
700
+ | `transientDispatchFailure` | planned | doc: docs/eval-surface-map.md |
701
+ | `userStoryScoreboard` | production | blueprint-agent:scripts/experiments/eval/launch-scoreboard/run.ts |
702
+ | `validateSearchLedgerEvent` | planned | consumer tests: discovery-lab:tools/search-ledger-planless.test.mjs |
703
+ | `verifyCodeSurface` | production | agent-runtime:src/candidate-execution/builder.ts |
704
+ | `verifyLoopProvenanceRecord` | none | only this package's tests: src/campaign/provenance-integrity.test.ts:5 |
705
+ | `verifySearchHistoryReceipt` | planned | doc: docs/search-history-receipts.md |
706
+ | `WorktreeAdapterError` | none | only this package's tests: tests/campaign/worktree.test.ts:16 |
707
+
708
+ ### `./contract`
709
+
710
+ 64 value exports — 40 production, 6 planned, 18 none.
711
+
712
+ | symbol | consumer | evidence |
713
+ | --- | --- | --- |
714
+ | `analyzeRuns` | production | agent-builder:eval/scripts/insight-report.ts |
715
+ | `buildDefaultAnalystRegistry` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
716
+ | `buildEvidenceVector` | planned | doc: docs/experiment.md |
717
+ | `campaignSplitDigest` | production | agent-runtime:bench/src/swe-arena/premeasured-from-cells.mts |
718
+ | `compareOptimizationMethods` | production | agent-app:src/eval-campaign/index.ts |
719
+ | `composeGate` | production | discovery-lab:tools/confirmation-gate.mjs |
720
+ | `createChatClient` | production | agent-dev-container:products/intelligence/api/src/lib/intent-audit-analyst.ts |
721
+ | `createReferenceEquivalenceJudge` | none | only this package's tests: src/reference-equivalence-judge.test.ts:13 |
722
+ | `defaultProductionGate` | production | agent-app:src/eval-campaign/index.ts |
723
+ | `defineAgentEval` | planned | example: examples/evaluate-a-change/index.ts:10 |
724
+ | `diffGenerations` | none | only this package's tests: tests/contract-diff.test.ts:2 |
725
+ | `diffRunBaselineToWinner` | none | only this package's tests: tests/contract-diff.test.ts:2 |
726
+ | `diffRuns` | none | only this package's tests: tests/contract-diff.test.ts:2 |
727
+ | `evalReportingSuite` | none | only this package's tests: tests/contract-eval-reporting-suite.test.ts:19 |
728
+ | `evaluatePairedMeasurements` | production | this package: src/contract/profile-measured-comparison.ts:26 |
729
+ | `externalTextOptimizationMethod` | production | agent-app:src/eval-campaign/index.ts |
730
+ | `FileSystemOutcomeStore` | production | phony:products/builder/api/src/eval/outcomes.ts |
731
+ | `fromClaudeCodeSession` | none | only this package's tests: tests/contract-code-agent-intake.test.ts:6 |
732
+ | `fromCodexSession` | none | only this package's tests: tests/contract-code-agent-intake.test.ts:6 |
733
+ | `fromFeedbackTable` | planned | example: examples/customer-feedback-loop/index.ts:12 |
734
+ | `fromKimiCodeSession` | none | only this package's tests: tests/contract-code-agent-intake.test.ts:6 |
735
+ | `fromOpenCodeSession` | none | only this package's tests: tests/contract-code-agent-intake.test.ts:6 |
736
+ | `fromOtelSpans` | production | agent-dev-container:packages/intelligence-kernel/src/deterministic.ts |
737
+ | `fromPiSession` | none | only this package's tests: tests/contract-code-agent-intake.test.ts:6 |
738
+ | `fromRunRecordDir` | production | this package: src/contract/eval-reporting-suite.ts:25 |
739
+ | `fsCampaignStorage` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
740
+ | `gepaOptimizationMethod` | production | agent-app:src/eval-campaign/index.ts |
741
+ | `heldOutGate` | production | agent-dev-container:products/intelligence/api/src/lib/optimization-gate.ts |
742
+ | `inMemoryCampaignStorage` | production | agent-dev-container:products/sandbox/evals/src/auto-optimization.ts |
743
+ | `InMemoryOutcomeStore` | planned | consumer tests: phony:products/builder/api/src/eval/outcomes-correlate.test.ts |
744
+ | `llmJudge` | production | agent-runtime:examples/agentic-data-creation/offline-fixtures.ts |
745
+ | `makeProposalFinding` | production | agent-runtime:bench/src/swe-arena/outer-loop.mts |
746
+ | `measuredComparisonFromAgentProfileImprovementExperiment` | production | agent-runtime:src/intelligence/authored-profile-improvement.ts |
747
+ | `measuredComparisonFromCandidateExperiment` | production | agent-runtime:src/intelligence/improvement-cycle.ts |
748
+ | `observeCodeAgentSession` | production | this package: src/contract/intake/code-agent-session.ts:13 |
749
+ | `paretoPolicy` | none | only this package's tests: src/campaign/gates/promotion-policy.test.ts:3 |
750
+ | `paretoSignificanceGate` | production | agent-app:src/eval-campaign/index.ts |
751
+ | `parseAgentTrace` | none | only this package's tests: tests/intake-agent-trace.test.ts:3 |
752
+ | `parseCodeAgentJsonl` | planned | named in discovery-lab (bind not in the import graph) |
753
+ | `parseCodeAgentJsonlFile` | none | only this package's tests: tests/contract-code-agent-intake.test.ts:6 |
754
+ | `partitionRunsByAuthoringModel` | none | only this package's tests: tests/intake-agent-trace.test.ts:3 |
755
+ | `REFERENCE_EQUIVALENCE_INPUT_LIMITS` | none | only this package's tests: src/reference-equivalence-judge.test.ts:13 |
756
+ | `REFERENCE_EQUIVALENCE_JUDGE_VERSION` | none | only this package's tests: src/reference-equivalence-judge.test.ts:13 |
757
+ | `runAgentProfileImprovementExperiment` | production | agent-runtime:src/intelligence/authored-profile-improvement.ts |
758
+ | `runCampaign` | production | agent-app:src/eval-campaign/index.ts |
759
+ | `runCandidateExperiment` | production | agent-runtime:src/intelligence/improvement-cycle.ts |
760
+ | `runEval` | production | ai-trading-blueprint:evals/src/sim/multishot-user-sim.ts |
761
+ | `runImprovementLoop` | production | blueprint-agent:scripts/experiments/vb-improve-trajectory.ts |
762
+ | `runReferenceEquivalenceJudge` | none | only this package's tests: src/reference-equivalence-judge.test.ts:13 |
763
+ | `sealAgentProfileImprovementExperiment` | production | agent-runtime:src/intelligence/authored-profile-improvement.ts |
764
+ | `sealAgentProfileImprovementSuite` | production | agent-runtime:src/intelligence/profile-improvement-experiment.ts |
765
+ | `sealAgentProfileImprovementTask` | production | agent-runtime:scripts/fixtures/packed-cohort-consumer.ts |
766
+ | `sealCandidateBenchmarkSuite` | production | agent-runtime:bench/scripts/verify-pier-agent.mts |
767
+ | `sealCandidateBenchmarkTask` | production | agent-runtime:bench/scripts/verify-pier-agent.mts |
768
+ | `sealCandidateExperiment` | production | agent-runtime:src/intelligence/improvement-cycle.ts |
769
+ | `selfImprove` | production | agent-app:src/eval-campaign/index.ts |
770
+ | `SelfImproveRunError` | production | agent-dev-container:products/intelligence/api/src/routes/optimizations.ts |
771
+ | `skillOptOptimizationMethod` | production | agent-app:src/eval-campaign/index.ts |
772
+ | `streamCodeAgentJsonlFile` | none | only this package's tests: tests/contract-code-agent-intake.test.ts:6 |
773
+ | `summarizeExecution` | production | traces:src/execution.ts |
774
+ | `transientDispatchFailure` | planned | doc: docs/eval-surface-map.md |
775
+ | `verifyAgentProfileImprovementExperimentComparison` | production | agent-runtime:src/intelligence/authored-profile-improvement.ts |
776
+ | `verifyCandidateExperiment` | production | agent-runtime:src/intelligence/improvement-cycle.ts |
777
+ | `verifyCandidateExperimentComparison` | production | agent-runtime:src/intelligence/improvement-cycle.ts |
778
+
779
+ ### `./experiment`
780
+
781
+ 83 value exports — 43 production, 22 planned, 18 none.
782
+
783
+ | symbol | consumer | evidence |
784
+ | --- | --- | --- |
785
+ | `amendExperiment` | planned | doc: docs/experiment.md |
786
+ | `assertDesignAdequate` | planned | doc: docs/experiment.md |
787
+ | `assertFunnelReconciles` | none | only this package's tests: tests/experiment/funnel.test.ts:7 |
788
+ | `assertMatchedBudgets` | planned | doc: docs/experiment.md |
789
+ | `benjaminiHochberg` | production | agent-runtime:bench/src/corpus-report.mts |
790
+ | `bonferroni` | planned | doc: docs/design/statistics-decisions.md |
791
+ | `BOOTSTRAP_GATE_MIN_N` | production | discovery-lab:tools/confirmation-activation.mjs |
792
+ | `buildEvidenceVector` | planned | doc: docs/experiment.md |
793
+ | `buildFunnel` | planned | doc: docs/experiment.md |
794
+ | `classifyReissue` | none | only this package's tests: tests/experiment/preregistration-acceptance.test.ts:11 |
795
+ | `clusteredPower` | planned | doc: docs/charter.md |
796
+ | `comparePairedArms` | production | agent-knowledge:src/memory/experiment/learning-metrics.ts |
797
+ | `composeFunnels` | planned | doc: docs/experiment.md |
798
+ | `computeEstimand` | production | this package: src/experiment/define.ts:21 |
799
+ | `computeInterval` | production | this package: src/experiment/define.ts:21 |
800
+ | `createEvidenceReceipt` | planned | consumer tests: agent-runtime:tests/integration/runtime-eval-pursuit-evidence.test.ts |
801
+ | `defineExperiment` | planned | doc: docs/charter.md |
802
+ | `DesignRefusalError` | none | only this package's tests: tests/experiment/power.test.ts:10 |
803
+ | `eProcess` | production | this package: src/campaign/gates/sequential.ts:35 |
804
+ | `evaluateCondition` | planned | named in agent-dev-container (bind not in the import graph) |
805
+ | `evaluateHaltRule` | production | this package: src/experiment/define.ts:21 |
806
+ | `evaluateHypothesis` | none | only this package's tests: tests/tier2.test.ts:8 |
807
+ | `evaluateIdentityGate` | production | this package: src/experiment/define.ts:21 |
808
+ | `evaluateOracleDeterminismGate` | production | this package: src/experiment/define.ts:21 |
809
+ | `evaluatePopulationReproducibilityGate` | production | this package: src/experiment/define.ts:21 |
810
+ | `evaluatePowerFloorGate` | production | this package: src/experiment/define.ts:21 |
811
+ | `evaluatePredicate` | production | this package: src/experiment/funnel.ts:17 |
812
+ | `evaluateProvenanceGate` | production | this package: src/experiment/define.ts:21 |
813
+ | `EVIDENCE_AUTHORITY_KINDS` | none | — |
814
+ | `EVIDENCE_RECEIPT_VERSION` | none | — |
815
+ | `EVIDENCE_STATES` | none | only this package's tests: src/experiment/evidence-record.test.ts:4 |
816
+ | `evidenceRegistryRecordSchema` | none | — |
817
+ | `executeAdmissionRule` | production | this package: src/experiment/define.ts:58 |
818
+ | `executeDecisionRule` | production | this package: src/experiment/define.ts:21 |
819
+ | `ExperimentTracker` | production | phony:products/builder/api/src/eval/evolution.ts |
820
+ | `fileExperimentStore` | production | phony:products/builder/api/src/eval/evolution.ts |
821
+ | `FunnelIntegrityError` | none | only this package's tests: tests/experiment/funnel.test.ts:7 |
822
+ | `hashJson` | production | agent-dev-container:products/intelligence/api/src/lib/ingest-optimization-mapping.ts |
823
+ | `heldoutSignificance` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
824
+ | `holm` | planned | doc: docs/design/statistics-decisions.md |
825
+ | `INDEPENDENT_EVIDENCE_AUTHORITY_KINDS` | none | — |
826
+ | `inMemoryExperimentStore` | none | only this package's tests: src/experiment-tracker.test.ts:2 |
827
+ | `isIndependentEvidence` | planned | consumer tests: agent-runtime:tests/integration/runtime-eval-pursuit-evidence.test.ts |
828
+ | `manifestContentDigest` | production | this package: src/campaign/gates/sequential.ts:34 |
829
+ | `MatchedBudgetError` | none | only this package's tests: tests/experiment/budget-and-seal.test.ts:8 |
830
+ | `mcnemar` | production | agent-runtime:bench/src/swe-arena/analyze.ts |
831
+ | `mcnemarPower` | production | supervisor-lab:bench/deepswe/headroom.ts |
832
+ | `mcnemarRequiredN` | production | blueprint-agent:scripts/experiments/analyze-convergence.ts |
833
+ | `mulberry32` | production | agent-knowledge:src/memory/holdout.ts |
834
+ | `openSealedExperiment` | planned | example: examples/sealed-experiment/index.ts:12 |
835
+ | `pairArms` | production | agent-knowledge:src/memory/experiment/learning-pairs.ts |
836
+ | `pairedBootstrap` | production | agent-builder:frontier/adapters/fleet-eval-runner.ts |
837
+ | `pairedEvalueSequence` | production | creative-agent:src/lib/experiments/ab-design.ts |
838
+ | `pairedMde` | production | discovery-lab:tools/design-gate.mjs |
839
+ | `pairedRiskDifference` | production | blueprint-agent:scripts/experiments/lib/validity-gates.ts |
840
+ | `pairedRiskDifferenceExact` | production | this package: src/paired-promotion-decision.ts:51 |
841
+ | `pairedRiskDifferenceScore` | production | this package: src/paired-promotion-decision.ts:51 |
842
+ | `pairHoldout` | production | agent-dev-container:products/intelligence/api/src/lib/run-provenance.ts |
843
+ | `pairRunRecords` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
844
+ | `paretoPolicy` | none | only this package's tests: src/campaign/gates/promotion-policy.test.ts:3 |
845
+ | `paretoSignificanceGate` | production | agent-app:src/eval-campaign/index.ts |
846
+ | `parseEvidenceRegistryRecord` | none | only this package's tests: src/experiment/evidence-record.test.ts:4 |
847
+ | `powerPreflight` | production | discovery-lab:tools/run-sequential-profile-improvement.mjs |
848
+ | `projectNLadderBudget` | production | this package: src/experiment/define.ts:21 |
849
+ | `readField` | planned | named in agent-dev-container (bind not in the import graph) |
850
+ | `renderEvidenceIndex` | none | only this package's tests: src/experiment/evidence-record.test.ts:4 |
851
+ | `renderFunnelTable` | planned | example: examples/sealed-experiment/index.ts:12 |
852
+ | `requiredPairedSampleSize` | production | discovery-lab:tools/design-gate.mjs |
853
+ | `requiredSampleSize` | planned | doc: docs/design/statistics-decisions.md |
854
+ | `runSelectionRule` | production | this package: src/experiment/define.ts:21 |
855
+ | `runUniformPassBudget` | production | this package: src/experiment/define.ts:21 |
856
+ | `sealExperiment` | planned | example: examples/sealed-experiment/index.ts:12 |
857
+ | `SealIntegrityError` | none | only this package's tests: tests/experiment/budget-and-seal.test.ts:8 |
858
+ | `sequentialCrossingHorizon` | none | only this package's tests: tests/sequential.test.ts:2 |
859
+ | `sequentialDecide` | planned | doc: docs/experiment.md |
860
+ | `sequentialPairedGate` | planned | consumer tests: legal-agent:tests/eval/self-improve.ts |
861
+ | `signManifest` | production | phony:products/builder/api/src/eval/champion-sign.ts |
862
+ | `validateEvidenceRegistry` | none | only this package's tests: src/experiment/evidence-record.test.ts:4 |
863
+ | `verifyEvidenceReceipt` | planned | consumer tests: agent-runtime:tests/integration/runtime-eval-pursuit-evidence.test.ts |
864
+ | `verifyManifest` | production | phony:products/builder/api/src/eval/champion-sign.ts |
865
+ | `verifyMatchedBudgets` | production | this package: src/experiment/define.ts:57 |
866
+ | `verifySealedExperiment` | planned | example: examples/sealed-experiment/index.ts:12 |
867
+ | `wilson` | production | agent-dev-container:products/intelligence/api/src/lib/project-analysis-coverage-recommendations.ts |
868
+
869
+ ### `./fuzz`
870
+
871
+ 14 value exports — 11 production, 1 planned, 2 none.
872
+
873
+ | symbol | consumer | evidence |
874
+ | --- | --- | --- |
875
+ | `adversarialObjective` | production | this package: src/fuzz/explorer.ts:22 |
876
+ | `BehaviorExplorer` | production | this package: src/fuzz/fuzz-agent.ts:9 |
877
+ | `buildCapsule` | production | this package: src/fuzz/explorer.ts:19 |
878
+ | `buildCoverage` | production | this package: src/fuzz/capsule.ts:12 |
879
+ | `cellId` | planned | consumer tests: tax-agent:tests/eval/lib/fuzz-space.ts |
880
+ | `composeGates` | production | creative-agent:eval/lib/fuzz-wiring.ts |
881
+ | `enumerateCells` | production | this package: src/fuzz/explorer.ts:21 |
882
+ | `fuzzAgent` | production | creative-agent:eval/fuzz.ts |
883
+ | `makeExploreTools` | none | only this package's tests: src/fuzz/fuzz-agent.test.ts:8 |
884
+ | `mutationProposer` | production | creative-agent:eval/fuzz.ts |
885
+ | `noveltyObjective` | none | only this package's tests: src/fuzz/fuzz-agent.test.ts:7 |
886
+ | `perturbationStabilityGate` | production | creative-agent:eval/lib/fuzz-wiring.ts |
887
+ | `renderCapsuleHtml` | production | creative-agent:eval/fuzz.ts |
888
+ | `severityFloorGate` | production | creative-agent:eval/lib/fuzz-wiring.ts |
889
+
890
+ ### `./hosted`
891
+
892
+ 11 value exports — 7 production, 2 planned, 2 none.
893
+
894
+ | symbol | consumer | evidence |
895
+ | --- | --- | --- |
896
+ | `createHostedClient` | production | traces:src/index.ts |
897
+ | `EvalRunEventSchema` | planned | example: examples/hosted-ingest-server/server.ts:33 |
898
+ | `HOSTED_WIRE_VERSION` | production | agent-dev-container:packages/sdk-data/src/sink/intelligence-sink.ts |
899
+ | `hostedClientFromEnv` | production | traces:src/index.ts |
900
+ | `hostedTenantFromEnv` | production | creative-agent:eval/self-improve.ts |
901
+ | `IngestEvalRunsRequestSchema` | production | this package: src/hosted/client.ts:19 |
902
+ | `IngestResponseSchema` | production | this package: src/hosted/client.ts:19 |
903
+ | `IngestTracesRequestSchema` | production | this package: src/hosted/client.ts:19 |
904
+ | `InsightReportSchema` | none | — |
905
+ | `TraceSpanEventSchema` | planned | example: examples/hosted-ingest-server/server.ts:33 |
906
+ | `UnixNanoTimestampSchema` | none | — |
907
+
908
+ ### `./ledger-core`
909
+
910
+ 16 value exports — 15 production, 0 planned, 1 none.
911
+
912
+ | symbol | consumer | evidence |
913
+ | --- | --- | --- |
914
+ | `appendLedgerLine` | production | this package: src/campaign/search-ledger-file.ts:3 |
915
+ | `AtomicFileLockError` | production | this package: src/ledger-core/journal-file.ts:16 |
916
+ | `canonicalString` | production | this package: src/agent-profile.ts:5 |
917
+ | `FileLedgerJournal` | production | this package: src/campaign/search-ledger.ts:22 |
918
+ | `hashCanonical` | production | discovery-lab:tools/provenance.mjs |
919
+ | `jsonDocument` | production | this package: src/analyst/benchmark-command-artifact.ts:2 |
920
+ | `LEDGER_HASH_PATTERN` | production | this package: src/ledger-core/trusted-head.ts:55 |
921
+ | `LedgerCanonicalizationError` | none | only this package's tests: src/ledger-core/canonical.test.ts:5 |
922
+ | `probeAtomicFileLock` | production | this package: src/campaign/single-run-lock.ts:17 |
923
+ | `readTrustedHeadFile` | production | this package: src/ledger-core/journal.ts:29 |
924
+ | `trustedHeadPathFor` | production | this package: src/ledger-core/journal.ts:29 |
925
+ | `tryAcquireAtomicFileLock` | production | this package: src/campaign/single-run-lock.ts:17 |
926
+ | `tryWithLedgerFileLock` | production | this package: src/campaign/search-ledger-file.ts:3 |
927
+ | `verifyEntriesAgainstTrustedHead` | production | this package: src/ledger-core/journal.ts:29 |
928
+ | `withLedgerFileLock` | production | this package: src/analyst/benchmark-response-cache.ts:6 |
929
+ | `writeLedgerFileAtomically` | production | this package: src/analyst/benchmark-response-cache.ts:6 |
930
+
931
+ ### `./matrix`
932
+
933
+ 3 value exports — 3 production, 0 planned, 0 none.
934
+
935
+ | symbol | consumer | evidence |
936
+ | --- | --- | --- |
937
+ | `readCellSpend` | production | this package: src/matrix/runner.ts:15 |
938
+ | `runAgentMatrix` | production | blueprint-agent:scripts/experiments/lib/vb-matrix.ts |
939
+ | `withCellSpend` | production | gtm-agent:eval/lib/multishot-graph.ts |
940
+
941
+ ### `./meta-eval`
942
+
943
+ 13 value exports — 10 production, 1 planned, 2 none.
944
+
945
+ | symbol | consumer | evidence |
946
+ | --- | --- | --- |
947
+ | `calibrationCurve` | none | only this package's tests: tests/meta-eval.test.ts:2 |
948
+ | `correlationStudy` | production | phony:products/builder/api/src/eval/outcomes.ts |
949
+ | `evalHealthStamp` | production | creative-agent:eval/lib/judge-sentinel.ts |
950
+ | `fileSentinelStore` | production | creative-agent:eval/lib/judge-sentinel.ts |
951
+ | `FileSystemOutcomeStore` | production | phony:products/builder/api/src/eval/outcomes.ts |
952
+ | `InMemoryOutcomeStore` | planned | consumer tests: phony:products/builder/api/src/eval/outcomes-correlate.test.ts |
953
+ | `inMemorySentinelStore` | none | only this package's tests: tests/judge-sentinel.test.ts:7 |
954
+ | `judgeSentinelReport` | production | creative-agent:eval/lib/judge-sentinel.ts |
955
+ | `rubricPredictiveValidity` | production | physim:apps/server/src/scripts/agent-eval/validate-rubrics.ts |
956
+ | `snapshotFromAgreement` | production | physim:apps/server/src/scripts/agent-eval/judge-sentinel.ts |
957
+ | `snapshotFromCalibration` | production | gtm-agent:eval/calibrate-judges.ts |
958
+ | `snapshotFromSentinelSet` | production | creative-agent:eval/lib/judge-sentinel.ts |
959
+ | `validateSentinelSnapshot` | production | creative-agent:eval/lib/judge-sentinel.ts |
960
+
961
+ ### `./multishot`
962
+
963
+ 18 value exports — 10 production, 4 planned, 4 none.
964
+
965
+ | symbol | consumer | evidence |
966
+ | --- | --- | --- |
967
+ | `assertMultishotShotResult` | production | this package: src/multishot/matrix.ts:17 |
968
+ | `computeCellComposite` | planned | consumer tests: tax-agent:tests/eval/multishot.ts |
969
+ | `DEFAULT_CODER_MODEL` | planned | named in blueprint-agent (bind not in the import graph) |
970
+ | `DEFAULT_JUDGE_MODEL` | planned | named in agent-builder (bind not in the import graph) |
971
+ | `defaultDelegationTools` | production | gtm-agent:eval/lib/multishot-graph.ts |
972
+ | `defaultMultishotDriverSystemPrompt` | none | only this package's tests: tests/multishot/shape-defaults.test.ts:10 |
973
+ | `defaultMultishotOpener` | none | only this package's tests: tests/multishot/shape-defaults.test.ts:10 |
974
+ | `defaultShapeFromProfile` | production | gtm-agent:eval/lib/multishot-graph.ts |
975
+ | `estimateMultishotCost` | production | this package: src/multishot/default-tools.ts:8 |
976
+ | `MultishotDriverEmptyError` | production | gtm-agent:eval/lib/multishot-graph.ts |
977
+ | `MultishotFatalToolError` | production | gtm-agent:eval/lib/multishot-graph.ts |
978
+ | `MultishotShotResultError` | none | only this package's tests: tests/multishot/matrix-shot-seam.test.ts:18 |
979
+ | `renderDimensions` | planned | consumer tests: tax-agent:tests/eval/multishot.ts |
980
+ | `renderJsonFooter` | production | gtm-agent:eval/scoring/multishot-judges.ts |
981
+ | `renderPersonaFacts` | none | only this package's tests: tests/multishot/shape-defaults.test.ts:10 |
982
+ | `runJudge` | production | gtm-agent:eval/monitoring/live-soak.ts |
983
+ | `runMultishot` | production | agent-runtime:examples/p1-parity/arms.ts |
984
+ | `runMultishotMatrix` | production | gtm-agent:eval/lib/multishot-matrix-graph.ts |
985
+
986
+ ### `./multishot/golden`
987
+
988
+ 21 value exports — 11 production, 8 planned, 2 none.
989
+
990
+ | symbol | consumer | evidence |
991
+ | --- | --- | --- |
992
+ | `assertMultishotGoldenScenario` | planned | consumer tests: gtm-agent:eval/lib/multishot-graph.test.ts |
993
+ | `assertMultishotMatrixGoldenScenario` | planned | consumer tests: gtm-agent:eval/lib/multishot-matrix-graph.test.ts |
994
+ | `checkMultishotGolden` | planned | doc: docs/multishot-golden-records.md |
995
+ | `checkMultishotGoldenScenario` | planned | doc: docs/multishot-golden-records.md |
996
+ | `checkMultishotMatrixGoldenScenario` | planned | doc: docs/multishot-golden-records.md |
997
+ | `compareJson` | production | this package: src/multishot/golden/harness.ts:10 |
998
+ | `CURRENT_MULTISHOT_GOLDEN_VERSION` | planned | doc: docs/multishot-golden-records.md |
999
+ | `goldenRecords` | production | this package: src/multishot/golden/harness.ts:23 |
1000
+ | `maskVolatileMarkdown` | none | only this package's tests: src/multishot/golden/golden.test.ts:20 |
1001
+ | `MultishotGoldenMismatchError` | planned | doc: docs/multishot-golden-records.md |
1002
+ | `multishotGoldenScenarios` | production | this package: src/multishot/golden/harness.ts:24 |
1003
+ | `multishotGoldenVersions` | none | — |
1004
+ | `multishotMatrixGoldenScenarios` | production | this package: src/multishot/golden/harness.ts:12 |
1005
+ | `readRunDir` | production | this package: src/multishot/golden/harness.ts:16 |
1006
+ | `recordError` | production | this package: src/multishot/golden/harness.ts:16 |
1007
+ | `recordJudgeRequest` | production | this package: src/multishot/golden/matrix-scenarios.ts:21 |
1008
+ | `recordMessage` | planned | named in agent-dev-container (bind not in the import graph) |
1009
+ | `recordRequest` | production | this package: src/multishot/golden/matrix-scenarios.ts:21 |
1010
+ | `recordResult` | production | this package: src/multishot/golden/harness.ts:16 |
1011
+ | `sortJudgeRequests` | production | this package: src/multishot/golden/harness.ts:16 |
1012
+ | `stripVolatile` | production | this package: src/multishot/golden/harness.ts:16 |
1013
+
1014
+ ### `./pipelines`
1015
+
1016
+ 8 value exports — 6 production, 0 planned, 2 none.
1017
+
1018
+ | symbol | consumer | evidence |
1019
+ | --- | --- | --- |
1020
+ | `budgetBreachView` | production | insurance-agent:scripts/analyze-agent-eval-evidence.ts |
1021
+ | `computeToolUseMetrics` | production | traces:src/pipelines.ts |
1022
+ | `failureClusterView` | production | insurance-agent:scripts/analyze-agent-eval-evidence.ts |
1023
+ | `firstDivergenceView` | none | only this package's tests: src/pipelines/first-divergence.test.ts:4 |
1024
+ | `judgeAgreementView` | production | insurance-agent:scripts/analyze-agent-eval-evidence.ts |
1025
+ | `regressionView` | none | only this package's tests: src/pipelines/regression.test.ts:4 |
1026
+ | `stuckLoopView` | production | agent-runtime:src/runtime/supervise/trajectory-recorder.ts |
1027
+ | `toolWasteView` | production | agent-runtime:src/runtime/supervise/trajectory-recorder.ts |
1028
+
1029
+ ### `./profile-cell`
1030
+
1031
+ 12 value exports — 5 production, 3 planned, 4 none.
1032
+
1033
+ | symbol | consumer | evidence |
1034
+ | --- | --- | --- |
1035
+ | `AGENT_PROFILE_KINDS` | production | blueprint-agent:scripts/experiments/lib/agent-profile-cell.ts |
1036
+ | `agentProfileCellHashMaterial` | planned | consumer tests: blueprint-agent:scripts/experiments/lib/__tests__/agent-profile-cell.test.ts |
1037
+ | `agentProfileCellKey` | planned | consumer tests: blueprint-agent:scripts/experiments/lib/__tests__/agent-profile-cell.test.ts |
1038
+ | `AgentProfileCellValidationError` | none | only this package's tests: tests/agent-profile-cell.test.ts:3 |
1039
+ | `assertRunAgentProfileCell` | none | only this package's tests: tests/agent-profile-cell.test.ts:3 |
1040
+ | `buildAgentInterfaceProfileCell` | none | only this package's tests: tests/agent-profile-cell.test.ts:3 |
1041
+ | `buildAgentProfileCell` | production | ai-trading-blueprint:evals/src/trading/agent-profile-cell.ts |
1042
+ | `groupRunsByAgentProfileCell` | planned | consumer tests: legal-agent:tests/eval/matrix-multi.ts |
1043
+ | `requireAgentProfileCell` | none | only this package's tests: tests/agent-profile-cell.test.ts:3 |
1044
+ | `toAgentProfileJson` | production | blueprint-agent:scripts/experiments/lib/agent-profile-cell.ts |
1045
+ | `validateAgentProfileCell` | production | blueprint-agent:apps/web/src/lib/.server/services/leaderboards/submission-profile-config.ts |
1046
+ | `verifyAgentProfileCell` | production | this package: src/eval-campaign.ts:39 |
1047
+
1048
+ ### `./reporting`
1049
+
1050
+ 15 value exports — 12 production, 2 planned, 1 none.
1051
+
1052
+ | symbol | consumer | evidence |
1053
+ | --- | --- | --- |
1054
+ | `assertReleaseConfidence` | none | only this package's tests: tests/release-confidence.test.ts:4 |
1055
+ | `benjaminiHochberg` | production | agent-runtime:bench/src/corpus-report.mts |
1056
+ | `bootstrapCi` | production | agent-dev-container:products/intelligence/api/src/lib/analysis-worker.ts |
1057
+ | `evaluateInterimReleaseConfidence` | production | agent-builder:scripts/eval.ts |
1058
+ | `evaluateReleaseConfidence` | production | agent-knowledge:src/release.ts |
1059
+ | `gainHistogram` | production | blueprint-agent:scripts/experiments/vb-runrecord-analysis.ts |
1060
+ | `judgeReplayGate` | planned | named in physim (bind not in the import graph) |
1061
+ | `pairedBootstrap` | production | agent-builder:frontier/adapters/fleet-eval-runner.ts |
1062
+ | `pairedEvalueSequence` | production | creative-agent:src/lib/experiments/ab-design.ts |
1063
+ | `paretoChart` | production | workcomp-agent:eval/benchmark/select.ts |
1064
+ | `RESEARCH_REPORT_HARD_PAIR_FLOOR` | planned | doc: docs/research-report-methodology.md |
1065
+ | `researchReport` | production | this package: src/eval-campaign.ts:58 |
1066
+ | `rubricPredictiveValidity` | production | physim:apps/server/src/scripts/agent-eval/validate-rubrics.ts |
1067
+ | `summaryTable` | production | blueprint-agent:scripts/experiments/vb-runrecord-analysis.ts |
1068
+ | `wilcoxonSignedRank` | production | agent-builder:src/lib/.server/eval/loops/differential-eval.ts |
1069
+
1070
+ ### `./rl`
1071
+
1072
+ 65 value exports — 18 production, 15 planned, 32 none.
1073
+
1074
+ | symbol | consumer | evidence |
1075
+ | --- | --- | --- |
1076
+ | `ABSENT_CATEGORY` | none | only this package's tests: src/rl/sim-fidelity.test.ts:3 |
1077
+ | `appendToCorpus` | none | only this package's tests: src/rl/corpus.test.ts:6 |
1078
+ | `applyEloUpdate` | none | only this package's tests: tests/rl-tournament.test.ts:2 |
1079
+ | `bestOfN` | planned | named in agent-runtime (bind not in the import graph) |
1080
+ | `bucketLabel` | planned | named in agent-dev-container (bind not in the import graph) |
1081
+ | `buildDatasetFromCorpus` | none | only this package's tests: src/rl/corpus.test.ts:6 |
1082
+ | `buildPairwiseFromCampaign` | none | only this package's tests: tests/rl-tournament.test.ts:2 |
1083
+ | `buildRlDataset` | production | this package: src/rl/corpus.ts:24 |
1084
+ | `buildVerifiedFindingRow` | none | only this package's tests: src/rl/verified-findings-dataset.test.ts:5 |
1085
+ | `campaignToRunRecords` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
1086
+ | `compareAdaptationCurves` | planned | doc: docs/design/statistics-decisions.md |
1087
+ | `datasheetToMarkdown` | none | only this package's tests: src/rl/dataset.test.ts:7 |
1088
+ | `defaultBehaviorFeatures` | none | only this package's tests: src/rl/sim-fidelity.test.ts:3 |
1089
+ | `detectRewardHacking` | production | agent-builder:scripts/eval.ts |
1090
+ | `doublyRobust` | planned | consumer tests: agent-knowledge:tests/memory-holdout.test.ts |
1091
+ | `easyModeCheck` | none | only this package's tests: src/rl/sim-fidelity.test.ts:3 |
1092
+ | `extractPreferences` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
1093
+ | `extractStepRewards` | none | only this package's tests: tests/rl-process-reward.test.ts:3 |
1094
+ | `extractVerifiableReward` | planned | named in gtm-agent (bind not in the import graph) |
1095
+ | `extractVerifiableRewardsFromRecords` | production | creative-agent:eval/canonical-export.ts |
1096
+ | `FileSystemOutcomeStore` | production | phony:products/builder/api/src/eval/outcomes.ts |
1097
+ | `filterDeterministicallyRewarded` | production | this package: src/rl/reward-hacking.ts:43 |
1098
+ | `firstPassK` | none | only this package's tests: tests/rl-adaptation-eval.test.ts:3 |
1099
+ | `fitBradleyTerry` | none | only this package's tests: tests/rl-tournament.test.ts:2 |
1100
+ | `injectIrrelevantClause` | none | only this package's tests: tests/rl-contamination.test.ts:2 |
1101
+ | `InMemoryOutcomeStore` | planned | consumer tests: phony:products/builder/api/src/eval/outcomes-correlate.test.ts |
1102
+ | `inverseProbabilityWeighting` | planned | consumer tests: agent-knowledge:tests/memory-holdout.test.ts |
1103
+ | `jsDivergence` | none | only this package's tests: src/rl/sim-fidelity.test.ts:3 |
1104
+ | `loadVerifiedFindingsDataset` | planned | doc: docs/verified-labels-flywheel.md |
1105
+ | `observationsFromRunRecords` | none | only this package's tests: tests/rl-active-curriculum.test.ts:3 |
1106
+ | `offPolicyEstimateAll` | none | only this package's tests: tests/rl-off-policy.test.ts:3 |
1107
+ | `paretoFrontier` | production | agent-dev-container:products/sandbox/evals/run.ts |
1108
+ | `PredictiveValidityResearcher` | none | only this package's tests: tests/rl-predictive-validity-researcher.test.ts:4 |
1109
+ | `prmTrainingPairs` | planned | doc: examples/fine-tune-with-prime-rl/README.md |
1110
+ | `quantileEdges` | none | only this package's tests: src/rl/sim-fidelity.test.ts:3 |
1111
+ | `readCorpus` | planned | named in agent-runtime (bind not in the import graph) |
1112
+ | `renameVariables` | none | only this package's tests: tests/rl-contamination.test.ts:2 |
1113
+ | `runAdaptationCurve` | none | only this package's tests: tests/rl-adaptation-eval.test.ts:3 |
1114
+ | `runComputeCurve` | none | only this package's tests: tests/rl-adversarial-and-compute.test.ts:2 |
1115
+ | `runContaminationProbe` | none | only this package's tests: tests/rl-contamination.test.ts:2 |
1116
+ | `runEvalCampaign` | production | agent-builder:scripts/eval.ts |
1117
+ | `runRLCampaign` | planned | named in creative-agent (bind not in the import graph) |
1118
+ | `runwiseStepRewardSummary` | none | only this package's tests: tests/rl-process-reward.test.ts:3 |
1119
+ | `selfConsistency` | none | only this package's tests: tests/rl-adversarial-and-compute.test.ts:2 |
1120
+ | `selfNormalizedImportanceWeighting` | planned | consumer tests: agent-knowledge:tests/memory-holdout.test.ts |
1121
+ | `shuffleOrder` | none | — |
1122
+ | `simFidelityReport` | none | only this package's tests: src/rl/sim-fidelity.test.ts:3 |
1123
+ | `stepRewardsToJsonl` | none | only this package's tests: src/rollout/gate-properties.test.ts:38 |
1124
+ | `summarizeVerifiedFindings` | none | only this package's tests: src/rl/verified-findings-dataset.test.ts:5 |
1125
+ | `thompsonCurriculum` | planned | doc: docs/design/statistics-decisions.md |
1126
+ | `toAnthropicFormat` | production | blueprint-agent:scripts/experiments/lib/steering-dataset.ts |
1127
+ | `toDpoJsonl` | production | this package: src/rl/dataset.ts:27 |
1128
+ | `toDpoRows` | production | this package: src/rl/dataset.ts:27 |
1129
+ | `toGrpoJsonl` | production | this package: src/rl/dataset.ts:27 |
1130
+ | `toGrpoRows` | production | this package: src/rl/dataset.ts:27 |
1131
+ | `toPrmJsonl` | none | only this package's tests: tests/rl-exporters.test.ts:2 |
1132
+ | `toPrmRows` | planned | doc: docs/rollout.md |
1133
+ | `toSftJsonl` | production | this package: src/rl/dataset.ts:27 |
1134
+ | `toSftRows` | production | this package: src/rl/dataset.ts:27 |
1135
+ | `toTRLFormat` | production | blueprint-agent:scripts/experiments/lib/steering-dataset.ts |
1136
+ | `validateDatasetFormats` | planned | example: examples/publish-rl-dataset/build-dataset.ts:30 |
1137
+ | `varianceBasedCurriculum` | production | this package: src/fuzz/explorer.ts:18 |
1138
+ | `verificationReportToRunRecord` | none | only this package's tests: tests/rl-adapters.test.ts:5 |
1139
+ | `VERIFIED_FINDING_SCHEMA` | none | only this package's tests: src/rl/verified-findings-dataset.test.ts:5 |
1140
+ | `verifiedFindingsToJsonl` | none | only this package's tests: src/rl/verified-findings-dataset.test.ts:5 |
1141
+
1142
+ ### `./rollout`
1143
+
1144
+ 67 value exports — 49 production, 8 planned, 10 none.
1145
+
1146
+ | symbol | consumer | evidence |
1147
+ | --- | --- | --- |
1148
+ | `addScrubCounts` | production | this package: src/rollout/release/hf-dataset.ts:34 |
1149
+ | `appendRolloutLines` | production | agent-runtime:bench/src/rollout-ledger/settle-capture.mts |
1150
+ | `assertGateReport` | production | this package: src/rollout/release/card.ts:12 |
1151
+ | `assertMinted` | production | blueprint-agent:scripts/experiments/lib/steering-dataset.ts |
1152
+ | `assertMintedLines` | planned | consumer tests: agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.test.mts |
1153
+ | `assertRolloutLine` | production | this package: src/rollout/interchange/harbor.ts:60 |
1154
+ | `ATIF_SCHEMA_VERSION` | none | only this package's tests: src/rollout/interchange/harbor.test.ts:6 |
1155
+ | `buildDatasetCard` | production | this package: src/rollout/release/hf-dataset.ts:25 |
1156
+ | `buildHfDataset` | production | supervisor-lab:bench/vertical-rollout-package.ts |
1157
+ | `claudeProjectSlug` | planned | consumer tests: agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.test.mts |
1158
+ | `DEFAULT_CLAUDE_PROJECTS_DIR` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1159
+ | `DEFAULT_OPENCODE_DB` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1160
+ | `defaultRolloutScrubber` | production | supervisor-lab:bench/vertical-rollout-package.ts |
1161
+ | `emptyScrubCounts` | production | this package: src/rollout/release/hf-dataset.ts:34 |
1162
+ | `findClaudeTranscripts` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1163
+ | `findOpencodeSessionsByDirectory` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1164
+ | `FORMAT_FILES` | production | this package: src/rollout/release/hf-dataset.ts:25 |
1165
+ | `FORMAT_GATE_DISPOSITION` | production | this package: src/rollout/release/card.ts:12 |
1166
+ | `fromHarborTrajectory` | planned | doc: docs/rollout.md |
1167
+ | `GATE_CHECK_IDS` | production | this package: src/rollout/release/gate-report.ts:47 |
1168
+ | `GATE_CHECKS` | planned | doc: docs/rollout.md |
1169
+ | `GATE_POLICIES` | production | this package: src/rollout/release/gate-report.ts:47 |
1170
+ | `gatedEvidenceOf` | production | this package: src/rollout/schema.ts:50 |
1171
+ | `gatedRolloutIds` | production | this package: src/rollout/release/hf-dataset.ts:26 |
1172
+ | `gateErrors` | production | this package: src/rollout/schema.ts:50 |
1173
+ | `gateGamedOutcome` | none | only this package's tests: src/rollout/gate-properties.test.ts:56 |
1174
+ | `HARBOR_IMPORT_GAP` | none | only this package's tests: src/rollout/interchange/harbor.test.ts:6 |
1175
+ | `isRealnessGated` | production | starter-foundry:src/lib/held-out-gate.ts |
1176
+ | `isRolloutLine` | production | this package: src/supervisor-run/rollout-nodes.ts:26 |
1177
+ | `isTrainableSplit` | production | this package: src/rollout/release/hf-dataset.ts:24 |
1178
+ | `measureFormatGate` | production | this package: src/rollout/release/hf-dataset.ts:26 |
1179
+ | `mintRolloutRows` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
1180
+ | `observedScore` | production | this package: src/rl/active-curriculum.ts:34 |
1181
+ | `observedSplitScore` | production | starter-foundry:src/lib/held-out-gate.ts |
1182
+ | `openOpencodeDb` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1183
+ | `parseRolloutReleaseArgs` | none | only this package's tests: src/rollout/release/hf-dataset.test.ts:9 |
1184
+ | `planPushCommand` | none | only this package's tests: src/rollout/release/hf-dataset.test.ts:9 |
1185
+ | `readClaudeTranscript` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1186
+ | `readOpencodeSessionMessages` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1187
+ | `readRolloutJournal` | none | only this package's tests: src/rollout/ledger.test.ts:6 |
1188
+ | `readRolloutLedger` | production | this package: src/rollout/release/hf-dataset.ts:23 |
1189
+ | `relabelImportedSplit` | planned | doc: docs/rollout.md |
1190
+ | `RELEASE_FORMATS` | production | this package: src/rollout/release/hf-dataset.ts:25 |
1191
+ | `releaseRowRefs` | production | this package: src/rollout/release/hf-dataset.ts:26 |
1192
+ | `ROLLOUT_CAPTURES` | production | this package: src/rollout/interchange/harbor.ts:60 |
1193
+ | `ROLLOUT_ROLES` | production | this package: src/rollout/interchange/harbor.ts:60 |
1194
+ | `ROLLOUT_SCHEMA` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1195
+ | `ROLLOUT_SPLITS` | production | this package: src/rollout/interchange/harbor.ts:60 |
1196
+ | `runRolloutReleaseCli` | production | this package: src/cli.ts:18 |
1197
+ | `scoreOrigin` | production | this package: src/rollout/mint.ts:38 |
1198
+ | `scrubLines` | production | supervisor-lab:bench/vertical-rollout-package.ts |
1199
+ | `scrubRolloutLine` | none | only this package's tests: src/rollout/release/scrub.test.ts:4 |
1200
+ | `scrubText` | none | only this package's tests: src/rollout/release/scrub.test.ts:4 |
1201
+ | `toHarborTrajectories` | planned | doc: docs/rollout.md |
1202
+ | `toHarborTrajectory` | planned | doc: docs/rollout.md |
1203
+ | `toJsonl` | production | supervisor-lab:bench/vertical-rollout-package.ts |
1204
+ | `toRewardRows` | planned | doc: docs/rollout.md |
1205
+ | `toRftItem` | none | only this package's tests: src/rollout/reward-invariant.test.ts:33 |
1206
+ | `toRftItems` | production | this package: src/rollout/release/hf-dataset.ts:22 |
1207
+ | `toSftRows` | production | this package: src/rl/dataset.ts:27 |
1208
+ | `toVerifiersRolloutOutput` | none | only this package's tests: src/rollout/exporters.test.ts:2 |
1209
+ | `toVerifiersRolloutOutputs` | production | this package: src/rollout/release/hf-dataset.ts:22 |
1210
+ | `trainingReward` | production | supervisor-lab:bench/vertical-gate-proof.ts |
1211
+ | `trainingScore` | production | this package: src/eval-trace-store.ts:20 |
1212
+ | `unmintableReasons` | production | supervisor-lab:bench/mint-audit.ts |
1213
+ | `validateRolloutLine` | production | this package: src/supervisor-run/integrity-tree.ts:1 |
1214
+ | `writeRolloutLedger` | production | agent-runtime:bench/src/rollout-ledger/backfill-swe-arena.mts |
1215
+
1216
+ ### `./storyboard`
1217
+
1218
+ 9 value exports — 5 production, 2 planned, 2 none.
1219
+
1220
+ | symbol | consumer | evidence |
1221
+ | --- | --- | --- |
1222
+ | `codeEditFromSpan` | none | only this package's tests: src/storyboard/code-edit.test.ts:4 |
1223
+ | `codeEditsForStoryboard` | none | only this package's tests: src/storyboard/code-edit.test.ts:4 |
1224
+ | `compileStoryboard` | production | run-capsule:src/index.ts |
1225
+ | `editAnimationText` | planned | named in run-capsule (bind not in the import graph) |
1226
+ | `extractCodeEdits` | production | run-capsule:src/index.ts |
1227
+ | `reduceToSemanticEvents` | production | run-capsule:src/artifacts.ts |
1228
+ | `renderCodeAnimationHtml` | planned | named in run-capsule (bind not in the import graph) |
1229
+ | `renderStoryboardHtml` | production | run-capsule:src/index.ts |
1230
+ | `renderStoryboardMarkdown` | production | run-capsule:src/index.ts |
1231
+
1232
+ ### `./supervisor-run`
1233
+
1234
+ 29 value exports — 21 production, 4 planned, 4 none.
1235
+
1236
+ | symbol | consumer | evidence |
1237
+ | --- | --- | --- |
1238
+ | `analyzeSupervisorRun` | production | discovery-lab:pursuits/agent-authored-runtime-control-deepseek-20260812h/workspaces/runtime-control-deps/runtime-control.mjs |
1239
+ | `analyzeSupervisorRunIntegrity` | production | discovery-lab:pursuits/final-agent-authored-campaign-20260810/materialize-candidates.mjs |
1240
+ | `analyzeSupervisorRunSources` | production | discovery-lab:experiments/0005-run-visibility/probe.mjs |
1241
+ | `claudeCodeSupervisorRunReader` | none | only this package's tests: src/supervisor-run/claude-code-reader.test.ts:6 |
1242
+ | `findSupervisorRunDirIn` | none | only this package's tests: src/supervisor-run/loops-reader.test.ts:11 |
1243
+ | `findSupervisorRunDirs` | production | traces:src/cli.ts |
1244
+ | `isRuntimeSupervisorRunDir` | production | discovery-lab:tools/disco.mjs |
1245
+ | `isUnavailable` | production | discovery-lab:tools/disco.mjs |
1246
+ | `loopsSupervisorRunReader` | planned | named in traces (bind not in the import graph) |
1247
+ | `NO_SOURCE_LIMITS` | production | traces:src/supervisor-run-context.ts |
1248
+ | `parsePatch` | planned | named in browser-agent-driver (bind not in the import graph) |
1249
+ | `parseSupervisorTree` | production | discovery-lab:experiments/0005-run-visibility/probe.mjs |
1250
+ | `readClaudeCodeSupervisorRun` | none | only this package's tests: src/supervisor-run/claude-code-reader.test.ts:6 |
1251
+ | `readLoopsSupervisorRun` | planned | named in discovery-lab (bind not in the import graph) |
1252
+ | `readRuntimeSupervisorRun` | production | discovery-lab:pursuits/final-agent-authored-campaign-20260810/materialize-candidates.mjs |
1253
+ | `renderSupervisorRollupMarkdown` | production | traces:src/cli.ts |
1254
+ | `renderSupervisorRunHeadline` | production | discovery-lab:experiments/0005-run-visibility/probe.mjs |
1255
+ | `renderSupervisorRunMarkdown` | production | agent-runtime:bench/src/swe-arena/run-report.mts |
1256
+ | `reportSupervisorRound` | production | agent-runtime:bench/src/swe-arena/outer-loop.mts |
1257
+ | `rollupSupervisorRuns` | production | discovery-lab:tools/disco.mjs |
1258
+ | `runtimeSupervisorRunReader` | none | only this package's tests: src/supervisor-run/index.test.ts:2 |
1259
+ | `showMeasured` | production | this package: src/supervisor-run/render.ts:8 |
1260
+ | `SUPERVISOR_RUN_INTEGRITY_SCHEMA` | production | this package: src/supervisor-run/integrity.ts:4 |
1261
+ | `SUPERVISOR_RUN_ROLLUP_SCHEMA` | production | this package: src/supervisor-run/analyze.ts:18 |
1262
+ | `SUPERVISOR_RUN_SCHEMA` | production | this package: src/supervisor-run/analyze.ts:18 |
1263
+ | `supervisorRunRolloutLines` | planned | consumer tests: discovery-lab:tools/runtime-journal-eval.test.mjs |
1264
+ | `unavailable` | production | this package: src/supervisor-run/analyze.ts:18 |
1265
+ | `writeSupervisorRunReport` | production | agent-runtime:bench/src/swe-arena/run-report.mts |
1266
+ | `writeSupervisorRunReportSafe` | production | agent-runtime:bench/src/swe-arena/outer-loop.mts |
1267
+
1268
+ ### `./trace-attributes`
1269
+
1270
+ 28 value exports — 26 production, 1 planned, 1 none.
1271
+
1272
+ | symbol | consumer | evidence |
1273
+ | --- | --- | --- |
1274
+ | `applyLlmSpanOtlpAttributes` | production | traces:src/adapters/claude.ts |
1275
+ | `asNumber` | planned | named in agent-app (bind not in the import graph) |
1276
+ | `contextInputTokens` | none | only this package's tests: tests/trace-otel.test.ts:3 |
1277
+ | `firstNumberAttr` | production | traces:src/file-export.ts |
1278
+ | `INPUT_VALUE` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1279
+ | `LLM_CACHE_WRITE_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1280
+ | `LLM_CACHE_WRITE_TOKENS` | production | traces:src/adapters/claude.ts |
1281
+ | `LLM_CACHED_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1282
+ | `LLM_CACHED_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1283
+ | `LLM_CONTEXT_TOKENS` | production | this package: src/trace-analyst/behavioral-metrics.ts:11 |
1284
+ | `LLM_COST_ATTR_KEYS` | production | traces:src/file-export.ts |
1285
+ | `LLM_COST_USD` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1286
+ | `LLM_INPUT_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1287
+ | `LLM_INPUT_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1288
+ | `LLM_MODEL_ATTR_KEYS` | production | this package: src/contract/intake/otel-spans.ts:44 |
1289
+ | `LLM_MODEL_NAME` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1290
+ | `LLM_OUTPUT_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1291
+ | `LLM_OUTPUT_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1292
+ | `LLM_REASONING_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1293
+ | `LLM_REASONING_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1294
+ | `OPENINFERENCE_SPAN_KIND` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1295
+ | `OUTPUT_VALUE` | production | agent-runtime:src/runtime/supervise-surface.ts |
1296
+ | `RUN_COST_ATTR_KEYS` | production | this package: src/trace/execution-measurements.ts:2 |
1297
+ | `SPAN_KIND_ATTR_KEYS` | production | traces:src/run-span-tree.ts |
1298
+ | `TOOL_ARGS_CAPTURED` | production | this package: src/trace/otlp-attributes.ts:3 |
1299
+ | `TOOL_LATENCY_MS` | production | this package: src/trace/otlp-attributes.ts:3 |
1300
+ | `TOOL_NAME` | production | supervisor-lab:bench/analyst/rewrite-views.mts |
1301
+ | `TOOL_NAME_ATTR_KEYS` | production | this package: src/trace/otlp-attributes.ts:3 |
1302
+
1303
+ ### `./trace-repair`
1304
+
1305
+ 96 value exports — 40 production, 31 planned, 25 none.
1306
+
1307
+ | symbol | consumer | evidence |
1308
+ | --- | --- | --- |
1309
+ | `ADMISSION_CONFIG_DEFAULTS` | planned | doc: docs/trace-repair-admission.md |
1310
+ | `ADMISSION_EXCLUSION_MEANING` | production | this package: src/trace-repair/admission-report.ts:11 |
1311
+ | `ADMISSION_EXCLUSION_ORDER` | production | this package: src/trace-repair/admission-report.ts:11 |
1312
+ | `ADMISSION_STRATA` | production | this package: src/trace-repair/admission-report.ts:11 |
1313
+ | `admissionArtifact` | planned | doc: docs/trace-repair-admission.md |
1314
+ | `AdmissionDenominatorError` | production | this package: src/trace-repair/admission.ts:36 |
1315
+ | `AdmissionIndependenceError` | none | only this package's tests: src/trace-repair/admission.test.ts:19 |
1316
+ | `admitRow` | production | this package: tests/trace-repair/fixtures.ts:11 |
1317
+ | `admittedCount` | planned | named in agent-knowledge (bind not in the import graph) |
1318
+ | `admittedRowIds` | none | only this package's tests: src/trace-repair/admission.test.ts:3 |
1319
+ | `askRepairArm` | planned | doc: docs/trace-repair-analyst-arms.md |
1320
+ | `assertAnalystIndependent` | production | this package: src/trace-repair/admission.ts:36 |
1321
+ | `assertArmSymmetry` | production | this package: src/trace-repair/admission.ts:52 |
1322
+ | `assertChainReconciles` | production | this package: src/trace-repair/admission-report.ts:11 |
1323
+ | `assertControlCalibrated` | production | this package: src/trace-repair/admission-contract.ts:39 |
1324
+ | `assertDenominatorIntact` | planned | doc: docs/design/statistics-decisions.md |
1325
+ | `assertDeterministicOracle` | planned | doc: docs/design/statistics-decisions.md |
1326
+ | `assertDspyRepairEngine` | none | only this package's tests: tests/trace-repair/arm-dspy.test.ts:19 |
1327
+ | `BLINDED_FIELDS` | none | only this package's tests: tests/trace-repair/admission.test.ts:3 |
1328
+ | `blindTrajectory` | production | this package: src/trace-repair/analyst-arm.ts:41 |
1329
+ | `buildDenominatorChain` | production | this package: src/trace-repair/admission.ts:36 |
1330
+ | `checkInterventionBudget` | production | this package: src/trace-repair/analyst-arm.ts:33 |
1331
+ | `classifyActionPayload` | none | only this package's tests: tests/trace-repair/action-budget.test.ts:2 |
1332
+ | `continuationPolicyDigest` | production | this package: src/trace-repair/admission.ts:52 |
1333
+ | `ContinuationPolicyViolationError` | none | only this package's tests: src/trace-repair/continuation-policy.test.ts:11 |
1334
+ | `continuationSeed` | planned | doc: docs/trace-repair-continuation.md |
1335
+ | `ContinuationSymmetryError` | none | only this package's tests: src/trace-repair/continuation-policy.test.ts:11 |
1336
+ | `CONTROL_SCREENING_MODES` | production | this package: src/trace-repair/admission.ts:58 |
1337
+ | `controlCanRescue` | none | only this package's tests: tests/trace-repair/control-policy.test.ts:4 |
1338
+ | `countFunnel` | production | this package: src/trace-repair/delta-repair.ts:27 |
1339
+ | `createCompletionRepairArm` | planned | doc: docs/trace-repair-analyst-arms.md |
1340
+ | `createDockerContinuationEnvironment` | planned | doc: docs/trace-repair-continuation.md |
1341
+ | `createDspyRepairArm` | planned | doc: docs/trace-repair-analyst-arms.md |
1342
+ | `CREDIT_TERMS` | none | only this package's tests: tests/trace-repair/degenerate-strategies.test.ts:8 |
1343
+ | `defineControlPolicy` | production | this package: tests/trace-repair/fixtures.ts:16 |
1344
+ | `definePinnedContinuationPolicy` | planned | doc: docs/trace-repair-admission.md |
1345
+ | `DEGENERATE_STRATEGIES` | planned | doc: docs/trace-repair-grader.md |
1346
+ | `deltaRepair` | planned | doc: README.md |
1347
+ | `dockerRunArgs` | none | only this package's tests: src/trace-repair/docker-environment.test.ts:3 |
1348
+ | `DSPY_REPAIR_SIGNATURE` | planned | doc: docs/trace-repair-analyst-arms.md |
1349
+ | `DSPY_REPAIR_TASK_TOKEN` | none | only this package's tests: tests/trace-repair/arm-dspy.test.ts:19 |
1350
+ | `dspyRepairInstructions` | none | only this package's tests: tests/trace-repair/arm-dspy.test.ts:19 |
1351
+ | `gradeRepairRow` | planned | doc: docs/trace-repair-analyst-arms.md |
1352
+ | `injectedTestOracle` | planned | doc: docs/trace-repair-grader.md |
1353
+ | `isRecordedTimeout` | production | this package: src/trace-repair/grade.ts:51 |
1354
+ | `MINI_SWE_SYSTEM_MESSAGE` | production | this package: src/trace-repair/continuation-policy.ts:28 |
1355
+ | `NO_DECISIVE_FAILURE` | none | only this package's tests: tests/trace-repair/analyst-response.test.ts:2 |
1356
+ | `NO_OP_ACTIONS` | none | — |
1357
+ | `nodeProcessRunner` | planned | doc: docs/trace-repair-continuation.md |
1358
+ | `NondeterministicOracleError` | none | only this package's tests: tests/trace-repair/oracle-determinism.test.ts:4 |
1359
+ | `noOpInjectionStep` | none | only this package's tests: src/trace-repair/admission.test.ts:3 |
1360
+ | `normalizeActionForComparison` | production | this package: src/trace-repair/grade.ts:31 |
1361
+ | `oracleDeterminism` | production | this package: tests/trace-repair/fixtures.ts:18 |
1362
+ | `OUTPUT_ELISION_THRESHOLD` | none | only this package's tests: src/trace-repair/mini-swe-scaffold.test.ts:2 |
1363
+ | `parseAction` | production | this package: src/trace-repair/continuation-policy.ts:28 |
1364
+ | `parseAnalystResponse` | planned | doc: docs/trace-repair-grader.md |
1365
+ | `parseTaskOracleRegistry` | planned | doc: docs/trace-repair-admission.md |
1366
+ | `renderAdmissionReport` | planned | doc: docs/trace-repair-admission.md |
1367
+ | `renderDeltaRepairReport` | planned | doc: docs/trace-repair-grader.md |
1368
+ | `renderFormatErrorObservation` | production | this package: src/trace-repair/continuation-policy.ts:28 |
1369
+ | `renderInstanceMessage` | production | this package: src/trace-repair/continuation-policy.ts:28 |
1370
+ | `renderObservation` | production | this package: src/trace-repair/continuation-policy.ts:28 |
1371
+ | `renderRepairTrajectory` | production | this package: src/trace-repair/arm-completion.ts:37 |
1372
+ | `renderTimeoutObservation` | production | this package: src/trace-repair/continuation-policy.ts:28 |
1373
+ | `REPAIR_CONTRACT_LINES` | production | this package: src/trace-repair/arm-completion.ts:37 |
1374
+ | `REPAIR_QUESTION` | production | this package: src/trace-repair/arm-completion.ts:37 |
1375
+ | `REPAIR_REPAIR_CONTRACT_LINES` | production | this package: src/trace-repair/arm-completion.ts:37 |
1376
+ | `repairArmAsymmetries` | planned | doc: docs/charter.md |
1377
+ | `repairArmPromptSha256` | production | this package: src/trace-repair/analyst-arm.ts:42 |
1378
+ | `repairArmResponse` | planned | doc: docs/trace-repair-analyst-arms.md |
1379
+ | `repairCredit` | production | this package: src/trace-repair/grade.ts:40 |
1380
+ | `repairFinding` | none | only this package's tests: tests/trace-repair/analyst-response.test.ts:2 |
1381
+ | `repairQuestionSha256` | production | this package: src/trace-repair/analyst-arm.ts:42 |
1382
+ | `repairTaskDefinition` | production | this package: src/trace-repair/arm-completion.ts:37 |
1383
+ | `repairTaskPolicy` | production | this package: src/trace-repair/arm-dspy.ts:48 |
1384
+ | `repairTrajectoryHeader` | production | this package: src/trace-repair/arm-completion.ts:37 |
1385
+ | `resolveAdmissionConfig` | none | only this package's tests: src/trace-repair/admission.test.ts:3 |
1386
+ | `rolloutDigest` | planned | doc: docs/trace-repair-continuation.md |
1387
+ | `rolloutRecordedSteps` | planned | doc: docs/trace-repair-continuation.md |
1388
+ | `runAdmission` | planned | doc: docs/trace-repair-admission.md |
1389
+ | `runContinuation` | planned | doc: docs/trace-repair-continuation.md |
1390
+ | `SCAFFOLD_INTERVENTION_BUDGET` | production | this package: src/trace-repair/arm-completion.ts:27 |
1391
+ | `scanShellAction` | none | only this package's tests: tests/trace-repair/action-budget.test.ts:2 |
1392
+ | `stratumOf` | production | this package: src/trace-repair/admission.ts:36 |
1393
+ | `submissionOf` | production | this package: src/trace-repair/continuation-policy.ts:28 |
1394
+ | `SUBMIT_SENTINEL` | production | this package: src/trace-repair/action-budget.ts:19 |
1395
+ | `SUITE_REWARD_UNIT` | none | only this package's tests: tests/trace-repair/oracle-determinism.test.ts:4 |
1396
+ | `taskOracleRegistry` | none | only this package's tests: src/trace-repair/admission-report.test.ts:7 |
1397
+ | `TB_REPAIR_ADMISSION_CRITERIA` | none | only this package's tests: tests/trace-repair/admission.test.ts:2 |
1398
+ | `TestOracleError` | none | only this package's tests: tests/trace-repair/test-oracle.test.ts:3 |
1399
+ | `testSuiteDigest` | production | this package: tests/trace-repair/fixtures.ts:34 |
1400
+ | `TestSuiteTamperedError` | planned | doc: docs/trace-repair-grader.md |
1401
+ | `toRecordedSteps` | none | only this package's tests: src/trace-repair/continuation-policy.test.ts:25 |
1402
+ | `totalCost` | planned | doc: examples/adapt-a-text-optimizer/README.md |
1403
+ | `totalUsage` | planned | named in agent-builder (bind not in the import graph) |
1404
+ | `UncalibratedControlError` | planned | doc: docs/trace-repair-admission.md |
1405
+
1406
+ ### `./traces`
1407
+
1408
+ 112 value exports — 94 production, 7 planned, 11 none.
1409
+
1410
+ | symbol | consumer | evidence |
1411
+ | --- | --- | --- |
1412
+ | `aggregateLlm` | production | this package: src/meta-eval/correlation-study.ts:14 |
1413
+ | `analyzeTraces` | production | agent-builder:src/lib/.server/eval/analysts/canonical-trace-analyst.ts |
1414
+ | `applyLlmSpanOtlpAttributes` | production | traces:src/adapters/claude.ts |
1415
+ | `applyToolSpanOtlpAttributes` | production | braid:src/adapters/analysis/trace-store.ts |
1416
+ | `argHash` | production | agent-runtime:src/runtime/supervise/detector-monitor.ts |
1417
+ | `asNumber` | planned | named in agent-app (bind not in the import graph) |
1418
+ | `assertRunCaptured` | production | creative-agent:eval/canonical-runner.ts |
1419
+ | `asString` | planned | named in agent-app (bind not in the import graph) |
1420
+ | `buildTraceAnalysisToolDescriptors` | production | this package: src/analyst/tool-groups.ts:14 |
1421
+ | `buildTraceInsightContext` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/report-html-sections.ts |
1422
+ | `buildTraceInsightPrompt` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/brief.ts |
1423
+ | `captureFetchToRawSink` | production | creative-agent:eval/canonical-runner.ts |
1424
+ | `classifyOtlpSpanRole` | production | this package: src/contract/intake/otel-spans.ts:44 |
1425
+ | `contextInputTokens` | none | only this package's tests: tests/trace-otel.test.ts:3 |
1426
+ | `convertTraceStoresToOtlp` | planned | consumer tests: legal-agent:tests/eval/lib/traces-to-otlp.ts |
1427
+ | `createBoundedTraceAnalysisStore` | production | braid:src/adapters/analysis/trace-store.ts |
1428
+ | `createOtelExporter` | planned | doc: docs/adapters-observability.md |
1429
+ | `createOtelTracingStore` | none | only this package's tests: tests/trace-contracts.test.ts:4 |
1430
+ | `DEFAULT_REDACTION_RULES` | production | agent-dev-container:products/intelligence/api/src/lib/redact.ts |
1431
+ | `DEFAULT_TRACE_ANALYST_BUDGETS` | production | agent-builder:src/lib/.server/eval/stores/d1-trace-analysis-store-adapter.ts |
1432
+ | `defaultProviderRedactor` | production | this package: src/llm-client.ts:34 |
1433
+ | `defaultTraceInsightPanel` | none | only this package's tests: src/trace-analyst/insights.test.ts:3 |
1434
+ | `describeTraceInsightScope` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/report-html.ts |
1435
+ | `domainEvidencePattern` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/session-loader.ts |
1436
+ | `exportRunAsOtlp` | production | agent-builder:src/lib/.server/eval/analysts/canonical-trace-analyst.ts |
1437
+ | `extractOtlpAttributes` | production | this package: src/trace-analyst/store-otlp.ts:40 |
1438
+ | `extractUsage` | production | this package: src/contract/intake/code-agent-session.ts:12 |
1439
+ | `extractUsageFromResponse` | none | only this package's tests: src/trace/extract-usage.test.ts:2 |
1440
+ | `extractUsageFromSse` | production | this package: src/trace/capture-fetch.ts:8 |
1441
+ | `FAILURE_CLASSES` | production | agent-dev-container:products/intelligence/api/src/lib/deep-trace-analyst.ts |
1442
+ | `FileSystemRawProviderSink` | production | agent-builder:scripts/eval.ts |
1443
+ | `FileSystemTraceStore` | production | agent-builder:scripts/eval.ts |
1444
+ | `firstNumberAttr` | production | traces:src/file-export.ts |
1445
+ | `firstStringAttr` | production | this package: src/trace-analyst/otlp-to-run-records.ts:63 |
1446
+ | `flattenOtlpExportToNdjson` | production | creative-agent:eval/trace-analyst-runner.ts |
1447
+ | `groupBy` | production | this package: src/tool-use-metrics.ts:10 |
1448
+ | `hasCapturedToolArgs` | production | this package: src/pipelines/failure-cluster.ts:10 |
1449
+ | `inferDomainKeywords` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/suite-projection.ts |
1450
+ | `InMemoryRawProviderSink` | production | agent-builder:src/lib/.server/eval/loops/canonical-campaign.ts |
1451
+ | `InMemoryTraceStore` | production | agent-builder:src/lib/.server/eval/loops/auto-research-runner.ts |
1452
+ | `INPUT_VALUE` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1453
+ | `isJudgeSpan` | production | this package: src/trace/query.ts:12 |
1454
+ | `isLlmSpan` | production | agent-runtime:src/candidate-execution/finalize.ts |
1455
+ | `isOtlpModelCall` | production | this package: src/contract/intake/otel-spans.ts:44 |
1456
+ | `isToolSpan` | production | agent-runtime:src/runtime/supervise/trace-evidence.ts |
1457
+ | `judgeSpans` | production | this package: src/builder-eval/three-layer-eval.ts:25 |
1458
+ | `LLM_CACHE_WRITE_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1459
+ | `LLM_CACHE_WRITE_TOKENS` | production | traces:src/adapters/claude.ts |
1460
+ | `LLM_CACHED_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1461
+ | `LLM_CACHED_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1462
+ | `LLM_CONTEXT_TOKENS` | production | this package: src/trace-analyst/behavioral-metrics.ts:11 |
1463
+ | `LLM_COST_ATTR_KEYS` | production | traces:src/file-export.ts |
1464
+ | `LLM_COST_USD` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1465
+ | `LLM_INPUT_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1466
+ | `LLM_INPUT_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1467
+ | `LLM_MODEL_ATTR_KEYS` | production | this package: src/contract/intake/otel-spans.ts:44 |
1468
+ | `LLM_MODEL_NAME` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1469
+ | `LLM_OUTPUT_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1470
+ | `LLM_OUTPUT_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1471
+ | `LLM_REASONING_TOKEN_ATTR_KEYS` | production | traces:src/file-export.ts |
1472
+ | `LLM_REASONING_TOKENS` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1473
+ | `llmSpans` | production | phony:products/builder/api/src/eval/replay.ts |
1474
+ | `NoopRawProviderSink` | production | agent-dev-container:products/intelligence/api/src/lib/eval-engine.ts |
1475
+ | `OPENINFERENCE_SPAN_KIND` | production | braid:src/adapters/analysis/trace-event-projection.ts |
1476
+ | `OTEL_AGENT_EVAL_SCOPE` | production | this package: src/trace/otel-export.ts:10 |
1477
+ | `OtlpFileTraceStore` | production | agent-builder:scripts/run-canonical-analyst-loop.ts |
1478
+ | `otlpRowsToRunRecords` | production | traces:src/execution.ts |
1479
+ | `otlpRowsToTraceRunRecords` | none | — |
1480
+ | `otlpTextToTraceAnalysisStore` | production | braid:src/adapters/analysis/trace-store.ts |
1481
+ | `otlpToRunRecords` | production | traces:src/execution.ts |
1482
+ | `otlpToTraceRunRecords` | none | only this package's tests: src/trace-analyst/otlp-to-run-records.test.ts:3 |
1483
+ | `OUTPUT_VALUE` | production | agent-runtime:src/runtime/supervise-surface.ts |
1484
+ | `planTraceInsightQuestions` | none | only this package's tests: src/trace-analyst/insights.test.ts:3 |
1485
+ | `projectOtlpFlatLine` | production | this package: src/trace-analyst/otlp-to-run-records.ts:63 |
1486
+ | `providerFromBaseUrl` | production | this package: src/llm-client.ts:34 |
1487
+ | `REDACTION_VERSION` | production | agent-dev-container:products/intelligence/api/src/lib/redact.ts |
1488
+ | `redactString` | production | agent-dev-container:products/intelligence/api/src/lib/redact.ts |
1489
+ | `redactValue` | production | traces:src/index.ts |
1490
+ | `RUN_COST_ATTR_KEYS` | production | this package: src/trace/execution-measurements.ts:2 |
1491
+ | `runFailureClass` | production | starter-foundry:registry/layers/agent-eval/redteam/files/src/eval/redteam/runner.ts |
1492
+ | `RunIntegrityError` | production | creative-agent:eval/canonical-runner.ts |
1493
+ | `runsForScenario` | production | phony:products/builder/api/src/eval/gate.ts |
1494
+ | `scoreTraceInsightReadiness` | production | blueprint-agent:scripts/experiments/lib/analyze-vb-run/report-html-sections.ts |
1495
+ | `SPAN_KIND_ATTR_KEYS` | production | traces:src/run-span-tree.ts |
1496
+ | `SpanNotFoundError` | production | this package: src/trace-analyst/store-boundary.ts:1 |
1497
+ | `stringField` | planned | named in agent-dev-container (bind not in the import graph) |
1498
+ | `throwIfRunIncomplete` | none | only this package's tests: tests/run-integrity.test.ts:3 |
1499
+ | `tokenizeDomainWords` | planned | consumer tests: blueprint-agent:scripts/experiments/__tests__/analyze-vb-run.test.ts |
1500
+ | `TOOL_ARGS_CAPTURED` | production | this package: src/trace/otlp-attributes.ts:3 |
1501
+ | `TOOL_LATENCY_MS` | production | this package: src/trace/otlp-attributes.ts:3 |
1502
+ | `TOOL_NAME` | production | supervisor-lab:bench/analyst/rewrite-views.mts |
1503
+ | `TOOL_NAME_ATTR_KEYS` | production | this package: src/trace/otlp-attributes.ts:3 |
1504
+ | `toolSpans` | production | this package: src/pipelines/failure-cluster.ts:10 |
1505
+ | `toolSpansToTraceAnalysisStore` | production | agent-runtime:src/runtime/supervise/trace-evidence.ts |
1506
+ | `ToolTraceMissingError` | none | only this package's tests: src/trace-analyst/store-tool-spans.test.ts:10 |
1507
+ | `TRACE_ANALYSIS_LIMITS` | production | this package: src/analyst/benchmark-evidence-validation.ts:1 |
1508
+ | `TRACE_ANALYST_ACTOR_DESCRIPTION` | production | this package: src/trace-analyst/analyst.ts:11 |
1509
+ | `TRACE_ANALYST_ACTOR_DESCRIPTION_VERSION` | production | this package: src/trace-analyst/analyst.ts:11 |
1510
+ | `TRACE_ANALYST_TOOL_NAMESPACE` | none | only this package's tests: src/trace-analyst/tools.test.ts:4 |
1511
+ | `TRACE_ANALYST_TRUNCATION_MARKER_PREFIX` | production | agent-runtime:src/analyst-loop/iterations-to-trace-store.ts |
1512
+ | `TRACE_SCHEMA_VERSION` | production | skeletal-os:src/eval/agentEvalBridge.ts |
1513
+ | `TraceAnalysisLimitError` | production | this package: src/trace-analyst/store-bounds.ts:3 |
1514
+ | `TraceAnalysisStoreContractError` | production | this package: src/trace-analyst/store-boundary.ts:1 |
1515
+ | `TraceAnalysisValidationError` | production | this package: src/trace-analyst/store-boundary.ts:1 |
1516
+ | `traceAnalystFunctionGroup` | none | only this package's tests: src/trace-analyst/tools.test.ts:4 |
1517
+ | `traceAnalystOnRunComplete` | planned | named in creative-agent (bind not in the import graph) |
1518
+ | `TraceEmitter` | production | agent-builder:src/lib/.server/runtime/trace-runtime.ts |
1519
+ | `TraceFileMalformedError` | production | this package: src/trace-analyst/store-otlp.ts:31 |
1520
+ | `TraceFileMissingError` | production | this package: src/trace-analyst/store-otlp.ts:31 |
1521
+ | `TraceFileTooLargeError` | production | this package: src/trace-analyst/store-otlp.ts:31 |
1522
+ | `TraceNotFoundError` | production | this package: src/trace-analyst/store-boundary.ts:1 |
1523
+ | `traceSpanKindToOpenInferenceKind` | production | this package: src/trace/otel-export.ts:11 |
1524
+
1525
+ ### `./trajectory-replay`
1526
+
1527
+ 49 value exports — 24 production, 19 planned, 6 none.
1528
+
1529
+ | symbol | consumer | evidence |
1530
+ | --- | --- | --- |
1531
+ | `assertReplayableTrajectory` | none | only this package's tests: tests/trajectory-replay/steps.test.ts:7 |
1532
+ | `buildFixPrompt` | production | this package: src/trajectory-replay/fix-loop.ts:17 |
1533
+ | `buildRetryFixPrompt` | production | this package: src/trajectory-replay/fix-loop.ts:17 |
1534
+ | `classifyObservation` | planned | doc: docs/trajectory-replay.md |
1535
+ | `classifyPrefixStep` | none | only this package's tests: tests/trajectory-replay/verify.test.ts:14 |
1536
+ | `classifyVerdict` | planned | named in agent-builder (bind not in the import graph) |
1537
+ | `clipText` | production | this package: src/trajectory-replay/fix-loop.ts:17 |
1538
+ | `countScriptCommands` | production | this package: src/trajectory-replay/fix-loop.ts:17 |
1539
+ | `decodeRecordedTurns` | planned | doc: docs/trajectory-replay.md |
1540
+ | `derivedImageTag` | planned | named in traces (bind not in the import graph) |
1541
+ | `deriveFailureSignature` | production | this package: src/trace-repair/grade.ts:25 |
1542
+ | `dockerImagePreparer` | production | this package: src/trajectory-replay/batch.ts:41 |
1543
+ | `enumerateReplayableCases` | production | this package: src/trajectory-replay/batch.ts:32 |
1544
+ | `extractFixCommand` | production | this package: src/trajectory-replay/fix-loop.ts:17 |
1545
+ | `finalRecordedOutcome` | planned | doc: docs/trajectory-replay.md |
1546
+ | `findingReplayStep` | planned | named in traces (bind not in the import graph) |
1547
+ | `findingTrajectoryId` | planned | named in traces (bind not in the import graph) |
1548
+ | `FORMAT_ERROR_OBSERVATION_PREFIX` | production | this package: src/trace-repair/mini-swe-scaffold.ts:17 |
1549
+ | `generateFixCommand` | production | this package: src/trajectory-replay/batch.ts:39 |
1550
+ | `goldIncorrectSteps` | planned | named in traces (bind not in the import graph) |
1551
+ | `ingestRecordedTrajectory` | production | this package: src/trajectory-replay/batch.ts:42 |
1552
+ | `isElidedField` | none | only this package's tests: tests/trajectory-replay/steps.test.ts:7 |
1553
+ | `isRecordedTimeout` | production | this package: src/trace-repair/grade.ts:51 |
1554
+ | `isSubmitAction` | production | this package: src/trajectory-replay/corpus.ts:19 |
1555
+ | `isSubmitOnlyAction` | none | only this package's tests: tests/trajectory-replay/steps.test.ts:7 |
1556
+ | `parseCorpusFlag` | planned | named in traces (bind not in the import graph) |
1557
+ | `parseIncorrectStepsSubject` | production | this package: src/trajectory-replay/findings.ts:41 |
1558
+ | `parseObservationOutput` | production | this package: src/trajectory-replay/corpus.ts:19 |
1559
+ | `parseRecordedReturncode` | production | this package: src/trace-repair/grade.ts:25 |
1560
+ | `PREFIX_DIVERGENCE_TOLERANCE_PCT` | production | this package: src/trajectory-replay/batch.ts:42 |
1561
+ | `readFindingsFile` | planned | named in traces (bind not in the import graph) |
1562
+ | `readLabelEntries` | planned | named in traces (bind not in the import graph) |
1563
+ | `renderBatchReport` | planned | named in traces (bind not in the import graph) |
1564
+ | `renderVerifiedFindingsSection` | planned | named in traces (bind not in the import graph) |
1565
+ | `replayVerify` | production | this package: src/trajectory-replay/batch.ts:42 |
1566
+ | `replayVerifyFinding` | planned | doc: docs/trajectory-replay.md |
1567
+ | `resolveCaseResources` | production | this package: src/trajectory-replay/findings.ts:36 |
1568
+ | `resolveFindingInvocation` | planned | named in traces (bind not in the import graph) |
1569
+ | `resolveFindingReplayability` | planned | named in traces (bind not in the import graph) |
1570
+ | `runFixLoop` | production | this package: src/trajectory-replay/batch.ts:40 |
1571
+ | `runReplayBatch` | planned | doc: docs/trajectory-replay.md |
1572
+ | `SandboxCounterfactualRunner` | production | this package: src/trajectory-replay/batch.ts:42 |
1573
+ | `seededSample` | planned | named in traces (bind not in the import graph) |
1574
+ | `SUBMIT_ACTION_SIGNATURE` | production | this package: src/trace-repair/mini-swe-scaffold.ts:17 |
1575
+ | `summarizePrefixReplay` | none | only this package's tests: tests/trajectory-replay/verify.test.ts:14 |
1576
+ | `TIMEOUT_OBSERVATION_MARKER` | production | this package: src/trace-repair/mini-swe-scaffold.ts:17 |
1577
+ | `unreadableExitCount` | none | only this package's tests: tests/trajectory-replay/steps.test.ts:7 |
1578
+ | `verifyFindings` | planned | doc: docs/trajectory-replay.md |
1579
+ | `wrapActionForExec` | production | this package: src/trace-repair/grade.ts:24 |
1580
+
1581
+ ### `./wire`
1582
+
1583
+ 30 value exports — 25 production, 3 planned, 2 none.
1584
+
1585
+ | symbol | consumer | evidence |
1586
+ | --- | --- | --- |
1587
+ | `buildOpenApi` | production | this package: src/cli.ts:20 |
1588
+ | `createApp` | planned | doc: docs/wire-protocol.md |
1589
+ | `dispatchRpc` | planned | doc: docs/wire-protocol.md |
1590
+ | `ErrorResponseSchema` | production | this package: src/wire/openapi.ts:15 |
1591
+ | `FeedbackIngestResponseSchema` | production | this package: src/wire/openapi.ts:15 |
1592
+ | `FeedbackTrajectorySchema` | production | insurance-agent:src/routes/api.feedback.ts |
1593
+ | `getBuiltinRubric` | production | this package: src/wire/handlers.ts:21 |
1594
+ | `handleFeedbackIngest` | production | this package: src/wire/server.ts:20 |
1595
+ | `handleJudge` | production | this package: src/wire/rpc.ts:15 |
1596
+ | `handleListRubrics` | production | this package: src/wire/rpc.ts:15 |
1597
+ | `handleTracesIngest` | production | this package: src/wire/server.ts:20 |
1598
+ | `handleVersion` | production | this package: src/cli.ts:19 |
1599
+ | `hashRubric` | production | this package: src/wire/handlers.ts:22 |
1600
+ | `HealthResponseSchema` | production | this package: src/wire/openapi.ts:15 |
1601
+ | `JudgeRequestSchema` | production | this package: src/wire/openapi.ts:15 |
1602
+ | `JudgeResultSchema` | production | this package: src/wire/openapi.ts:15 |
1603
+ | `listBuiltinRubrics` | production | this package: src/wire/handlers.ts:21 |
1604
+ | `ListRubricsResponseSchema` | production | this package: src/wire/openapi.ts:15 |
1605
+ | `RubricDimensionSchema` | none | only this package's tests: tests/wire/schemas.test.ts:11 |
1606
+ | `RubricSchema` | none | only this package's tests: tests/wire/schemas.test.ts:11 |
1607
+ | `runRpcBatch` | production | this package: src/cli.ts:21 |
1608
+ | `runRpcOnce` | production | this package: src/cli.ts:21 |
1609
+ | `startServer` | production | creative-agent:eval/ingestion-server.ts |
1610
+ | `startServerAsync` | production | this package: src/cli.ts:22 |
1611
+ | `TraceEventSchema` | planned | named in insurance-agent (bind not in the import graph) |
1612
+ | `TracesIngestRequestSchema` | production | this package: src/wire/openapi.ts:15 |
1613
+ | `TracesIngestResponseSchema` | production | this package: src/wire/openapi.ts:15 |
1614
+ | `VersionResponseSchema` | production | this package: src/wire/openapi.ts:15 |
1615
+ | `WIRE_VERSION` | production | this package: src/wire/handlers.ts:22 |
1616
+ | `WireError` | production | this package: src/wire/rpc.ts:15 |