dsh-aris-panel 0.1.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (442) hide show
  1. package/LICENSE +21 -0
  2. package/README.md +98 -0
  3. package/README_CN.md +87 -0
  4. package/dsh/checkout.patch.yml +38 -0
  5. package/dsh/client.js +634 -0
  6. package/dsh/cordis.patch.yml +44 -0
  7. package/dsh/index.mjs +76 -0
  8. package/dsh/run-status.mjs +182 -0
  9. package/dsh/scope-limits.mjs +50 -0
  10. package/dsh/workbench.mjs +291 -0
  11. package/mcp-servers/claude-review/README.md +93 -0
  12. package/mcp-servers/claude-review/run_with_claude_aws.sh +49 -0
  13. package/mcp-servers/claude-review/server.py +718 -0
  14. package/mcp-servers/codex-image2/README.md +65 -0
  15. package/mcp-servers/codex-image2/server.py +893 -0
  16. package/mcp-servers/feishu-bridge/requirements.txt +1 -0
  17. package/mcp-servers/feishu-bridge/server.py +240 -0
  18. package/mcp-servers/gemini-review/README.md +171 -0
  19. package/mcp-servers/gemini-review/server.py +1856 -0
  20. package/mcp-servers/llm-chat/requirements.txt +1 -0
  21. package/mcp-servers/llm-chat/server.py +664 -0
  22. package/mcp-servers/manual-review/README.md +133 -0
  23. package/mcp-servers/manual-review/server.py +910 -0
  24. package/mcp-servers/manual-review/ui.html +279 -0
  25. package/mcp-servers/minimax-chat/requirements.txt +1 -0
  26. package/mcp-servers/minimax-chat/server.py +381 -0
  27. package/package.json +51 -0
  28. package/skills/ablation-planner/SKILL.md +123 -0
  29. package/skills/alphaxiv/SKILL.md +196 -0
  30. package/skills/analyze-results/SKILL.md +46 -0
  31. package/skills/arxiv/SKILL.md +248 -0
  32. package/skills/auto-paper-improvement-loop/SKILL.md +651 -0
  33. package/skills/auto-review-loop/SKILL.md +1137 -0
  34. package/skills/auto-review-loop-llm/SKILL.md +259 -0
  35. package/skills/auto-review-loop-minimax/SKILL.md +302 -0
  36. package/skills/citation-audit/SKILL.md +502 -0
  37. package/skills/claims-drafting/SKILL.md +227 -0
  38. package/skills/comm-lit-review/SKILL.md +297 -0
  39. package/skills/deepxiv/SKILL.md +263 -0
  40. package/skills/dse-loop/SKILL.md +296 -0
  41. package/skills/embodiment-description/SKILL.md +129 -0
  42. package/skills/exa-search/SKILL.md +205 -0
  43. package/skills/experiment-audit/SKILL.md +311 -0
  44. package/skills/experiment-bridge/SKILL.md +376 -0
  45. package/skills/experiment-plan/SKILL.md +249 -0
  46. package/skills/experiment-queue/SKILL.md +431 -0
  47. package/skills/experiment-queue/scripts/build_manifest.py +142 -0
  48. package/skills/experiment-queue/scripts/queue_manager.py +433 -0
  49. package/skills/feishu-notify/SKILL.md +156 -0
  50. package/skills/figure-description/SKILL.md +138 -0
  51. package/skills/figure-spec/SKILL.md +262 -0
  52. package/skills/figure-spec/scripts/figure_renderer.py +799 -0
  53. package/skills/formula-derivation/SKILL.md +280 -0
  54. package/skills/gemini-search/SKILL.md +231 -0
  55. package/skills/grant-proposal/SKILL.md +698 -0
  56. package/skills/idea-creator/SKILL.md +542 -0
  57. package/skills/idea-discovery/SKILL.md +521 -0
  58. package/skills/idea-discovery-robot/SKILL.md +363 -0
  59. package/skills/integrity-forensics/SKILL.md +284 -0
  60. package/skills/interview-cheatsheet/SKILL.md +245 -0
  61. package/skills/invention-structuring/SKILL.md +188 -0
  62. package/skills/jurisdiction-format/SKILL.md +192 -0
  63. package/skills/kill-argument/SKILL.md +437 -0
  64. package/skills/mermaid-diagram/SKILL.md +419 -0
  65. package/skills/meta-apply/SKILL.md +141 -0
  66. package/skills/meta-optimize/SKILL.md +437 -0
  67. package/skills/monitor-experiment/SKILL.md +140 -0
  68. package/skills/novelty-check/SKILL.md +101 -0
  69. package/skills/openalex/SKILL.md +237 -0
  70. package/skills/overleaf-sync/SKILL.md +220 -0
  71. package/skills/paper-claim-audit/SKILL.md +348 -0
  72. package/skills/paper-compile/SKILL.md +266 -0
  73. package/skills/paper-figure/SKILL.md +312 -0
  74. package/skills/paper-illustration/SKILL.md +736 -0
  75. package/skills/paper-illustration-image2/SKILL.md +391 -0
  76. package/skills/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  77. package/skills/paper-plan/SKILL.md +386 -0
  78. package/skills/paper-poster/SKILL.md +19 -0
  79. package/skills/paper-poster-html/DESIGN_FINAL.md +176 -0
  80. package/skills/paper-poster-html/IMPLEMENTATION_CONVENTIONS.md +161 -0
  81. package/skills/paper-poster-html/LICENSES/posterly-MIT.txt +21 -0
  82. package/skills/paper-poster-html/NOTICE.md +57 -0
  83. package/skills/paper-poster-html/SKILL.md +323 -0
  84. package/skills/paper-poster-html/scripts/_posterly/__init__.py +0 -0
  85. package/skills/paper-poster-html/scripts/_posterly/canvas.py +200 -0
  86. package/skills/paper-poster-html/scripts/_posterly/measure.py +588 -0
  87. package/skills/paper-poster-html/scripts/_posterly/polish.py +498 -0
  88. package/skills/paper-poster-html/scripts/_posterly/preflight.py +489 -0
  89. package/skills/paper-poster-html/scripts/_posterly/render.py +215 -0
  90. package/skills/paper-poster-html/scripts/_posterly/textutil.py +16 -0
  91. package/skills/paper-poster-html/scripts/_posterly/verify_final.py +171 -0
  92. package/skills/paper-poster-html/scripts/asset_check.py +897 -0
  93. package/skills/paper-poster-html/scripts/extract_pdf_figures.py +666 -0
  94. package/skills/paper-poster-html/scripts/poster_check.py +251 -0
  95. package/skills/paper-poster-html/scripts/preprocess_figures.py +238 -0
  96. package/skills/paper-poster-html/scripts/render_preview.py +217 -0
  97. package/skills/paper-poster-html/scripts/run_gates.py +556 -0
  98. package/skills/paper-poster-html/scripts/style_check.py +1324 -0
  99. package/skills/paper-poster-html/templates/COMPONENTS.md +462 -0
  100. package/skills/paper-poster-html/templates/README.md +170 -0
  101. package/skills/paper-poster-html/templates/landscape_4col.html +1032 -0
  102. package/skills/paper-poster-html/templates/landscape_hero.html +1046 -0
  103. package/skills/paper-poster-html/templates/portrait_2col.html +947 -0
  104. package/skills/paper-poster-html/templates/tokens/acl.json +9 -0
  105. package/skills/paper-poster-html/templates/tokens/cvpr.json +9 -0
  106. package/skills/paper-poster-html/templates/tokens/generic.json +9 -0
  107. package/skills/paper-poster-html/templates/tokens/iclr.json +9 -0
  108. package/skills/paper-poster-html/templates/tokens/icml.json +9 -0
  109. package/skills/paper-poster-html/templates/tokens/neurips.json +9 -0
  110. package/skills/paper-slides/SKILL.md +635 -0
  111. package/skills/paper-talk/SKILL.md +381 -0
  112. package/skills/paper-write/SKILL.md +604 -0
  113. package/skills/paper-write/templates/IEEEtran.bst +2409 -0
  114. package/skills/paper-write/templates/IEEEtran.cls +6347 -0
  115. package/skills/paper-write/templates/iclr2026.tex +84 -0
  116. package/skills/paper-write/templates/icml2025.tex +87 -0
  117. package/skills/paper-write/templates/ieee_conference.tex +89 -0
  118. package/skills/paper-write/templates/ieee_journal.tex +93 -0
  119. package/skills/paper-write/templates/math_commands.tex +48 -0
  120. package/skills/paper-write/templates/neurips2025.tex +80 -0
  121. package/skills/paper-writing/SKILL.md +916 -0
  122. package/skills/patent-novelty-check/SKILL.md +153 -0
  123. package/skills/patent-pipeline/SKILL.md +344 -0
  124. package/skills/patent-review/SKILL.md +203 -0
  125. package/skills/pixel-art/SKILL.md +137 -0
  126. package/skills/prior-art-search/SKILL.md +146 -0
  127. package/skills/proof-checker/SKILL.md +866 -0
  128. package/skills/proof-orchestrator/NOTICE.md +24 -0
  129. package/skills/proof-orchestrator/SKILL.md +254 -0
  130. package/skills/proof-orchestrator/references/audit-output-contract.md +126 -0
  131. package/skills/proof-orchestrator/references/deepseek-routing.md +74 -0
  132. package/skills/proof-orchestrator/references/dispatch-prompts.md +227 -0
  133. package/skills/proof-orchestrator/references/notation-audit.md +135 -0
  134. package/skills/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  135. package/skills/proof-orchestrator/references/stress-tests.md +38 -0
  136. package/skills/proof-writer/SKILL.md +223 -0
  137. package/skills/qzcli/SKILL.md +324 -0
  138. package/skills/rebuttal/SKILL.md +376 -0
  139. package/skills/render-html/SKILL.md +316 -0
  140. package/skills/render-html/scripts/render_html.py +1006 -0
  141. package/skills/render-html/scripts/templates/academic.html +703 -0
  142. package/skills/render-html/scripts/templates/dashboard.html +333 -0
  143. package/skills/research-lit/SKILL.md +756 -0
  144. package/skills/research-pipeline/SKILL.md +384 -0
  145. package/skills/research-refine/SKILL.md +770 -0
  146. package/skills/research-refine-pipeline/SKILL.md +186 -0
  147. package/skills/research-review/SKILL.md +198 -0
  148. package/skills/research-wiki/SKILL.md +461 -0
  149. package/skills/resubmit-pipeline/SKILL.md +447 -0
  150. package/skills/result-to-claim/SKILL.md +311 -0
  151. package/skills/run-experiment/SKILL.md +313 -0
  152. package/skills/semantic-scholar/SKILL.md +236 -0
  153. package/skills/serverless-modal/SKILL.md +335 -0
  154. package/skills/shared-references/acceptance-gate.md +324 -0
  155. package/skills/shared-references/assurance-contract.md +248 -0
  156. package/skills/shared-references/capture-antipatterns.md +78 -0
  157. package/skills/shared-references/citation-discipline.md +583 -0
  158. package/skills/shared-references/compute-env-contract.md +163 -0
  159. package/skills/shared-references/effort-contract.md +183 -0
  160. package/skills/shared-references/evidence-precheck.md +65 -0
  161. package/skills/shared-references/experiment-integrity.md +49 -0
  162. package/skills/shared-references/external-cadence.md +326 -0
  163. package/skills/shared-references/fan-out-pattern.md +366 -0
  164. package/skills/shared-references/injection-hygiene.md +127 -0
  165. package/skills/shared-references/integration-contract.md +461 -0
  166. package/skills/shared-references/output-composition.md +93 -0
  167. package/skills/shared-references/output-language.md +45 -0
  168. package/skills/shared-references/output-manifest.md +49 -0
  169. package/skills/shared-references/output-versioning.md +111 -0
  170. package/skills/shared-references/patent-format-cn.md +199 -0
  171. package/skills/shared-references/patent-format-ep.md +173 -0
  172. package/skills/shared-references/patent-format-us.md +161 -0
  173. package/skills/shared-references/patent-writing-principles.md +197 -0
  174. package/skills/shared-references/prior-art-databases.md +141 -0
  175. package/skills/shared-references/resumable-runs.md +109 -0
  176. package/skills/shared-references/review-scope-limits.md +81 -0
  177. package/skills/shared-references/review-tracing.md +391 -0
  178. package/skills/shared-references/reviewer-independence.md +79 -0
  179. package/skills/shared-references/reviewer-routing.md +852 -0
  180. package/skills/shared-references/skill-governance.md +104 -0
  181. package/skills/shared-references/taste-calibration.md +85 -0
  182. package/skills/shared-references/venue-checklists.md +114 -0
  183. package/skills/shared-references/wiki-helper-resolution.md +134 -0
  184. package/skills/shared-references/writing-principles.md +525 -0
  185. package/skills/skills-codex/README.md +102 -0
  186. package/skills/skills-codex/README_CN.md +100 -0
  187. package/skills/skills-codex/ablation-planner/SKILL.md +126 -0
  188. package/skills/skills-codex/alphaxiv/SKILL.md +186 -0
  189. package/skills/skills-codex/analyze-results/SKILL.md +45 -0
  190. package/skills/skills-codex/arxiv/SKILL.md +210 -0
  191. package/skills/skills-codex/auto-paper-improvement-loop/SKILL.md +574 -0
  192. package/skills/skills-codex/auto-review-loop/SKILL.md +500 -0
  193. package/skills/skills-codex/auto-review-loop-llm/SKILL.md +247 -0
  194. package/skills/skills-codex/auto-review-loop-minimax/SKILL.md +290 -0
  195. package/skills/skills-codex/citation-audit/SKILL.md +504 -0
  196. package/skills/skills-codex/claims-drafting/SKILL.md +239 -0
  197. package/skills/skills-codex/comm-lit-review/SKILL.md +299 -0
  198. package/skills/skills-codex/comm-lit-review/references/domain-taxonomy.md +57 -0
  199. package/skills/skills-codex/comm-lit-review/references/output-template.md +37 -0
  200. package/skills/skills-codex/comm-lit-review/references/source-policy.md +99 -0
  201. package/skills/skills-codex/comm-lit-review/references/venue-tiering.md +112 -0
  202. package/skills/skills-codex/deepxiv/SKILL.md +142 -0
  203. package/skills/skills-codex/dse-loop/SKILL.md +285 -0
  204. package/skills/skills-codex/embodiment-description/SKILL.md +129 -0
  205. package/skills/skills-codex/exa-search/SKILL.md +192 -0
  206. package/skills/skills-codex/experiment-audit/SKILL.md +286 -0
  207. package/skills/skills-codex/experiment-bridge/SKILL.md +356 -0
  208. package/skills/skills-codex/experiment-plan/SKILL.md +249 -0
  209. package/skills/skills-codex/experiment-queue/SKILL.md +401 -0
  210. package/skills/skills-codex/feishu-notify/SKILL.md +155 -0
  211. package/skills/skills-codex/figure-description/SKILL.md +138 -0
  212. package/skills/skills-codex/figure-spec/SKILL.md +252 -0
  213. package/skills/skills-codex/formula-derivation/SKILL.md +280 -0
  214. package/skills/skills-codex/gemini-search/SKILL.md +205 -0
  215. package/skills/skills-codex/grant-proposal/SKILL.md +626 -0
  216. package/skills/skills-codex/idea-creator/SKILL.md +405 -0
  217. package/skills/skills-codex/idea-discovery/SKILL.md +475 -0
  218. package/skills/skills-codex/idea-discovery-robot/SKILL.md +362 -0
  219. package/skills/skills-codex/integrity-forensics/SKILL.md +106 -0
  220. package/skills/skills-codex/interview-cheatsheet/SKILL.md +245 -0
  221. package/skills/skills-codex/invention-structuring/SKILL.md +188 -0
  222. package/skills/skills-codex/jurisdiction-format/SKILL.md +192 -0
  223. package/skills/skills-codex/kill-argument/SKILL.md +403 -0
  224. package/skills/skills-codex/mermaid-diagram/SKILL.md +379 -0
  225. package/skills/skills-codex/meta-apply/SKILL.md +154 -0
  226. package/skills/skills-codex/meta-optimize/SKILL.md +348 -0
  227. package/skills/skills-codex/monitor-experiment/SKILL.md +98 -0
  228. package/skills/skills-codex/novelty-check/SKILL.md +89 -0
  229. package/skills/skills-codex/openalex/SKILL.md +228 -0
  230. package/skills/skills-codex/overleaf-sync/SKILL.md +220 -0
  231. package/skills/skills-codex/paper-claim-audit/SKILL.md +350 -0
  232. package/skills/skills-codex/paper-compile/SKILL.md +253 -0
  233. package/skills/skills-codex/paper-figure/SKILL.md +311 -0
  234. package/skills/skills-codex/paper-illustration/SKILL.md +690 -0
  235. package/skills/skills-codex/paper-illustration-image2/SKILL.md +383 -0
  236. package/skills/skills-codex/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  237. package/skills/skills-codex/paper-plan/SKILL.md +278 -0
  238. package/skills/skills-codex/paper-poster/SKILL.md +19 -0
  239. package/skills/skills-codex/paper-poster-html/SKILL.md +377 -0
  240. package/skills/skills-codex/paper-slides/SKILL.md +571 -0
  241. package/skills/skills-codex/paper-talk/SKILL.md +381 -0
  242. package/skills/skills-codex/paper-write/SKILL.md +411 -0
  243. package/skills/skills-codex/paper-write/templates/IEEEtran.bst +2409 -0
  244. package/skills/skills-codex/paper-write/templates/IEEEtran.cls +6347 -0
  245. package/skills/skills-codex/paper-write/templates/aaai2026.bst +1493 -0
  246. package/skills/skills-codex/paper-write/templates/aaai2026.sty +315 -0
  247. package/skills/skills-codex/paper-write/templates/aaai2026.tex +952 -0
  248. package/skills/skills-codex/paper-write/templates/acl.sty +312 -0
  249. package/skills/skills-codex/paper-write/templates/acl2026.tex +377 -0
  250. package/skills/skills-codex/paper-write/templates/acl_natbib.bst +1940 -0
  251. package/skills/skills-codex/paper-write/templates/acm.bst +3081 -0
  252. package/skills/skills-codex/paper-write/templates/acm_mm2026.tex +204 -0
  253. package/skills/skills-codex/paper-write/templates/acmart.cls +3520 -0
  254. package/skills/skills-codex/paper-write/templates/cvpr.bst +1448 -0
  255. package/skills/skills-codex/paper-write/templates/cvpr.sty +508 -0
  256. package/skills/skills-codex/paper-write/templates/cvpr2026.tex +63 -0
  257. package/skills/skills-codex/paper-write/templates/iclr2026.tex +84 -0
  258. package/skills/skills-codex/paper-write/templates/iclr2026_conference.bst +1440 -0
  259. package/skills/skills-codex/paper-write/templates/iclr2026_conference.sty +246 -0
  260. package/skills/skills-codex/paper-write/templates/icml2026.sty +767 -0
  261. package/skills/skills-codex/paper-write/templates/icml2026.tex +662 -0
  262. package/skills/skills-codex/paper-write/templates/ieee_conference.tex +89 -0
  263. package/skills/skills-codex/paper-write/templates/ieee_journal.tex +93 -0
  264. package/skills/skills-codex/paper-write/templates/math_commands.tex +48 -0
  265. package/skills/skills-codex/paper-write/templates/neurips2026.tex +493 -0
  266. package/skills/skills-codex/paper-write/templates/neurips_2026.sty +437 -0
  267. package/skills/skills-codex/paper-writing/SKILL.md +731 -0
  268. package/skills/skills-codex/patent-novelty-check/SKILL.md +153 -0
  269. package/skills/skills-codex/patent-pipeline/SKILL.md +344 -0
  270. package/skills/skills-codex/patent-review/SKILL.md +202 -0
  271. package/skills/skills-codex/pixel-art/SKILL.md +139 -0
  272. package/skills/skills-codex/prior-art-search/SKILL.md +146 -0
  273. package/skills/skills-codex/proof-checker/SKILL.md +554 -0
  274. package/skills/skills-codex/proof-orchestrator/SKILL.md +260 -0
  275. package/skills/skills-codex/proof-orchestrator/references/audit-output-contract.md +126 -0
  276. package/skills/skills-codex/proof-orchestrator/references/deepseek-routing.md +76 -0
  277. package/skills/skills-codex/proof-orchestrator/references/dispatch-prompts.md +227 -0
  278. package/skills/skills-codex/proof-orchestrator/references/notation-audit.md +135 -0
  279. package/skills/skills-codex/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  280. package/skills/skills-codex/proof-orchestrator/references/stress-tests.md +38 -0
  281. package/skills/skills-codex/proof-writer/SKILL.md +222 -0
  282. package/skills/skills-codex/qzcli/SKILL.md +324 -0
  283. package/skills/skills-codex/rebuttal/SKILL.md +305 -0
  284. package/skills/skills-codex/render-html/SKILL.md +305 -0
  285. package/skills/skills-codex/render-html/scripts/__pycache__/render_html.cpython-314.pyc +0 -0
  286. package/skills/skills-codex/render-html/scripts/render_html.py +909 -0
  287. package/skills/skills-codex/render-html/scripts/templates/academic.html +342 -0
  288. package/skills/skills-codex/render-html/scripts/templates/dashboard.html +333 -0
  289. package/skills/skills-codex/research-lit/SKILL.md +464 -0
  290. package/skills/skills-codex/research-pipeline/SKILL.md +340 -0
  291. package/skills/skills-codex/research-refine/SKILL.md +721 -0
  292. package/skills/skills-codex/research-refine-pipeline/SKILL.md +186 -0
  293. package/skills/skills-codex/research-review/SKILL.md +135 -0
  294. package/skills/skills-codex/research-wiki/SKILL.md +421 -0
  295. package/skills/skills-codex/resubmit-pipeline/SKILL.md +444 -0
  296. package/skills/skills-codex/result-to-claim/SKILL.md +246 -0
  297. package/skills/skills-codex/run-experiment/SKILL.md +236 -0
  298. package/skills/skills-codex/semantic-scholar/SKILL.md +219 -0
  299. package/skills/skills-codex/serverless-modal/SKILL.md +335 -0
  300. package/skills/skills-codex/shared-references/acceptance-gate.md +336 -0
  301. package/skills/skills-codex/shared-references/assurance-contract.md +139 -0
  302. package/skills/skills-codex/shared-references/capture-antipatterns.md +84 -0
  303. package/skills/skills-codex/shared-references/citation-discipline.md +452 -0
  304. package/skills/skills-codex/shared-references/compute-env-contract.md +163 -0
  305. package/skills/skills-codex/shared-references/effort-contract.md +143 -0
  306. package/skills/skills-codex/shared-references/evidence-precheck.md +73 -0
  307. package/skills/skills-codex/shared-references/experiment-integrity.md +49 -0
  308. package/skills/skills-codex/shared-references/external-cadence.md +334 -0
  309. package/skills/skills-codex/shared-references/fan-out-pattern.md +375 -0
  310. package/skills/skills-codex/shared-references/injection-hygiene.md +132 -0
  311. package/skills/skills-codex/shared-references/integration-contract.md +372 -0
  312. package/skills/skills-codex/shared-references/output-composition.md +98 -0
  313. package/skills/skills-codex/shared-references/output-language.md +45 -0
  314. package/skills/skills-codex/shared-references/output-manifest.md +40 -0
  315. package/skills/skills-codex/shared-references/output-versioning.md +111 -0
  316. package/skills/skills-codex/shared-references/patent-format-cn.md +199 -0
  317. package/skills/skills-codex/shared-references/patent-format-ep.md +173 -0
  318. package/skills/skills-codex/shared-references/patent-format-us.md +161 -0
  319. package/skills/skills-codex/shared-references/patent-writing-principles.md +197 -0
  320. package/skills/skills-codex/shared-references/prior-art-databases.md +141 -0
  321. package/skills/skills-codex/shared-references/resumable-runs.md +125 -0
  322. package/skills/skills-codex/shared-references/review-scope-limits.md +81 -0
  323. package/skills/skills-codex/shared-references/review-tracing.md +144 -0
  324. package/skills/skills-codex/shared-references/reviewer-independence.md +66 -0
  325. package/skills/skills-codex/shared-references/reviewer-routing.md +128 -0
  326. package/skills/skills-codex/shared-references/skill-governance.md +119 -0
  327. package/skills/skills-codex/shared-references/taste-calibration.md +90 -0
  328. package/skills/skills-codex/shared-references/venue-checklists.md +73 -0
  329. package/skills/skills-codex/shared-references/wiki-helper-resolution.md +69 -0
  330. package/skills/skills-codex/shared-references/writing-principles.md +525 -0
  331. package/skills/skills-codex/slides-polish/SKILL.md +563 -0
  332. package/skills/skills-codex/specification-writing/SKILL.md +211 -0
  333. package/skills/skills-codex/system-profile/SKILL.md +103 -0
  334. package/skills/skills-codex/training-check/SKILL.md +83 -0
  335. package/skills/skills-codex/vast-gpu/SKILL.md +394 -0
  336. package/skills/skills-codex/web-debug-search/SKILL.md +334 -0
  337. package/skills/skills-codex/wiki-enrich/SKILL.md +255 -0
  338. package/skills/skills-codex/writing-systems-papers/SKILL.md +184 -0
  339. package/skills/skills-codex-claude-review/README.md +79 -0
  340. package/skills/skills-codex-claude-review/README_CN.md +78 -0
  341. package/skills/skills-codex-claude-review/auto-paper-improvement-loop/SKILL.md +581 -0
  342. package/skills/skills-codex-claude-review/auto-review-loop/SKILL.md +510 -0
  343. package/skills/skills-codex-claude-review/novelty-check/SKILL.md +102 -0
  344. package/skills/skills-codex-claude-review/paper-figure/SKILL.md +319 -0
  345. package/skills/skills-codex-claude-review/paper-plan/SKILL.md +287 -0
  346. package/skills/skills-codex-claude-review/paper-write/SKILL.md +420 -0
  347. package/skills/skills-codex-claude-review/research-refine/SKILL.md +732 -0
  348. package/skills/skills-codex-claude-review/research-review/SKILL.md +149 -0
  349. package/skills/skills-codex-gemini-review/README.md +176 -0
  350. package/skills/skills-codex-gemini-review/README_CN.md +175 -0
  351. package/skills/skills-codex-gemini-review/auto-paper-improvement-loop/SKILL.md +331 -0
  352. package/skills/skills-codex-gemini-review/auto-review-loop/SKILL.md +304 -0
  353. package/skills/skills-codex-gemini-review/grant-proposal/SKILL.md +630 -0
  354. package/skills/skills-codex-gemini-review/idea-creator/SKILL.md +263 -0
  355. package/skills/skills-codex-gemini-review/idea-discovery/SKILL.md +275 -0
  356. package/skills/skills-codex-gemini-review/idea-discovery-robot/SKILL.md +365 -0
  357. package/skills/skills-codex-gemini-review/novelty-check/SKILL.md +92 -0
  358. package/skills/skills-codex-gemini-review/paper-figure/SKILL.md +289 -0
  359. package/skills/skills-codex-gemini-review/paper-plan/SKILL.md +265 -0
  360. package/skills/skills-codex-gemini-review/paper-poster-html/SKILL.md +102 -0
  361. package/skills/skills-codex-gemini-review/paper-slides/SKILL.md +582 -0
  362. package/skills/skills-codex-gemini-review/paper-write/SKILL.md +346 -0
  363. package/skills/skills-codex-gemini-review/paper-writing/SKILL.md +312 -0
  364. package/skills/skills-codex-gemini-review/research-refine/SKILL.md +674 -0
  365. package/skills/skills-codex-gemini-review/research-review/SKILL.md +112 -0
  366. package/skills/slides-polish/SKILL.md +565 -0
  367. package/skills/specification-writing/SKILL.md +211 -0
  368. package/skills/system-profile/SKILL.md +103 -0
  369. package/skills/training-check/SKILL.md +132 -0
  370. package/skills/vast-gpu/SKILL.md +394 -0
  371. package/skills/web-debug-search/SKILL.md +334 -0
  372. package/skills/wiki-enrich/SKILL.md +257 -0
  373. package/skills/writing-systems-papers/SKILL.md +184 -0
  374. package/templates/CLAUDE_MD_TEMPLATE.md +29 -0
  375. package/templates/EXPERIMENT_LOG_TEMPLATE.md +47 -0
  376. package/templates/EXPERIMENT_PLAN_TEMPLATE.md +51 -0
  377. package/templates/EXPERIMENT_PLAN_TEMPLATE_CN.md +53 -0
  378. package/templates/FINDINGS_TEMPLATE.md +52 -0
  379. package/templates/IDEA_CANDIDATES_TEMPLATE.md +47 -0
  380. package/templates/IDEA_CANDIDATES_TEMPLATE_CN.md +47 -0
  381. package/templates/INVENTION_BRIEF_TEMPLATE.md +87 -0
  382. package/templates/MANIFEST_TEMPLATE.md +7 -0
  383. package/templates/NARRATIVE_REPORT_TEMPLATE.md +49 -0
  384. package/templates/PAPER_PLAN_TEMPLATE.md +47 -0
  385. package/templates/PATENT_CLAIMS_TEMPLATE.md +78 -0
  386. package/templates/PATENT_SPECIFICATION_TEMPLATE.md +68 -0
  387. package/templates/README.md +57 -0
  388. package/templates/RESEARCH_BRIEF_TEMPLATE.md +35 -0
  389. package/templates/RESEARCH_BRIEF_TEMPLATE_CN.md +41 -0
  390. package/templates/RESEARCH_CONTRACT_TEMPLATE.md +60 -0
  391. package/templates/claude-hooks/corpus_write_guard.json +16 -0
  392. package/templates/claude-hooks/corpus_write_guard.py +85 -0
  393. package/templates/claude-hooks/meta_logging.json +74 -0
  394. package/templates/gitignore-trace.txt +3 -0
  395. package/tools/__pycache__/check_skills_inventory.cpython-314.pyc +0 -0
  396. package/tools/arxiv_fetch.py +311 -0
  397. package/tools/capture_filter.py +126 -0
  398. package/tools/check_skills_inventory.py +273 -0
  399. package/tools/convert_skills_to_llm_chat.py +282 -0
  400. package/tools/copilot_native_evidence.py +818 -0
  401. package/tools/deepxiv_fetch.py +213 -0
  402. package/tools/evidence_check.py +212 -0
  403. package/tools/exa_search.py +425 -0
  404. package/tools/experiment_queue/README.md +118 -0
  405. package/tools/experiment_queue/build_manifest.py +44 -0
  406. package/tools/experiment_queue/queue_manager.py +44 -0
  407. package/tools/extract_paper_style.py +560 -0
  408. package/tools/figure_renderer.py +69 -0
  409. package/tools/forensics_gate.py +669 -0
  410. package/tools/generate_codex_claude_review_overrides.py +299 -0
  411. package/tools/idea_discovery_gate.py +256 -0
  412. package/tools/install_aris.ps1 +1372 -0
  413. package/tools/install_aris.sh +1370 -0
  414. package/tools/install_aris_codex.sh +1023 -0
  415. package/tools/install_aris_copilot.sh +1052 -0
  416. package/tools/iteration_log.py +143 -0
  417. package/tools/lint_skills_helpers.sh +84 -0
  418. package/tools/meta_opt/check_ready.sh +80 -0
  419. package/tools/meta_opt/log_event.sh +91 -0
  420. package/tools/meta_opt/trigger_eval.py +280 -0
  421. package/tools/meta_opt/trigger_evals.sample.json +28 -0
  422. package/tools/openalex_fetch.py +326 -0
  423. package/tools/overleaf_audit.sh +104 -0
  424. package/tools/overleaf_setup.sh +150 -0
  425. package/tools/paper_illustration_image2.py +62 -0
  426. package/tools/provenance.py +294 -0
  427. package/tools/research_wiki.py +1720 -0
  428. package/tools/review_gate.py +502 -0
  429. package/tools/run_state.py +399 -0
  430. package/tools/save_trace.sh +477 -0
  431. package/tools/semantic_scholar_fetch.py +438 -0
  432. package/tools/skill-groups.tsv +116 -0
  433. package/tools/skill_picker.py +238 -0
  434. package/tools/smart_update.ps1 +521 -0
  435. package/tools/smart_update.sh +591 -0
  436. package/tools/smart_update_codex.sh +419 -0
  437. package/tools/smart_update_copilot.sh +605 -0
  438. package/tools/threat_scan.py +222 -0
  439. package/tools/verify_paper_audits.sh +487 -0
  440. package/tools/verify_papers.py +613 -0
  441. package/tools/verify_wiki_coverage.sh +176 -0
  442. package/tools/watchdog.py +485 -0
@@ -0,0 +1,391 @@
1
+ # Review Tracing Protocol
2
+
3
+ ## Purpose
4
+
5
+ Save full prompt/response pairs for every cross-model reviewer call, enabling:
6
+ - **Reviewer-independence audit**: verify the executor only passed file paths, not summaries
7
+ - **Reproducibility**: threadId preservation allows conversation continuation
8
+ - **Meta-optimize input**: richer data for harness improvement analysis
9
+
10
+ ## When to Trace
11
+
12
+ After **every** native `task(agent_type=rubber-duck)` (`copilot-native`),
13
+ `mcp__codex__codex`, `mcp__codex__codex-reply`, compatibility
14
+ `copilot --agent`, `mcp__manual_review__review`, or
15
+ `mcp__manual_review__review_reply` call that serves a reviewer/critique
16
+ function. This includes review scoring, experiment auditing, claim
17
+ verification, idea critique, and patch gating.
18
+
19
+ Do NOT trace: purely informational LLM calls (e.g., `codex exec` for code generation that is not a review).
20
+
21
+ ## Trace Directory
22
+
23
+ ```
24
+ .aris/traces/<skill-name>/<YYYY-MM-DD>_run<NN>/
25
+ ├── run.meta.json # Run-level metadata
26
+ ├── 001-<purpose>.request.json # Request snapshot
27
+ ├── 001-<purpose>.response.md # Full response text
28
+ ├── 001-<purpose>.meta.json # Response metadata
29
+ ├── 002-<purpose>.request.json # Second call (e.g., reply)
30
+ └── ...
31
+ ```
32
+
33
+ - `<skill-name>`: the ARIS skill that triggered this call (e.g., `auto-review-loop`)
34
+ - `<YYYY-MM-DD>_run<NN>`: date + sequential run number (start from `01`)
35
+ - `<purpose>`: short kebab-case label (e.g., `round-1-review`, `critique`, `ideation`, `audit`, `patch-gate`)
36
+
37
+ ## How to Trace
38
+
39
+ After each reviewer call — including every FAILED attempt in a
40
+ capability-fallback chain (one trace entry per attempt: `--status error` +
41
+ `--fallback-reason`; the successful entry records the RESOLVED pair) — save the trace using `save_trace.sh`,
42
+ resolved through the canonical helper chain (see
43
+ `integration-contract.md` §2 — failure policy C, "forensic helper").
44
+ The full invocation:
45
+
46
+ ```bash
47
+ # Resolve $TRACE_HELPER (canonical strict-safe chain; see integration-contract.md §2).
48
+ cd "$(git rev-parse --show-toplevel 2>/dev/null || pwd)" || exit 1
49
+ if [ -z "${ARIS_REPO:-}" ] && [ -f .aris/installed-skills.txt ]; then
50
+ ARIS_REPO=$(awk -F'\t' '$1=="repo_root"{print $2; exit}' .aris/installed-skills.txt 2>/dev/null) || true
51
+ fi
52
+ if [ -z "${ARIS_REPO:-}" ] && [ -f "$HOME/.aris/repo" ]; then
53
+ ARIS_REPO=$(cat "$HOME/.aris/repo" 2>/dev/null) || true
54
+ fi
55
+ TRACE_HELPER=".aris/tools/save_trace.sh"
56
+ [ -f "$TRACE_HELPER" ] || TRACE_HELPER="tools/save_trace.sh"
57
+ [ -f "$TRACE_HELPER" ] || { [ -n "${ARIS_REPO:-}" ] && TRACE_HELPER="$ARIS_REPO/tools/save_trace.sh"; }
58
+ [ -f "$TRACE_HELPER" ] || TRACE_HELPER=""
59
+
60
+ if [ -n "$TRACE_HELPER" ]; then
61
+ bash "$TRACE_HELPER" \
62
+ --skill "<skill-name>" \
63
+ --purpose "<purpose>" \
64
+ --model "<model that actually ran — the RESOLVED pair, not the target>" \
65
+ --effort "<effort that actually ran>" \
66
+ --fallback-reason "<why the capability chain stepped down; empty when it didn't>" \
67
+ --status "<ok | fallback_used | error>" \
68
+ --thread-id "<threadId from response>" \
69
+ --backend "<codex | copilot-native | copilot | manual | oracle-pro | agy>" \
70
+ --tool "<mcp__codex__codex | task(agent_type=rubber-duck) | copilot --agent | ...>" \
71
+ --executor "<claude-code | copilot | codex>" \
72
+ --executor-model "<from --executor-model; omit this flag if not set>" \
73
+ --executor-family "<legacy consistency hint; helper re-derives from executor-model>" \
74
+ --reviewer-profile "<profile name for copilot backend; empty for others>" \
75
+ --reviewer-family "<legacy consistency hint; helper re-derives from reviewer model>" \
76
+ --requested-reviewer-model "<model originally requested>" \
77
+ --reported-reviewer-model "<model the backend reports it used>" \
78
+ --memory-hash "<sha256 of memory artifact if available; empty otherwise>" \
79
+ --native-evidence "<required for copilot-native; omit otherwise>" \
80
+ --independence-verified "<legacy consistency hint; helper ignores and re-derives>" \
81
+ --prompt "<full prompt as sent>" \
82
+ --response "<full response content>"
83
+ else
84
+ # Required fallback: the resolver exhausted all four layers and
85
+ # save_trace.sh is unreachable, but trace artifacts are still
86
+ # required (unless `--- trace: off` was explicitly set on this
87
+ # SKILL invocation). Write the four files below directly per the
88
+ # schemas in "File Schemas", into:
89
+ # .aris/traces/<skill-name>/<YYYY-MM-DD>_run<NN>/
90
+ # run.meta.json
91
+ # <NNN>-<purpose>.request.json
92
+ # <NNN>-<purpose>.response.md
93
+ # <NNN>-<purpose>.meta.json
94
+ # Do NOT silently skip — trace_path is load-bearing for any
95
+ # mandatory audit emitting `trace_path` in its artifact (see
96
+ # assurance-contract.md §"Required Audit Artifact Schema").
97
+ echo "WARN: save_trace.sh not resolved; writing trace files directly per review-tracing.md schema." >&2
98
+ fi
99
+ ```
100
+
101
+ The helper, when present, handles directory creation, run numbering, file
102
+ writing, and provenance classification. For a successful `copilot-native`
103
+ call, it first revalidates `--native-evidence`, overrides caller model/response
104
+ fields with the host-bound values, and records the evidence ID/path; missing or
105
+ invalid evidence is an error. A native dispatch that failed before evidence
106
+ could exist may be traced only with `--status error` and no
107
+ `--native-evidence`; that record has no verified model provenance or
108
+ independence and can never enter the stop gate. For other backends it derives
109
+ families from model strings
110
+ (`reported_reviewer_model`, otherwise `requested_reviewer_model`, otherwise
111
+ `model`) and ignores contradictory caller family/independence claims. It also
112
+ records that `--executor-model` is `caller-declared`; a different family pair is
113
+ `family_relation: "different"` but remains `independence_verified: "unverified"`.
114
+ Validated native evidence instead records both sources as
115
+ `host-session-event`, `family_relation: "different"`, and
116
+ `independence_verified: true`.
117
+
118
+ The native-specific required pair for a completed review is:
119
+
120
+ ```bash
121
+ bash "$TRACE_HELPER" ... --backend copilot-native \
122
+ --native-evidence "review-stage/COPILOT_NATIVE_${RUN_ID}_ROUND_${ROUND}_REVIEW.evidence.json"
123
+ ```
124
+
125
+ Omit `--native-evidence` for every other backend.
126
+ For a failed native dispatch, instead use `--backend copilot-native --status
127
+ error --fallback-reason "<host error>"` without evidence, then write a separate
128
+ trace for any fallback backend that actually reviews the artifacts.
129
+ The fallback branch above documents what to do
130
+ when the helper is unreachable — the trace is forensic evidence, so
131
+ "helper missing" never means "skip the trace." A direct fallback writer MUST
132
+ apply the same model-string derivation and evidence-source rules; it may not
133
+ copy caller family labels or promote caller-declared identity to verification.
134
+
135
+ ## File Schemas
136
+
137
+ ### `run.meta.json`
138
+ ```json
139
+ {
140
+ "skill": "auto-review-loop",
141
+ "run_id": "2026-04-15_run01",
142
+ "started_at": "2026-04-15T14:30:00+08:00",
143
+ "executor": "claude-code",
144
+ "executor_model": "claude-sonnet-4-5",
145
+ "executor_model_source": "caller-declared",
146
+ "executor_family": "anthropic",
147
+ "reviewer_model_source": "requested",
148
+ "reviewer_family": "openai",
149
+ "reviewer_backend": "codex",
150
+ "native_evidence_id": null,
151
+ "native_evidence_path": null,
152
+ "family_relation": "different",
153
+ "independence_verified": "unverified",
154
+ "project_dir": "/path/to/project"
155
+ }
156
+ ```
157
+
158
+ - `executor`: the name of the running executor (from `--executor` parameter; defaults to `"claude-code"`). Dynamic — set by the caller, not hardcoded.
159
+ - `executor_model`: the declared model running this ARIS invocation (from `--executor-model` when available; otherwise `null`).
160
+ - `executor_model_source`: `"host-session-event"` for validated native evidence,
161
+ `"caller-declared"` when a non-native caller passes `--executor-model`, or
162
+ `"unavailable"`.
163
+ - `executor_family`: derived from `executor_model` (`openai` / `anthropic` / `google` / `unknown`).
164
+ - `reviewer_model_source`: `"host-session-event"`, `"backend-reported"`,
165
+ `"requested"`, or `"unavailable"`.
166
+ - `reviewer_family`: derived from the backend-reported model when available, otherwise the requested model.
167
+ - `family_relation`: `different`, `same`, or `unknown`, derived from model strings.
168
+ - `independence_verified`: `true` only for a revalidated native cross-family
169
+ event chain, `false` for known same-family identities, otherwise
170
+ `"unverified"`.
171
+ - `native_evidence_id` / `native_evidence_path`: set only for
172
+ `copilot-native`; `null` for other backends.
173
+ - `reviewer_backend`: `codex` / `copilot-native` / `copilot` / `manual` /
174
+ `oracle-pro` / `agy`.
175
+
176
+ ### `NNN-<purpose>.request.json`
177
+ ```json
178
+ {
179
+ "call_number": 1,
180
+ "purpose": "round-1-review",
181
+ "timestamp": "2026-04-15T14:31:00+08:00",
182
+ "tool": "mcp__codex__codex",
183
+ "backend": "codex",
184
+ "model": "gpt-5.6-sol",
185
+ "config": {"model_reasoning_effort": "xhigh"},
186
+ "reviewer_profile": null,
187
+ "files_referenced": ["paper/sections/3_method.tex", "results/table1.csv"],
188
+ "prompt": "<full prompt text>"
189
+ }
190
+ ```
191
+
192
+ For native Copilot backend (no reviewer override in a bound Copilot session):
193
+ ```json
194
+ {
195
+ "call_number": 1,
196
+ "purpose": "round-1-review",
197
+ "tool": "task(agent_type=rubber-duck)",
198
+ "backend": "copilot-native",
199
+ "model": "gpt-5.5",
200
+ "executor_model": "claude-sonnet-4.6",
201
+ "executor_model_source": "host-session-event",
202
+ "executor_family": "anthropic",
203
+ "reported_reviewer_model": "gpt-5.5",
204
+ "reviewer_model_source": "host-session-event",
205
+ "reviewer_family": "openai",
206
+ "family_relation": "different",
207
+ "independence_verified": true,
208
+ "native_evidence_id": "cne_0123456789abcdef0123456789abcdef",
209
+ "native_evidence_path": "/project/review-stage/COPILOT_NATIVE_run_20260715_a1b2c3d4_ROUND_1_REVIEW.evidence.json",
210
+ "prompt": "<full nonce-bound native task prompt>"
211
+ }
212
+ ```
213
+
214
+ For compatibility copilot backend (`--reviewer: copilot`):
215
+ ```json
216
+ {
217
+ "call_number": 1,
218
+ "purpose": "round-1-review",
219
+ "timestamp": "2026-04-15T14:31:00+08:00",
220
+ "tool": "copilot --agent",
221
+ "backend": "copilot",
222
+ "model": "gpt-5.4",
223
+ "effort": "xhigh",
224
+ "effort_unpinned": false,
225
+ "reviewer_profile": "aris-reviewer-openai",
226
+ "requested_reviewer_model": "gpt-5.4",
227
+ "reported_reviewer_model": null,
228
+ "executor_model": "claude-sonnet-4-5",
229
+ "executor_model_source": "caller-declared",
230
+ "executor_family": "anthropic",
231
+ "reviewer_model_source": "requested",
232
+ "reviewer_family": "openai",
233
+ "family_relation": "different",
234
+ "independence_verified": "unverified",
235
+ "files_referenced": ["paper/sections/3_method.tex", "results/table1.csv"],
236
+ "prompt": "<full prompt text>"
237
+ }
238
+ ```
239
+
240
+ Fields:
241
+ - `tool`: the tool name used (`task(agent_type=rubber-duck)`,
242
+ `mcp__codex__codex`, `copilot --agent`, etc.).
243
+ - `backend`: the logical backend (`codex`, `copilot-native`, compatibility
244
+ `copilot`, `manual`, `oracle-pro`, `agy`).
245
+ - `effort_unpinned`: applies only to compatibility `copilot`; native
246
+ complementary dispatch does not accept an ARIS model/effort pin.
247
+ - `reviewer_profile`: custom profile for compatibility copilot; `rubber-duck`
248
+ may be used as a descriptive value for native traces; `null` otherwise.
249
+ - `requested_reviewer_model`: the model parsed from profile frontmatter and repeated through subprocess `--model`; `null` when unavailable.
250
+ - `reported_reviewer_model`: the model the tool reports actually using (for example, captured Copilot `gen_ai.response.model` telemetry); `null` when unavailable.
251
+ - `executor_model`: native host-event value, or from `--executor-model` on
252
+ legacy/non-native paths; `null` if unavailable.
253
+ - `executor_model_source`: `host-session-event`, `caller-declared`, or
254
+ `unavailable`.
255
+ - `executor_family`: derived from `executor_model`.
256
+ - `reviewer_model_source`: `backend-reported`, `requested`, or `unavailable`.
257
+ - `reviewer_family`: derived by the helper from the reported/requested/actual reviewer model, never trusted from the caller.
258
+ - `family_relation`: string-derived relation (`different` / `same` / `unknown`).
259
+ - `independence_verified`: `true` only for validated native cross-family
260
+ evidence; `false` for known same-family; `"unverified"` for advisory pairs.
261
+ - `native_evidence_id` / `native_evidence_path`: bind a native trace to the
262
+ revalidated session-event artifact; `null` otherwise.
263
+
264
+ ### `NNN-<purpose>.response.md`
265
+ The reviewer's full response, verbatim. No truncation, no summarization.
266
+
267
+ ### `NNN-<purpose>.meta.json`
268
+ ```json
269
+ {
270
+ "call_number": 1,
271
+ "purpose": "round-1-review",
272
+ "timestamp": "2026-04-15T14:33:00+08:00",
273
+ "thread_id": "019d8fe0-b25d-...",
274
+ "model": "gpt-5.6-sol",
275
+ "model_family": "openai",
276
+ "executor_model": null,
277
+ "executor_model_source": "unavailable",
278
+ "executor_family": "unknown",
279
+ "reviewer_model_source": "requested",
280
+ "family_relation": "unknown",
281
+ "independence_verified": "unverified",
282
+ "reviewer_profile": null,
283
+ "duration_ms": 142000,
284
+ "status": "ok"
285
+ }
286
+ ```
287
+
288
+ For compatibility copilot backend:
289
+ ```json
290
+ {
291
+ "call_number": 1,
292
+ "purpose": "round-1-review",
293
+ "timestamp": "2026-04-15T14:33:00+08:00",
294
+ "thread_id": null,
295
+ "model": "gpt-5.4",
296
+ "model_family": "openai",
297
+ "effort": "xhigh",
298
+ "effort_unpinned": false,
299
+ "executor_model": "claude-sonnet-4-5",
300
+ "executor_model_source": "caller-declared",
301
+ "executor_family": "anthropic",
302
+ "requested_reviewer_model": "gpt-5.4",
303
+ "reported_reviewer_model": null,
304
+ "reviewer_model_source": "requested",
305
+ "family_relation": "different",
306
+ "independence_verified": "unverified",
307
+ "reviewer_profile": "aris-reviewer-openai",
308
+ "duration_ms": 142000,
309
+ "status": "ok"
310
+ }
311
+ ```
312
+
313
+ Fields new per this fix:
314
+ - `model_family`: `openai` / `anthropic` / `google` / `unknown` — derived from the model that actually ran.
315
+ - `effort_unpinned`: whether a compatibility Copilot subprocess lacked the
316
+ required explicit `xhigh` pin; it does not apply to native dispatch.
317
+ - `executor_model` / `executor_model_source`: native calls record the actual
318
+ host event model/source; compatibility routing records `caller-declared`.
319
+ - `executor_family`: derived from `executor_model`; `unknown` if not known.
320
+ - `requested_reviewer_model`: the profile model also passed through subprocess `--model`; `null` when not available.
321
+ - `reported_reviewer_model` / `reviewer_model_source`: backend output when available, otherwise the requested model and source.
322
+ - `family_relation`: the model-string relation, separate from evidence assurance.
323
+ - `independence_verified`: never `true` solely because caller-declared and requested model strings differ; use `"unverified"` in that case.
324
+ - `reviewer_profile`: for compatibility copilot, the custom agent profile;
325
+ native may record `rubber-duck`; `null` for other backends.
326
+ - `native_evidence_id` / `native_evidence_path`: present only on validated
327
+ `copilot-native` traces.
328
+ - `memory_hash`: SHA-256 of `review-stage/REVIEWER_MEMORY.md` at trace time; `null` when no memory file exists.
329
+
330
+ ## Configuration
331
+
332
+ Tracing respects three modes, set via inline parameter `--- trace: off | meta | full`:
333
+ - **`full`** (default): save full prompt + full response
334
+ - **`meta`**: save metadata only (no prompt/response text), useful for sensitive projects
335
+ - **`off`**: disable tracing entirely
336
+
337
+ ## Integration with events.jsonl
338
+
339
+ After writing a trace, append a compact summary event to `.aris/meta/events.jsonl`:
340
+
341
+ ```json
342
+ {"event":"review_trace","skill":"auto-review-loop","purpose":"round-1-review","thread_id":null,"trace_path":".aris/traces/auto-review-loop/2026-04-15_run01/","backend":"copilot-native","tool":"task(agent_type=rubber-duck)","executor_model":"claude-sonnet-4.6","executor_model_source":"host-session-event","executor_family":"anthropic","reviewer_model_source":"host-session-event","reviewer_family":"openai","family_relation":"different","independence_verified":true,"native_evidence_id":"cne_0123456789abcdef0123456789abcdef","native_evidence_path":"/project/review-stage/COPILOT_NATIVE_run_20260715_a1b2c3d4_ROUND_1_REVIEW.evidence.json","status":"ok"}
343
+ ```
344
+
345
+ This allows `/meta-optimize` to discover traces without reading the full trace files.
346
+
347
+ ## Debugging With Traces
348
+
349
+ Traces are not only audit evidence — they are the **first place to look when a
350
+ verdict is surprising**: a score regresses round-to-round, two reviewer backends
351
+ disagree, or `/result-to-claim` contradicts an earlier claim. Before re-invoking
352
+ the reviewer for "a better answer", read the raw transcript and find the moment
353
+ its judgment actually changed:
354
+
355
+ ```bash
356
+ # Diff the raw response bodies across the two calls in question
357
+ skill=auto-review-loop run=2026-04-15_run01
358
+ diff ".aris/traces/$skill/$run/002-round-2.response.md" \
359
+ ".aris/traces/$skill/$run/003-round-3.response.md"
360
+
361
+ # Grep for the sentence where the assessment turned
362
+ grep -En 'however|but|concern|missing|cannot' \
363
+ ".aris/traces/$skill/$run/003-round-3.response.md"
364
+ ```
365
+
366
+ The paragraph where the assessment changed **is** the causal explanation for the
367
+ divergence — cite it, don't guess. Re-running the reviewer without reading the
368
+ trace is tuning by vibe: you get a new opinion, not an explanation.
369
+
370
+ This is the same muscle ARIS already applies to code failures (the "**Read the
371
+ error** — parse traceback, stderr, and log files" step in `/experiment-bridge`'s
372
+ auto-debug sequence, and `/codex:rescue` reading tracebacks before a retry) —
373
+ applied to saved AI-judgment transcripts instead of stderr. The trace is written
374
+ in English and most of it is the reviewer talking to itself; the discipline is
375
+ identical: read the primary artifact first, then act on the exact divergence
376
+ point rather than re-rolling the dice.
377
+
378
+ Practical triggers:
379
+
380
+ | Surprise | Trace move |
381
+ |---|---|
382
+ | Score dropped after a "fix" round | diff the two rounds' `.response.md`; find which criterion flipped |
383
+ | Two backends disagree (codex vs gemini/manual) | grep both responses for the SAME artifact path; compare what each actually read |
384
+ | Reviewer "forgot" an earlier concern | grep prior rounds for the concern keyword; if present-then-absent, cite it in the next prompt instead of restating from memory |
385
+ | Verdict contradicts a deterministic checker | read the request `.md` — was the checker's output actually in the files the reviewer was pointed at? |
386
+
387
+ ## Privacy
388
+
389
+ - `.aris/traces/` should be in `.gitignore` — traces are project-local, never committed
390
+ - Traces may contain sensitive research content; treat them as confidential
391
+ - Use `--- trace: off` for projects with strict confidentiality requirements
@@ -0,0 +1,79 @@
1
+ # Reviewer Independence Protocol
2
+
3
+ ## Core Principle
4
+
5
+ **Content must reach the reviewer unfiltered. The executor points to files and sets the review task; the reviewer reads and judges independently.**
6
+
7
+ Cross-model adversarial collaboration only works if the reviewer forms its own assessment from primary artifacts. If the executor pre-digests, summarizes, or interprets content before passing it to the reviewer, the reviewer is evaluating the executor's framing — not the actual work. This re-introduces the correlated blind spots that heterogeneous review is designed to avoid.
8
+
9
+ ## What CAN be passed to the reviewer
10
+
11
+ - **Role/persona** — e.g., "Review as a NeurIPS-level reviewer"
12
+ - **Review objective** — e.g., "Evaluate publishability", "Check code correctness", "Score 1-10 on clarity"
13
+ - **File paths** — let the reviewer read file contents directly
14
+ - **Structural metadata** — e.g., "The paper has 8 sections", "Experiments are in experiments/"
15
+ - **Venue constraints** — e.g., "ICLR format, 9-page limit"
16
+
17
+ ## What CANNOT be passed (counts as "subjective interference")
18
+
19
+ - ❌ Executor's summary or paraphrase of file contents
20
+ - ❌ Executor's interpretation of results (e.g., "I think the problem is...", "This suggests...")
21
+ - ❌ Executor's recommendations or conclusions (e.g., "I suggest changing...", "The likely cause is...")
22
+ - ❌ Key findings or bullet points extracted by the executor
23
+ - ❌ Leading questions (e.g., "Is this publishable?", "Is this trade-off reasonable?")
24
+ - ❌ Previous review rounds' feedback or critique (let the reviewer assess the current state fresh)
25
+ - ❌ Executor's description of what was changed since last round (e.g., "I fixed X, Y, Z")
26
+ - ❌ Statements asserting the current approach's strengths
27
+
28
+ ## Why this matters
29
+
30
+ | With filtering | Without filtering |
31
+ |---|---|
32
+ | Reviewer sees executor's framing | Reviewer sees raw artifacts |
33
+ | Correlated blind spots persist | Genuinely independent assessment |
34
+ | Executor can "coach" favorable review | Review probes real weaknesses |
35
+ | Defeats the purpose of cross-model | Achieves adversarial collaboration |
36
+
37
+ ## Correct pattern
38
+
39
+ ```
40
+ mcp__codex__codex:
41
+ prompt: |
42
+ Review the following research project as a senior ML reviewer.
43
+
44
+ Files to read:
45
+ - Proposal: /path/to/PROPOSAL.md
46
+ - Experiment results: /path/to/EXPERIMENT_LOG.md
47
+ - Paper draft: /path/to/paper/main.tex
48
+ - Code: /path/to/src/
49
+
50
+ Please read all files yourself and provide a complete review.
51
+ Score 1-10 on: novelty, soundness, clarity, significance.
52
+ ```
53
+
54
+ ## Incorrect pattern
55
+
56
+ ```
57
+ mcp__codex__codex:
58
+ prompt: |
59
+ The main contribution is a new loss function that improves by 15%.
60
+ However, I noticed the ablation is incomplete.
61
+ Here's my summary of the key results: [...]
62
+ Please review whether this is publishable.
63
+ ```
64
+
65
+ ## When to apply
66
+
67
+ This protocol applies to ALL cross-model review calls in ARIS:
68
+ - `/research-review` — paper review
69
+ - `/auto-review-loop` — iterative review
70
+ - `/paper-plan` — outline review
71
+ - `/paper-write` — section review
72
+ - `/paper-figure` — figure quality review
73
+ - `/rebuttal` — stress test
74
+ - `/meta-optimize` — patch review
75
+ - Any skill that sends artifacts to `mcp__codex__codex` or `mcp__codex__codex-reply`
76
+
77
+ ## Exception
78
+
79
+ Multi-round review within the SAME thread (`codex-reply`) may reference the reviewer's own previous feedback to check resolution — but still must not include executor interpretations of that feedback.