dsh-aris-panel 0.1.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (442) hide show
  1. package/LICENSE +21 -0
  2. package/README.md +98 -0
  3. package/README_CN.md +87 -0
  4. package/dsh/checkout.patch.yml +38 -0
  5. package/dsh/client.js +634 -0
  6. package/dsh/cordis.patch.yml +44 -0
  7. package/dsh/index.mjs +76 -0
  8. package/dsh/run-status.mjs +182 -0
  9. package/dsh/scope-limits.mjs +50 -0
  10. package/dsh/workbench.mjs +291 -0
  11. package/mcp-servers/claude-review/README.md +93 -0
  12. package/mcp-servers/claude-review/run_with_claude_aws.sh +49 -0
  13. package/mcp-servers/claude-review/server.py +718 -0
  14. package/mcp-servers/codex-image2/README.md +65 -0
  15. package/mcp-servers/codex-image2/server.py +893 -0
  16. package/mcp-servers/feishu-bridge/requirements.txt +1 -0
  17. package/mcp-servers/feishu-bridge/server.py +240 -0
  18. package/mcp-servers/gemini-review/README.md +171 -0
  19. package/mcp-servers/gemini-review/server.py +1856 -0
  20. package/mcp-servers/llm-chat/requirements.txt +1 -0
  21. package/mcp-servers/llm-chat/server.py +664 -0
  22. package/mcp-servers/manual-review/README.md +133 -0
  23. package/mcp-servers/manual-review/server.py +910 -0
  24. package/mcp-servers/manual-review/ui.html +279 -0
  25. package/mcp-servers/minimax-chat/requirements.txt +1 -0
  26. package/mcp-servers/minimax-chat/server.py +381 -0
  27. package/package.json +51 -0
  28. package/skills/ablation-planner/SKILL.md +123 -0
  29. package/skills/alphaxiv/SKILL.md +196 -0
  30. package/skills/analyze-results/SKILL.md +46 -0
  31. package/skills/arxiv/SKILL.md +248 -0
  32. package/skills/auto-paper-improvement-loop/SKILL.md +651 -0
  33. package/skills/auto-review-loop/SKILL.md +1137 -0
  34. package/skills/auto-review-loop-llm/SKILL.md +259 -0
  35. package/skills/auto-review-loop-minimax/SKILL.md +302 -0
  36. package/skills/citation-audit/SKILL.md +502 -0
  37. package/skills/claims-drafting/SKILL.md +227 -0
  38. package/skills/comm-lit-review/SKILL.md +297 -0
  39. package/skills/deepxiv/SKILL.md +263 -0
  40. package/skills/dse-loop/SKILL.md +296 -0
  41. package/skills/embodiment-description/SKILL.md +129 -0
  42. package/skills/exa-search/SKILL.md +205 -0
  43. package/skills/experiment-audit/SKILL.md +311 -0
  44. package/skills/experiment-bridge/SKILL.md +376 -0
  45. package/skills/experiment-plan/SKILL.md +249 -0
  46. package/skills/experiment-queue/SKILL.md +431 -0
  47. package/skills/experiment-queue/scripts/build_manifest.py +142 -0
  48. package/skills/experiment-queue/scripts/queue_manager.py +433 -0
  49. package/skills/feishu-notify/SKILL.md +156 -0
  50. package/skills/figure-description/SKILL.md +138 -0
  51. package/skills/figure-spec/SKILL.md +262 -0
  52. package/skills/figure-spec/scripts/figure_renderer.py +799 -0
  53. package/skills/formula-derivation/SKILL.md +280 -0
  54. package/skills/gemini-search/SKILL.md +231 -0
  55. package/skills/grant-proposal/SKILL.md +698 -0
  56. package/skills/idea-creator/SKILL.md +542 -0
  57. package/skills/idea-discovery/SKILL.md +521 -0
  58. package/skills/idea-discovery-robot/SKILL.md +363 -0
  59. package/skills/integrity-forensics/SKILL.md +284 -0
  60. package/skills/interview-cheatsheet/SKILL.md +245 -0
  61. package/skills/invention-structuring/SKILL.md +188 -0
  62. package/skills/jurisdiction-format/SKILL.md +192 -0
  63. package/skills/kill-argument/SKILL.md +437 -0
  64. package/skills/mermaid-diagram/SKILL.md +419 -0
  65. package/skills/meta-apply/SKILL.md +141 -0
  66. package/skills/meta-optimize/SKILL.md +437 -0
  67. package/skills/monitor-experiment/SKILL.md +140 -0
  68. package/skills/novelty-check/SKILL.md +101 -0
  69. package/skills/openalex/SKILL.md +237 -0
  70. package/skills/overleaf-sync/SKILL.md +220 -0
  71. package/skills/paper-claim-audit/SKILL.md +348 -0
  72. package/skills/paper-compile/SKILL.md +266 -0
  73. package/skills/paper-figure/SKILL.md +312 -0
  74. package/skills/paper-illustration/SKILL.md +736 -0
  75. package/skills/paper-illustration-image2/SKILL.md +391 -0
  76. package/skills/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  77. package/skills/paper-plan/SKILL.md +386 -0
  78. package/skills/paper-poster/SKILL.md +19 -0
  79. package/skills/paper-poster-html/DESIGN_FINAL.md +176 -0
  80. package/skills/paper-poster-html/IMPLEMENTATION_CONVENTIONS.md +161 -0
  81. package/skills/paper-poster-html/LICENSES/posterly-MIT.txt +21 -0
  82. package/skills/paper-poster-html/NOTICE.md +57 -0
  83. package/skills/paper-poster-html/SKILL.md +323 -0
  84. package/skills/paper-poster-html/scripts/_posterly/__init__.py +0 -0
  85. package/skills/paper-poster-html/scripts/_posterly/canvas.py +200 -0
  86. package/skills/paper-poster-html/scripts/_posterly/measure.py +588 -0
  87. package/skills/paper-poster-html/scripts/_posterly/polish.py +498 -0
  88. package/skills/paper-poster-html/scripts/_posterly/preflight.py +489 -0
  89. package/skills/paper-poster-html/scripts/_posterly/render.py +215 -0
  90. package/skills/paper-poster-html/scripts/_posterly/textutil.py +16 -0
  91. package/skills/paper-poster-html/scripts/_posterly/verify_final.py +171 -0
  92. package/skills/paper-poster-html/scripts/asset_check.py +897 -0
  93. package/skills/paper-poster-html/scripts/extract_pdf_figures.py +666 -0
  94. package/skills/paper-poster-html/scripts/poster_check.py +251 -0
  95. package/skills/paper-poster-html/scripts/preprocess_figures.py +238 -0
  96. package/skills/paper-poster-html/scripts/render_preview.py +217 -0
  97. package/skills/paper-poster-html/scripts/run_gates.py +556 -0
  98. package/skills/paper-poster-html/scripts/style_check.py +1324 -0
  99. package/skills/paper-poster-html/templates/COMPONENTS.md +462 -0
  100. package/skills/paper-poster-html/templates/README.md +170 -0
  101. package/skills/paper-poster-html/templates/landscape_4col.html +1032 -0
  102. package/skills/paper-poster-html/templates/landscape_hero.html +1046 -0
  103. package/skills/paper-poster-html/templates/portrait_2col.html +947 -0
  104. package/skills/paper-poster-html/templates/tokens/acl.json +9 -0
  105. package/skills/paper-poster-html/templates/tokens/cvpr.json +9 -0
  106. package/skills/paper-poster-html/templates/tokens/generic.json +9 -0
  107. package/skills/paper-poster-html/templates/tokens/iclr.json +9 -0
  108. package/skills/paper-poster-html/templates/tokens/icml.json +9 -0
  109. package/skills/paper-poster-html/templates/tokens/neurips.json +9 -0
  110. package/skills/paper-slides/SKILL.md +635 -0
  111. package/skills/paper-talk/SKILL.md +381 -0
  112. package/skills/paper-write/SKILL.md +604 -0
  113. package/skills/paper-write/templates/IEEEtran.bst +2409 -0
  114. package/skills/paper-write/templates/IEEEtran.cls +6347 -0
  115. package/skills/paper-write/templates/iclr2026.tex +84 -0
  116. package/skills/paper-write/templates/icml2025.tex +87 -0
  117. package/skills/paper-write/templates/ieee_conference.tex +89 -0
  118. package/skills/paper-write/templates/ieee_journal.tex +93 -0
  119. package/skills/paper-write/templates/math_commands.tex +48 -0
  120. package/skills/paper-write/templates/neurips2025.tex +80 -0
  121. package/skills/paper-writing/SKILL.md +916 -0
  122. package/skills/patent-novelty-check/SKILL.md +153 -0
  123. package/skills/patent-pipeline/SKILL.md +344 -0
  124. package/skills/patent-review/SKILL.md +203 -0
  125. package/skills/pixel-art/SKILL.md +137 -0
  126. package/skills/prior-art-search/SKILL.md +146 -0
  127. package/skills/proof-checker/SKILL.md +866 -0
  128. package/skills/proof-orchestrator/NOTICE.md +24 -0
  129. package/skills/proof-orchestrator/SKILL.md +254 -0
  130. package/skills/proof-orchestrator/references/audit-output-contract.md +126 -0
  131. package/skills/proof-orchestrator/references/deepseek-routing.md +74 -0
  132. package/skills/proof-orchestrator/references/dispatch-prompts.md +227 -0
  133. package/skills/proof-orchestrator/references/notation-audit.md +135 -0
  134. package/skills/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  135. package/skills/proof-orchestrator/references/stress-tests.md +38 -0
  136. package/skills/proof-writer/SKILL.md +223 -0
  137. package/skills/qzcli/SKILL.md +324 -0
  138. package/skills/rebuttal/SKILL.md +376 -0
  139. package/skills/render-html/SKILL.md +316 -0
  140. package/skills/render-html/scripts/render_html.py +1006 -0
  141. package/skills/render-html/scripts/templates/academic.html +703 -0
  142. package/skills/render-html/scripts/templates/dashboard.html +333 -0
  143. package/skills/research-lit/SKILL.md +756 -0
  144. package/skills/research-pipeline/SKILL.md +384 -0
  145. package/skills/research-refine/SKILL.md +770 -0
  146. package/skills/research-refine-pipeline/SKILL.md +186 -0
  147. package/skills/research-review/SKILL.md +198 -0
  148. package/skills/research-wiki/SKILL.md +461 -0
  149. package/skills/resubmit-pipeline/SKILL.md +447 -0
  150. package/skills/result-to-claim/SKILL.md +311 -0
  151. package/skills/run-experiment/SKILL.md +313 -0
  152. package/skills/semantic-scholar/SKILL.md +236 -0
  153. package/skills/serverless-modal/SKILL.md +335 -0
  154. package/skills/shared-references/acceptance-gate.md +324 -0
  155. package/skills/shared-references/assurance-contract.md +248 -0
  156. package/skills/shared-references/capture-antipatterns.md +78 -0
  157. package/skills/shared-references/citation-discipline.md +583 -0
  158. package/skills/shared-references/compute-env-contract.md +163 -0
  159. package/skills/shared-references/effort-contract.md +183 -0
  160. package/skills/shared-references/evidence-precheck.md +65 -0
  161. package/skills/shared-references/experiment-integrity.md +49 -0
  162. package/skills/shared-references/external-cadence.md +326 -0
  163. package/skills/shared-references/fan-out-pattern.md +366 -0
  164. package/skills/shared-references/injection-hygiene.md +127 -0
  165. package/skills/shared-references/integration-contract.md +461 -0
  166. package/skills/shared-references/output-composition.md +93 -0
  167. package/skills/shared-references/output-language.md +45 -0
  168. package/skills/shared-references/output-manifest.md +49 -0
  169. package/skills/shared-references/output-versioning.md +111 -0
  170. package/skills/shared-references/patent-format-cn.md +199 -0
  171. package/skills/shared-references/patent-format-ep.md +173 -0
  172. package/skills/shared-references/patent-format-us.md +161 -0
  173. package/skills/shared-references/patent-writing-principles.md +197 -0
  174. package/skills/shared-references/prior-art-databases.md +141 -0
  175. package/skills/shared-references/resumable-runs.md +109 -0
  176. package/skills/shared-references/review-scope-limits.md +81 -0
  177. package/skills/shared-references/review-tracing.md +391 -0
  178. package/skills/shared-references/reviewer-independence.md +79 -0
  179. package/skills/shared-references/reviewer-routing.md +852 -0
  180. package/skills/shared-references/skill-governance.md +104 -0
  181. package/skills/shared-references/taste-calibration.md +85 -0
  182. package/skills/shared-references/venue-checklists.md +114 -0
  183. package/skills/shared-references/wiki-helper-resolution.md +134 -0
  184. package/skills/shared-references/writing-principles.md +525 -0
  185. package/skills/skills-codex/README.md +102 -0
  186. package/skills/skills-codex/README_CN.md +100 -0
  187. package/skills/skills-codex/ablation-planner/SKILL.md +126 -0
  188. package/skills/skills-codex/alphaxiv/SKILL.md +186 -0
  189. package/skills/skills-codex/analyze-results/SKILL.md +45 -0
  190. package/skills/skills-codex/arxiv/SKILL.md +210 -0
  191. package/skills/skills-codex/auto-paper-improvement-loop/SKILL.md +574 -0
  192. package/skills/skills-codex/auto-review-loop/SKILL.md +500 -0
  193. package/skills/skills-codex/auto-review-loop-llm/SKILL.md +247 -0
  194. package/skills/skills-codex/auto-review-loop-minimax/SKILL.md +290 -0
  195. package/skills/skills-codex/citation-audit/SKILL.md +504 -0
  196. package/skills/skills-codex/claims-drafting/SKILL.md +239 -0
  197. package/skills/skills-codex/comm-lit-review/SKILL.md +299 -0
  198. package/skills/skills-codex/comm-lit-review/references/domain-taxonomy.md +57 -0
  199. package/skills/skills-codex/comm-lit-review/references/output-template.md +37 -0
  200. package/skills/skills-codex/comm-lit-review/references/source-policy.md +99 -0
  201. package/skills/skills-codex/comm-lit-review/references/venue-tiering.md +112 -0
  202. package/skills/skills-codex/deepxiv/SKILL.md +142 -0
  203. package/skills/skills-codex/dse-loop/SKILL.md +285 -0
  204. package/skills/skills-codex/embodiment-description/SKILL.md +129 -0
  205. package/skills/skills-codex/exa-search/SKILL.md +192 -0
  206. package/skills/skills-codex/experiment-audit/SKILL.md +286 -0
  207. package/skills/skills-codex/experiment-bridge/SKILL.md +356 -0
  208. package/skills/skills-codex/experiment-plan/SKILL.md +249 -0
  209. package/skills/skills-codex/experiment-queue/SKILL.md +401 -0
  210. package/skills/skills-codex/feishu-notify/SKILL.md +155 -0
  211. package/skills/skills-codex/figure-description/SKILL.md +138 -0
  212. package/skills/skills-codex/figure-spec/SKILL.md +252 -0
  213. package/skills/skills-codex/formula-derivation/SKILL.md +280 -0
  214. package/skills/skills-codex/gemini-search/SKILL.md +205 -0
  215. package/skills/skills-codex/grant-proposal/SKILL.md +626 -0
  216. package/skills/skills-codex/idea-creator/SKILL.md +405 -0
  217. package/skills/skills-codex/idea-discovery/SKILL.md +475 -0
  218. package/skills/skills-codex/idea-discovery-robot/SKILL.md +362 -0
  219. package/skills/skills-codex/integrity-forensics/SKILL.md +106 -0
  220. package/skills/skills-codex/interview-cheatsheet/SKILL.md +245 -0
  221. package/skills/skills-codex/invention-structuring/SKILL.md +188 -0
  222. package/skills/skills-codex/jurisdiction-format/SKILL.md +192 -0
  223. package/skills/skills-codex/kill-argument/SKILL.md +403 -0
  224. package/skills/skills-codex/mermaid-diagram/SKILL.md +379 -0
  225. package/skills/skills-codex/meta-apply/SKILL.md +154 -0
  226. package/skills/skills-codex/meta-optimize/SKILL.md +348 -0
  227. package/skills/skills-codex/monitor-experiment/SKILL.md +98 -0
  228. package/skills/skills-codex/novelty-check/SKILL.md +89 -0
  229. package/skills/skills-codex/openalex/SKILL.md +228 -0
  230. package/skills/skills-codex/overleaf-sync/SKILL.md +220 -0
  231. package/skills/skills-codex/paper-claim-audit/SKILL.md +350 -0
  232. package/skills/skills-codex/paper-compile/SKILL.md +253 -0
  233. package/skills/skills-codex/paper-figure/SKILL.md +311 -0
  234. package/skills/skills-codex/paper-illustration/SKILL.md +690 -0
  235. package/skills/skills-codex/paper-illustration-image2/SKILL.md +383 -0
  236. package/skills/skills-codex/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  237. package/skills/skills-codex/paper-plan/SKILL.md +278 -0
  238. package/skills/skills-codex/paper-poster/SKILL.md +19 -0
  239. package/skills/skills-codex/paper-poster-html/SKILL.md +377 -0
  240. package/skills/skills-codex/paper-slides/SKILL.md +571 -0
  241. package/skills/skills-codex/paper-talk/SKILL.md +381 -0
  242. package/skills/skills-codex/paper-write/SKILL.md +411 -0
  243. package/skills/skills-codex/paper-write/templates/IEEEtran.bst +2409 -0
  244. package/skills/skills-codex/paper-write/templates/IEEEtran.cls +6347 -0
  245. package/skills/skills-codex/paper-write/templates/aaai2026.bst +1493 -0
  246. package/skills/skills-codex/paper-write/templates/aaai2026.sty +315 -0
  247. package/skills/skills-codex/paper-write/templates/aaai2026.tex +952 -0
  248. package/skills/skills-codex/paper-write/templates/acl.sty +312 -0
  249. package/skills/skills-codex/paper-write/templates/acl2026.tex +377 -0
  250. package/skills/skills-codex/paper-write/templates/acl_natbib.bst +1940 -0
  251. package/skills/skills-codex/paper-write/templates/acm.bst +3081 -0
  252. package/skills/skills-codex/paper-write/templates/acm_mm2026.tex +204 -0
  253. package/skills/skills-codex/paper-write/templates/acmart.cls +3520 -0
  254. package/skills/skills-codex/paper-write/templates/cvpr.bst +1448 -0
  255. package/skills/skills-codex/paper-write/templates/cvpr.sty +508 -0
  256. package/skills/skills-codex/paper-write/templates/cvpr2026.tex +63 -0
  257. package/skills/skills-codex/paper-write/templates/iclr2026.tex +84 -0
  258. package/skills/skills-codex/paper-write/templates/iclr2026_conference.bst +1440 -0
  259. package/skills/skills-codex/paper-write/templates/iclr2026_conference.sty +246 -0
  260. package/skills/skills-codex/paper-write/templates/icml2026.sty +767 -0
  261. package/skills/skills-codex/paper-write/templates/icml2026.tex +662 -0
  262. package/skills/skills-codex/paper-write/templates/ieee_conference.tex +89 -0
  263. package/skills/skills-codex/paper-write/templates/ieee_journal.tex +93 -0
  264. package/skills/skills-codex/paper-write/templates/math_commands.tex +48 -0
  265. package/skills/skills-codex/paper-write/templates/neurips2026.tex +493 -0
  266. package/skills/skills-codex/paper-write/templates/neurips_2026.sty +437 -0
  267. package/skills/skills-codex/paper-writing/SKILL.md +731 -0
  268. package/skills/skills-codex/patent-novelty-check/SKILL.md +153 -0
  269. package/skills/skills-codex/patent-pipeline/SKILL.md +344 -0
  270. package/skills/skills-codex/patent-review/SKILL.md +202 -0
  271. package/skills/skills-codex/pixel-art/SKILL.md +139 -0
  272. package/skills/skills-codex/prior-art-search/SKILL.md +146 -0
  273. package/skills/skills-codex/proof-checker/SKILL.md +554 -0
  274. package/skills/skills-codex/proof-orchestrator/SKILL.md +260 -0
  275. package/skills/skills-codex/proof-orchestrator/references/audit-output-contract.md +126 -0
  276. package/skills/skills-codex/proof-orchestrator/references/deepseek-routing.md +76 -0
  277. package/skills/skills-codex/proof-orchestrator/references/dispatch-prompts.md +227 -0
  278. package/skills/skills-codex/proof-orchestrator/references/notation-audit.md +135 -0
  279. package/skills/skills-codex/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  280. package/skills/skills-codex/proof-orchestrator/references/stress-tests.md +38 -0
  281. package/skills/skills-codex/proof-writer/SKILL.md +222 -0
  282. package/skills/skills-codex/qzcli/SKILL.md +324 -0
  283. package/skills/skills-codex/rebuttal/SKILL.md +305 -0
  284. package/skills/skills-codex/render-html/SKILL.md +305 -0
  285. package/skills/skills-codex/render-html/scripts/__pycache__/render_html.cpython-314.pyc +0 -0
  286. package/skills/skills-codex/render-html/scripts/render_html.py +909 -0
  287. package/skills/skills-codex/render-html/scripts/templates/academic.html +342 -0
  288. package/skills/skills-codex/render-html/scripts/templates/dashboard.html +333 -0
  289. package/skills/skills-codex/research-lit/SKILL.md +464 -0
  290. package/skills/skills-codex/research-pipeline/SKILL.md +340 -0
  291. package/skills/skills-codex/research-refine/SKILL.md +721 -0
  292. package/skills/skills-codex/research-refine-pipeline/SKILL.md +186 -0
  293. package/skills/skills-codex/research-review/SKILL.md +135 -0
  294. package/skills/skills-codex/research-wiki/SKILL.md +421 -0
  295. package/skills/skills-codex/resubmit-pipeline/SKILL.md +444 -0
  296. package/skills/skills-codex/result-to-claim/SKILL.md +246 -0
  297. package/skills/skills-codex/run-experiment/SKILL.md +236 -0
  298. package/skills/skills-codex/semantic-scholar/SKILL.md +219 -0
  299. package/skills/skills-codex/serverless-modal/SKILL.md +335 -0
  300. package/skills/skills-codex/shared-references/acceptance-gate.md +336 -0
  301. package/skills/skills-codex/shared-references/assurance-contract.md +139 -0
  302. package/skills/skills-codex/shared-references/capture-antipatterns.md +84 -0
  303. package/skills/skills-codex/shared-references/citation-discipline.md +452 -0
  304. package/skills/skills-codex/shared-references/compute-env-contract.md +163 -0
  305. package/skills/skills-codex/shared-references/effort-contract.md +143 -0
  306. package/skills/skills-codex/shared-references/evidence-precheck.md +73 -0
  307. package/skills/skills-codex/shared-references/experiment-integrity.md +49 -0
  308. package/skills/skills-codex/shared-references/external-cadence.md +334 -0
  309. package/skills/skills-codex/shared-references/fan-out-pattern.md +375 -0
  310. package/skills/skills-codex/shared-references/injection-hygiene.md +132 -0
  311. package/skills/skills-codex/shared-references/integration-contract.md +372 -0
  312. package/skills/skills-codex/shared-references/output-composition.md +98 -0
  313. package/skills/skills-codex/shared-references/output-language.md +45 -0
  314. package/skills/skills-codex/shared-references/output-manifest.md +40 -0
  315. package/skills/skills-codex/shared-references/output-versioning.md +111 -0
  316. package/skills/skills-codex/shared-references/patent-format-cn.md +199 -0
  317. package/skills/skills-codex/shared-references/patent-format-ep.md +173 -0
  318. package/skills/skills-codex/shared-references/patent-format-us.md +161 -0
  319. package/skills/skills-codex/shared-references/patent-writing-principles.md +197 -0
  320. package/skills/skills-codex/shared-references/prior-art-databases.md +141 -0
  321. package/skills/skills-codex/shared-references/resumable-runs.md +125 -0
  322. package/skills/skills-codex/shared-references/review-scope-limits.md +81 -0
  323. package/skills/skills-codex/shared-references/review-tracing.md +144 -0
  324. package/skills/skills-codex/shared-references/reviewer-independence.md +66 -0
  325. package/skills/skills-codex/shared-references/reviewer-routing.md +128 -0
  326. package/skills/skills-codex/shared-references/skill-governance.md +119 -0
  327. package/skills/skills-codex/shared-references/taste-calibration.md +90 -0
  328. package/skills/skills-codex/shared-references/venue-checklists.md +73 -0
  329. package/skills/skills-codex/shared-references/wiki-helper-resolution.md +69 -0
  330. package/skills/skills-codex/shared-references/writing-principles.md +525 -0
  331. package/skills/skills-codex/slides-polish/SKILL.md +563 -0
  332. package/skills/skills-codex/specification-writing/SKILL.md +211 -0
  333. package/skills/skills-codex/system-profile/SKILL.md +103 -0
  334. package/skills/skills-codex/training-check/SKILL.md +83 -0
  335. package/skills/skills-codex/vast-gpu/SKILL.md +394 -0
  336. package/skills/skills-codex/web-debug-search/SKILL.md +334 -0
  337. package/skills/skills-codex/wiki-enrich/SKILL.md +255 -0
  338. package/skills/skills-codex/writing-systems-papers/SKILL.md +184 -0
  339. package/skills/skills-codex-claude-review/README.md +79 -0
  340. package/skills/skills-codex-claude-review/README_CN.md +78 -0
  341. package/skills/skills-codex-claude-review/auto-paper-improvement-loop/SKILL.md +581 -0
  342. package/skills/skills-codex-claude-review/auto-review-loop/SKILL.md +510 -0
  343. package/skills/skills-codex-claude-review/novelty-check/SKILL.md +102 -0
  344. package/skills/skills-codex-claude-review/paper-figure/SKILL.md +319 -0
  345. package/skills/skills-codex-claude-review/paper-plan/SKILL.md +287 -0
  346. package/skills/skills-codex-claude-review/paper-write/SKILL.md +420 -0
  347. package/skills/skills-codex-claude-review/research-refine/SKILL.md +732 -0
  348. package/skills/skills-codex-claude-review/research-review/SKILL.md +149 -0
  349. package/skills/skills-codex-gemini-review/README.md +176 -0
  350. package/skills/skills-codex-gemini-review/README_CN.md +175 -0
  351. package/skills/skills-codex-gemini-review/auto-paper-improvement-loop/SKILL.md +331 -0
  352. package/skills/skills-codex-gemini-review/auto-review-loop/SKILL.md +304 -0
  353. package/skills/skills-codex-gemini-review/grant-proposal/SKILL.md +630 -0
  354. package/skills/skills-codex-gemini-review/idea-creator/SKILL.md +263 -0
  355. package/skills/skills-codex-gemini-review/idea-discovery/SKILL.md +275 -0
  356. package/skills/skills-codex-gemini-review/idea-discovery-robot/SKILL.md +365 -0
  357. package/skills/skills-codex-gemini-review/novelty-check/SKILL.md +92 -0
  358. package/skills/skills-codex-gemini-review/paper-figure/SKILL.md +289 -0
  359. package/skills/skills-codex-gemini-review/paper-plan/SKILL.md +265 -0
  360. package/skills/skills-codex-gemini-review/paper-poster-html/SKILL.md +102 -0
  361. package/skills/skills-codex-gemini-review/paper-slides/SKILL.md +582 -0
  362. package/skills/skills-codex-gemini-review/paper-write/SKILL.md +346 -0
  363. package/skills/skills-codex-gemini-review/paper-writing/SKILL.md +312 -0
  364. package/skills/skills-codex-gemini-review/research-refine/SKILL.md +674 -0
  365. package/skills/skills-codex-gemini-review/research-review/SKILL.md +112 -0
  366. package/skills/slides-polish/SKILL.md +565 -0
  367. package/skills/specification-writing/SKILL.md +211 -0
  368. package/skills/system-profile/SKILL.md +103 -0
  369. package/skills/training-check/SKILL.md +132 -0
  370. package/skills/vast-gpu/SKILL.md +394 -0
  371. package/skills/web-debug-search/SKILL.md +334 -0
  372. package/skills/wiki-enrich/SKILL.md +257 -0
  373. package/skills/writing-systems-papers/SKILL.md +184 -0
  374. package/templates/CLAUDE_MD_TEMPLATE.md +29 -0
  375. package/templates/EXPERIMENT_LOG_TEMPLATE.md +47 -0
  376. package/templates/EXPERIMENT_PLAN_TEMPLATE.md +51 -0
  377. package/templates/EXPERIMENT_PLAN_TEMPLATE_CN.md +53 -0
  378. package/templates/FINDINGS_TEMPLATE.md +52 -0
  379. package/templates/IDEA_CANDIDATES_TEMPLATE.md +47 -0
  380. package/templates/IDEA_CANDIDATES_TEMPLATE_CN.md +47 -0
  381. package/templates/INVENTION_BRIEF_TEMPLATE.md +87 -0
  382. package/templates/MANIFEST_TEMPLATE.md +7 -0
  383. package/templates/NARRATIVE_REPORT_TEMPLATE.md +49 -0
  384. package/templates/PAPER_PLAN_TEMPLATE.md +47 -0
  385. package/templates/PATENT_CLAIMS_TEMPLATE.md +78 -0
  386. package/templates/PATENT_SPECIFICATION_TEMPLATE.md +68 -0
  387. package/templates/README.md +57 -0
  388. package/templates/RESEARCH_BRIEF_TEMPLATE.md +35 -0
  389. package/templates/RESEARCH_BRIEF_TEMPLATE_CN.md +41 -0
  390. package/templates/RESEARCH_CONTRACT_TEMPLATE.md +60 -0
  391. package/templates/claude-hooks/corpus_write_guard.json +16 -0
  392. package/templates/claude-hooks/corpus_write_guard.py +85 -0
  393. package/templates/claude-hooks/meta_logging.json +74 -0
  394. package/templates/gitignore-trace.txt +3 -0
  395. package/tools/__pycache__/check_skills_inventory.cpython-314.pyc +0 -0
  396. package/tools/arxiv_fetch.py +311 -0
  397. package/tools/capture_filter.py +126 -0
  398. package/tools/check_skills_inventory.py +273 -0
  399. package/tools/convert_skills_to_llm_chat.py +282 -0
  400. package/tools/copilot_native_evidence.py +818 -0
  401. package/tools/deepxiv_fetch.py +213 -0
  402. package/tools/evidence_check.py +212 -0
  403. package/tools/exa_search.py +425 -0
  404. package/tools/experiment_queue/README.md +118 -0
  405. package/tools/experiment_queue/build_manifest.py +44 -0
  406. package/tools/experiment_queue/queue_manager.py +44 -0
  407. package/tools/extract_paper_style.py +560 -0
  408. package/tools/figure_renderer.py +69 -0
  409. package/tools/forensics_gate.py +669 -0
  410. package/tools/generate_codex_claude_review_overrides.py +299 -0
  411. package/tools/idea_discovery_gate.py +256 -0
  412. package/tools/install_aris.ps1 +1372 -0
  413. package/tools/install_aris.sh +1370 -0
  414. package/tools/install_aris_codex.sh +1023 -0
  415. package/tools/install_aris_copilot.sh +1052 -0
  416. package/tools/iteration_log.py +143 -0
  417. package/tools/lint_skills_helpers.sh +84 -0
  418. package/tools/meta_opt/check_ready.sh +80 -0
  419. package/tools/meta_opt/log_event.sh +91 -0
  420. package/tools/meta_opt/trigger_eval.py +280 -0
  421. package/tools/meta_opt/trigger_evals.sample.json +28 -0
  422. package/tools/openalex_fetch.py +326 -0
  423. package/tools/overleaf_audit.sh +104 -0
  424. package/tools/overleaf_setup.sh +150 -0
  425. package/tools/paper_illustration_image2.py +62 -0
  426. package/tools/provenance.py +294 -0
  427. package/tools/research_wiki.py +1720 -0
  428. package/tools/review_gate.py +502 -0
  429. package/tools/run_state.py +399 -0
  430. package/tools/save_trace.sh +477 -0
  431. package/tools/semantic_scholar_fetch.py +438 -0
  432. package/tools/skill-groups.tsv +116 -0
  433. package/tools/skill_picker.py +238 -0
  434. package/tools/smart_update.ps1 +521 -0
  435. package/tools/smart_update.sh +591 -0
  436. package/tools/smart_update_codex.sh +419 -0
  437. package/tools/smart_update_copilot.sh +605 -0
  438. package/tools/threat_scan.py +222 -0
  439. package/tools/verify_paper_audits.sh +487 -0
  440. package/tools/verify_papers.py +613 -0
  441. package/tools/verify_wiki_coverage.sh +176 -0
  442. package/tools/watchdog.py +485 -0
@@ -0,0 +1,554 @@
1
+ ---
2
+ name: proof-checker
3
+ description: Rigorous mathematical proof verification and fixing workflow. Reads a LaTeX proof, identifies gaps via fresh-agent Codex GPT-5.6-Sol ultra review, fixes each gap with full derivations, re-reviews, and generates an audit report. Base review is same-family provisional. Use when user says "检查证明", "verify proof", "proof check", "审证明", "check this proof", or wants rigorous mathematical verification of a theory paper.
4
+ argument-hint: "[path-to-tex-file or proof-description]"
5
+ allowed-tools: Bash(*), Read, Grep, Glob, Write, Edit
6
+ ---
7
+
8
+ # Proof Checker: Rigorous Mathematical Verification & Fixing
9
+
10
+ > **Codex assurance:** base Codex proof judgments are
11
+ > `review_independence: same-family` and `acceptance_status: provisional`.
12
+ > Deterministic compilation/algebra checks may be accepted; a semantic proof
13
+ > acceptance requires a cross-family overlay. Reviewer failure emits BLOCKED.
14
+
15
+ Systematically verify a mathematical proof via fresh-agent adversarial review, fix identified gaps, re-review until convergence, and generate a detailed audit report with proof-obligation accounting.
16
+
17
+ ## Context: $ARGUMENTS
18
+
19
+ ## Constants
20
+
21
+ - MAX_REVIEW_ROUNDS = 3
22
+ - REVIEWER_MODEL = `gpt-5.6-sol` via Codex reviewer agent, reasoning effort
23
+ `ultra` for this deep-audit skill (capability fallback never below `xhigh`)
24
+ - **REVIEWER_BACKEND = `codex`** — Default: Codex reviewer agent (`spawn_agent`, ultra — deep-audit tier). Override with `— reviewer: oracle-pro` for GPT-5.5 Pro via Oracle MCP. See `shared-references/reviewer-routing.md`.
25
+ - AUDIT_DOC: `PROOF_AUDIT.md` at the paper directory root, alongside `main.tex` (cumulative log; when invoked via `/paper-writing`, this is `paper/PROOF_AUDIT.md`)
26
+ - REPORT_TEX: `proof_audit_report.tex` (formal before/after PDF)
27
+ - STATE_FILE: `PROOF_CHECK_STATE.json` (for recovery)
28
+ - SKELETON_DOC: `PROOF_SKELETON.md` (micro-claim inventory)
29
+ - **RENDER_HTML = true** — When `true` (default), auto-render `PROOF_AUDIT.md` to HTML at workflow end via `/render-html`. Uses **full review gate** (audit-class, math-heavy — render-fidelity check protects against MathJax breakage). Set `false` to skip, or pass `— render html: false`. **Non-blocking**: failures don't invalidate the proof audit.
30
+
31
+ ### Acceptance Gate (objective, replaces subjective scoring)
32
+
33
+ The proof passes when ALL of the following hold:
34
+ 1. Zero open FATAL or CRITICAL issues
35
+ 2. Every theorem/lemma has: (i) explicit hypotheses, (ii) proof with all interchanges justified, (iii) every application discharges hypotheses in the ledger
36
+ 3. All big-O/Θ/o statements have declared parameter dependence and uniformity scope
37
+ 4. Counterexample pass executed on all key lemmas (log candidates even if none found)
38
+
39
+ ## Issue Taxonomy (20 categories, 4 groups)
40
+
41
+ ### Group A: Logic & Proof Structure
42
+
43
+ | Category | Description | Example |
44
+ |----------|-------------|---------|
45
+ | **UNJUSTIFIED_ASSERTION** | Claim stated without proof or reference | "The Hessian splits into Gram blocks" |
46
+ | **UNPROVEN_SUBCLAIM** | "Clearly" / "it follows" hides a nontrivial lemma | "By symmetry, the cross-terms vanish" without checking |
47
+ | **QUANTIFIER_ERROR** | Wrong order ∀/∃, missing "for sufficiently small κ" | "For all π, there exists ε" vs "there exists ε for all π" |
48
+ | **IMPLICATION_REVERSAL** | Uses (A⇒B) as (B⇒A), or claims equivalence with only one direction | |
49
+ | **CASE_INCOMPLETE** | Misses boundary/degenerate cases | Singular covariance, zero weight, non-unique argmin |
50
+ | **CIRCULAR_DEPENDENCY** | Lemma uses theorem that depends on it | |
51
+ | **LOGICAL_GAP** | A step is not justified by what precedes it | B=Θ(1) → β_K=0 without analyzing W |
52
+
53
+ ### Group B: Analysis & Measure Theory
54
+
55
+ | Category | Description | Example |
56
+ |----------|-------------|---------|
57
+ | **ILLEGAL_INTERCHANGE** | Swaps limit/expectation/derivative/integral without DCT/MCT/Fubini | Differentiating under E without domination |
58
+ | **NONUNIFORM_CONVERGENCE** | Pointwise convergence used as uniform | sup and limit swapped |
59
+ | **MISSING_DOMINATION** | DCT cited but no dominating function given | |
60
+ | **INTEGRABILITY_GAP** | Uses E|X|^p without proving/assuming finite moments | |
61
+ | **REGULARITY_GAP** | Differentiability/Lipschitz/convexity used but not established | |
62
+ | **STOCHASTIC_MODE_CONFUSION** | Mixes a.s./in prob./in L²/in expectation | |
63
+
64
+ ### Group C: Model & Parameter Tracking
65
+
66
+ | Category | Description | Example |
67
+ |----------|-------------|---------|
68
+ | **MISSING_DERIVATION** | A quantity is used but never derived from the model | Risk functional with undefined B, W |
69
+ | **HIDDEN_ASSUMPTION** | Proof silently uses a condition not in the theorem | Gaussianity assumed but not stated |
70
+ | **INSUFFICIENT_ASSUMPTION** | Hypotheses too weak for proof (counterexample exists) | Moment conditions admitting 2-point distributions |
71
+ | **DIMENSION_TRACKING** | Parameter dependence (d, n, K, ...) not explicit | d enters only through κ |
72
+ | **NORMALIZATION_MISMATCH** | Coordinate/scaling conventions inconsistent | Rescaled vs raw coordinates |
73
+ | **CONSTANT_DEPENDENCE_HIDDEN** | "C" depends on d,n,K but treated as universal | |
74
+
75
+ ### Group D: Scope & Claims
76
+
77
+ | Category | Description | Example |
78
+ |----------|-------------|---------|
79
+ | **SCOPE_OVERCLAIM** | Conclusion stated more broadly than proof supports | "β_K=0" with only generic overlap |
80
+ | **REFERENCE_MISMATCH** | Cited theorem's hypotheses not verified at point of use | |
81
+
82
+ ## Two-Axis Severity System
83
+
84
+ ### Axis A — Proof Status (what is wrong)
85
+
86
+ | Status | Meaning |
87
+ |--------|---------|
88
+ | **INVALID** | Statement false as written (counterexample exists or contradiction) |
89
+ | **UNJUSTIFIED** | Could be true, but current proof does not establish it |
90
+ | **UNDERSTATED** | True only after strengthening assumptions |
91
+ | **OVERSTATED** | True only after weakening conclusion / adding qualifiers |
92
+ | **UNCLEAR** | Ambiguous notation / definition drift (not wrong per se) |
93
+
94
+ ### Axis B — Impact (how much breaks)
95
+
96
+ | Impact | Meaning |
97
+ |--------|---------|
98
+ | **GLOBAL** | Breaks main theorem or core dependency chain |
99
+ | **LOCAL** | Affects a side result but not the main theorem |
100
+ | **COSMETIC** | Exposition only |
101
+
102
+ ### Severity Labels (derived)
103
+
104
+ | Label | Definition |
105
+ |-------|------------|
106
+ | **FATAL** | INVALID + GLOBAL |
107
+ | **CRITICAL** | (INVALID + LOCAL) or (UNJUSTIFIED + GLOBAL) |
108
+ | **MAJOR** | (UNJUSTIFIED + LOCAL) or (UNDERSTATED/OVERSTATED + GLOBAL) |
109
+ | **MINOR** | Clarity / notation / dimension bookkeeping that doesn't change claims |
110
+
111
+ ## Side-Condition Checklists for Common Theorems
112
+
113
+ When the proof invokes any of the following, require explicit verification of ALL listed conditions:
114
+
115
+ | Theorem | Required Conditions |
116
+ |---------|-------------------|
117
+ | **DCT** (Dominated Convergence) | Pointwise a.e. convergence + integrable dominating function |
118
+ | **MCT** (Monotone Convergence) | Monotone increasing + non-negative |
119
+ | **Fubini/Tonelli** | Product measurability + integrability (Fubini) or non-negative (Tonelli) |
120
+ | **Leibniz integral rule** | Continuity of integrand + dominating function for derivative |
121
+ | **Implicit Function Theorem** | Continuous differentiability + non-singular Jacobian |
122
+ | **Taylor with remainder** | Sufficient differentiability + remainder form (Lagrange/integral) |
123
+ | **Jensen's inequality** | Convexity of function + integrability |
124
+ | **Cauchy-Schwarz** | Correct inner product space + integrability of both factors |
125
+ | **Weyl/Davis-Kahan** | Symmetry/Hermiticity + perturbation bound conditions |
126
+ | **Analytic continuation** | Domain connectivity + identity theorem conditions |
127
+ | **WLOG reduction** | Invariance under claimed symmetry + reduction is reversible |
128
+
129
+ ## Workflow
130
+
131
+ ### Proof-obligation fan-out
132
+
133
+ Independent sections/theorems may be extracted by fresh read-only
134
+ `spawn_agent` shards, with a sequential fresh-context fallback. Each shard
135
+ returns `{"shard_id": ..., "entries": [{"payload": ..., "dedup_key":
136
+ "<theorem-or-obligation-id>"}]}` and must not declare the proof valid. The
137
+ parent mechanically merges obligations; the fresh Codex review that evaluates
138
+ them records `review_independence: same-family` and
139
+ `acceptance_status: provisional`. See
140
+ [`fan-out-pattern.md`](../shared-references/fan-out-pattern.md).
141
+
142
+ ### Phase 0: Preparation
143
+
144
+ 1. **Locate the proof**: Find the main `.tex` file(s).
145
+ 2. **Read the entire proof**: Extract list of all theorems/lemmas/propositions/corollaries/definitions/assumptions.
146
+ 3. **Read reference materials**: Reference papers, prior results.
147
+ 4. **Build a section map**: Structured list with line numbers and key claims.
148
+ 5. **Identify the main theorem**: Central result, assumptions, claims.
149
+
150
+ ### Phase 0.5: Proof-Obligation Ledger
151
+
152
+ Build formal accounting artifacts. Save to `PROOF_SKELETON.md`:
153
+
154
+ #### 1. Dependency DAG
155
+ Nodes = Definitions / Assumptions / Lemmas / Theorems. Edges = "uses". **Detect cycles** (including semantic circularity where Lemma A uses a corollary that quietly depends on A).
156
+
157
+ #### 2. Assumption Ledger
158
+ For each theorem/lemma, list every hypothesis with WHERE each is verified (or mark "UNVERIFIED"). Track **usage-minimal assumption sets** — which assumptions were actually used vs merely stated.
159
+
160
+ #### 3. Typed Symbol Table
161
+ Each symbol must have a **type signature**:
162
+ ```
163
+ κ : scalar ∈ (0,1), depends on (d, α_t, Σ, μ)
164
+ u* : vector ∈ ℝ^d, u* = C^{-1}m
165
+ B^even : matrix ∈ ℝ^{(L+1)×(L+1)}, symmetric PSD
166
+ Ψ_v : function ℝ → ℝ, analytic in (ζ,κ), parity determined by v
167
+ ```
168
+ Flag any symbol whose meaning changes or whose type is inconsistent across uses.
169
+
170
+ #### 4. Canonical Quantified Statements
171
+ For each theorem/lemma, rewrite the statement with **explicit quantifiers, domains, and limit order**:
172
+ ```
173
+ ∀K ≥ 3, ∀π ∈ Π_K^{ms,∘} \ E_K, ∃κ_0 > 0 such that ∀κ ∈ (0, κ_0):
174
+ h_act^{(K,π)} = Θ(κ^{α_K^act}) [uniform in π on compact subsets]
175
+ ```
176
+ If you cannot restate a theorem this precisely, mark it **UNCLEAR — needs disambiguation**.
177
+
178
+ #### 5. Micro-Claim Inventory
179
+ Every nontrivial step becomes a numbered micro-claim in **sequent form**:
180
+ ```
181
+ MC-17: Context: [Lemma 3.1, κ < κ_0, Z_κ has bounded moments up to order 2m+2]
182
+ ⊢ Goal: P̂_0 is positive definite
183
+ Rule: monomials linearly independent on support of continuous distribution
184
+ Side-conditions: positive density near origin ✓ (by GMM weak convergence)
185
+ ```
186
+ Each micro-claim has: justification rule name + required conditions + where conditions are proven.
187
+
188
+ #### 6. Limit-Order Map
189
+ Track every asymptotic statement's **limit order and uniformity scope**:
190
+ ```
191
+ h_act = Θ(κ^α) [as κ→0, uniform in π on compact subsets of Π_K, for fixed K]
192
+ τ_act ~ (b/a)n [as n→∞, for fixed κ,K,π with x_K ≪ 1]
193
+ ```
194
+ Flag any statement where limit order is ambiguous or uniformity is unclear.
195
+
196
+ ### Phase 1: First Review (Codex GPT-5.6-Sol ultra)
197
+
198
+ Submit the **complete proof content** with the following **mandatory reviewer checklist** in the prompt:
199
+
200
+ ```text
201
+ spawn_agent:
202
+ model: gpt-5.6-sol
203
+ reasoning_effort: ultra
204
+ message: |
205
+ You are performing a rigorous mathematical proof review. For EVERY theorem,
206
+ lemma, and proposition, check ALL of the following:
207
+
208
+ ## MANDATORY CHECKS
209
+
210
+ A. DEFINITIONS: List any symbol whose meaning is ambiguous or changes.
211
+ B. HYPOTHESIS DISCHARGE: For each lemma/theorem APPLICATION (not statement),
212
+ list each hypothesis and whether it was verified, with location.
213
+ C. INEQUALITY AUDIT: For each inequality chain, verify direction, missing
214
+ absolute values, missing conditions (convexity, PSD, integrability).
215
+ D. INTERCHANGE AUDIT: Flag every limit/derivative/expectation/integral
216
+ interchange. State which theorem justifies it (DCT/MCT/Fubini/Leibniz)
217
+ and which conditions are verified/missing.
218
+ E. PROBABILITY MODE: Track whether claims are a.s./in prob./in expectation/
219
+ w.h.p. Ensure transitions are justified.
220
+ F. UNIFORMITY & CONSTANTS: For every O(·), o(·), Θ(·), ≲, state whether
221
+ it is uniform over all parameters. List hidden parameter dependence.
222
+ G. EDGE/DEGENERATE CASES: Attempt to break each key lemma with a 1D,
223
+ low-rank, or extreme-parameter construction.
224
+ H. DEPENDENCY CONSISTENCY: Detect cycles or forward references to unproven
225
+ results.
226
+
227
+ ## OUTPUT FORMAT (per issue)
228
+ For each issue found, provide:
229
+ - id: sequential number
230
+ - status: INVALID / UNJUSTIFIED / UNDERSTATED / OVERSTATED / UNCLEAR
231
+ - impact: GLOBAL / LOCAL / COSMETIC
232
+ - category: [from taxonomy]
233
+ - location: section/equation/line
234
+ - statement: what the proof claims
235
+ - why_invalid: why this is wrong or unjustified
236
+ - counterexample: YES (describe) / NO / CANDIDATE (describe attempt)
237
+ - affects: which downstream results break if this is wrong
238
+ - minimal_fix: how to fix it
239
+
240
+ [FULL PROOF CONTENT HERE]
241
+ ```
242
+
243
+ **Save the reviewer `agent_id`.** Parse into structured issue list. Write to `PROOF_AUDIT.md`.
244
+
245
+ ### Phase 1.5: Counterexample Red Team
246
+
247
+ For each CRITICAL or MAJOR issue, and for every key lemma that introduces:
248
+ - a new inequality bound
249
+ - an identifiability/uniqueness claim
250
+ - a curvature/PSD/strong convexity assertion
251
+ - a uniform-in-parameter claim
252
+ - a convergence mode upgrade (pointwise → uniform, in prob → w.h.p.)
253
+
254
+ Systematically attempt to construct counterexamples using:
255
+
256
+ | Strategy | Description |
257
+ |----------|-------------|
258
+ | **Dimensional collapse** | Set d=1 or 2, K=2, n small |
259
+ | **Degeneracy** | Singular covariance, tiny weight, overlapping means, identical components |
260
+ | **Extremal distributions** | Two-point ±a, bounded non-subGaussian, heavy tails |
261
+ | **Adversarial parameter scaling** | Pick parameters making neglected terms dominate |
262
+ | **Numeric falsification** | Translate lemma to a function, brute-force optimize over small domain |
263
+
264
+ **Rule**: Label "counterexample found" ONLY if algebraically verified. Otherwise log as "candidate counterexample — needs verification."
265
+
266
+ Record all attempts (successful or not) in `PROOF_AUDIT.md`.
267
+
268
+ ### Phase 2: Fix Implementation
269
+
270
+ For each issue, ordered by severity (FATAL → CRITICAL → MAJOR → MINOR):
271
+
272
+ #### Step 2a: Choose fix strategy
273
+ For each issue, explicitly choose one of:
274
+ - **ADD_DERIVATION**: Write missing proof steps
275
+ - **STRENGTHEN_ASSUMPTION**: Add conditions to theorem statement
276
+ - **WEAKEN_CLAIM**: Reduce conclusion scope
277
+ - **ADD_REFERENCE**: Cite known result + verify its conditions apply
278
+
279
+ Log this choice — it is a scope-changing decision when it alters theorem statements.
280
+
281
+ #### Step 2b: Derive the fix mathematically
282
+ - Complete mathematical derivation, not just a claim
283
+ - If new proposition/lemma needed, write in full theorem-proof style
284
+
285
+ #### Step 2c: Implement in LaTeX
286
+ - Edit the `.tex` file
287
+ - Preserve existing `\label` references where possible
288
+
289
+ #### Step 2d: Record the fix
290
+ ```markdown
291
+ ### Fix N: [SHORT TITLE]
292
+ **Issue**: [id] [CATEGORY] — [description]
293
+ **Severity**: FATAL / CRITICAL / MAJOR / MINOR
294
+ **Status**: INVALID / UNJUSTIFIED / UNDERSTATED / OVERSTATED
295
+ **Impact**: GLOBAL / LOCAL / COSMETIC
296
+ **Fix strategy**: ADD_DERIVATION / STRENGTHEN_ASSUMPTION / WEAKEN_CLAIM / ADD_REFERENCE
297
+ **Location**: Section X, Lines Y-Z
298
+
299
+ **BEFORE**: [what the proof originally did]
300
+ **WHY WRONG**: [mathematical problem, with counterexample if applicable]
301
+ **AFTER**: [what the fix does]
302
+ **KEY EQUATION**: [central new equation]
303
+ **PROOF OBLIGATIONS ADDED**: [new conditions/lemmas introduced]
304
+ **DOWNSTREAM EFFECTS**: [which results now need re-checking]
305
+ ```
306
+
307
+ #### Step 2e: Compile check
308
+ ```bash
309
+ pdflatex -interaction=nonstopmode <file>.tex 2>&1 | grep -E "Error|Warning|undefined"
310
+ ```
311
+
312
+ ### Phase 3: Re-Review (Codex GPT-5.6-Sol ultra)
313
+
314
+ Launch a fresh reviewer agent for the next review round. Do not use `send_input` here; proof-checker keeps each round independent. Request the same mandatory checklist.
315
+
316
+ Check acceptance gate. If not met, repeat Phases 2-3 (up to MAX_REVIEW_ROUNDS).
317
+
318
+ ### Phase 3.5: Global Closure & Independent Verification
319
+
320
+ #### Global closure checks
321
+ After all fixes, verify the proof as a whole:
322
+ - **Statement–conclusion match**: Does the proof end with EXACTLY what the theorem claims (quantifiers, constants, uniformity)?
323
+ - **All obligations discharged**: Every node in the obligation DAG is proven or explicitly assumed (and the theorem statement includes it).
324
+ - **Case analysis coverage**: Cases partition the domain AND include boundary/degenerate cases.
325
+ - **Induction correctness** (if applicable): Base case, inductive step, correct use of IH, induction measure strictly decreases.
326
+ - **WLOG reductions**: Each "without loss of generality" spawns a micro-claim proving the reduction is lossless.
327
+ - **No silent assumption strengthening**: Any fix that strengthened assumptions has propagated to the main theorem statement.
328
+
329
+ #### Independent second review for FATAL/CRITICAL fixes
330
+ For any fix that resolved a FATAL or CRITICAL issue, submit the **fixed section alone** (without showing the previous critique) to a **fresh Codex thread**:
331
+
332
+ ```text
333
+ spawn_agent:
334
+ model: gpt-5.6-sol
335
+ reasoning_effort: ultra
336
+ message: |
337
+ Blind review of the following proof section. You have NOT seen any prior
338
+ review or discussion. Check every step for correctness, hidden assumptions,
339
+ illegal interchanges, and counterexamples.
340
+ [FIXED SECTION ONLY]
341
+ ```
342
+
343
+ If the blind reviewer finds new issues, re-enter Phase 2.
344
+
345
+ #### Regression proof-audit
346
+ After fixes, re-run:
347
+ - DAG acyclicity check (no new cycles introduced)
348
+ - Counterexample suite on all DOWNSTREAM lemmas of modified results
349
+ - Assumption-delta report: what became stronger/weaker due to fixes?
350
+
351
+ ### Phase 3.9: Unrecoverable Proof Protocol
352
+
353
+ If acceptance gate is not met after MAX_REVIEW_ROUNDS, output a **Proof Unrecoverable Report**:
354
+ 1. Minimal set of blocking FATAL/CRITICAL issues that could not be resolved
355
+ 2. Salvage options ranked: (a) weaken claim, (b) strengthen assumptions, (c) add missing lemmas, (d) restructure argument
356
+ 3. Which parts of the proof are likely still reusable
357
+ 4. Recommended next steps for the author
358
+
359
+ Do NOT silently declare success. The report must be honest.
360
+
361
+ ### Phase 4: Audit Report Generation
362
+
363
+ Generate `proof_audit_report.tex` with:
364
+
365
+ 1. **Overview table**: All issues with two-axis severity, category, fix strategy, status
366
+ 2. **Before/After logic chain**: Red (BEFORE) → Green (AFTER) comparison
367
+ 3. **For each fix**: original proof → why wrong → counterexample (if any) → complete derivation → remaining subtleties
368
+ 4. **Proof-obligation diff**: What was unverified before, what is verified now
369
+ 5. **Summary**: Now proven / still assumed / open problems
370
+ 6. **Colored boxes**: BEFORE (red), AFTER (green), WHY WRONG (orange), KEY INSIGHT (blue), WARNING (yellow)
371
+
372
+ Compile: `pdflatex proof_audit_report.tex && pdflatex proof_audit_report.tex`
373
+
374
+ ### Phase 5: State Persistence
375
+
376
+ Write `PROOF_CHECK_STATE.json`:
377
+ ```json
378
+ {
379
+ "status": "completed",
380
+ "rounds": 2,
381
+ "review_agent_ids": ["..."],
382
+ "fatal_fixed": 0,
383
+ "critical_fixed": 3,
384
+ "major_fixed": 2,
385
+ "minor_fixed": 1,
386
+ "counterexamples_found": 1,
387
+ "counterexample_candidates": 2,
388
+ "acceptance_gate": "PASS",
389
+ "timestamp": "..."
390
+ }
391
+ ```
392
+
393
+ ### Phase 5.5: Research Wiki Claim Ledger (additive; only if a wiki is active)
394
+
395
+ If — and only if — a `research-wiki/` exists, persist each top-level theorem/headline
396
+ as a **claim node** (the wiki's PROVE/JUDGE ledger). This is the **birth point** for wiki
397
+ claim nodes. It is a **detect-only record, never a verdict**: it never changes the audit's
398
+ `verdict`/`reason_code`, never blocks, and is skipped when `verdict == NOT_APPLICABLE` or no
399
+ wiki is found. The claim's `status` is the **PROOF axis only** ({drafted, unproven,
400
+ sound-modulo-imports, verified, refuted, retracted}); empirical experiment support is a
401
+ separate axis carried by edges (`/result-to-claim`), never written into `status`.
402
+
403
+ Resolve the helper via the Codex-side chain (skip cleanly if unavailable; the audit is
404
+ already complete):
405
+ ```
406
+ ARIS_REPO="${ARIS_REPO:-$(awk -F'\t' '$1=="repo_root"{print $2; exit}' .aris/installed-skills-codex.txt 2>/dev/null)}"
407
+ WIKI_SCRIPT=""
408
+ [ -n "$ARIS_REPO" ] && [ -f "$ARIS_REPO/tools/research_wiki.py" ] && WIKI_SCRIPT="$ARIS_REPO/tools/research_wiki.py"
409
+ [ -z "$WIKI_SCRIPT" ] && [ -f tools/research_wiki.py ] && WIKI_SCRIPT="tools/research_wiki.py"
410
+ [ -z "$WIKI_SCRIPT" ] && [ -f ~/.codex/skills/research-wiki/research_wiki.py ] && WIKI_SCRIPT="$HOME/.codex/skills/research-wiki/research_wiki.py"
411
+ ```
412
+
413
+ If `research-wiki/` exists and `WIKI_SCRIPT` is available and `verdict != NOT_APPLICABLE`,
414
+ for each top-level theorem map the audit outcome to an honest status — `PASS`/all proofs
415
+ complete → `verified`; closes modulo flagged imports → `sound-modulo-imports`; counterexample
416
+ found or statement judged false → `refuted`; open gap (UNJUSTIFIED, no counterexample) →
417
+ `unproven` (never fake a gap as `refuted`/`verified`) — then record it (idempotent):
418
+ ```
419
+ python3 "$WIKI_SCRIPT" add_claim research-wiki/ --slug "<stable-theorem-id>" \
420
+ --name "<theorem headline>" --status "<mapped status>" \
421
+ --provenance "<trace_path from PROOF_AUDIT.json>" --statement "<canonical statement>" \
422
+ --scope "<what it does NOT say; flagged imports>" --update-on-exist
423
+ ```
424
+ `add_claim` failure is non-fatal (warn and continue; the audit is unaffected).
425
+
426
+ ## Key Rules
427
+
428
+ ### Mathematical rigor
429
+ - **Never accept a proof step on faith**. "Clearly" / "it follows" / "by standard arguments" are red flags — each must spawn a micro-claim.
430
+ - **Hypothesis discharge**: Every time a lemma is APPLIED, verify EACH of its hypotheses at that point. Use the side-condition checklists above.
431
+ - **Interchange discipline**: Every swap of limit/expectation/derivative/integral must cite a theorem (DCT/MCT/Fubini/Leibniz) and verify its conditions with explicit dominating function or integrability proof.
432
+ - **Uniformity discipline**: Every O(·)/Θ(·) must declare what parameters it is uniform over. "O(1)" that secretly depends on d,n,K is a CONSTANT_DEPENDENCE_HIDDEN issue.
433
+ - **Quantifier discipline**: Check ∀/∃ order. "For sufficiently small κ" must specify: does κ₀ depend on K? On π? On d?
434
+ - **Counterexample-first**: Before trying to fix a gap, first try to break it.
435
+ - **WLOG prohibition**: Every "without loss of generality" must have an explicit micro-claim proving the reduction. No free WLOGs.
436
+ - **No silent assumption strengthening**: Any fix that adds conditions must propagate to the theorem statement.
437
+
438
+ ### Review-independence protocol
439
+ - **Codex executor analyzes and implements; a fresh Codex reviewer provides adversarial review.** Base review remains same-family/provisional.
440
+ - **Codex reasoning always ultra** (deep-audit tier): never below `xhigh` — only the capability fallback in `reviewer-routing.md` may step down, and only on explicit capability errors.
441
+ - **Send full content**: Don't summarize — send actual math for line-by-line checking.
442
+ - **Fresh reviewer agents**: Save each returned `agent_id` for traceability, but launch a new `spawn_agent` for each review round. Do not use `send_input` across proof-checker rounds.
443
+
444
+ ### Fix quality
445
+ - **Minimal fixes**: Fix exactly what's broken, nothing more.
446
+ - **Full derivation**: Every fix includes complete mathematical argument.
447
+ - **Explicit scope decisions**: Each fix is tagged ADD_DERIVATION / STRENGTHEN_ASSUMPTION / WEAKEN_CLAIM / ADD_REFERENCE.
448
+ - **Compile after each fix**: LaTeX must compile cleanly.
449
+
450
+ ### Scope honesty
451
+ - **Don't overclaim**: If a fix makes a result conditional, say so.
452
+ - **Separate "proven" from "assumed"**: The audit report has an explicit section for this.
453
+ - **Log open problems**: Issues requiring future work are listed, not hidden.
454
+
455
+ ## Output Files
456
+
457
+ | File | Content | When |
458
+ |------|---------|------|
459
+ | `PROOF_SKELETON.md` | Dependency DAG + assumption ledger + micro-claims | Phase 0.5 |
460
+ | `PROOF_AUDIT.md` | Cumulative round-by-round audit log | Updated each round |
461
+ | `PROOF_AUDIT.json` | Machine-readable submission verdict (see below) | Always emitted |
462
+ | `proof_audit_report.tex/.pdf` | Formal before/after report | Phase 4 |
463
+ | `PROOF_CHECK_STATE.json` | State for recovery | Phase 5 |
464
+ | `PROOF_AUDIT.html` (+ `.review.json` sidecar) | Single-file HTML view auto-rendered via `/render-html "PROOF_AUDIT.md" --json "PROOF_AUDIT.json"`. **Non-blocking** — if `/render-html` fails the audit still counts as complete. | Workflow end (when `RENDER_HTML = true`, default) |
465
+
466
+ ## Submission Artifact Emission
467
+
468
+ This skill **always** writes `PROOF_AUDIT.json` at the paper directory
469
+ root (i.e. `paper/PROOF_AUDIT.json` when invoked from `/paper-writing`
470
+ with paper-dir `paper/`; `<your-paper-dir>/PROOF_AUDIT.json` when invoked
471
+ standalone), regardless of caller or whether the paper contains theorems.
472
+ A paper with no `\begin{theorem}` / `\begin{lemma}` / `\begin{proof}` emits
473
+ verdict `NOT_APPLICABLE`; silent skip is forbidden. `paper-writing`
474
+ Phase 6 and `verify_paper_audits.sh` both rely on this artifact
475
+ existing at `<paper-dir>/PROOF_AUDIT.json`.
476
+
477
+ The artifact conforms to the schema in `shared-references/assurance-contract.md`:
478
+
479
+ ```json
480
+ {
481
+ "audit_skill": "proof-checker",
482
+ "verdict": "PASS | WARN | FAIL | NOT_APPLICABLE | BLOCKED | ERROR",
483
+ "reason_code": "all_proofs_complete | minor_gaps | critical_gap | no_theorems | ...",
484
+ "summary": "One-line human-readable verdict summary.",
485
+ "audited_input_hashes": {
486
+ "main.tex": "sha256:...",
487
+ "sections/4.theory.tex": "sha256:..."
488
+ },
489
+ "trace_path": ".aris/traces/proof-checker/<date>_run<NN>/",
490
+ "thread_id": "<codex mcp thread id>",
491
+ "executor_model": "codex-gpt-5.6-sol",
492
+ "executor_family": "openai",
493
+ "reviewer_model": "gpt-5.6-sol",
494
+ "reviewer_family": "openai",
495
+ "review_independence": "same-family",
496
+ "acceptance_status": "provisional",
497
+ "reviewer_reasoning": "ultra",
498
+ "generated_at": "<UTC ISO-8601>",
499
+ "details": {
500
+ "theorems_audited": <int>,
501
+ "issues": [ { "id": "T1-H3", "severity": "FATAL|CRITICAL|MAJOR|MINOR",
502
+ "category": "quantifier|domination|...",
503
+ "location": "sections/4.theory.tex:L182",
504
+ "note": "..." }, ... ]
505
+ }
506
+ }
507
+ ```
508
+
509
+ ### `audited_input_hashes` scope
510
+
511
+ Hash the **declared input set** actually reviewed — the theorem-bearing
512
+ `.tex` files passed into this invocation — not a repo-wide union and not
513
+ the reviewer's self-reported opened subset. The external verifier rehashes
514
+ these entries; any mismatch flags `STALE`.
515
+
516
+ **Path convention** (must match `verify_paper_audits.sh`): keys are
517
+ **paths relative to the paper directory** (no `paper/` prefix — the
518
+ verifier resolves relative to the paper dir; prefixing produces
519
+ `paper/paper/...` and false-fails as STALE). Use **absolute paths** for
520
+ files outside the paper dir.
521
+
522
+ ### Verdict decision table
523
+
524
+ | Input state | Verdict | `reason_code` example |
525
+ |-------------------------------------------------------|------------------|-----------------------|
526
+ | No theorems / lemmas / proofs in paper | `NOT_APPLICABLE` | `no_theorems` |
527
+ | Theorems present but referenced files unreadable | `BLOCKED` | `source_unreadable` |
528
+ | All proof obligations discharged, no gaps | `PASS` | `all_proofs_complete` |
529
+ | Only MINOR issues (notation / exposition) | `WARN` | `minor_gaps` |
530
+ | Any FATAL or CRITICAL issue (logic gap, wrong claim) | `FAIL` | `critical_gap` |
531
+ | Reviewer invocation failed (network / malformed) | `ERROR` | `reviewer_error` |
532
+
533
+ MAJOR issues alone map to `WARN` or `FAIL` at the reviewer's discretion and
534
+ must carry an explicit justification in `summary` + `details.issues`.
535
+
536
+ ### Thread independence
537
+
538
+ Every invocation uses a fresh reviewer agent. Never use `send_input` across
539
+ proof-checker runs. Do not accept prior audit outputs
540
+ (PAPER_CLAIM_AUDIT, CITATION_AUDIT, EXPERIMENT_LOG) as input — the fresh
541
+ thread preserves reviewer independence per
542
+ `shared-references/reviewer-independence.md`.
543
+
544
+ This skill never blocks by itself; `paper-writing` Phase 6 plus the
545
+ verifier decide whether the verdict blocks finalization based on the
546
+ `assurance` level.
547
+
548
+ ## Example Invocations
549
+
550
+ ```
551
+ /proof-checker "neurips_2025.tex"
552
+ /proof-checker "check the GMM generalization proof, focus on dimension dependence"
553
+ /proof-checker "verify proof in paper.tex — difficulty: nightmare"
554
+ ```