@furongjun1999/dsh-memory 0.4.11 → 0.5.1

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (580) hide show
  1. package/README.md +552 -465
  2. package/codebuddy/CODEBUDDY.md +11 -3
  3. package/codebuddy/README.md +92 -90
  4. package/codebuddy/mcp.json +27 -27
  5. package/docs/GBrain/345/217/257/345/200/237/351/211/264/347/202/271_/347/201/265/346/236/242/350/220/275/347/202/271/344/272/244/346/216/245_20260919.md +169 -169
  6. package/docs/Pi/345/217/257/345/255/246/344/271/240/344/274/230/347/202/271_/347/201/265/346/236/242/345/244/247/350/204/221/346/224/271/350/277/233/344/272/244/346/216/245_20260915.md +144 -144
  7. package/docs/README.md +143 -111
  8. package/docs/discipline/harnesses.yaml +244 -226
  9. package/docs/discipline/templates/full.md.tmpl +61 -61
  10. package/docs/discipline/templates/rules.mdc.tmpl +68 -0
  11. package/docs/discipline/templates/skill.md.tmpl +23 -23
  12. package/docs/eval/AGI/344/270/203/347/273/264/350/257/204/345/210/206/346/212/245/345/221/212_md_cg_v1.0.md +299 -299
  13. package/docs/eval/AGI/344/270/203/347/273/264/350/257/204/345/210/206/346/212/245/345/221/212_md_cg_v2.0.md +239 -239
  14. package/docs/eval/bench_lingshu_self/bench_self.py +140 -0
  15. package/docs/eval/bench_lingshu_self/self_bench_result.json +404 -0
  16. package/docs/eval/bench_lingshu_self//347/201/265/346/236/242/350/207/252/345/272/223/347/253/257/345/210/260/347/253/257/346/243/200/347/264/242/345/256/236/346/265/213_v1.0.md +38 -0
  17. package/docs/eval//344/270/215/345/217/257/351/235/240/346/200/247/350/220/275/345/234/260_P0_v1.0.md +185 -0
  18. package/docs/eval//345/256/236/351/252/214/346/226/271/346/241/210_/345/255/246/344/271/240/351/227/255/347/216/257AB/344/270/216/346/250/252/350/257/204_v1.0.md +174 -174
  19. package/docs/eval//346/225/205/351/232/234/346/263/250/345/205/245/345/256/236/346/265/213_v1.0.md +422 -0
  20. package/docs/eval//346/250/252/350/257/204_/345/205/255/345/256/266100/351/242/230/344/270/255/350/213/261/345/217/214/346/237/245_v1.0.md +223 -223
  21. package/docs/eval//347/253/257/345/210/260/347/253/257LoCoMoQA/345/220/214/345/217/243/345/276/204/345/257/271/347/205/247_v1.0.md +100 -0
  22. package/docs/eval//347/253/257/345/210/260/347/253/257/345/271/262/346/211/260/346/261/240/350/257/204/346/265/213_/347/241/256/345/256/232/346/200/247/350/243/201/345/206/263vsLLM_judge_v1.1.md +197 -0
  23. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v1.md +156 -0
  24. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v10.md +210 -0
  25. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v11.md +227 -0
  26. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v12.md +203 -0
  27. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v13.md +233 -0
  28. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v14.md +191 -0
  29. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v15.md +213 -0
  30. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v16.md +214 -0
  31. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v17.md +199 -0
  32. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v2.md +156 -0
  33. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v3.md +152 -0
  34. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v4.md +128 -0
  35. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v5.md +114 -0
  36. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v6.md +192 -0
  37. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v7.md +187 -0
  38. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v8.md +207 -0
  39. package/docs/eval//347/274/272/351/231/267/346/214/226/346/216/230_/350/207/252/344/270/273/350/277/255/344/273/243_v9.md +203 -0
  40. package/docs/{mdcg → hive}//344/273/244/347/211/214/344/270/216/350/247/222/350/211/262/346/235/203/350/201/214/345/210/206/347/246/273_v0.1.md +165 -160
  41. package/docs/hive//344/273/244/347/211/214/350/257/255/344/271/211/344/277/256/346/255/243_/344/273/262/350/243/201/344/275/215_v0.1.md +76 -0
  42. package/docs/{mdcg → hive}//345/255/220/344/273/243/347/220/206/351/205/215/347/275/256/346/240/207/345/207/206_v0.5.md +236 -236
  43. package/docs/hive//345/256/211/345/205/250/345/256/241/350/256/241/345/256/236/351/224/232_v0.1.md +202 -0
  44. package/docs/hive//346/243/200/347/264/242/346/224/266/346/225/233/345/256/236/346/265/213/344/270/216S1b/350/256/276/350/256/241_v0.1.md +43 -43
  45. package/docs/hive//346/243/200/347/264/242/347/256/227/346/263/225/345/217/243/345/276/204/345/257/271/347/205/247_v0.1.md +85 -0
  46. package/docs/hive//346/243/200/347/264/242/350/267/257/345/276/204/344/270/216/350/256/244/347/237/245/347/273/223/346/236/204/345/245/221/347/272/246_v0.1.md +102 -102
  47. package/docs/hive//350/234/202/345/267/242M6_ingest/345/256/236/346/226/275/350/256/241/345/210/222_v0.1.md +47 -0
  48. package/docs/hive//350/234/202/345/267/242/345/217/214/345/256/236/344/276/213/344/272/222/351/252/214_/350/256/276/350/256/241/345/256/232/347/250/277.md +519 -503
  49. package/docs/hive//350/234/202/345/267/242/345/267/245/344/275/234/350/256/260/345/277/206_/351/241/271/347/233/256/350/256/241/345/210/222.md +85 -85
  50. package/docs/hive//350/234/202/345/267/242/350/256/276/350/256/241_/347/220/206/350/256/272/345/257/271/351/275/220_v0.1.md +191 -0
  51. package/docs/hive//350/234/202/345/267/242/350/277/255/344/273/243_/345/256/217/350/247/202/344/270/216/347/276/244/344/275/223/350/260/203/345/272/246_v0.1.md +90 -0
  52. package/docs/images/lingshu-moonlight-covenant-poster-preview.jpg +0 -0
  53. package/docs/images/lingshu-moonlight-covenant-poster.png +0 -0
  54. package/docs/mdcg/D_meta_/345/267/245/347/250/213/345/214/226/346/226/271/346/241/210_v0.2.md +216 -216
  55. package/docs/mdcg/README/350/257/246/347/273/206/347/211/210_v0.4.10.md +646 -646
  56. package/docs/mdcg/lingshu_tutorial.html +14449 -14449
  57. package/docs/mdcg/release_v0.4.11.md +49 -0
  58. package/docs/mdcg/release_v0.4.5.md +55 -55
  59. package/docs/mdcg/tool_table_v0.3.0.md +117 -117
  60. package/docs/mdcg//345/205/250/345/272/223/344/273/243/347/240/201/350/257/204/345/256/241/344/270/216/346/235/241/344/273/266/345/214/226/346/263/250/351/207/212_/350/256/241/345/210/222_v0.1.md +600 -600
  61. package/docs/mdcg//345/212/237/350/203/275/350/260/203/347/224/250/346/230/240/345/260/204/350/241/250_v0.1.md +40 -40
  62. package/docs/mdcg//345/215/225/345/205/203/350/207/252/346/210/221/351/224/232/347/202/271_/347/263/273/347/273/237/346/217/220/347/244/272/350/257/215/346/240/207/345/207/206_v0.3.md +275 -275
  63. package/docs/mdcg//345/217/221/345/270/203/351/227/250/347/246/201/351/223/276_v0.1.md +54 -0
  64. package/docs/mdcg//346/272/220/347/240/201/347/272/247/346/236/266/346/236/204/345/256/241/350/256/241_GPT/346/211/271/350/257/204/345/257/271/347/205/247_v1.0.md +130 -130
  65. package/docs/mdcg//347/201/265/346/236/24282/345/267/245/345/205/267_/345/212/237/350/203/275/346/225/264/347/220/206/344/270/216/350/277/201/347/247/273/346/230/240/345/260/204_v0.1.md +278 -278
  66. package/docs/mdcg//347/201/265/346/236/242/350/256/260/345/277/206/345/212/250/350/257/215/345/215/217/350/256/256_v1.0-draft.md +172 -172
  67. package/docs/mdcg//347/274/272/345/217/243/345/215/225_P0/346/224/266/345/217/243_v0.1.md +270 -270
  68. package/docs/mdcg//350/256/244/347/237/245/345/233/276_G4-G8/347/274/272/345/217/243/350/243/201/345/256/232/345/215/225_v0.1.md +491 -491
  69. package/docs/mdcg//350/256/244/347/237/245/345/233/276_/346/235/241/344/273/266/347/251/272/351/227/264/345/220/210/346/210/220/344/270/216/347/224/237/346/225/210/346/235/241/344/273/266/345/217/243/345/276/204_v0.1.md +340 -340
  70. package/docs/mdcg//350/256/244/347/237/245/345/233/276_/347/264/242/345/274/225/344/270/216/345/267/245/347/250/213/350/247/204/350/214/203/345/214/226_/350/256/241/345/210/222_v0.1.md +588 -588
  71. package/docs/plans/GridWorld/346/234/200/345/260/217/351/227/255/347/216/257/350/247/204/346/240/274_v0.1.md +176 -176
  72. package/docs/plans//345/221/275/345/220/215/346/262/273/347/220/206_/351/241/271/347/233/256/350/256/241/345/210/222.md +81 -81
  73. package/docs/swarm//350/234/202/347/276/244/344/272/222/350/201/224_v0.1.md +704 -704
  74. package/docs/swarm//350/234/202/347/276/244/345/220/214/351/224/231/346/243/200/346/265/213/345/256/236/351/252/214/345/215/217/350/256/256_v0.1.md +242 -242
  75. package/docs/theory//344/270/215/345/217/257/351/235/240/345/256/232/347/220/206/344/270/216/345/244/261/346/225/210/344/274/230/345/205/210/346/241/206/346/236/266_v0.3.md +365 -0
  76. package/docs/theory//345/215/225/347/272/277/347/250/213/344/270/216/346/263/250/346/204/217/345/212/233/351/233/206/344/270/255_/346/227/240/344/272/211/350/256/256/347/220/206/350/256/272/346/226/207/346/241/243_v1.0.md +129 -129
  77. package/docs/theory//345/271/266/345/217/221/345/277/205/347/204/266/346/200/247/347/220/206/350/256/272_v0.2.md +134 -134
  78. package/docs/theory//346/231/272/350/203/275/350/256/2723.4.md +5260 -5260
  79. package/docs/theory//346/246/202/345/277/265/345/210/206/345/261/202/345/257/271/351/275/220/350/241/250_v0.1.md +145 -145
  80. package/docs/theory//347/220/206/350/256/272_/346/234/272/345/210/266_/344/273/243/347/240/201_/345/256/236/351/252/214_/347/274/272/345/217/243/347/237/251/351/230/265_v0.1.md +144 -144
  81. package/docs/theory//347/220/206/350/256/272/344/273/223/346/213/206/345/210/206/344/270/216/347/231/275/347/256/261/347/237/245/350/257/206/345/272/223/345/206/205/350/277/201_v0.1.md +198 -198
  82. package/docs/theory//350/256/244/347/237/245/344/273/243/347/220/206/344/270/216/346/224/266/346/225/233/347/273/223/346/236/204_/346/235/241/344/273/266/350/256/272/351/207/215/346/236/204_v0.1.md +221 -221
  83. package/docs/world_model//344/270/226/347/225/214/346/250/241/345/236/213_/347/245/236/347/273/217/347/275/221/347/273/234/345/272/225/345/261/202/346/236/266/346/236/204_v1.0.md +237 -237
  84. package/docs//344/270/215/345/217/257/351/235/240/346/200/247/347/220/206/350/256/272_v0.1.md +343 -0
  85. package/docs//345/217/221/345/270/203/344/273/266/345/233/236/346/272/257/350/257/264/346/230/216_v0.1.md +82 -82
  86. package/docs//345/267/245/344/275/234/347/272/252/345/276/213_/350/256/244/347/237/245/345/233/276/346/235/241/347/233/256_v1.1.json +28 -2
  87. package/dsh/README.md +82 -82
  88. package/dsh/cordis.yml.example +139 -139
  89. package/dsh/update-lingshu.bat +11 -11
  90. package/lib/bridge.d.ts +9 -0
  91. package/lib/bridge.js +35 -0
  92. package/lib/hooks.js +36 -2
  93. package/lib/index.js +7 -1
  94. package/lib/lib/roleplay_web.js +116 -29
  95. package/lib/lib/token_store.d.ts +7 -1
  96. package/lib/lib/token_store.js +12 -3
  97. package/md_cg/__init__.py +7 -7
  98. package/md_cg/audit.py +379 -368
  99. package/md_cg/autonomy.py +287 -287
  100. package/md_cg/backfill.py +1328 -1327
  101. package/md_cg/backfill_bigdomain.py +34 -34
  102. package/md_cg/backfill_bucket_zh.py +35 -0
  103. package/md_cg/bench6_arms.py +410 -410
  104. package/md_cg/bench6_common.py +230 -230
  105. package/md_cg/bench6_competitors.py +212 -212
  106. package/md_cg/bench_axis_domain.py +257 -257
  107. package/md_cg/bench_blind_comp.py +308 -308
  108. package/md_cg/bench_e2e_judge.py +532 -0
  109. package/md_cg/bench_e2e_locomo_qa.py +368 -0
  110. package/md_cg/bench_e2e_qa.py +256 -0
  111. package/md_cg/bench_en_atoms_public.py +230 -230
  112. package/md_cg/bench_governance.py +348 -348
  113. package/md_cg/bench_lme_zh.py +410 -410
  114. package/md_cg/bench_locomo.py +121 -121
  115. package/md_cg/bench_locomo_zh.py +450 -450
  116. package/md_cg/bench_locomo_zh_public.py +147 -147
  117. package/md_cg/bench_longmem.py +112 -112
  118. package/md_cg/bench_membench.py +632 -632
  119. package/md_cg/bench_p0.py +149 -149
  120. package/md_cg/bench_progressive.py +287 -287
  121. package/md_cg/bench_role_views.py +238 -238
  122. package/md_cg/bench_task_ab.py +243 -243
  123. package/md_cg/bench_task_ab_llm.py +408 -408
  124. package/md_cg/bench_unified_en.py +204 -204
  125. package/md_cg/bench_zh_mad.py +601 -601
  126. package/md_cg/blindspot_tickets.py +123 -123
  127. package/md_cg/branches.py +301 -285
  128. package/md_cg/build_postings.py +73 -73
  129. package/md_cg/ccgc.py +1006 -948
  130. package/md_cg/census.py +132 -132
  131. package/md_cg/chain.py +315 -300
  132. package/md_cg/codeindex.py +531 -531
  133. package/md_cg/coldverify.py +292 -292
  134. package/md_cg/comment_gate.py +337 -337
  135. package/md_cg/cond_compose.py +190 -190
  136. package/md_cg/cond_facts.py +154 -154
  137. package/md_cg/cond_template.json +106 -106
  138. package/md_cg/condition_anchor.py +142 -142
  139. package/md_cg/conformance.py +726 -726
  140. package/md_cg/consistency.py +717 -717
  141. package/md_cg/consolidate.py +1537 -1439
  142. package/md_cg/corpus.py +110 -110
  143. package/md_cg/crosscheck.py +1098 -1097
  144. package/md_cg/crypto.py +3 -1
  145. package/md_cg/d_meta.py +310 -310
  146. package/md_cg/datapath.py +78 -18
  147. package/md_cg/docindex.py +473 -473
  148. package/md_cg/eval_common.py +575 -575
  149. package/md_cg/evidence.py +4 -2
  150. package/md_cg/evolution.py +477 -477
  151. package/md_cg/export.py +222 -220
  152. package/md_cg/forgetting.py +581 -581
  153. package/md_cg/fsutil.py +377 -329
  154. package/md_cg/hotcache.py +48 -7
  155. package/md_cg/hyperedge.py +251 -251
  156. package/md_cg/identity.py +390 -390
  157. package/md_cg/insight.py +500 -500
  158. package/md_cg/interop.py +338 -0
  159. package/md_cg/judgment_manifest.py +177 -0
  160. package/md_cg/lexicon/build_cedict_en_zh.py +329 -329
  161. package/md_cg/lexicon/build_standard_en.py +171 -171
  162. package/md_cg/lexicon/expand_en_zh.py +211 -211
  163. package/md_cg/lifecycle.py +272 -272
  164. package/md_cg/linkref.py +280 -280
  165. package/md_cg/links.py +140 -107
  166. package/md_cg/mcp_server.py +403 -55
  167. package/md_cg/md_whitebox.py +345 -345
  168. package/md_cg/mdcg.py +783 -233
  169. package/md_cg/mdcos.py +500 -70
  170. package/md_cg/metacognition.py +591 -591
  171. package/md_cg/migrate.py +119 -119
  172. package/md_cg/migrate_aeis.py +221 -221
  173. package/md_cg/migrate_roleplay.py +293 -293
  174. package/md_cg/migrate_wisdom_graph.py +360 -360
  175. package/md_cg/mreview/__init__.py +25 -25
  176. package/md_cg/mreview/__main__.py +110 -110
  177. package/md_cg/mreview/bundle.py +178 -178
  178. package/md_cg/mreview/candidates.py +262 -262
  179. package/md_cg/mreview/govern.py +694 -693
  180. package/md_cg/mreview/locate.py +939 -939
  181. package/md_cg/mreview/pipeline.py +728 -728
  182. package/md_cg/mreview/rules/duplication.json +21 -21
  183. package/md_cg/mreview/rules/field_coverage.json +54 -54
  184. package/md_cg/mreview/rules/source_license.json +21 -21
  185. package/md_cg/mreview/rules/template_flow.json +21 -21
  186. package/md_cg/mreview/ruleset.py +252 -252
  187. package/md_cg/nodefile.py +575 -575
  188. package/md_cg/pooling.py +484 -472
  189. package/md_cg/postings.py +300 -298
  190. package/md_cg/predict.py +1100 -1100
  191. package/md_cg/progressive.py +123 -123
  192. package/md_cg/protect.py +272 -272
  193. package/md_cg/protocol/md_cg_gate.proto +33 -33
  194. package/md_cg/protocol.py +372 -372
  195. package/md_cg/provenance.py +582 -582
  196. package/md_cg/reach.py +453 -453
  197. package/md_cg/readcache.py +143 -0
  198. package/md_cg/reconcile.py +228 -0
  199. package/md_cg/refindex.py +833 -833
  200. package/md_cg/refine.py +604 -604
  201. package/md_cg/review_cli.py +170 -0
  202. package/md_cg/roleviews.py +89 -89
  203. package/md_cg/routing.py +393 -365
  204. package/md_cg/run_tests.py +211 -0
  205. package/md_cg/scrub.py +13 -3
  206. package/md_cg/security.py +128 -18
  207. package/md_cg/self_state.py +1029 -1029
  208. package/md_cg/selfreport.py +152 -151
  209. package/md_cg/semantic/__init__.py +10 -10
  210. package/md_cg/semantic/canonical.py +122 -122
  211. package/md_cg/semantic/en_normalizer.py +364 -364
  212. package/md_cg/semantic/en_zh_map.json +28694 -0
  213. package/md_cg/semantic/export_en_zh_map.py +64 -0
  214. package/md_cg/semantic/unify.py +45 -0
  215. package/md_cg/semantic/zh_en_atoms.py +139 -139
  216. package/md_cg/signer.py +7 -4
  217. package/md_cg/sources.py +816 -582
  218. package/md_cg/statushdr.py +179 -179
  219. package/md_cg/stg.py +54 -37
  220. package/md_cg/subgraph.py +729 -729
  221. package/md_cg/sustain.py +35 -5
  222. package/md_cg/tasks.py +470 -470
  223. package/md_cg/test_access_hints.py +147 -0
  224. package/md_cg/test_action_derive.py +203 -203
  225. package/md_cg/test_audit_rotate.py +270 -270
  226. package/md_cg/test_autonomy.py +143 -143
  227. package/md_cg/test_bench_governance.py +102 -102
  228. package/md_cg/test_blindspot_tickets.py +166 -166
  229. package/md_cg/test_branch_discard_tombstone.py +136 -0
  230. package/md_cg/test_branches.py +13 -3
  231. package/md_cg/test_ccg_perturb.py +184 -184
  232. package/md_cg/test_ccgc.py +433 -433
  233. package/md_cg/test_census_prune.py +81 -81
  234. package/md_cg/test_chain_read_isolate.py +168 -0
  235. package/md_cg/test_cond_compose_anchors.py +76 -76
  236. package/md_cg/test_cond_match.py +165 -165
  237. package/md_cg/test_condition_anchor.py +81 -81
  238. package/md_cg/test_d_meta.py +412 -412
  239. package/md_cg/test_datapath_device_name.py +203 -0
  240. package/md_cg/test_datapath_root.py +199 -199
  241. package/md_cg/test_emit_negtail_cache.py +156 -0
  242. package/md_cg/test_en_pipeline.py +22 -2
  243. package/md_cg/test_gain_gate.py +212 -212
  244. package/md_cg/test_govern_directread.py +421 -0
  245. package/md_cg/test_health_scale.py +173 -173
  246. package/md_cg/test_hive_ingest.py +285 -0
  247. package/md_cg/test_hot_cold.py +215 -215
  248. package/md_cg/test_hyperedge.py +245 -245
  249. package/md_cg/test_i26_empty_first_write.py +116 -0
  250. package/md_cg/test_i27_e041_identity.py +128 -0
  251. package/md_cg/test_i28_hotcache_prodpath.py +122 -0
  252. package/md_cg/test_i32_hotcache_env_key.py +218 -0
  253. package/md_cg/test_identity_attribution.py +96 -15
  254. package/md_cg/test_index_durability.py +17 -3
  255. package/md_cg/test_interop.py +95 -0
  256. package/md_cg/test_interop_judgment.py +228 -0
  257. package/md_cg/test_issue39_utf8_stdio.py +273 -0
  258. package/md_cg/test_lifecycle.py +309 -309
  259. package/md_cg/test_linkref.py +306 -306
  260. package/md_cg/test_links_concurrent_write.py +188 -0
  261. package/md_cg/test_lock.py +43 -43
  262. package/md_cg/test_md_access_parity.py +255 -255
  263. package/md_cg/test_md_writepath.py +345 -345
  264. package/md_cg/test_mdstore_search_parity.py +160 -0
  265. package/md_cg/test_merge_upsert.py +168 -0
  266. package/md_cg/test_mr_m2.py +587 -587
  267. package/md_cg/test_mr_m3.py +710 -710
  268. package/md_cg/test_mr_m4.py +485 -485
  269. package/md_cg/test_n123_derive_expiry_chain.py +205 -0
  270. package/md_cg/test_n130_verify_falsified_protect.py +185 -0
  271. package/md_cg/test_n131_merge_gate.py +205 -0
  272. package/md_cg/test_p0.py +250 -250
  273. package/md_cg/test_p1.py +316 -316
  274. package/md_cg/test_p10_identity.py +173 -173
  275. package/md_cg/test_p11_consistency.py +233 -233
  276. package/md_cg/test_p12_metacognition.py +212 -212
  277. package/md_cg/test_p13_encryption.py +241 -241
  278. package/md_cg/test_p14_sustain.py +249 -249
  279. package/md_cg/test_p15_scrub.py +280 -280
  280. package/md_cg/test_p16_self_state.py +301 -301
  281. package/md_cg/test_p17_predict.py +354 -354
  282. package/md_cg/test_p18_whitebox.py +171 -171
  283. package/md_cg/test_p19_migrate_roleplay.py +149 -149
  284. package/md_cg/test_p1x_ref_root.py +160 -0
  285. package/md_cg/test_p20_evolution.py +315 -315
  286. package/md_cg/test_p21_tokens.py +293 -270
  287. package/md_cg/test_p22_theory.py +175 -175
  288. package/md_cg/test_p23_links.py +311 -311
  289. package/md_cg/test_p24_evidence.py +227 -227
  290. package/md_cg/test_p25_weights.py +156 -156
  291. package/md_cg/test_p26_refindex.py +416 -416
  292. package/md_cg/test_p27_docindex.py +16 -7
  293. package/md_cg/test_p28_refcheck.py +305 -305
  294. package/md_cg/test_p29_session_ingest_export.py +354 -333
  295. package/md_cg/test_p2_mcp.py +3 -0
  296. package/md_cg/test_p3.py +11 -2
  297. package/md_cg/test_p30_maintain.py +330 -330
  298. package/md_cg/test_p31_insight.py +534 -534
  299. package/md_cg/test_p32_backfill.py +7 -1
  300. package/md_cg/test_p33_ccg_wiring.py +293 -293
  301. package/md_cg/test_p34_crosscheck.py +331 -331
  302. package/md_cg/test_p35_conditioned_claim.py +252 -252
  303. package/md_cg/test_p36_kp_align.py +230 -230
  304. package/md_cg/test_p37_condition_space.py +248 -248
  305. package/md_cg/test_p38_concurrent_flush.py +102 -0
  306. package/md_cg/test_p38_contextualize.py +273 -273
  307. package/md_cg/test_p39_verify_flow.py +153 -0
  308. package/md_cg/test_p39_vision_evidence.py +369 -369
  309. package/md_cg/test_p40_refine_worklist.py +241 -241
  310. package/md_cg/test_p41_evolve_patrol.py +224 -224
  311. package/md_cg/test_p42_provenance.py +269 -269
  312. package/md_cg/test_p43_pooling.py +412 -398
  313. package/md_cg/test_p44_md_whitebox.py +231 -231
  314. package/md_cg/test_p45_session_identity.py +219 -219
  315. package/md_cg/test_p46_unit_scope.py +272 -272
  316. package/md_cg/test_p47_session_view.py +316 -0
  317. package/md_cg/test_p4_fuzzy.py +223 -223
  318. package/md_cg/test_p5_semantic.py +226 -226
  319. package/md_cg/test_p6_consolidate.py +440 -387
  320. package/md_cg/test_p7_goals_recent.py +202 -202
  321. package/md_cg/test_p8_subgraph_chain.py +200 -200
  322. package/md_cg/test_p9_forget_protect.py +231 -231
  323. package/md_cg/test_predict_beta.py +135 -135
  324. package/md_cg/test_preflight_failclosed.py +100 -100
  325. package/md_cg/test_progressive.py +146 -146
  326. package/md_cg/test_propose_tail_index.py +157 -0
  327. package/md_cg/test_protocol.py +243 -243
  328. package/md_cg/test_reach.py +378 -378
  329. package/md_cg/test_reach_keys.py +201 -201
  330. package/md_cg/test_read_clip.py +141 -141
  331. package/md_cg/test_read_scope_b27.py +277 -0
  332. package/md_cg/test_readcache_default_on.py +168 -0
  333. package/md_cg/test_readcache_precise_inval.py +270 -0
  334. package/md_cg/test_readcache_prodpath.py +203 -0
  335. package/md_cg/test_reconcile_v0.py +294 -0
  336. package/md_cg/test_retr_gates_prodpath.py +140 -0
  337. package/md_cg/test_retr_s1.py +344 -340
  338. package/md_cg/test_retr_s1b.py +276 -209
  339. package/md_cg/test_retr_s3.py +194 -194
  340. package/md_cg/test_retr_s4.py +163 -163
  341. package/md_cg/test_retr_s5.py +200 -200
  342. package/md_cg/test_retr_s6.py +157 -157
  343. package/md_cg/test_retr_s7.py +392 -384
  344. package/md_cg/test_retr_s8_time.py +369 -316
  345. package/md_cg/test_retr_s9_edges.py +286 -286
  346. package/md_cg/test_retr_s9_entity_ctx.py +9 -3
  347. package/md_cg/test_retr_score_once.py +208 -0
  348. package/md_cg/test_review_conformance.py +367 -367
  349. package/md_cg/test_review_onepass.py +170 -0
  350. package/md_cg/test_role_views.py +354 -354
  351. package/md_cg/test_rrf_graph_seed_cache.py +154 -0
  352. package/md_cg/test_security_audit.py +155 -0
  353. package/md_cg/test_security_audit_b26.py +161 -0
  354. package/md_cg/test_security_audit_v21.py +250 -0
  355. package/md_cg/test_sem_noise.py +242 -242
  356. package/md_cg/test_semantic_canonical.py +16 -2
  357. package/md_cg/test_session_isolation.py +168 -0
  358. package/md_cg/test_snapshot_autoclose.py +187 -0
  359. package/md_cg/test_subproc_encoding.py +192 -192
  360. package/md_cg/test_sustain_mutual.py +153 -153
  361. package/md_cg/test_tail_watermark_race.py +208 -0
  362. package/md_cg/test_tasks.py +409 -409
  363. package/md_cg/test_tenant_env_override_warn.py +139 -0
  364. package/md_cg/test_tenant_registry_corrupt_warn.py +151 -0
  365. package/md_cg/test_tool_face.py +189 -189
  366. package/md_cg/test_transfer.py +180 -180
  367. package/md_cg/test_trust.py +361 -361
  368. package/md_cg/test_twophase.py +286 -286
  369. package/md_cg/test_v14_fixes.py +38 -20
  370. package/md_cg/test_validity_filter.py +280 -280
  371. package/md_cg/test_verify_answer.py +138 -138
  372. package/md_cg/test_verify_dirty_reconcile.py +157 -0
  373. package/md_cg/test_wisdom_md_store.py +292 -292
  374. package/md_cg/test_writelimit.py +197 -197
  375. package/md_cg/test_writepipe.py +214 -214
  376. package/md_cg/theory.py +6 -3
  377. package/md_cg/tokens.py +85 -14
  378. package/md_cg/tool_face.py +260 -260
  379. package/md_cg/trust.py +986 -950
  380. package/md_cg/twophase.py +231 -231
  381. package/md_cg/units.py +668 -667
  382. package/md_cg/vision_evidence.py +667 -666
  383. package/md_cg/weights.py +624 -624
  384. package/md_cg/whitebox.py +527 -527
  385. package/md_cg/whitebox_kb/__init__.py +37 -37
  386. package/md_cg/whitebox_kb/aeis_core/__init__.py +42 -42
  387. package/md_cg/whitebox_kb/aeis_core/semantic.py +280 -280
  388. package/md_cg/whitebox_kb/aeis_core/textutil.py +13 -13
  389. package/md_cg/whitebox_kb/engine.py +310 -310
  390. package/md_cg/whitebox_kb/seed_knowledge//346/231/272/350/203/275/350/256/2723.4.md +5260 -5260
  391. package/md_cg/whitebox_kb/wisdom/browser_units.py +2631 -2631
  392. package/md_cg/whitebox_kb/wisdom/causal_discover.py +432 -432
  393. package/md_cg/whitebox_kb/wisdom/chat_engine.py +1506 -1506
  394. package/md_cg/whitebox_kb/wisdom/compiler_code_units.py +3033 -3033
  395. package/md_cg/whitebox_kb/wisdom/condition_algebra.py +112 -112
  396. package/md_cg/whitebox_kb/wisdom/condition_frame.py +315 -315
  397. package/md_cg/whitebox_kb/wisdom/condition_kb.py +108 -108
  398. package/md_cg/whitebox_kb/wisdom/conflict_map.json +4445 -4445
  399. package/md_cg/whitebox_kb/wisdom/core/lexer.py +512 -512
  400. package/md_cg/whitebox_kb/wisdom/core/name_checker.py +1023 -1023
  401. package/md_cg/whitebox_kb/wisdom/cspmn.py +258 -258
  402. package/md_cg/whitebox_kb/wisdom/csre.py +264 -264
  403. package/md_cg/whitebox_kb/wisdom/danmaku_audit.py +252 -252
  404. package/md_cg/whitebox_kb/wisdom/distilled_condition_units.json +6417 -6417
  405. package/md_cg/whitebox_kb/wisdom/docs/WB-EVAL-20260902b.json +1950 -1950
  406. package/md_cg/whitebox_kb/wisdom/docs/WB-EVAL-20260902c.json +1296 -1296
  407. package/md_cg/whitebox_kb/wisdom/docs/whitebox_capability_graph_demo.json +59 -59
  408. package/md_cg/whitebox_kb/wisdom/graph_db_units.py +3089 -3089
  409. package/md_cg/whitebox_kb/wisdom/knowledge_points.py +318 -318
  410. package/md_cg/whitebox_kb/wisdom/md_access.py +470 -470
  411. package/md_cg/whitebox_kb/wisdom/md_store.py +276 -251
  412. package/md_cg/whitebox_kb/wisdom/migrate_wisdom.py +476 -476
  413. package/md_cg/whitebox_kb/wisdom/multilang_ir.py +128 -124
  414. package/md_cg/whitebox_kb/wisdom/navigate.py +248 -248
  415. package/md_cg/whitebox_kb/wisdom/neural_retrieve.py +224 -224
  416. package/md_cg/whitebox_kb/wisdom/os_units.py +2735 -2735
  417. package/md_cg/whitebox_kb/wisdom/pattern_separation.py +376 -376
  418. package/md_cg/whitebox_kb/wisdom/prereq_map.json +364 -364
  419. package/md_cg/whitebox_kb/wisdom/python_code_units.py +2766 -2766
  420. package/md_cg/whitebox_kb/wisdom/role_solidified.json +7 -7
  421. package/md_cg/whitebox_kb/wisdom/route_memory.py +244 -244
  422. package/md_cg/whitebox_kb/wisdom/scene_reconstruction.py +148 -148
  423. package/md_cg/whitebox_kb/wisdom/snr_report.json +40 -40
  424. package/md_cg/whitebox_kb/wisdom/test_code_compose_domains.py +7541 -7541
  425. package/md_cg/whitebox_kb/wisdom/test_compiler_self_bootstrap.py +354 -354
  426. package/md_cg/whitebox_kb/wisdom/test_ecosystem_assembly.py +158 -158
  427. package/md_cg/whitebox_kb/wisdom/test_ecosystem_demos.py +193 -193
  428. package/md_cg/whitebox_kb/wisdom/test_graph_db.py +292 -292
  429. package/md_cg/whitebox_kb/wisdom/test_python_self_bootstrap.py +160 -160
  430. package/md_cg/whitebox_kb/wisdom/trigger_words_index.json +4366 -4366
  431. package/md_cg/whitebox_kb/wisdom/verifier.py +1297 -1297
  432. package/md_cg/writelimit.py +356 -356
  433. package/md_cg/writepipe.py +20 -8
  434. package/package.json +101 -96
  435. package/skills/plugin.json +54 -54
  436. package/skills/skills/designer-perspective/SKILL.md +158 -158
  437. package/skills/skills/designer-perspective/references/01-observation-position.md +66 -66
  438. package/skills/skills/designer-perspective/references/02-structure-recognition.md +62 -62
  439. package/skills/skills/designer-perspective/references/03-direction-judgment.md +55 -55
  440. package/skills/skills/designer-perspective/references/04-qualification-verdict.md +72 -72
  441. package/skills/skills/designer-perspective/references/05-condition-attribution.md +74 -74
  442. package/skills/skills/designer-perspective/scripts/designer.py +545 -545
  443. package/skills/skills/designer-perspective/tests/cases.jsonl +17 -17
  444. package/skills/skills/designer-perspective/tests/selftest.py +61 -61
  445. package/skills/skills/lingshu-browser/SKILL.md +60 -60
  446. package/skills/skills/lingshu-compiler/SKILL.md +56 -56
  447. package/skills/skills/lingshu-compiler/units/analyze-type-infer/SKILL.md +45 -45
  448. package/skills/skills/lingshu-compiler/units/check-name-real/SKILL.md +45 -45
  449. package/skills/skills/lingshu-compiler/units/compile-assign/SKILL.md +45 -45
  450. package/skills/skills/lingshu-compiler/units/compile-expr-tree/SKILL.md +45 -45
  451. package/skills/skills/lingshu-compiler/units/compile-full-pipeline/SKILL.md +45 -45
  452. package/skills/skills/lingshu-compiler/units/compile-func-def/SKILL.md +45 -45
  453. package/skills/skills/lingshu-compiler/units/compile-if-then/SKILL.md +45 -45
  454. package/skills/skills/lingshu-compiler/units/compile-logic-expr/SKILL.md +45 -45
  455. package/skills/skills/lingshu-compiler/units/compile-recursive/SKILL.md +45 -45
  456. package/skills/skills/lingshu-compiler/units/compile-scope/SKILL.md +45 -45
  457. package/skills/skills/lingshu-compiler/units/compile-type-check/SKILL.md +45 -45
  458. package/skills/skills/lingshu-compiler/units/compile-while/SKILL.md +45 -45
  459. package/skills/skills/lingshu-compiler/units/compiler-0010c4bf/SKILL.md +45 -45
  460. package/skills/skills/lingshu-compiler/units/compiler-0355bffb/SKILL.md +45 -45
  461. package/skills/skills/lingshu-compiler/units/compiler-0361708a/SKILL.md +45 -45
  462. package/skills/skills/lingshu-compiler/units/compiler-054a0414/SKILL.md +45 -45
  463. package/skills/skills/lingshu-compiler/units/compiler-05a1691a/SKILL.md +45 -45
  464. package/skills/skills/lingshu-compiler/units/compiler-05eeed1e/SKILL.md +45 -45
  465. package/skills/skills/lingshu-compiler/units/compiler-0622a1f6/SKILL.md +45 -45
  466. package/skills/skills/lingshu-compiler/units/compiler-08b54217/SKILL.md +45 -45
  467. package/skills/skills/lingshu-compiler/units/compiler-0a62b70c/SKILL.md +45 -45
  468. package/skills/skills/lingshu-compiler/units/compiler-0ab24d00/SKILL.md +45 -45
  469. package/skills/skills/lingshu-compiler/units/compiler-0e093688/SKILL.md +45 -45
  470. package/skills/skills/lingshu-compiler/units/compiler-0e82b966/SKILL.md +45 -45
  471. package/skills/skills/lingshu-compiler/units/compiler-0ee9b9b9/SKILL.md +45 -45
  472. package/skills/skills/lingshu-compiler/units/compiler-0f3787e6/SKILL.md +45 -45
  473. package/skills/skills/lingshu-compiler/units/compiler-1028685f/SKILL.md +45 -45
  474. package/skills/skills/lingshu-compiler/units/compiler-11897630/SKILL.md +45 -45
  475. package/skills/skills/lingshu-compiler/units/compiler-16661b9b/SKILL.md +45 -45
  476. package/skills/skills/lingshu-compiler/units/compiler-1f722303/SKILL.md +45 -45
  477. package/skills/skills/lingshu-compiler/units/compiler-25be1262/SKILL.md +45 -45
  478. package/skills/skills/lingshu-compiler/units/compiler-2a76ba07/SKILL.md +45 -45
  479. package/skills/skills/lingshu-compiler/units/compiler-2dbea54a/SKILL.md +45 -45
  480. package/skills/skills/lingshu-compiler/units/compiler-2df52f16/SKILL.md +45 -45
  481. package/skills/skills/lingshu-compiler/units/compiler-2ec2c9d2/SKILL.md +45 -45
  482. package/skills/skills/lingshu-compiler/units/compiler-2f8c8f39/SKILL.md +45 -45
  483. package/skills/skills/lingshu-compiler/units/compiler-38378ac7/SKILL.md +45 -45
  484. package/skills/skills/lingshu-compiler/units/compiler-39457d2e/SKILL.md +45 -45
  485. package/skills/skills/lingshu-compiler/units/compiler-39fb5926/SKILL.md +45 -45
  486. package/skills/skills/lingshu-compiler/units/compiler-47f4fbfa/SKILL.md +45 -45
  487. package/skills/skills/lingshu-compiler/units/compiler-4a8cd1f1/SKILL.md +45 -45
  488. package/skills/skills/lingshu-compiler/units/compiler-4b230b6b/SKILL.md +45 -45
  489. package/skills/skills/lingshu-compiler/units/compiler-4cbbda95/SKILL.md +45 -45
  490. package/skills/skills/lingshu-compiler/units/compiler-4d5a68ff/SKILL.md +45 -45
  491. package/skills/skills/lingshu-compiler/units/compiler-56dc9bdd/SKILL.md +45 -45
  492. package/skills/skills/lingshu-compiler/units/compiler-57e76ebe/SKILL.md +45 -45
  493. package/skills/skills/lingshu-compiler/units/compiler-61ff016b/SKILL.md +45 -45
  494. package/skills/skills/lingshu-compiler/units/compiler-63eae588/SKILL.md +45 -45
  495. package/skills/skills/lingshu-compiler/units/compiler-64c4224b/SKILL.md +45 -45
  496. package/skills/skills/lingshu-compiler/units/compiler-66377955/SKILL.md +45 -45
  497. package/skills/skills/lingshu-compiler/units/compiler-674ab4f6/SKILL.md +45 -45
  498. package/skills/skills/lingshu-compiler/units/compiler-6af76fd6/SKILL.md +45 -45
  499. package/skills/skills/lingshu-compiler/units/compiler-6d979db2/SKILL.md +45 -45
  500. package/skills/skills/lingshu-compiler/units/compiler-6de326c0/SKILL.md +45 -45
  501. package/skills/skills/lingshu-compiler/units/compiler-6f1c0eea/SKILL.md +45 -45
  502. package/skills/skills/lingshu-compiler/units/compiler-724d6c9b/SKILL.md +45 -45
  503. package/skills/skills/lingshu-compiler/units/compiler-7690177f/SKILL.md +45 -45
  504. package/skills/skills/lingshu-compiler/units/compiler-7b229a7d/SKILL.md +45 -45
  505. package/skills/skills/lingshu-compiler/units/compiler-7f471b3a/SKILL.md +45 -45
  506. package/skills/skills/lingshu-compiler/units/compiler-815cad08/SKILL.md +45 -45
  507. package/skills/skills/lingshu-compiler/units/compiler-83c61634/SKILL.md +45 -45
  508. package/skills/skills/lingshu-compiler/units/compiler-8c798c81/SKILL.md +45 -45
  509. package/skills/skills/lingshu-compiler/units/compiler-8dd747ac/SKILL.md +45 -45
  510. package/skills/skills/lingshu-compiler/units/compiler-90324d0a/SKILL.md +45 -45
  511. package/skills/skills/lingshu-compiler/units/compiler-94b8d72d/SKILL.md +45 -45
  512. package/skills/skills/lingshu-compiler/units/compiler-94f12231/SKILL.md +45 -45
  513. package/skills/skills/lingshu-compiler/units/compiler-95937c16/SKILL.md +45 -45
  514. package/skills/skills/lingshu-compiler/units/compiler-98a5625b/SKILL.md +45 -45
  515. package/skills/skills/lingshu-compiler/units/compiler-98b4c42f/SKILL.md +45 -45
  516. package/skills/skills/lingshu-compiler/units/compiler-98e3894a/SKILL.md +45 -45
  517. package/skills/skills/lingshu-compiler/units/compiler-9bdfe4b8/SKILL.md +45 -45
  518. package/skills/skills/lingshu-compiler/units/compiler-9d9b4e83/SKILL.md +45 -45
  519. package/skills/skills/lingshu-compiler/units/compiler-9f7be5ad/SKILL.md +45 -45
  520. package/skills/skills/lingshu-compiler/units/compiler-a6daf076/SKILL.md +45 -45
  521. package/skills/skills/lingshu-compiler/units/compiler-a7995a0c/SKILL.md +45 -45
  522. package/skills/skills/lingshu-compiler/units/compiler-a8399248/SKILL.md +45 -45
  523. package/skills/skills/lingshu-compiler/units/compiler-aabbd099/SKILL.md +45 -45
  524. package/skills/skills/lingshu-compiler/units/compiler-afe169d8/SKILL.md +45 -45
  525. package/skills/skills/lingshu-compiler/units/compiler-b0678bda/SKILL.md +45 -45
  526. package/skills/skills/lingshu-compiler/units/compiler-b09bd196/SKILL.md +45 -45
  527. package/skills/skills/lingshu-compiler/units/compiler-b1396e23/SKILL.md +45 -45
  528. package/skills/skills/lingshu-compiler/units/compiler-b9b31ce0/SKILL.md +45 -45
  529. package/skills/skills/lingshu-compiler/units/compiler-c10264a7/SKILL.md +45 -45
  530. package/skills/skills/lingshu-compiler/units/compiler-cb1e8e4b/SKILL.md +45 -45
  531. package/skills/skills/lingshu-compiler/units/compiler-ccafd438/SKILL.md +45 -45
  532. package/skills/skills/lingshu-compiler/units/compiler-cda9c262/SKILL.md +45 -45
  533. package/skills/skills/lingshu-compiler/units/compiler-ce648068/SKILL.md +45 -45
  534. package/skills/skills/lingshu-compiler/units/compiler-cf5776a4/SKILL.md +45 -45
  535. package/skills/skills/lingshu-compiler/units/compiler-d974e5d3/SKILL.md +45 -45
  536. package/skills/skills/lingshu-compiler/units/compiler-e3979fd3/SKILL.md +45 -45
  537. package/skills/skills/lingshu-compiler/units/compiler-eb1cf2b5/SKILL.md +45 -45
  538. package/skills/skills/lingshu-compiler/units/compiler-ecb30d5b/SKILL.md +45 -45
  539. package/skills/skills/lingshu-compiler/units/compiler-f8c8b24b/SKILL.md +45 -45
  540. package/skills/skills/lingshu-compiler/units/compiler-f99fedbe/SKILL.md +45 -45
  541. package/skills/skills/lingshu-compiler/units/compiler-fa8ff5f7/SKILL.md +45 -45
  542. package/skills/skills/lingshu-compiler/units/compiler-fe1b058d/SKILL.md +45 -45
  543. package/skills/skills/lingshu-compiler/units/lex-chinese-program/SKILL.md +45 -45
  544. package/skills/skills/lingshu-compiler/units/lex-dao-de-jing/SKILL.md +45 -45
  545. package/skills/skills/lingshu-compiler/units/lex-nine-chapters/SKILL.md +45 -45
  546. package/skills/skills/lingshu-compiler/units/vm-arithmetic/SKILL.md +45 -45
  547. package/skills/skills/lingshu-compiler/units/vm-array-ops/SKILL.md +45 -45
  548. package/skills/skills/lingshu-compiler/units/vm-closure-call/SKILL.md +45 -45
  549. package/skills/skills/lingshu-compiler/units/vm-closure-create/SKILL.md +45 -45
  550. package/skills/skills/lingshu-compiler/units/vm-compare/SKILL.md +45 -45
  551. package/skills/skills/lingshu-compiler/units/vm-cond-jump/SKILL.md +45 -45
  552. package/skills/skills/lingshu-compiler/units/vm-cond-space/SKILL.md +45 -45
  553. package/skills/skills/lingshu-compiler/units/vm-exception/SKILL.md +45 -45
  554. package/skills/skills/lingshu-compiler/units/vm-func-call/SKILL.md +45 -45
  555. package/skills/skills/lingshu-compiler/units/vm-loop-run/SKILL.md +45 -45
  556. package/skills/skills/lingshu-compiler/units/vm-profiling/SKILL.md +45 -45
  557. package/skills/skills/lingshu-compiler/units/vm-refcount/SKILL.md +45 -45
  558. package/skills/skills/lingshu-compiler/units/vm-run-loop/SKILL.md +45 -45
  559. package/skills/skills/lingshu-compiler/units/vm-short-circuit/SKILL.md +45 -45
  560. package/skills/skills/lingshu-compiler/units/vm-stack-guard/SKILL.md +45 -45
  561. package/skills/skills/lingshu-compiler/units/vm-stack-ops/SKILL.md +45 -45
  562. package/skills/skills/lingshu-compiler/units/vm-trust-accum/SKILL.md +45 -45
  563. package/skills/skills/lingshu-graph/SKILL.md +63 -63
  564. package/skills/skills/lingshu-net/SKILL.md +48 -48
  565. package/skills/skills/lingshu-os/SKILL.md +64 -64
  566. package/skills/skills/lingshu-pylang/SKILL.md +71 -71
  567. package/src/bridge.ts +33 -0
  568. package/src/hooks.ts +38 -2
  569. package/src/index.ts +526 -518
  570. package/src/lib/datapath.ts +326 -326
  571. package/src/lib/mdcg_client.ts +413 -413
  572. package/src/lib/mutual.ts +428 -428
  573. package/src/lib/prompt_safety.ts +62 -62
  574. package/src/lib/python_path.ts +71 -71
  575. package/src/lib/roleplay_web.ts +116 -29
  576. package/src/lib/token_store.ts +13 -3
  577. package/src/tools.ts +212 -212
  578. package/zcode/AGENTS.md +11 -3
  579. package/zcode/README.md +41 -41
  580. /package/docs/{mdcg → hive}//344/270/273/344/273/243/347/220/206/345/255/220/344/273/243/347/220/206/350/256/260/345/277/206/346/236/266/346/236/204/350/256/276/350/256/241.md" +0 -0
@@ -1,710 +1,710 @@
1
- # -*- coding: utf-8 -*-
2
- """M3 定位自测(D1 字段级定位):六类定位器 + 契约四键 + 批量/包 + 零写入。
3
-
4
- 真源:docs/记忆评审系统_立项设计与施工交接_20260915.md §5 D1
5
- `locate(node_id, issue_hint) -> [{field, span, issue_kind, evidence}]`
6
-
7
- 纪律(同 test_mr_m1/m2):
8
- · **自备数据源**——合成节点 + tempfile 合成库,零依赖真源库(外部 clone 全绿)。
9
- · **零写入实锤**——全部相位跑完后认知图指纹逐字节不变(M3 是只读模块)。
10
- · **确定性**——同一输入两次调用逐字节一致;`now` 显式传入,不靠墙钟。
11
- · **不猜**——语义级矛盾归 blindspot(`contradiction_semantic`),不编造字符区间。
12
-
13
- 覆盖:A 基础工具(纯函数) B 问题面定位器(含字段层门限、指纹不一致成因)+ 观测面
14
- (observation_aged:观测时刻不是失效声明,**不进告警面**)
15
- C 提示过滤与 blindspot(含 stale 的**依赖存在性**判据:载体消失/漂移)
16
- D 契约与确定性 E 批量与包 F 零写入
17
- 运行:python -m md_cg.test_mr_m3
18
- """
19
- from __future__ import annotations
20
-
21
- import hashlib
22
- import json
23
- import os
24
- import shutil
25
- import sys
26
- import tempfile
27
-
28
- from . import codeindex as CI
29
- from . import conformance as CF
30
- from . import nodefile as NF
31
- from . import writelimit as WL
32
- from .mreview import locate as LC
33
-
34
- PASS = FAIL = SKIP = 0
35
- FAILS = []
36
-
37
-
38
- def ok(cond, label):
39
- global PASS, FAIL
40
- if cond:
41
- PASS += 1
42
- else:
43
- FAIL += 1
44
- FAILS.append(label)
45
- print(" FAIL %s" % label)
46
-
47
-
48
- def skip(label):
49
- global SKIP
50
- SKIP += 1
51
- print(" SKIP %s" % label)
52
-
53
-
54
- # ---------------------------- 合成库 ----------------------------
55
-
56
- def ccg(fn="示例节点", *, 生效="载体/位置:本地仓;时间:全时窗(任意时刻成立);方法:test;约束:无",
57
- sub="a/b", exe="python -m md_cg.demo", ver="test", neg="无", extra=""):
58
- """六要素齐全的正文(防 `_loc_missing_field` 的 CCG 缺行噪声干扰其它判据)。"""
59
- return ("# 功能名:%s\n# 生效条件:%s\n# 子功能:%s\n# 执行:%s\n"
60
- "# 验证方式:%s\n# 不适用条件:%s\n%s" % (fn, 生效, sub, exe, ver, neg, extra))
61
-
62
-
63
- def full_cs(*, pos="本地仓", tw=(0.0, 9999999999.0), tool="test", con="无"):
64
- return {"observation_position": pos, "time_window": list(tw),
65
- "observation_tool": tool, "existence_constraint": con}
66
-
67
-
68
- def mkroot(root, nodes):
69
- """合成认知图根:`_index.json` + 节点盘文件(文件文本由 NF.dumps 生成)。"""
70
- os.makedirs(root, exist_ok=True)
71
- idx = {}
72
- for nid, spec in nodes.items():
73
- fm = dict(spec.get("fm") or {})
74
- fm.setdefault("id", nid)
75
- content = spec.get("content") or ""
76
- rel = spec.get("path") or ("knowledge/%s.md" % nid)
77
- p = os.path.join(root, rel)
78
- os.makedirs(os.path.dirname(p), exist_ok=True)
79
- with open(p, "w", encoding="utf-8") as f:
80
- f.write(NF.dumps(fm, content))
81
- meta = dict(spec.get("meta") or {})
82
- meta.setdefault("path", rel)
83
- meta.setdefault("layer", spec.get("layer") or "knowledge")
84
- meta.setdefault("content_hash", NF.content_hash(content))
85
- meta.update({"id": nid, "path": rel})
86
- meta["path"] = rel
87
- idx[nid] = meta
88
- with open(os.path.join(root, CF.INDEX_FILE), "w", encoding="utf-8") as f:
89
- json.dump({"nodes": idx}, f, ensure_ascii=False)
90
- return root
91
-
92
-
93
- def snapshot(root):
94
- out = {}
95
- for dp, _dns, fns in os.walk(root):
96
- for fn in fns:
97
- p = os.path.join(dp, fn)
98
- with open(p, "rb") as f:
99
- out[os.path.relpath(p, root)] = hashlib.md5(f.read()).hexdigest()
100
- return out
101
-
102
-
103
- def kinds_of(hits):
104
- return sorted({h["issue_kind"] for h in hits})
105
-
106
-
107
- def fields_of(hits, kind=None):
108
- return sorted({h["field"] for h in hits if kind is None or h["issue_kind"] == kind})
109
-
110
-
111
- NOW = 1789000000.0 # 固定「当前时间」(2026-09 量级),不靠墙钟
112
- EXPIRED = (1000.0, 2000.0)
113
-
114
-
115
- # ---------------------------- A 组:基础工具 ----------------------------
116
-
117
- def phase_a(tmp):
118
- print("[A] 基础工具(纯函数)")
119
-
120
- c = "第一句。第二句!第三句?第四句;"
121
- sp = LC.sentence_spans(c)
122
- ok(len(sp) == 4 and [i for i, _s, _e, _t in sp] == [0, 1, 2, 3],
123
- "A1 sentence_spans 句索引连续(%s)" % [i for i, _s, _e, _t in sp])
124
- ok(all(c[s:e] == t for _i, s, e, t in sp), "A2 span 与原文切片逐字对应")
125
- ok(LC.sentence_spans("") == [] and LC.sentence_spans(" ") == [],
126
- "A3 空/纯空白正文零句(不编号空句)")
127
- ok([t for _i, _s, _e, t in LC.sentence_spans("甲。\n乙。")] == ["甲。", "乙。"],
128
- "A4 换行是句尾且空句不编号(分隔符不残留在句首)")
129
-
130
- body = ccg(extra="尾句。")
131
- ms = LC.mark_spans(body)
132
- ok(set(ms) == set(NF.CCG_MARKS), "A5 mark_spans 六要素全提(实得 %s)" % sorted(ms))
133
- ok(all(body[s:e].startswith("#") for s, e in ms.values()),
134
- "A6 要素行 span 覆盖整行")
135
- ok(LC.mark_spans(ccg() + "# 功能名:第二个\n").get("功能名")
136
- == LC.mark_spans(ccg()).get("功能名"), "A7 同要素取首次出现(确定性)")
137
-
138
- text = NF.dumps({"id": "n1", "path": "knowledge/n1.md", "tags": []}, ccg())
139
- ks = LC.key_line_spans(text)
140
- ok(ks.get("id", (None, None))[0] == 2, "A8 key_line_spans 行号 1-based(实得 %s)"
141
- % (ks.get("id") or (None,))[0])
142
- ok("功能名" not in ks and "---" not in ks,
143
- "A9 正文 `#` 行与 `---` 分隔线都排除在 frontmatter 之外")
144
- ok(ks.get("id") and text[ks["id"][1][0]:ks["id"][1][1]].startswith('id:'),
145
- "A10 键行 span 切片以键名开头")
146
-
147
- root = mkroot(os.path.join(tmp, "a"), {
148
- "n1": {"content": ccg(), "meta": {"layer": "knowledge", "tags": ["a"]}},
149
- "n2": {"content": "短", "path": "", "meta": {"layer": "knowledge"}},
150
- })
151
- nd = LC.load_node("n1", root)
152
- ok(nd and nd["meta"].get("layer") == "knowledge" and "# 功能名:" in (nd["content"] or ""),
153
- "A11 load_node 取索引 meta + 文件正文")
154
- ok(nd and nd["fm"].get("id") == "n1", "A12 loads 解析出的 fm 是文件真源")
155
- ok(LC.load_node("ghost", root) is None, "A13 索引无此节点 → None(不猜路径)")
156
- ok(LC.load_node("n1", root, index={"nodes": {"n1": {"path": "knowledge/n1.md"}}})
157
- is not None, "A14 index 可显式注入(不读盘 index)")
158
- ok(LC._index(root).get("n1") is not None, "A15 _index 兼容 {nodes:…} 形态")
159
- ok(LC._index(root, index={"n1": {"path": "x"}}) == {"n1": {"path": "x"}},
160
- "A16 裸 dict 索引原样透传")
161
-
162
- ok(LC._snippet("甲" * 200, [0, 200]).endswith("…"), "A17 超长片段截断加省略号")
163
- ok(LC._line_of(text, nd["content"], [0, 5]) == 6,
164
- "A18 _line_of 定位到正文首行(实得 %s)" % LC._line_of(text, nd["content"], [0, 5]))
165
- ok(LC._line_of(None, "x", [0, 1]) is None and LC._line_of(text, "不存在", [0, 1]) is None,
166
- "A19 无 text / 正文不在文件内 → line=None(不编造)")
167
- ok(LC.canonical_kind("dup_content") == "dup"
168
- and LC.canonical_kind("template_flow_digits_only") == "template_flow",
169
- "A20 M1/D1 用词归并到 D1 规范名")
170
- ok(LC.canonical_kind("天外飞仙") == "天外飞仙", "A21 未知类别原样返回(不假装认路)")
171
-
172
-
173
- # ---------------------------- B 组:六类定位器 ----------------------------
174
-
175
- def phase_b(tmp):
176
- print("[B] 六类定位器")
177
-
178
- # B1-B6 missing_field
179
- m = {"id": "b1", "layer": "knowledge", "content_hash": "h1",
180
- "tags": ["a"], "role": "", "importance": 0.5, "evidence_count": 3,
181
- "lifecycle_state": "active", "condition_space": full_cs()}
182
- hits = LC.locate_ex("b1", "missing_field", meta=m, content=ccg(),
183
- text="", fm={}, peers=[], now=NOW)["hits"]
184
- ok(fields_of(hits) == ["role"] and all(h["span"] is None for h in hits),
185
- "B1 字段两处皆空 → missing_field 且 span=None(实得 %s)" % fields_of(hits))
186
-
187
- m2 = dict(m, evidence_count=0, role="knowledge-card")
188
- hits = LC.locate_ex("b1", "missing_field", meta=m2, content=ccg(),
189
- text="", fm={}, peers=[], now=NOW)["hits"]
190
- ok("evidence_count" in fields_of(hits) and "evidence_zero" in {h["rule"] for h in hits},
191
- "B2 evidence_count=0 → 专项命中(rule=evidence_zero)")
192
-
193
- m3 = dict(m, importance=1.7, role="k")
194
- hits = LC.locate_ex("b1", "missing_field", meta=m3, content=ccg(),
195
- text="", fm={}, peers=[], now=NOW)["hits"]
196
- ok("importance" in fields_of(hits) and "field_invalid" in {h["rule"] for h in hits},
197
- "B3 importance 越界 → 字段存在但不可用")
198
-
199
- m4 = dict(m, condition_space={"observation_position": "本地"},
200
- role="k", tags=["a"])
201
- body = ccg()
202
- hits = LC.locate_ex("b1", "missing_field", meta=m4, content=body,
203
- text="", fm={}, peers=[], now=NOW)["hits"]
204
- h = [x for x in hits if x["field"] == "condition_space"]
205
- ok(bool(h) and h[0]["span"] is not None
206
- and body[h[0]["span"][0]:h[0]["span"][1]].startswith("# 生效条件"),
207
- "B4 四槽不全 → 指向正文「# 生效条件」行(%d/4)"
208
- % (len(NF.CONDITION_SLOTS) - 3))
209
-
210
- cut = "# 功能名:只有一行\n正文没有其它要素。\n"
211
- hits = LC.locate_ex("b1", "missing_field", meta=dict(m, role="k", tags=["a"]),
212
- content=cut, text="", fm={}, peers=[], now=NOW)["hits"]
213
- hm = [x for x in hits if x["rule"] == "ccg_incomplete"]
214
- ok(bool(hm) and hm[0]["field"] == "content" and hm[0]["span"] is None
215
- and "生效条件" in hm[0]["evidence"],
216
- "B5 正文缺 CCG 要素行 → field=content 且 span=None(行不存在,不编造区间)")
217
- ok(not [x for x in LC.locate_ex("b1", "missing_field", meta=m3, content=ccg(),
218
- text="", fm={}, peers=[], now=NOW)["hits"]
219
- if x["rule"] == "ccg_incomplete"],
220
- "B6 要素齐全 → 无 ccg_incomplete 噪声")
221
-
222
- # B7-B10 weak_source
223
- hits = LC.locate_ex("b1", "weak_source",
224
- meta={"verification_basis": "", "tags": ["计算机"]},
225
- content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
226
- ok(len(hits) == 1 and hits[0]["rule"] == "basis_absent", "B7 基底为空 → basis_absent")
227
-
228
- hits = LC.locate_ex("b1", "weak_source",
229
- meta={"verification_basis": "self", "tags": ["计算机"]},
230
- content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
231
- ok(len(hits) == 1 and hits[0]["rule"] == "basis_enum", "B8 基底越枚举 → basis_enum")
232
-
233
- hits = LC.locate_ex("b1", "weak_source",
234
- meta={"verification_basis": "textbook", "tags": ["计算机"]},
235
- content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
236
- ok(len(hits) == 1 and hits[0]["rule"] == "basis_licensed"
237
- and "science" in hits[0]["evidence"], "B9 理科×textbook → 赛道不相容(实得 %s)"
238
- % (hits[0]["evidence"] if hits else "无"))
239
-
240
- hits = LC.locate_ex("b1", "weak_source",
241
- meta={"verification_basis": "textbook", "tags": ["语文"]},
242
- content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
243
- ok(hits == [], "B10 文科×textbook 合规 → 零命中(不误报)")
244
-
245
- # B11-B12c observation_aged(原「stale」的时间窗口径:**观测时刻不是失效声明**)
246
- m5 = dict(m, role="k", tags=["a"], condition_space=full_cs(tw=EXPIRED))
247
- body = ccg()
248
- hits = LC.locate_ex("b1", "observation_aged", meta=m5, content=body, text="",
249
- fm={}, peers=[], now=NOW)["hits"]
250
- ok(len(hits) == 1 and hits[0]["field"] == "condition_space"
251
- and hits[0]["severity"] == "info" and "观测时刻" in hits[0]["evidence"],
252
- "B11 时间窗过期 → observation_aged(观测面·info,不冒充失效)")
253
-
254
- m6 = dict(m5, condition_space=full_cs(tw=(0.0, NF.FULL_TIME_WINDOW_MAX)))
255
- ok(LC.locate_ex("b1", "observation_aged", meta=m6, content=body, text="", fm={},
256
- peers=[], now=NOW)["hits"] == [],
257
- "B12 全时窗是合法声明 → 不判")
258
-
259
- # B12b-B12c 时间窗来源链(真库口径:索引快照只带 time_window,**无 condition_space 键**)
260
- m_nocs = {k: v for k, v in m.items() if k != "condition_space"}
261
- hits = LC.locate_ex("b1", "observation_aged", meta=dict(m_nocs, role="k"),
262
- content=body, text="",
263
- fm={"condition_space": full_cs(tw=EXPIRED)}, peers=[],
264
- now=NOW)["hits"]
265
- ok(len(hits) == 1 and hits[0]["rule"] == "observation_window_passed",
266
- "B12b 快照无条件空间 → 回退 fm 真源仍判(实得 %d 条)" % len(hits))
267
-
268
- hits = LC.locate_ex("b1", "observation_aged",
269
- meta=dict(m_nocs, role="k", time_window=list(EXPIRED)),
270
- content=body, text="", fm={}, peers=[], now=NOW)["hits"]
271
- ok(len(hits) == 1 and hits[0]["rule"] == "observation_window_passed",
272
- "B12c 快照只带 time_window → 真库口径下仍判")
273
-
274
- # B12d-B12e **观测面不进评审告警面**(B+C 修正的关键隔离断言)
275
- hits = LC.locate_ex("b1", None, meta=m5, content=body, text="", fm={},
276
- peers=[], now=NOW)["hits"]
277
- ok(all(h["issue_kind"] != "observation_aged" for h in hits),
278
- "B12d 默认全量定位不含 observation_aged(不进评审告警面)")
279
- ok("observation_aged" in LC.ADVISORY_KINDS
280
- and "observation_aged" not in LC.ISSUE_KINDS,
281
- "B12e observation_aged 归观测面(ADVISORY_KINDS),不占问题面 D1 六类")
282
-
283
- # C stale —— **依赖存在性**(B+C 修正:时效判定看载体是否还在,不看观测时刻)
284
- src = tempfile.mkdtemp(prefix="m3src_")
285
- slines = ["def f():", " return 1", "", "def g():", " return 2"]
286
- with open(os.path.join(src, "mod.py"), "w", encoding="utf-8") as f:
287
- f.write("\n".join(slines))
288
- ref_ok = {"path": "mod.py", "name": "f", "kind": "def", "lineno": 1, "end": 2,
289
- "lang": "py", "precise": True,
290
- "hash": CI.region_hash(slines, 1, 2), "root": src}
291
- hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
292
- fm={"code_ref": dict(ref_ok)}, peers=[], now=NOW)["hits"]
293
- ok(hits == [], "C1 依赖源文件在且区间哈希吻合 → 零命中(不误报)")
294
-
295
- hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
296
- fm={"code_ref": dict(ref_ok, hash="000000000000")},
297
- peers=[], now=NOW)["hits"]
298
- ok(len(hits) == 1 and hits[0]["rule"] == "ref_stale"
299
- and hits[0]["field"] == "code_ref",
300
- "C2 源文件在但区间哈希不符(已漂移)→ stale")
301
-
302
- hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
303
- fm={"code_ref": dict(ref_ok, path="gone.py")},
304
- peers=[], now=NOW)["hits"]
305
- ok(len(hits) == 1 and hits[0]["rule"] == "ref_dangling",
306
- "C3 依赖源文件不存在(悬空)→ stale(载体消失)")
307
-
308
- hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
309
- fm={"doc_ref": {"path": "x.md", "lineno": 1, "end": 2,
310
- "hash": "000000000000"}},
311
- peers=[], now=NOW)["hits"]
312
- ok(hits == [], "C4 ref 无 root(判不了)→ 零命中(观测手段不足不冒充失效)")
313
-
314
- hits = LC.locate_ex("c1", None, meta=m5, content=body, text="",
315
- fm={"code_ref": dict(ref_ok, path="gone.py")},
316
- peers=[], now=NOW)["hits"]
317
- ok(any(h["issue_kind"] == "stale" for h in hits),
318
- "C5 默认全量定位会跑 stale(依赖存在性属问题面)")
319
- shutil.rmtree(src, ignore_errors=True)
320
-
321
- # B13-B14 dup
322
- same = ccg(fn="重复节点")
323
- peers = LC.build_peers([{"node_id": "b1", "content": same},
324
- {"node_id": "b2", "content": same}])
325
- hits = LC.locate_ex("b1", "dup",
326
- meta=dict(m, role="k", content_hash=NF.content_hash(same)),
327
- content=same, text="", fm={}, peers=peers.get("b1"),
328
- now=NOW)["hits"]
329
- ok(len(hits) == 1 and hits[0]["field"] == "content_hash"
330
- and hits[0]["span"] == [0, len(same)] and hits[0]["peer"] == "b2",
331
- "B13 同内容指纹 → dup 指整篇正文(peer=%s)"
332
- % (hits[0]["peer"] if hits else "无"))
333
-
334
- hits = LC.locate_ex("b1", "dup", meta=dict(m, role="k", tags=["a"]),
335
- content=ccg(fn="甲"), text="", fm={},
336
- peers=LC.build_peers([{"node_id": "b1", "content": ccg(fn="甲")},
337
- {"node_id": "b2", "content": ccg(fn="乙")}]
338
- ).get("b1"), now=NOW)["hits"]
339
- ok(hits == [], "B14 内容不同 → 不判 dup(不误报)")
340
-
341
- # B15 template_flow:逐句骨架相同、仅数值不同
342
- t1 = ccg(fn="批次 1 收官", extra="本批处理 100 条记录,耗用 12 秒。第二句写 200 条。")
343
- t2 = ccg(fn="批次 2 收官", extra="本批处理 300 条记录,耗用 45 秒。第二句写 400 条。")
344
- peers2 = LC.build_peers([{"node_id": "b1", "content": t1},
345
- {"node_id": "b2", "content": t2}])
346
- hits = LC.locate_ex("b1", "template_flow", meta=dict(m, role="k", tags=["a"]),
347
- content=t1, text="", fm={}, peers=peers2.get("b1"), now=NOW)["hits"]
348
- ok(len(hits) >= 2 and all(h["field"] == "content" and h["span"] is not None
349
- for h in hits),
350
- "B15 同模板流水 → 逐句给出 span(%d 句命中)" % len(hits))
351
- ok(all(t1[h["span"][0]:h["span"][1]].strip()[:LC.SNIPPET_MAX] == h["snippet"]
352
- for h in hits),
353
- "B16 片段=span 切片去空白截断(与实现同口径,可肉眼复核)")
354
- ok(hits and hits[0]["sentence"] is not None and hits[0]["peer"] == "b2",
355
- "B17 携带句索引与对照节点(人工可跳行)")
356
-
357
- hits = LC.locate_ex("b1", "template_flow", meta=dict(m, role="k", tags=["a"]),
358
- content=t1, text="", fm={},
359
- peers=LC.build_peers([{"node_id": "b1", "content": t1}]).get("b1"),
360
- now=NOW)["hits"]
361
- ok(hits == [], "B18 无对照节点 → 不判流水")
362
-
363
- # B19-B21 contradiction(确定性)
364
- body = ccg()
365
- hits = LC.locate_ex("b1", "contradiction",
366
- meta={"content_hash": "declared!", "role": "k", "tags": ["a"]},
367
- content=body, text="", fm={}, peers=[], now=NOW)["hits"]
368
- ok(len(hits) == 1 and hits[0]["rule"] == "hash_declared_vs_actual"
369
- and NF.content_hash(body) in hits[0]["evidence"],
370
- "B19 索引声明指纹 ≠ 正文实算 → contradiction")
371
-
372
- hits = LC.locate_ex("b1", "contradiction",
373
- meta={"content_hash": NF.content_hash(body), "role": "k"},
374
- content=body, text="", fm={"id": "别的id"}, peers=[],
375
- now=NOW)["hits"]
376
- ok(len(hits) == 1 and hits[0]["rule"] == "id_declared_vs_index"
377
- and hits[0]["field"] == "id", "B20 文件 id ≠ 索引键 → contradiction")
378
-
379
- hits = LC.locate_ex("b1", "contradiction",
380
- meta={"content_hash": NF.content_hash(body), "role": "k"},
381
- content=body, text="", fm={"id": "b1"}, peers=[], now=NOW)["hits"]
382
- ok(hits == [], "B21 声明与事实一致 → 零命中")
383
-
384
- # B22-B25 字段层门限(判据同源:派生自 M1 规则库,不另立一份)
385
- scope = LC.field_layer_scope()
386
- ok(scope.get("role") == ["knowledge"]
387
- and scope.get("evidence_count") == ["knowledge"],
388
- "B22 层门限派生自 rules/*.json 的 matcher.layer(实得 %s)" % scope)
389
- ok("verification_basis" not in scope,
390
- "B23 未限层的字段不入表 = 全层适用(M1 R-BASIS-MISSING 无 layer 门限)")
391
-
392
- hits = LC.locate_ex("b1", "missing_field",
393
- meta=dict(m, layer="contextual", role="", evidence_count=0,
394
- tags=[]),
395
- content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
396
- fld = fields_of(hits)
397
- ok("role" not in fld and "evidence_count" not in fld,
398
- "B24 非 knowledge 层 → 限层字段越层不报(与 M1 `_scope` 同口径,实得 %s)" % fld)
399
- ok("tags" in fld,
400
- "B25 未限层字段不受门限影响(tags 仍全层检查,实得 %s)" % fld)
401
-
402
- # B26-B28 指纹不一致的**成因**(同一条命中,两种成因,处置完全不同)
403
- cause, note = LC.hash_mismatch_cause({"content_hash": "x"}, root=None, path=None)
404
- ok(cause == "unknown" and "成因未判定" in note,
405
- "B26 无盘上证据 → cause=unknown(不假装知道成因)")
406
-
407
- root = mkroot(os.path.join(tmp, "b_lag"), {
408
- "n1": {"content": body, "meta": {"content_hash": "declared!"}},
409
- })
410
- node_f = os.path.join(root, "knowledge", "n1.md")
411
- idx_f = os.path.join(root, CF.INDEX_FILE)
412
- mt = os.path.getmtime(node_f)
413
- os.utime(idx_f, (mt - 60.0, mt - 60.0)) # 索引快照比节点文件旧 60s
414
- hc = [h for h in LC.locate_ex("n1", "contradiction", root=root, now=NOW)["hits"]
415
- if h["rule"] == "hash_declared_vs_actual"]
416
- ok(hc and hc[0]["cause"] == "index_lag" and "快照滞后" in hc[0]["evidence"],
417
- "B27 文件比索引快照新 → cause=index_lag(正常写路径现象,实得 %s)"
418
- % (hc[0]["cause"] if hc else "无命中"))
419
-
420
- os.utime(idx_f, (mt + 60.0, mt + 60.0)) # 索引快照不旧于节点文件
421
- hc = [h for h in LC.locate_ex("n1", "contradiction", root=root, now=NOW)["hits"]
422
- if h["rule"] == "hash_declared_vs_actual"]
423
- ok(hc and hc[0]["cause"] == "true_mismatch" and "真源相抵触" in hc[0]["evidence"],
424
- "B28 索引快照不旧于文件 → cause=true_mismatch(需查,实得 %s)"
425
- % (hc[0]["cause"] if hc else "无命中"))
426
-
427
-
428
- # ---------------------------- C 组:提示过滤与 blindspot ----------------------------
429
-
430
- def _mixed():
431
- """同时命中多类的节点:role 空 + 四槽不全(missing_field)+ 指纹不符(contradiction)。
432
-
433
- `condition_space` 刻意只声明 1 槽——让 D 组同时存在「可指区间」(`# 生效条件` 行)
434
- 与「无区间」(frontmatter 声明类)两种命中,契约两侧都被覆盖。
435
- 基底取 `textbook` + `语文`(文科档)——让 `weak_source` 真的干净,
436
- C 组才能验证「该类无问题就返回空、不借机报别的类」。
437
- """
438
- body = ccg()
439
- return dict({"content_hash": "declared!", "role": "", "tags": ["语文"],
440
- "layer": "knowledge", "verification_basis": "textbook",
441
- "importance": 0.5, "evidence_count": 3, "lifecycle_state": "active",
442
- "condition_space": {"observation_position": "本地仓"}}), body
443
-
444
-
445
- def phase_c(tmp):
446
- print("[C] 提示过滤与 blindspot")
447
- m, body = _mixed()
448
- kw = dict(meta=m, content=body, text="", fm={}, peers=[], now=NOW)
449
-
450
- all_hits = LC.locate_ex("c1", None, **kw)["hits"]
451
- ok({"missing_field", "contradiction"} <= set(kinds_of(all_hits)),
452
- "C1 hint=None → 全量定位(实得 %s)" % kinds_of(all_hits))
453
-
454
- only = LC.locate_ex("c1", "weak_source", **kw)["hits"]
455
- ok(all(h["issue_kind"] == "weak_source" for h in only),
456
- "C2 hint=类别 → 只跑该类(实得 %s)" % kinds_of(only))
457
- ok(only == [], "C3 该类无问题 → 空(不借机报别的类)")
458
-
459
- by_field = LC.locate_ex("c1", "role", **kw)["hits"]
460
- ok(by_field and all(h["field"] == "role" for h in by_field),
461
- "C4 hint=字段名 → 按字段过滤全量结果(实得 %s)" % fields_of(by_field))
462
-
463
- both = LC.locate_ex("c1", {"issue_kind": "missing_field", "field": "role"}, **kw)["hits"]
464
- ok(both and len(both) == len(by_field), "C5 dict 形态 hint 同时收类别与字段")
465
-
466
- lst = LC.locate_ex("c1", ["contradiction", "missing_field"], **kw)["hits"]
467
- ok(len(lst) == len(all_hits)
468
- and set(kinds_of(lst)) == {"contradiction", "missing_field"}
469
- and LC.locate_ex("c1", ["contradiction", {"field": "role"}], **kw)["hits"] == [],
470
- "C6 list 形态 hint 收集多个类别;类别与字段是收窄关系(交集空即空,不退回全量)"
471
- "(实得 %d 条 %s)" % (len(lst), kinds_of(lst)))
472
-
473
- ex = LC.locate_ex("c1", "contradiction_semantic", **kw)
474
- ok(ex["hits"] == [] and ex["blindspot"] and "语义" in ex["blindspot"][0],
475
- "C7 语义级矛盾 → blindspot 且零 hits(不猜、不编造区间)")
476
- ok(ex["blindspot"] == LC.locate_ex("c1", "contradiction_semantic", **kw)["blindspot"],
477
- "C8 blindspot 文本确定(可断言)")
478
-
479
- ex2 = LC.locate_ex("c1", ["contradiction_semantic", "missing_field"], **kw)
480
- ok(ex2["blindspot"] and ex2["hits"]
481
- and all(h["issue_kind"] == "missing_field" for h in ex2["hits"]),
482
- "C9 blindspot 与可定位类别同批共存(互不吞没)")
483
-
484
- unk = LC.locate_ex("c1", "天外飞仙", **kw)
485
- ok(unk["hits"] == [] and unk["blindspot"] == [],
486
- "C10 未知提示 → 当字段过滤后为空,不炸也不假装认路")
487
-
488
- alias = LC.locate_ex("c1", "dup_content", **kw)["hits"]
489
- ok(all(h["issue_kind"] == "dup" for h in alias) and bool(alias) is False,
490
- "C11 M1 用词 dup_content 归并为 dup(无重复故空)")
491
-
492
- ok(LC.blindspot_reason("contradiction") == "" and
493
- LC.blindspot_reason("contradiction_semantic") != "",
494
- "C12 blindspot_reason 单一归口(可定位类别返回空串)")
495
-
496
-
497
- # ---------------------------- D 组:契约与确定性 ----------------------------
498
-
499
- def phase_d(tmp):
500
- print("[D] 契约与确定性")
501
- m, body = _mixed()
502
- kw = dict(meta=m, content=body, text="", fm={}, peers=[], now=NOW)
503
- hits = LC.locate("d1", None, **kw)
504
-
505
- ok(isinstance(hits, list), "D1 locate() 契约入口返回 list")
506
- ok(all({"field", "span", "issue_kind", "evidence"} <= set(h) for h in hits),
507
- "D2 四键齐备(实得 %s)" % (sorted(hits[0]) if hits else "无命中"))
508
- ok(all(h["issue_kind"] in LC.ISSUE_KINDS for h in hits),
509
- "D3 issue_kind 全落 D1 枚举(实得 %s)" % kinds_of(hits))
510
- ok(all(str(h["evidence"]).strip() for h in hits),
511
- "D4 evidence 非空(白箱判据:为何算问题)")
512
-
513
- sp = [h for h in hits if h["span"] is not None]
514
- ok(all(isinstance(h["span"], list) and len(h["span"]) == 2
515
- and h["span"][0] < h["span"][1] and h["span"][1] <= len(body) for h in sp),
516
- "D5 span 是正文内的半开区间 [start,end)")
517
- ok(all(isinstance(h["span"][0], int) and isinstance(h["span"][1], int) for h in sp),
518
- "D6 span 端点为整数(可复算切片)")
519
- ok(all(body[h["span"][0]:h["span"][1]].strip() != "" for h in sp),
520
- "D7 span 指向非空片段")
521
- ok(all(h.get("snippet") != "" for h in sp),
522
- "D8 有 span 即有片段(人工核对可肉眼确认)")
523
- ok(all(not h.get("snippet") for h in hits if h["span"] is None),
524
- "D9 声明类命中无区间 → 片段为空(不编造区间)")
525
-
526
- h2 = LC.locate("d1", None, **kw)
527
- ok(json.dumps(hits, ensure_ascii=False) == json.dumps(h2, ensure_ascii=False),
528
- "D10 同一输入两次调用逐字节一致(确定性)")
529
- ok(LC.locate_ex("d1", None, **kw)["load"]["content_len"] == len(body),
530
- "D11 load 回报正文长度(审计留痕)")
531
-
532
- key = [(h["issue_kind"], str(h["field"])) for h in hits]
533
- ok(key == sorted(key), "D12 命中按 (issue_kind, field, span, peer) 稳定排序(实得 %s)" % key)
534
-
535
- try:
536
- LC.locate_ex("d1", None, meta=m, content=None, text="", fm={}, peers=[], now=NOW)
537
- raised = False
538
- except ValueError:
539
- raised = True
540
- ok(raised, "D13 无正文且无 root → fail-closed 报错(不假装能定位)")
541
-
542
-
543
- # ---------------------------- E 组:批量与包 ----------------------------
544
-
545
- def _pkg(entries, bid="bE"):
546
- return {"bundle_id": bid, "group_kind": "batch", "group_key": "g",
547
- "size": len(entries), "entries": entries}
548
-
549
-
550
- def _ent(nid, *, excerpt="", h="", **kw):
551
- e = {"ref": nid, "node_id": nid, "excerpt": excerpt, "content_hash": h,
552
- "layer": "knowledge", "tags": ["a"]}
553
- e.update(kw)
554
- return e
555
-
556
-
557
- def _clean_meta(**kw):
558
- """合规 meta(理科档 × test 基底)——「干净」必须是真干净,否则 clean 断言无意义。
559
-
560
- `layer="knowledge"`:层门限(真源 = M1 规则库 `matcher.layer`)下,
561
- `role`/`evidence_count` 只在本层检查;缺层节点根本不进判据,clean 断言会空转。
562
- """
563
- m = {"role": "knowledge-card", "tags": ["计算机"], "layer": "knowledge",
564
- "verification_basis": "test",
565
- "importance": 0.5, "evidence_count": 2, "lifecycle_state": "active",
566
- "condition_space": full_cs()}
567
- m.update(kw)
568
- return m
569
-
570
-
571
- def phase_e(tmp):
572
- print("[E] 批量与包")
573
- same = ccg(fn="重复的")
574
- diff = ccg(fn="独一无二的甲")
575
- peers = LC.build_peers([{"node_id": "e1", "content": same},
576
- {"node_id": "e2", "content": same},
577
- {"node_id": "e3", "content": diff}])
578
- ok([p["node_id"] for p in peers["e1"]] == ["e2"]
579
- and [p["node_id"] for p in peers["e2"]] == ["e1"],
580
- "E1 build_peers 同内容指纹互为对照(双向)")
581
- ok(peers["e3"] == [], "E2 内容不同 → 不同组(不是「同批即同组」)")
582
-
583
- t1 = ccg(fn="批次 1 收官", extra="处理 100 条。")
584
- t2 = ccg(fn="批次 2 收官", extra="处理 200 条。")
585
- p2 = LC.build_peers([{"node_id": "e1", "content": t1},
586
- {"node_id": "e2", "content": t2}])
587
- ok([p["node_id"] for p in p2["e1"]] == ["e2"],
588
- "E3 同标题模板骨架 → 互为对照(指纹不同也入组)")
589
- ok(all(p["sk"] for p in p2["e1"]), "E4 对照项携带模板骨架(M1 同源口径)")
590
-
591
- ok(LC.build_peers([]) == {}, "E5 空批 → 空映射(不炸)")
592
- ok("e1" in LC.build_peers([{"node_id": "e1", "content": same}]),
593
- "E6 单条批仍回填键(调用方不必判空)")
594
-
595
- items = [{"node_id": "e1", "content": ccg(fn="干净的"), "meta": _clean_meta()},
596
- {"node_id": "e2", "content": ccg(fn="有问题的"),
597
- "meta": _clean_meta(role="")}]
598
- rep = LC.locate_many(items=items, now=NOW)
599
- ok(rep["nodes"] == 2 and rep["clean"] == ["e1"] and rep["missing"] == [],
600
- "E7 locate_many 报「干净」条(没问题≠没看,clean=%s)" % rep["clean"])
601
- ok(rep["by_kind"].get("missing_field") == 1 and rep["by_field"].get("role") == 1,
602
- "E8 by_kind/by_field 汇总正确(%s / %s)" % (rep["by_kind"], rep["by_field"]))
603
- ok(all(h["node_id"] == "e2" for h in rep["hits"]), "E9 命中归属到正确节点")
604
-
605
- root = mkroot(os.path.join(tmp, "e"), {
606
- "n1": {"content": ccg(fn="盘上节点"), "meta": _clean_meta(role="")},
607
- "n2": {"content": ccg(fn="盘上无问题"), "meta": _clean_meta()},
608
- })
609
- rep2 = LC.locate_many(["n1", "n2", "ghost"], root=root, now=NOW)
610
- ok(rep2["nodes"] == 2 and rep2["missing"] == ["ghost"],
611
- "E10 读不到的节点单列 missing(「没看」≠「没问题」,missing=%s)" % rep2["missing"])
612
- ok(rep2["clean"] == ["n2"], "E11 读盘形态同样分流 clean")
613
- ok(rep2["hits"] and rep2["hits"][0]["field"] == "role",
614
- "E12 读盘形态命中与显式 items 同判据")
615
-
616
- exc9 = ccg(fn="只在摘录里的节点")
617
- pkg = _pkg([_ent("n1", excerpt="摘录里没有特征码", h="过期指纹"),
618
- _ent("p9", excerpt=exc9, h=NF.content_hash(exc9),
619
- **_clean_meta(condition_space={"observation_position": "本地仓",
620
- "time_window": [0.0, 9999999999.0],
621
- "observation_tool": "test"}))])
622
- rep3 = LC.locate_package(pkg, root=root, now=NOW)
623
- ok(rep3["bundle_id"] == "bE" and rep3["entries"] == 2 and rep3["nodes"] == 2,
624
- "E13 locate_package 带包标识与条目计数")
625
- ok(any(h["node_id"] == "n1" for h in rep3["hits"]),
626
- "E14 包内可读节点走读盘正文(准确)")
627
- p9 = [h for h in rep3["hits"] if h["node_id"] == "p9"]
628
- ok(p9 and all(h["field"] == "condition_space" for h in p9)
629
- and exc9[p9[0]["span"][0]:p9[0]["span"][1]].startswith("# 生效条件"),
630
- "E15 文件不可读 → 用包内 excerpt 仍给出正文区间(降级但仍可指,实得 %s)"
631
- % kinds_of(p9))
632
- ok(not [h for h in rep3["hits"] if h["node_id"] == "n1"
633
- and h["issue_kind"] == "dup"],
634
- "E16 excerpt 不冒充正文做重复判定(诚实降级)")
635
-
636
- s = LC.summary(rep3["hits"])
637
- ok(s["total"] == len(rep3["hits"]) and s["nodes"] == len(s["node_ids"])
638
- and s["by_kind"], "E17 summary 汇总口径自洽")
639
- ok(LC.summary([])["total"] == 0 and LC.summary([])["node_ids"] == [],
640
- "E18 空命中 summary 不炸")
641
-
642
- md = LC.markdown_table(rep3["hits"])
643
- ok(md.count("\n") >= len(rep3["hits"]) + 1 and "人工判定" in md,
644
- "E19 markdown_table 逐条一行且留人工判定列")
645
- ok(LC.markdown_table([{"node_id": "x", "issue_kind": "dup", "field": "c",
646
- "span": None, "evidence": "含|竖线"}]).count("\\|") == 1,
647
- "E20 表格竖线转义(不破坏表格结构)")
648
-
649
- code = LC.main(["--root", root, "--node", "n1", "--json"])
650
- ok(code == 0, "E21 CLI --json 退出码 0")
651
- code2 = LC.main(["--root", root, "--node", "n1", "--kind", "missing_field",
652
- "--markdown"])
653
- ok(code2 == 0, "E22 CLI --kind + --markdown 退出码 0")
654
- # 本用例断言「缺 root → fail-closed」,而 locate 的 --root 缺省读环境变量 MDCG_ROOT
655
- # → 环境里存在该变量时用例必假失败(非 hermetic)。故用例内显式清除、用完还原。
656
- _saved_root = os.environ.pop("MDCG_ROOT", None)
657
- try:
658
- ok(LC.main(["--node", "n1"]) == 2, "E23 缺 root → 退出码 2(fail-closed)")
659
- finally:
660
- if _saved_root is not None:
661
- os.environ["MDCG_ROOT"] = _saved_root
662
- ok(LC.main(["--root", root]) == 2, "E24 缺 node → 退出码 2(不静默空跑)")
663
-
664
-
665
- # ---------------------------- F 组:零写入 ----------------------------
666
-
667
- def phase_f(tmp):
668
- print("[F] 零写入")
669
- root = mkroot(os.path.join(tmp, "f"), {
670
- "n1": {"content": ccg(fn="批次 1 收官", extra="处理 100 条。"),
671
- "meta": {"role": "", "tags": ["a"], "importance": 0.5,
672
- "evidence_count": 0, "lifecycle_state": "active",
673
- "condition_space": full_cs(tw=EXPIRED)}},
674
- "n2": {"content": ccg(fn="批次 2 收官", extra="处理 200 条。"),
675
- "meta": {"role": "knowledge-card", "tags": ["语文"],
676
- "importance": 0.5, "evidence_count": 1,
677
- "lifecycle_state": "active",
678
- "condition_space": full_cs()}},
679
- })
680
- before = snapshot(root)
681
- LC.locate_many(["n1", "n2"], root=root, now=NOW)
682
- LC.locate_package(_pkg([_ent("n1"), _ent("n2")]), root=root, now=NOW)
683
- LC.locate("n1", None, root=root, now=NOW)
684
- LC.main(["--root", root, "--node", "n1", "--node", "n2"])
685
- ok(snapshot(root) == before,
686
- "F1 全相位跑完认知图指纹逐字节不变(M3 只读)")
687
- ok(not os.path.exists(os.path.join(root, "_mreview")),
688
- "F2 定位不另立状态目录(无残留)")
689
-
690
-
691
- # ---------------------------- main ----------------------------
692
-
693
- def main(argv=None):
694
- with tempfile.TemporaryDirectory(prefix="mrev_m3_") as tmp:
695
- phase_a(tmp)
696
- phase_b(tmp)
697
- phase_c(tmp)
698
- phase_d(tmp)
699
- phase_e(tmp)
700
- phase_f(tmp)
701
- print("\nM3 自测:%d 通过 / %d 失败 / %d 跳过" % (PASS, FAIL, SKIP))
702
- if FAILS:
703
- print("失败项:")
704
- for f in FAILS:
705
- print(" - %s" % f)
706
- return 1 if FAIL else 0
707
-
708
-
709
- if __name__ == "__main__":
710
- sys.exit(main())
1
+ # -*- coding: utf-8 -*-
2
+ """M3 定位自测(D1 字段级定位):六类定位器 + 契约四键 + 批量/包 + 零写入。
3
+
4
+ 真源:docs/记忆评审系统_立项设计与施工交接_20260915.md §5 D1
5
+ `locate(node_id, issue_hint) -> [{field, span, issue_kind, evidence}]`
6
+
7
+ 纪律(同 test_mr_m1/m2):
8
+ · **自备数据源**——合成节点 + tempfile 合成库,零依赖真源库(外部 clone 全绿)。
9
+ · **零写入实锤**——全部相位跑完后认知图指纹逐字节不变(M3 是只读模块)。
10
+ · **确定性**——同一输入两次调用逐字节一致;`now` 显式传入,不靠墙钟。
11
+ · **不猜**——语义级矛盾归 blindspot(`contradiction_semantic`),不编造字符区间。
12
+
13
+ 覆盖:A 基础工具(纯函数) B 问题面定位器(含字段层门限、指纹不一致成因)+ 观测面
14
+ (observation_aged:观测时刻不是失效声明,**不进告警面**)
15
+ C 提示过滤与 blindspot(含 stale 的**依赖存在性**判据:载体消失/漂移)
16
+ D 契约与确定性 E 批量与包 F 零写入
17
+ 运行:python -m md_cg.test_mr_m3
18
+ """
19
+ from __future__ import annotations
20
+
21
+ import hashlib
22
+ import json
23
+ import os
24
+ import shutil
25
+ import sys
26
+ import tempfile
27
+
28
+ from . import codeindex as CI
29
+ from . import conformance as CF
30
+ from . import nodefile as NF
31
+ from . import writelimit as WL
32
+ from .mreview import locate as LC
33
+
34
+ PASS = FAIL = SKIP = 0
35
+ FAILS = []
36
+
37
+
38
+ def ok(cond, label):
39
+ global PASS, FAIL
40
+ if cond:
41
+ PASS += 1
42
+ else:
43
+ FAIL += 1
44
+ FAILS.append(label)
45
+ print(" FAIL %s" % label)
46
+
47
+
48
+ def skip(label):
49
+ global SKIP
50
+ SKIP += 1
51
+ print(" SKIP %s" % label)
52
+
53
+
54
+ # ---------------------------- 合成库 ----------------------------
55
+
56
+ def ccg(fn="示例节点", *, 生效="载体/位置:本地仓;时间:全时窗(任意时刻成立);方法:test;约束:无",
57
+ sub="a/b", exe="python -m md_cg.demo", ver="test", neg="无", extra=""):
58
+ """六要素齐全的正文(防 `_loc_missing_field` 的 CCG 缺行噪声干扰其它判据)。"""
59
+ return ("# 功能名:%s\n# 生效条件:%s\n# 子功能:%s\n# 执行:%s\n"
60
+ "# 验证方式:%s\n# 不适用条件:%s\n%s" % (fn, 生效, sub, exe, ver, neg, extra))
61
+
62
+
63
+ def full_cs(*, pos="本地仓", tw=(0.0, 9999999999.0), tool="test", con="无"):
64
+ return {"observation_position": pos, "time_window": list(tw),
65
+ "observation_tool": tool, "existence_constraint": con}
66
+
67
+
68
+ def mkroot(root, nodes):
69
+ """合成认知图根:`_index.json` + 节点盘文件(文件文本由 NF.dumps 生成)。"""
70
+ os.makedirs(root, exist_ok=True)
71
+ idx = {}
72
+ for nid, spec in nodes.items():
73
+ fm = dict(spec.get("fm") or {})
74
+ fm.setdefault("id", nid)
75
+ content = spec.get("content") or ""
76
+ rel = spec.get("path") or ("knowledge/%s.md" % nid)
77
+ p = os.path.join(root, rel)
78
+ os.makedirs(os.path.dirname(p), exist_ok=True)
79
+ with open(p, "w", encoding="utf-8") as f:
80
+ f.write(NF.dumps(fm, content))
81
+ meta = dict(spec.get("meta") or {})
82
+ meta.setdefault("path", rel)
83
+ meta.setdefault("layer", spec.get("layer") or "knowledge")
84
+ meta.setdefault("content_hash", NF.content_hash(content))
85
+ meta.update({"id": nid, "path": rel})
86
+ meta["path"] = rel
87
+ idx[nid] = meta
88
+ with open(os.path.join(root, CF.INDEX_FILE), "w", encoding="utf-8") as f:
89
+ json.dump({"nodes": idx}, f, ensure_ascii=False)
90
+ return root
91
+
92
+
93
+ def snapshot(root):
94
+ out = {}
95
+ for dp, _dns, fns in os.walk(root):
96
+ for fn in fns:
97
+ p = os.path.join(dp, fn)
98
+ with open(p, "rb") as f:
99
+ out[os.path.relpath(p, root)] = hashlib.md5(f.read()).hexdigest()
100
+ return out
101
+
102
+
103
+ def kinds_of(hits):
104
+ return sorted({h["issue_kind"] for h in hits})
105
+
106
+
107
+ def fields_of(hits, kind=None):
108
+ return sorted({h["field"] for h in hits if kind is None or h["issue_kind"] == kind})
109
+
110
+
111
+ NOW = 1789000000.0 # 固定「当前时间」(2026-09 量级),不靠墙钟
112
+ EXPIRED = (1000.0, 2000.0)
113
+
114
+
115
+ # ---------------------------- A 组:基础工具 ----------------------------
116
+
117
+ def phase_a(tmp):
118
+ print("[A] 基础工具(纯函数)")
119
+
120
+ c = "第一句。第二句!第三句?第四句;"
121
+ sp = LC.sentence_spans(c)
122
+ ok(len(sp) == 4 and [i for i, _s, _e, _t in sp] == [0, 1, 2, 3],
123
+ "A1 sentence_spans 句索引连续(%s)" % [i for i, _s, _e, _t in sp])
124
+ ok(all(c[s:e] == t for _i, s, e, t in sp), "A2 span 与原文切片逐字对应")
125
+ ok(LC.sentence_spans("") == [] and LC.sentence_spans(" ") == [],
126
+ "A3 空/纯空白正文零句(不编号空句)")
127
+ ok([t for _i, _s, _e, t in LC.sentence_spans("甲。\n乙。")] == ["甲。", "乙。"],
128
+ "A4 换行是句尾且空句不编号(分隔符不残留在句首)")
129
+
130
+ body = ccg(extra="尾句。")
131
+ ms = LC.mark_spans(body)
132
+ ok(set(ms) == set(NF.CCG_MARKS), "A5 mark_spans 六要素全提(实得 %s)" % sorted(ms))
133
+ ok(all(body[s:e].startswith("#") for s, e in ms.values()),
134
+ "A6 要素行 span 覆盖整行")
135
+ ok(LC.mark_spans(ccg() + "# 功能名:第二个\n").get("功能名")
136
+ == LC.mark_spans(ccg()).get("功能名"), "A7 同要素取首次出现(确定性)")
137
+
138
+ text = NF.dumps({"id": "n1", "path": "knowledge/n1.md", "tags": []}, ccg())
139
+ ks = LC.key_line_spans(text)
140
+ ok(ks.get("id", (None, None))[0] == 2, "A8 key_line_spans 行号 1-based(实得 %s)"
141
+ % (ks.get("id") or (None,))[0])
142
+ ok("功能名" not in ks and "---" not in ks,
143
+ "A9 正文 `#` 行与 `---` 分隔线都排除在 frontmatter 之外")
144
+ ok(ks.get("id") and text[ks["id"][1][0]:ks["id"][1][1]].startswith('id:'),
145
+ "A10 键行 span 切片以键名开头")
146
+
147
+ root = mkroot(os.path.join(tmp, "a"), {
148
+ "n1": {"content": ccg(), "meta": {"layer": "knowledge", "tags": ["a"]}},
149
+ "n2": {"content": "短", "path": "", "meta": {"layer": "knowledge"}},
150
+ })
151
+ nd = LC.load_node("n1", root)
152
+ ok(nd and nd["meta"].get("layer") == "knowledge" and "# 功能名:" in (nd["content"] or ""),
153
+ "A11 load_node 取索引 meta + 文件正文")
154
+ ok(nd and nd["fm"].get("id") == "n1", "A12 loads 解析出的 fm 是文件真源")
155
+ ok(LC.load_node("ghost", root) is None, "A13 索引无此节点 → None(不猜路径)")
156
+ ok(LC.load_node("n1", root, index={"nodes": {"n1": {"path": "knowledge/n1.md"}}})
157
+ is not None, "A14 index 可显式注入(不读盘 index)")
158
+ ok(LC._index(root).get("n1") is not None, "A15 _index 兼容 {nodes:…} 形态")
159
+ ok(LC._index(root, index={"n1": {"path": "x"}}) == {"n1": {"path": "x"}},
160
+ "A16 裸 dict 索引原样透传")
161
+
162
+ ok(LC._snippet("甲" * 200, [0, 200]).endswith("…"), "A17 超长片段截断加省略号")
163
+ ok(LC._line_of(text, nd["content"], [0, 5]) == 6,
164
+ "A18 _line_of 定位到正文首行(实得 %s)" % LC._line_of(text, nd["content"], [0, 5]))
165
+ ok(LC._line_of(None, "x", [0, 1]) is None and LC._line_of(text, "不存在", [0, 1]) is None,
166
+ "A19 无 text / 正文不在文件内 → line=None(不编造)")
167
+ ok(LC.canonical_kind("dup_content") == "dup"
168
+ and LC.canonical_kind("template_flow_digits_only") == "template_flow",
169
+ "A20 M1/D1 用词归并到 D1 规范名")
170
+ ok(LC.canonical_kind("天外飞仙") == "天外飞仙", "A21 未知类别原样返回(不假装认路)")
171
+
172
+
173
+ # ---------------------------- B 组:六类定位器 ----------------------------
174
+
175
+ def phase_b(tmp):
176
+ print("[B] 六类定位器")
177
+
178
+ # B1-B6 missing_field
179
+ m = {"id": "b1", "layer": "knowledge", "content_hash": "h1",
180
+ "tags": ["a"], "role": "", "importance": 0.5, "evidence_count": 3,
181
+ "lifecycle_state": "active", "condition_space": full_cs()}
182
+ hits = LC.locate_ex("b1", "missing_field", meta=m, content=ccg(),
183
+ text="", fm={}, peers=[], now=NOW)["hits"]
184
+ ok(fields_of(hits) == ["role"] and all(h["span"] is None for h in hits),
185
+ "B1 字段两处皆空 → missing_field 且 span=None(实得 %s)" % fields_of(hits))
186
+
187
+ m2 = dict(m, evidence_count=0, role="knowledge-card")
188
+ hits = LC.locate_ex("b1", "missing_field", meta=m2, content=ccg(),
189
+ text="", fm={}, peers=[], now=NOW)["hits"]
190
+ ok("evidence_count" in fields_of(hits) and "evidence_zero" in {h["rule"] for h in hits},
191
+ "B2 evidence_count=0 → 专项命中(rule=evidence_zero)")
192
+
193
+ m3 = dict(m, importance=1.7, role="k")
194
+ hits = LC.locate_ex("b1", "missing_field", meta=m3, content=ccg(),
195
+ text="", fm={}, peers=[], now=NOW)["hits"]
196
+ ok("importance" in fields_of(hits) and "field_invalid" in {h["rule"] for h in hits},
197
+ "B3 importance 越界 → 字段存在但不可用")
198
+
199
+ m4 = dict(m, condition_space={"observation_position": "本地"},
200
+ role="k", tags=["a"])
201
+ body = ccg()
202
+ hits = LC.locate_ex("b1", "missing_field", meta=m4, content=body,
203
+ text="", fm={}, peers=[], now=NOW)["hits"]
204
+ h = [x for x in hits if x["field"] == "condition_space"]
205
+ ok(bool(h) and h[0]["span"] is not None
206
+ and body[h[0]["span"][0]:h[0]["span"][1]].startswith("# 生效条件"),
207
+ "B4 四槽不全 → 指向正文「# 生效条件」行(%d/4)"
208
+ % (len(NF.CONDITION_SLOTS) - 3))
209
+
210
+ cut = "# 功能名:只有一行\n正文没有其它要素。\n"
211
+ hits = LC.locate_ex("b1", "missing_field", meta=dict(m, role="k", tags=["a"]),
212
+ content=cut, text="", fm={}, peers=[], now=NOW)["hits"]
213
+ hm = [x for x in hits if x["rule"] == "ccg_incomplete"]
214
+ ok(bool(hm) and hm[0]["field"] == "content" and hm[0]["span"] is None
215
+ and "生效条件" in hm[0]["evidence"],
216
+ "B5 正文缺 CCG 要素行 → field=content 且 span=None(行不存在,不编造区间)")
217
+ ok(not [x for x in LC.locate_ex("b1", "missing_field", meta=m3, content=ccg(),
218
+ text="", fm={}, peers=[], now=NOW)["hits"]
219
+ if x["rule"] == "ccg_incomplete"],
220
+ "B6 要素齐全 → 无 ccg_incomplete 噪声")
221
+
222
+ # B7-B10 weak_source
223
+ hits = LC.locate_ex("b1", "weak_source",
224
+ meta={"verification_basis": "", "tags": ["计算机"]},
225
+ content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
226
+ ok(len(hits) == 1 and hits[0]["rule"] == "basis_absent", "B7 基底为空 → basis_absent")
227
+
228
+ hits = LC.locate_ex("b1", "weak_source",
229
+ meta={"verification_basis": "self", "tags": ["计算机"]},
230
+ content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
231
+ ok(len(hits) == 1 and hits[0]["rule"] == "basis_enum", "B8 基底越枚举 → basis_enum")
232
+
233
+ hits = LC.locate_ex("b1", "weak_source",
234
+ meta={"verification_basis": "textbook", "tags": ["计算机"]},
235
+ content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
236
+ ok(len(hits) == 1 and hits[0]["rule"] == "basis_licensed"
237
+ and "science" in hits[0]["evidence"], "B9 理科×textbook → 赛道不相容(实得 %s)"
238
+ % (hits[0]["evidence"] if hits else "无"))
239
+
240
+ hits = LC.locate_ex("b1", "weak_source",
241
+ meta={"verification_basis": "textbook", "tags": ["语文"]},
242
+ content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
243
+ ok(hits == [], "B10 文科×textbook 合规 → 零命中(不误报)")
244
+
245
+ # B11-B12c observation_aged(原「stale」的时间窗口径:**观测时刻不是失效声明**)
246
+ m5 = dict(m, role="k", tags=["a"], condition_space=full_cs(tw=EXPIRED))
247
+ body = ccg()
248
+ hits = LC.locate_ex("b1", "observation_aged", meta=m5, content=body, text="",
249
+ fm={}, peers=[], now=NOW)["hits"]
250
+ ok(len(hits) == 1 and hits[0]["field"] == "condition_space"
251
+ and hits[0]["severity"] == "info" and "观测时刻" in hits[0]["evidence"],
252
+ "B11 时间窗过期 → observation_aged(观测面·info,不冒充失效)")
253
+
254
+ m6 = dict(m5, condition_space=full_cs(tw=(0.0, NF.FULL_TIME_WINDOW_MAX)))
255
+ ok(LC.locate_ex("b1", "observation_aged", meta=m6, content=body, text="", fm={},
256
+ peers=[], now=NOW)["hits"] == [],
257
+ "B12 全时窗是合法声明 → 不判")
258
+
259
+ # B12b-B12c 时间窗来源链(真库口径:索引快照只带 time_window,**无 condition_space 键**)
260
+ m_nocs = {k: v for k, v in m.items() if k != "condition_space"}
261
+ hits = LC.locate_ex("b1", "observation_aged", meta=dict(m_nocs, role="k"),
262
+ content=body, text="",
263
+ fm={"condition_space": full_cs(tw=EXPIRED)}, peers=[],
264
+ now=NOW)["hits"]
265
+ ok(len(hits) == 1 and hits[0]["rule"] == "observation_window_passed",
266
+ "B12b 快照无条件空间 → 回退 fm 真源仍判(实得 %d 条)" % len(hits))
267
+
268
+ hits = LC.locate_ex("b1", "observation_aged",
269
+ meta=dict(m_nocs, role="k", time_window=list(EXPIRED)),
270
+ content=body, text="", fm={}, peers=[], now=NOW)["hits"]
271
+ ok(len(hits) == 1 and hits[0]["rule"] == "observation_window_passed",
272
+ "B12c 快照只带 time_window → 真库口径下仍判")
273
+
274
+ # B12d-B12e **观测面不进评审告警面**(B+C 修正的关键隔离断言)
275
+ hits = LC.locate_ex("b1", None, meta=m5, content=body, text="", fm={},
276
+ peers=[], now=NOW)["hits"]
277
+ ok(all(h["issue_kind"] != "observation_aged" for h in hits),
278
+ "B12d 默认全量定位不含 observation_aged(不进评审告警面)")
279
+ ok("observation_aged" in LC.ADVISORY_KINDS
280
+ and "observation_aged" not in LC.ISSUE_KINDS,
281
+ "B12e observation_aged 归观测面(ADVISORY_KINDS),不占问题面 D1 六类")
282
+
283
+ # C stale —— **依赖存在性**(B+C 修正:时效判定看载体是否还在,不看观测时刻)
284
+ src = tempfile.mkdtemp(prefix="m3src_")
285
+ slines = ["def f():", " return 1", "", "def g():", " return 2"]
286
+ with open(os.path.join(src, "mod.py"), "w", encoding="utf-8") as f:
287
+ f.write("\n".join(slines))
288
+ ref_ok = {"path": "mod.py", "name": "f", "kind": "def", "lineno": 1, "end": 2,
289
+ "lang": "py", "precise": True,
290
+ "hash": CI.region_hash(slines, 1, 2), "root": src}
291
+ hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
292
+ fm={"code_ref": dict(ref_ok)}, peers=[], now=NOW)["hits"]
293
+ ok(hits == [], "C1 依赖源文件在且区间哈希吻合 → 零命中(不误报)")
294
+
295
+ hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
296
+ fm={"code_ref": dict(ref_ok, hash="000000000000")},
297
+ peers=[], now=NOW)["hits"]
298
+ ok(len(hits) == 1 and hits[0]["rule"] == "ref_stale"
299
+ and hits[0]["field"] == "code_ref",
300
+ "C2 源文件在但区间哈希不符(已漂移)→ stale")
301
+
302
+ hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
303
+ fm={"code_ref": dict(ref_ok, path="gone.py")},
304
+ peers=[], now=NOW)["hits"]
305
+ ok(len(hits) == 1 and hits[0]["rule"] == "ref_dangling",
306
+ "C3 依赖源文件不存在(悬空)→ stale(载体消失)")
307
+
308
+ hits = LC.locate_ex("c1", "stale", meta={}, content=body, text="",
309
+ fm={"doc_ref": {"path": "x.md", "lineno": 1, "end": 2,
310
+ "hash": "000000000000"}},
311
+ peers=[], now=NOW)["hits"]
312
+ ok(hits == [], "C4 ref 无 root(判不了)→ 零命中(观测手段不足不冒充失效)")
313
+
314
+ hits = LC.locate_ex("c1", None, meta=m5, content=body, text="",
315
+ fm={"code_ref": dict(ref_ok, path="gone.py")},
316
+ peers=[], now=NOW)["hits"]
317
+ ok(any(h["issue_kind"] == "stale" for h in hits),
318
+ "C5 默认全量定位会跑 stale(依赖存在性属问题面)")
319
+ shutil.rmtree(src, ignore_errors=True)
320
+
321
+ # B13-B14 dup
322
+ same = ccg(fn="重复节点")
323
+ peers = LC.build_peers([{"node_id": "b1", "content": same},
324
+ {"node_id": "b2", "content": same}])
325
+ hits = LC.locate_ex("b1", "dup",
326
+ meta=dict(m, role="k", content_hash=NF.content_hash(same)),
327
+ content=same, text="", fm={}, peers=peers.get("b1"),
328
+ now=NOW)["hits"]
329
+ ok(len(hits) == 1 and hits[0]["field"] == "content_hash"
330
+ and hits[0]["span"] == [0, len(same)] and hits[0]["peer"] == "b2",
331
+ "B13 同内容指纹 → dup 指整篇正文(peer=%s)"
332
+ % (hits[0]["peer"] if hits else "无"))
333
+
334
+ hits = LC.locate_ex("b1", "dup", meta=dict(m, role="k", tags=["a"]),
335
+ content=ccg(fn="甲"), text="", fm={},
336
+ peers=LC.build_peers([{"node_id": "b1", "content": ccg(fn="甲")},
337
+ {"node_id": "b2", "content": ccg(fn="乙")}]
338
+ ).get("b1"), now=NOW)["hits"]
339
+ ok(hits == [], "B14 内容不同 → 不判 dup(不误报)")
340
+
341
+ # B15 template_flow:逐句骨架相同、仅数值不同
342
+ t1 = ccg(fn="批次 1 收官", extra="本批处理 100 条记录,耗用 12 秒。第二句写 200 条。")
343
+ t2 = ccg(fn="批次 2 收官", extra="本批处理 300 条记录,耗用 45 秒。第二句写 400 条。")
344
+ peers2 = LC.build_peers([{"node_id": "b1", "content": t1},
345
+ {"node_id": "b2", "content": t2}])
346
+ hits = LC.locate_ex("b1", "template_flow", meta=dict(m, role="k", tags=["a"]),
347
+ content=t1, text="", fm={}, peers=peers2.get("b1"), now=NOW)["hits"]
348
+ ok(len(hits) >= 2 and all(h["field"] == "content" and h["span"] is not None
349
+ for h in hits),
350
+ "B15 同模板流水 → 逐句给出 span(%d 句命中)" % len(hits))
351
+ ok(all(t1[h["span"][0]:h["span"][1]].strip()[:LC.SNIPPET_MAX] == h["snippet"]
352
+ for h in hits),
353
+ "B16 片段=span 切片去空白截断(与实现同口径,可肉眼复核)")
354
+ ok(hits and hits[0]["sentence"] is not None and hits[0]["peer"] == "b2",
355
+ "B17 携带句索引与对照节点(人工可跳行)")
356
+
357
+ hits = LC.locate_ex("b1", "template_flow", meta=dict(m, role="k", tags=["a"]),
358
+ content=t1, text="", fm={},
359
+ peers=LC.build_peers([{"node_id": "b1", "content": t1}]).get("b1"),
360
+ now=NOW)["hits"]
361
+ ok(hits == [], "B18 无对照节点 → 不判流水")
362
+
363
+ # B19-B21 contradiction(确定性)
364
+ body = ccg()
365
+ hits = LC.locate_ex("b1", "contradiction",
366
+ meta={"content_hash": "declared!", "role": "k", "tags": ["a"]},
367
+ content=body, text="", fm={}, peers=[], now=NOW)["hits"]
368
+ ok(len(hits) == 1 and hits[0]["rule"] == "hash_declared_vs_actual"
369
+ and NF.content_hash(body) in hits[0]["evidence"],
370
+ "B19 索引声明指纹 ≠ 正文实算 → contradiction")
371
+
372
+ hits = LC.locate_ex("b1", "contradiction",
373
+ meta={"content_hash": NF.content_hash(body), "role": "k"},
374
+ content=body, text="", fm={"id": "别的id"}, peers=[],
375
+ now=NOW)["hits"]
376
+ ok(len(hits) == 1 and hits[0]["rule"] == "id_declared_vs_index"
377
+ and hits[0]["field"] == "id", "B20 文件 id ≠ 索引键 → contradiction")
378
+
379
+ hits = LC.locate_ex("b1", "contradiction",
380
+ meta={"content_hash": NF.content_hash(body), "role": "k"},
381
+ content=body, text="", fm={"id": "b1"}, peers=[], now=NOW)["hits"]
382
+ ok(hits == [], "B21 声明与事实一致 → 零命中")
383
+
384
+ # B22-B25 字段层门限(判据同源:派生自 M1 规则库,不另立一份)
385
+ scope = LC.field_layer_scope()
386
+ ok(scope.get("role") == ["knowledge"]
387
+ and scope.get("evidence_count") == ["knowledge"],
388
+ "B22 层门限派生自 rules/*.json 的 matcher.layer(实得 %s)" % scope)
389
+ ok("verification_basis" not in scope,
390
+ "B23 未限层的字段不入表 = 全层适用(M1 R-BASIS-MISSING 无 layer 门限)")
391
+
392
+ hits = LC.locate_ex("b1", "missing_field",
393
+ meta=dict(m, layer="contextual", role="", evidence_count=0,
394
+ tags=[]),
395
+ content=ccg(), text="", fm={}, peers=[], now=NOW)["hits"]
396
+ fld = fields_of(hits)
397
+ ok("role" not in fld and "evidence_count" not in fld,
398
+ "B24 非 knowledge 层 → 限层字段越层不报(与 M1 `_scope` 同口径,实得 %s)" % fld)
399
+ ok("tags" in fld,
400
+ "B25 未限层字段不受门限影响(tags 仍全层检查,实得 %s)" % fld)
401
+
402
+ # B26-B28 指纹不一致的**成因**(同一条命中,两种成因,处置完全不同)
403
+ cause, note = LC.hash_mismatch_cause({"content_hash": "x"}, root=None, path=None)
404
+ ok(cause == "unknown" and "成因未判定" in note,
405
+ "B26 无盘上证据 → cause=unknown(不假装知道成因)")
406
+
407
+ root = mkroot(os.path.join(tmp, "b_lag"), {
408
+ "n1": {"content": body, "meta": {"content_hash": "declared!"}},
409
+ })
410
+ node_f = os.path.join(root, "knowledge", "n1.md")
411
+ idx_f = os.path.join(root, CF.INDEX_FILE)
412
+ mt = os.path.getmtime(node_f)
413
+ os.utime(idx_f, (mt - 60.0, mt - 60.0)) # 索引快照比节点文件旧 60s
414
+ hc = [h for h in LC.locate_ex("n1", "contradiction", root=root, now=NOW)["hits"]
415
+ if h["rule"] == "hash_declared_vs_actual"]
416
+ ok(hc and hc[0]["cause"] == "index_lag" and "快照滞后" in hc[0]["evidence"],
417
+ "B27 文件比索引快照新 → cause=index_lag(正常写路径现象,实得 %s)"
418
+ % (hc[0]["cause"] if hc else "无命中"))
419
+
420
+ os.utime(idx_f, (mt + 60.0, mt + 60.0)) # 索引快照不旧于节点文件
421
+ hc = [h for h in LC.locate_ex("n1", "contradiction", root=root, now=NOW)["hits"]
422
+ if h["rule"] == "hash_declared_vs_actual"]
423
+ ok(hc and hc[0]["cause"] == "true_mismatch" and "真源相抵触" in hc[0]["evidence"],
424
+ "B28 索引快照不旧于文件 → cause=true_mismatch(需查,实得 %s)"
425
+ % (hc[0]["cause"] if hc else "无命中"))
426
+
427
+
428
+ # ---------------------------- C 组:提示过滤与 blindspot ----------------------------
429
+
430
+ def _mixed():
431
+ """同时命中多类的节点:role 空 + 四槽不全(missing_field)+ 指纹不符(contradiction)。
432
+
433
+ `condition_space` 刻意只声明 1 槽——让 D 组同时存在「可指区间」(`# 生效条件` 行)
434
+ 与「无区间」(frontmatter 声明类)两种命中,契约两侧都被覆盖。
435
+ 基底取 `textbook` + `语文`(文科档)——让 `weak_source` 真的干净,
436
+ C 组才能验证「该类无问题就返回空、不借机报别的类」。
437
+ """
438
+ body = ccg()
439
+ return dict({"content_hash": "declared!", "role": "", "tags": ["语文"],
440
+ "layer": "knowledge", "verification_basis": "textbook",
441
+ "importance": 0.5, "evidence_count": 3, "lifecycle_state": "active",
442
+ "condition_space": {"observation_position": "本地仓"}}), body
443
+
444
+
445
+ def phase_c(tmp):
446
+ print("[C] 提示过滤与 blindspot")
447
+ m, body = _mixed()
448
+ kw = dict(meta=m, content=body, text="", fm={}, peers=[], now=NOW)
449
+
450
+ all_hits = LC.locate_ex("c1", None, **kw)["hits"]
451
+ ok({"missing_field", "contradiction"} <= set(kinds_of(all_hits)),
452
+ "C1 hint=None → 全量定位(实得 %s)" % kinds_of(all_hits))
453
+
454
+ only = LC.locate_ex("c1", "weak_source", **kw)["hits"]
455
+ ok(all(h["issue_kind"] == "weak_source" for h in only),
456
+ "C2 hint=类别 → 只跑该类(实得 %s)" % kinds_of(only))
457
+ ok(only == [], "C3 该类无问题 → 空(不借机报别的类)")
458
+
459
+ by_field = LC.locate_ex("c1", "role", **kw)["hits"]
460
+ ok(by_field and all(h["field"] == "role" for h in by_field),
461
+ "C4 hint=字段名 → 按字段过滤全量结果(实得 %s)" % fields_of(by_field))
462
+
463
+ both = LC.locate_ex("c1", {"issue_kind": "missing_field", "field": "role"}, **kw)["hits"]
464
+ ok(both and len(both) == len(by_field), "C5 dict 形态 hint 同时收类别与字段")
465
+
466
+ lst = LC.locate_ex("c1", ["contradiction", "missing_field"], **kw)["hits"]
467
+ ok(len(lst) == len(all_hits)
468
+ and set(kinds_of(lst)) == {"contradiction", "missing_field"}
469
+ and LC.locate_ex("c1", ["contradiction", {"field": "role"}], **kw)["hits"] == [],
470
+ "C6 list 形态 hint 收集多个类别;类别与字段是收窄关系(交集空即空,不退回全量)"
471
+ "(实得 %d 条 %s)" % (len(lst), kinds_of(lst)))
472
+
473
+ ex = LC.locate_ex("c1", "contradiction_semantic", **kw)
474
+ ok(ex["hits"] == [] and ex["blindspot"] and "语义" in ex["blindspot"][0],
475
+ "C7 语义级矛盾 → blindspot 且零 hits(不猜、不编造区间)")
476
+ ok(ex["blindspot"] == LC.locate_ex("c1", "contradiction_semantic", **kw)["blindspot"],
477
+ "C8 blindspot 文本确定(可断言)")
478
+
479
+ ex2 = LC.locate_ex("c1", ["contradiction_semantic", "missing_field"], **kw)
480
+ ok(ex2["blindspot"] and ex2["hits"]
481
+ and all(h["issue_kind"] == "missing_field" for h in ex2["hits"]),
482
+ "C9 blindspot 与可定位类别同批共存(互不吞没)")
483
+
484
+ unk = LC.locate_ex("c1", "天外飞仙", **kw)
485
+ ok(unk["hits"] == [] and unk["blindspot"] == [],
486
+ "C10 未知提示 → 当字段过滤后为空,不炸也不假装认路")
487
+
488
+ alias = LC.locate_ex("c1", "dup_content", **kw)["hits"]
489
+ ok(all(h["issue_kind"] == "dup" for h in alias) and bool(alias) is False,
490
+ "C11 M1 用词 dup_content 归并为 dup(无重复故空)")
491
+
492
+ ok(LC.blindspot_reason("contradiction") == "" and
493
+ LC.blindspot_reason("contradiction_semantic") != "",
494
+ "C12 blindspot_reason 单一归口(可定位类别返回空串)")
495
+
496
+
497
+ # ---------------------------- D 组:契约与确定性 ----------------------------
498
+
499
+ def phase_d(tmp):
500
+ print("[D] 契约与确定性")
501
+ m, body = _mixed()
502
+ kw = dict(meta=m, content=body, text="", fm={}, peers=[], now=NOW)
503
+ hits = LC.locate("d1", None, **kw)
504
+
505
+ ok(isinstance(hits, list), "D1 locate() 契约入口返回 list")
506
+ ok(all({"field", "span", "issue_kind", "evidence"} <= set(h) for h in hits),
507
+ "D2 四键齐备(实得 %s)" % (sorted(hits[0]) if hits else "无命中"))
508
+ ok(all(h["issue_kind"] in LC.ISSUE_KINDS for h in hits),
509
+ "D3 issue_kind 全落 D1 枚举(实得 %s)" % kinds_of(hits))
510
+ ok(all(str(h["evidence"]).strip() for h in hits),
511
+ "D4 evidence 非空(白箱判据:为何算问题)")
512
+
513
+ sp = [h for h in hits if h["span"] is not None]
514
+ ok(all(isinstance(h["span"], list) and len(h["span"]) == 2
515
+ and h["span"][0] < h["span"][1] and h["span"][1] <= len(body) for h in sp),
516
+ "D5 span 是正文内的半开区间 [start,end)")
517
+ ok(all(isinstance(h["span"][0], int) and isinstance(h["span"][1], int) for h in sp),
518
+ "D6 span 端点为整数(可复算切片)")
519
+ ok(all(body[h["span"][0]:h["span"][1]].strip() != "" for h in sp),
520
+ "D7 span 指向非空片段")
521
+ ok(all(h.get("snippet") != "" for h in sp),
522
+ "D8 有 span 即有片段(人工核对可肉眼确认)")
523
+ ok(all(not h.get("snippet") for h in hits if h["span"] is None),
524
+ "D9 声明类命中无区间 → 片段为空(不编造区间)")
525
+
526
+ h2 = LC.locate("d1", None, **kw)
527
+ ok(json.dumps(hits, ensure_ascii=False) == json.dumps(h2, ensure_ascii=False),
528
+ "D10 同一输入两次调用逐字节一致(确定性)")
529
+ ok(LC.locate_ex("d1", None, **kw)["load"]["content_len"] == len(body),
530
+ "D11 load 回报正文长度(审计留痕)")
531
+
532
+ key = [(h["issue_kind"], str(h["field"])) for h in hits]
533
+ ok(key == sorted(key), "D12 命中按 (issue_kind, field, span, peer) 稳定排序(实得 %s)" % key)
534
+
535
+ try:
536
+ LC.locate_ex("d1", None, meta=m, content=None, text="", fm={}, peers=[], now=NOW)
537
+ raised = False
538
+ except ValueError:
539
+ raised = True
540
+ ok(raised, "D13 无正文且无 root → fail-closed 报错(不假装能定位)")
541
+
542
+
543
+ # ---------------------------- E 组:批量与包 ----------------------------
544
+
545
+ def _pkg(entries, bid="bE"):
546
+ return {"bundle_id": bid, "group_kind": "batch", "group_key": "g",
547
+ "size": len(entries), "entries": entries}
548
+
549
+
550
+ def _ent(nid, *, excerpt="", h="", **kw):
551
+ e = {"ref": nid, "node_id": nid, "excerpt": excerpt, "content_hash": h,
552
+ "layer": "knowledge", "tags": ["a"]}
553
+ e.update(kw)
554
+ return e
555
+
556
+
557
+ def _clean_meta(**kw):
558
+ """合规 meta(理科档 × test 基底)——「干净」必须是真干净,否则 clean 断言无意义。
559
+
560
+ `layer="knowledge"`:层门限(真源 = M1 规则库 `matcher.layer`)下,
561
+ `role`/`evidence_count` 只在本层检查;缺层节点根本不进判据,clean 断言会空转。
562
+ """
563
+ m = {"role": "knowledge-card", "tags": ["计算机"], "layer": "knowledge",
564
+ "verification_basis": "test",
565
+ "importance": 0.5, "evidence_count": 2, "lifecycle_state": "active",
566
+ "condition_space": full_cs()}
567
+ m.update(kw)
568
+ return m
569
+
570
+
571
+ def phase_e(tmp):
572
+ print("[E] 批量与包")
573
+ same = ccg(fn="重复的")
574
+ diff = ccg(fn="独一无二的甲")
575
+ peers = LC.build_peers([{"node_id": "e1", "content": same},
576
+ {"node_id": "e2", "content": same},
577
+ {"node_id": "e3", "content": diff}])
578
+ ok([p["node_id"] for p in peers["e1"]] == ["e2"]
579
+ and [p["node_id"] for p in peers["e2"]] == ["e1"],
580
+ "E1 build_peers 同内容指纹互为对照(双向)")
581
+ ok(peers["e3"] == [], "E2 内容不同 → 不同组(不是「同批即同组」)")
582
+
583
+ t1 = ccg(fn="批次 1 收官", extra="处理 100 条。")
584
+ t2 = ccg(fn="批次 2 收官", extra="处理 200 条。")
585
+ p2 = LC.build_peers([{"node_id": "e1", "content": t1},
586
+ {"node_id": "e2", "content": t2}])
587
+ ok([p["node_id"] for p in p2["e1"]] == ["e2"],
588
+ "E3 同标题模板骨架 → 互为对照(指纹不同也入组)")
589
+ ok(all(p["sk"] for p in p2["e1"]), "E4 对照项携带模板骨架(M1 同源口径)")
590
+
591
+ ok(LC.build_peers([]) == {}, "E5 空批 → 空映射(不炸)")
592
+ ok("e1" in LC.build_peers([{"node_id": "e1", "content": same}]),
593
+ "E6 单条批仍回填键(调用方不必判空)")
594
+
595
+ items = [{"node_id": "e1", "content": ccg(fn="干净的"), "meta": _clean_meta()},
596
+ {"node_id": "e2", "content": ccg(fn="有问题的"),
597
+ "meta": _clean_meta(role="")}]
598
+ rep = LC.locate_many(items=items, now=NOW)
599
+ ok(rep["nodes"] == 2 and rep["clean"] == ["e1"] and rep["missing"] == [],
600
+ "E7 locate_many 报「干净」条(没问题≠没看,clean=%s)" % rep["clean"])
601
+ ok(rep["by_kind"].get("missing_field") == 1 and rep["by_field"].get("role") == 1,
602
+ "E8 by_kind/by_field 汇总正确(%s / %s)" % (rep["by_kind"], rep["by_field"]))
603
+ ok(all(h["node_id"] == "e2" for h in rep["hits"]), "E9 命中归属到正确节点")
604
+
605
+ root = mkroot(os.path.join(tmp, "e"), {
606
+ "n1": {"content": ccg(fn="盘上节点"), "meta": _clean_meta(role="")},
607
+ "n2": {"content": ccg(fn="盘上无问题"), "meta": _clean_meta()},
608
+ })
609
+ rep2 = LC.locate_many(["n1", "n2", "ghost"], root=root, now=NOW)
610
+ ok(rep2["nodes"] == 2 and rep2["missing"] == ["ghost"],
611
+ "E10 读不到的节点单列 missing(「没看」≠「没问题」,missing=%s)" % rep2["missing"])
612
+ ok(rep2["clean"] == ["n2"], "E11 读盘形态同样分流 clean")
613
+ ok(rep2["hits"] and rep2["hits"][0]["field"] == "role",
614
+ "E12 读盘形态命中与显式 items 同判据")
615
+
616
+ exc9 = ccg(fn="只在摘录里的节点")
617
+ pkg = _pkg([_ent("n1", excerpt="摘录里没有特征码", h="过期指纹"),
618
+ _ent("p9", excerpt=exc9, h=NF.content_hash(exc9),
619
+ **_clean_meta(condition_space={"observation_position": "本地仓",
620
+ "time_window": [0.0, 9999999999.0],
621
+ "observation_tool": "test"}))])
622
+ rep3 = LC.locate_package(pkg, root=root, now=NOW)
623
+ ok(rep3["bundle_id"] == "bE" and rep3["entries"] == 2 and rep3["nodes"] == 2,
624
+ "E13 locate_package 带包标识与条目计数")
625
+ ok(any(h["node_id"] == "n1" for h in rep3["hits"]),
626
+ "E14 包内可读节点走读盘正文(准确)")
627
+ p9 = [h for h in rep3["hits"] if h["node_id"] == "p9"]
628
+ ok(p9 and all(h["field"] == "condition_space" for h in p9)
629
+ and exc9[p9[0]["span"][0]:p9[0]["span"][1]].startswith("# 生效条件"),
630
+ "E15 文件不可读 → 用包内 excerpt 仍给出正文区间(降级但仍可指,实得 %s)"
631
+ % kinds_of(p9))
632
+ ok(not [h for h in rep3["hits"] if h["node_id"] == "n1"
633
+ and h["issue_kind"] == "dup"],
634
+ "E16 excerpt 不冒充正文做重复判定(诚实降级)")
635
+
636
+ s = LC.summary(rep3["hits"])
637
+ ok(s["total"] == len(rep3["hits"]) and s["nodes"] == len(s["node_ids"])
638
+ and s["by_kind"], "E17 summary 汇总口径自洽")
639
+ ok(LC.summary([])["total"] == 0 and LC.summary([])["node_ids"] == [],
640
+ "E18 空命中 summary 不炸")
641
+
642
+ md = LC.markdown_table(rep3["hits"])
643
+ ok(md.count("\n") >= len(rep3["hits"]) + 1 and "人工判定" in md,
644
+ "E19 markdown_table 逐条一行且留人工判定列")
645
+ ok(LC.markdown_table([{"node_id": "x", "issue_kind": "dup", "field": "c",
646
+ "span": None, "evidence": "含|竖线"}]).count("\\|") == 1,
647
+ "E20 表格竖线转义(不破坏表格结构)")
648
+
649
+ code = LC.main(["--root", root, "--node", "n1", "--json"])
650
+ ok(code == 0, "E21 CLI --json 退出码 0")
651
+ code2 = LC.main(["--root", root, "--node", "n1", "--kind", "missing_field",
652
+ "--markdown"])
653
+ ok(code2 == 0, "E22 CLI --kind + --markdown 退出码 0")
654
+ # 本用例断言「缺 root → fail-closed」,而 locate 的 --root 缺省读环境变量 MDCG_ROOT
655
+ # → 环境里存在该变量时用例必假失败(非 hermetic)。故用例内显式清除、用完还原。
656
+ _saved_root = os.environ.pop("MDCG_ROOT", None)
657
+ try:
658
+ ok(LC.main(["--node", "n1"]) == 2, "E23 缺 root → 退出码 2(fail-closed)")
659
+ finally:
660
+ if _saved_root is not None:
661
+ os.environ["MDCG_ROOT"] = _saved_root
662
+ ok(LC.main(["--root", root]) == 2, "E24 缺 node → 退出码 2(不静默空跑)")
663
+
664
+
665
+ # ---------------------------- F 组:零写入 ----------------------------
666
+
667
+ def phase_f(tmp):
668
+ print("[F] 零写入")
669
+ root = mkroot(os.path.join(tmp, "f"), {
670
+ "n1": {"content": ccg(fn="批次 1 收官", extra="处理 100 条。"),
671
+ "meta": {"role": "", "tags": ["a"], "importance": 0.5,
672
+ "evidence_count": 0, "lifecycle_state": "active",
673
+ "condition_space": full_cs(tw=EXPIRED)}},
674
+ "n2": {"content": ccg(fn="批次 2 收官", extra="处理 200 条。"),
675
+ "meta": {"role": "knowledge-card", "tags": ["语文"],
676
+ "importance": 0.5, "evidence_count": 1,
677
+ "lifecycle_state": "active",
678
+ "condition_space": full_cs()}},
679
+ })
680
+ before = snapshot(root)
681
+ LC.locate_many(["n1", "n2"], root=root, now=NOW)
682
+ LC.locate_package(_pkg([_ent("n1"), _ent("n2")]), root=root, now=NOW)
683
+ LC.locate("n1", None, root=root, now=NOW)
684
+ LC.main(["--root", root, "--node", "n1", "--node", "n2"])
685
+ ok(snapshot(root) == before,
686
+ "F1 全相位跑完认知图指纹逐字节不变(M3 只读)")
687
+ ok(not os.path.exists(os.path.join(root, "_mreview")),
688
+ "F2 定位不另立状态目录(无残留)")
689
+
690
+
691
+ # ---------------------------- main ----------------------------
692
+
693
+ def main(argv=None):
694
+ with tempfile.TemporaryDirectory(prefix="mrev_m3_") as tmp:
695
+ phase_a(tmp)
696
+ phase_b(tmp)
697
+ phase_c(tmp)
698
+ phase_d(tmp)
699
+ phase_e(tmp)
700
+ phase_f(tmp)
701
+ print("\nM3 自测:%d 通过 / %d 失败 / %d 跳过" % (PASS, FAIL, SKIP))
702
+ if FAILS:
703
+ print("失败项:")
704
+ for f in FAILS:
705
+ print(" - %s" % f)
706
+ return 1 if FAIL else 0
707
+
708
+
709
+ if __name__ == "__main__":
710
+ sys.exit(main())