ltcai 11.2.0 → 11.5.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (253) hide show
  1. package/README.md +50 -53
  2. package/docs/CHANGELOG.md +87 -0
  3. package/docs/COMMUNITY_AND_PLUGINS.md +1 -1
  4. package/docs/DEVELOPMENT.md +1 -1
  5. package/docs/MULTI_AGENT_RUNTIME.md +1 -1
  6. package/docs/ONBOARDING.md +1 -1
  7. package/docs/OPERATIONS.md +6 -2
  8. package/docs/PERMISSION_MODE.md +1 -1
  9. package/docs/TRUST_MODEL.md +1 -1
  10. package/docs/WHY_LATTICE.md +1 -1
  11. package/docs/kg-schema.md +2 -2
  12. package/docs/v11.3.0_PLAN.md +202 -0
  13. package/docs/v11.4.0_RUST_FOUNDATION_PLAN.md +181 -0
  14. package/docs/v11.5.0_RUST_COMPLETE_PLAN.md +145 -0
  15. package/lattice_brain/__init__.py +1 -1
  16. package/lattice_brain/graph/_kg_common/__init__.py +287 -0
  17. package/lattice_brain/graph/_kg_common/extraction.py +516 -0
  18. package/lattice_brain/graph/_kg_common/relations.py +161 -0
  19. package/lattice_brain/graph/_kg_common/text.py +479 -0
  20. package/lattice_brain/graph/discovery_index/__init__.py +35 -0
  21. package/lattice_brain/graph/discovery_index/cleanup.py +182 -0
  22. package/lattice_brain/graph/discovery_index/extract.py +137 -0
  23. package/lattice_brain/graph/discovery_index/scan.py +411 -0
  24. package/lattice_brain/graph/discovery_index/upsert.py +495 -0
  25. package/lattice_brain/graph/projection/__init__.py +42 -0
  26. package/lattice_brain/graph/projection/curation.py +500 -0
  27. package/lattice_brain/graph/{projection.py → projection/v2_schema.py} +15 -477
  28. package/lattice_brain/graph/retrieval/__init__.py +54 -0
  29. package/lattice_brain/graph/retrieval/context.py +197 -0
  30. package/lattice_brain/graph/retrieval/graph_view.py +319 -0
  31. package/lattice_brain/graph/retrieval/hybrid.py +488 -0
  32. package/lattice_brain/graph/retrieval/maintenance.py +121 -0
  33. package/lattice_brain/graph/retrieval/signals.py +95 -0
  34. package/lattice_brain/graph/retrieval_vector/__init__.py +42 -0
  35. package/lattice_brain/graph/retrieval_vector/fingerprint.py +97 -0
  36. package/lattice_brain/graph/retrieval_vector/indexing.py +347 -0
  37. package/lattice_brain/graph/retrieval_vector/search.py +560 -0
  38. package/lattice_brain/graph/retrieval_vector/status.py +374 -0
  39. package/lattice_brain/ingestion/__init__.py +130 -0
  40. package/lattice_brain/ingestion/_contract.py +90 -0
  41. package/lattice_brain/ingestion/constants.py +127 -0
  42. package/lattice_brain/ingestion/folder_scan.py +57 -0
  43. package/lattice_brain/ingestion/folders.py +258 -0
  44. package/lattice_brain/ingestion/hashing.py +26 -0
  45. package/lattice_brain/ingestion/jobs_api.py +107 -0
  46. package/lattice_brain/ingestion/models.py +80 -0
  47. package/lattice_brain/ingestion/pipeline.py +486 -0
  48. package/lattice_brain/ingestion/quality.py +209 -0
  49. package/lattice_brain/ingestion/routing.py +295 -0
  50. package/lattice_brain/multimodal/__init__.py +164 -0
  51. package/lattice_brain/multimodal/audio.py +77 -0
  52. package/lattice_brain/multimodal/common.py +118 -0
  53. package/lattice_brain/multimodal/images.py +498 -0
  54. package/lattice_brain/multimodal/ports.py +169 -0
  55. package/lattice_brain/multimodal/video.py +410 -0
  56. package/lattice_brain/portability/__init__.py +90 -0
  57. package/lattice_brain/portability/_contract.py +42 -0
  58. package/lattice_brain/portability/backups.py +338 -0
  59. package/lattice_brain/portability/bundles.py +136 -0
  60. package/lattice_brain/portability/constants.py +93 -0
  61. package/lattice_brain/portability/fsops.py +138 -0
  62. package/lattice_brain/portability/service.py +41 -0
  63. package/lattice_brain/{portability.py → portability/sharing.py} +44 -677
  64. package/lattice_brain/runtime/__init__.py +1 -1
  65. package/lattice_brain/runtime/multi_agent.py +1 -1
  66. package/latticeai/__init__.py +1 -1
  67. package/latticeai/api/chronicle.py +63 -0
  68. package/latticeai/api/index_jobs.py +145 -0
  69. package/latticeai/core/agent/__init__.py +93 -0
  70. package/latticeai/core/agent/_contract.py +79 -0
  71. package/latticeai/core/agent/context.py +57 -0
  72. package/latticeai/core/agent/deps.py +125 -0
  73. package/latticeai/core/agent/execution.py +622 -0
  74. package/latticeai/core/agent/planning.py +145 -0
  75. package/latticeai/core/agent/recovery.py +157 -0
  76. package/latticeai/core/agent/runtime.py +210 -0
  77. package/latticeai/core/agent/verification.py +231 -0
  78. package/latticeai/core/embedding_providers/__init__.py +151 -0
  79. package/latticeai/core/embedding_providers/base.py +199 -0
  80. package/latticeai/core/embedding_providers/captions.py +162 -0
  81. package/latticeai/core/embedding_providers/profiles.py +126 -0
  82. package/latticeai/core/embedding_providers/text.py +350 -0
  83. package/latticeai/core/embedding_providers/vision.py +352 -0
  84. package/latticeai/core/file_generation/__init__.py +115 -0
  85. package/latticeai/core/file_generation/bundles.py +76 -0
  86. package/latticeai/core/file_generation/extraction.py +154 -0
  87. package/latticeai/core/file_generation/inference.py +235 -0
  88. package/latticeai/core/file_generation/orchestration.py +152 -0
  89. package/latticeai/core/file_generation/prompting.py +117 -0
  90. package/latticeai/core/file_generation/repair.py +114 -0
  91. package/latticeai/core/file_generation/sanitize.py +61 -0
  92. package/latticeai/core/file_generation/validation.py +201 -0
  93. package/latticeai/core/legacy_compatibility.py +1 -1
  94. package/latticeai/core/marketplace.py +1 -1
  95. package/latticeai/core/messages.py +14 -0
  96. package/latticeai/core/workspace_os_constants.py +1 -1
  97. package/latticeai/integrations/telegram_bot/__init__.py +123 -0
  98. package/latticeai/integrations/telegram_bot/__main__.py +17 -0
  99. package/latticeai/integrations/telegram_bot/config.py +86 -0
  100. package/latticeai/integrations/telegram_bot/dispatch.py +311 -0
  101. package/latticeai/integrations/telegram_bot/flows.py +478 -0
  102. package/latticeai/integrations/telegram_bot/helpers.py +322 -0
  103. package/latticeai/integrations/telegram_bot/screens.py +394 -0
  104. package/latticeai/models/router/__init__.py +88 -0
  105. package/latticeai/models/router/_contract.py +66 -0
  106. package/latticeai/models/router/branding.py +56 -0
  107. package/latticeai/models/router/catalog.py +69 -0
  108. package/latticeai/models/router/documents.py +199 -0
  109. package/latticeai/models/router/errors.py +37 -0
  110. package/latticeai/models/router/generation.py +258 -0
  111. package/latticeai/models/router/loading.py +291 -0
  112. package/latticeai/models/router/local_models.py +85 -0
  113. package/latticeai/models/router/registry.py +147 -0
  114. package/latticeai/runtime/build_phases/__init__.py +82 -0
  115. package/latticeai/runtime/build_phases/features.py +421 -0
  116. package/latticeai/runtime/build_phases/foundation.py +555 -0
  117. package/latticeai/runtime/build_phases/web.py +492 -0
  118. package/latticeai/runtime/runtime_context.py +1 -0
  119. package/latticeai/services/architecture_readiness.py +48 -19
  120. package/latticeai/services/brain_intelligence/__init__.py +58 -0
  121. package/latticeai/services/brain_intelligence/_contract.py +71 -0
  122. package/latticeai/services/brain_intelligence/consistency.py +193 -0
  123. package/latticeai/services/brain_intelligence/constants.py +47 -0
  124. package/latticeai/services/brain_intelligence/digest.py +258 -0
  125. package/latticeai/services/brain_intelligence/health.py +331 -0
  126. package/latticeai/services/brain_intelligence/proposals.py +264 -0
  127. package/latticeai/services/brain_intelligence/sampling.py +84 -0
  128. package/latticeai/services/brain_intelligence/service.py +48 -0
  129. package/latticeai/services/chronicle.py +557 -0
  130. package/latticeai/services/memory_service/__init__.py +52 -0
  131. package/latticeai/services/memory_service/_contract.py +100 -0
  132. package/latticeai/services/memory_service/brief.py +431 -0
  133. package/latticeai/services/memory_service/constants.py +57 -0
  134. package/latticeai/services/memory_service/maintenance.py +138 -0
  135. package/latticeai/services/memory_service/manager.py +186 -0
  136. package/latticeai/services/memory_service/proof.py +136 -0
  137. package/latticeai/services/memory_service/recall.py +225 -0
  138. package/latticeai/services/memory_service/service.py +48 -0
  139. package/latticeai/services/memory_service/stores.py +110 -0
  140. package/latticeai/services/model_runtime/__init__.py +322 -0
  141. package/latticeai/services/model_runtime/cloud.py +87 -0
  142. package/latticeai/services/model_runtime/download.py +282 -0
  143. package/latticeai/services/model_runtime/engines.py +341 -0
  144. package/latticeai/services/model_runtime/loading.py +178 -0
  145. package/latticeai/services/model_runtime/service.py +129 -0
  146. package/latticeai/services/model_runtime/state.py +131 -0
  147. package/latticeai/services/model_runtime/status.py +255 -0
  148. package/latticeai/services/product_readiness.py +15 -7
  149. package/latticeai/setup/wizard/__init__.py +126 -0
  150. package/latticeai/setup/wizard/catalog.py +172 -0
  151. package/latticeai/setup/wizard/detect.py +323 -0
  152. package/latticeai/setup/wizard/install.py +348 -0
  153. package/latticeai/setup/wizard/paths.py +168 -0
  154. package/latticeai/setup/wizard/plans.py +74 -0
  155. package/latticeai/setup/wizard/recommend.py +320 -0
  156. package/package.json +6 -2
  157. package/scripts/bump_version.py +14 -0
  158. package/scripts/capture_release_evidence.mjs +33 -21
  159. package/scripts/check_current_release_docs.mjs +1 -1
  160. package/scripts/check_i18n_namespace_coverage.mjs +41 -4
  161. package/scripts/check_max_file_lines.mjs +102 -0
  162. package/scripts/check_release_evidence_bound.mjs +30 -15
  163. package/scripts/check_screenshot_pixel_delta.py +34 -4
  164. package/scripts/check_server_i18n.mjs +2 -0
  165. package/scripts/chunking_parity_corpus.py +449 -0
  166. package/scripts/generate_agent_parity_fixtures.py +752 -0
  167. package/scripts/generate_chunking_parity_fixtures.py +259 -0
  168. package/scripts/generate_rust_parity_fixtures.py +997 -0
  169. package/scripts/lib/mock_server_fingerprint.mjs +94 -0
  170. package/scripts/release_screen_claims.json +42 -2
  171. package/src-tauri/Cargo.lock +404 -3
  172. package/src-tauri/Cargo.toml +13 -1
  173. package/src-tauri/src/backend.rs +460 -0
  174. package/src-tauri/src/folder.rs +33 -0
  175. package/src-tauri/src/main.rs +109 -399
  176. package/src-tauri/src/topology.rs +356 -0
  177. package/src-tauri/tauri.conf.json +1 -1
  178. package/static/app/asset-manifest.json +41 -37
  179. package/static/app/assets/Act-CWnxSCgN.js +1 -0
  180. package/static/app/assets/AdminConsole-BEQYU6kF.js +1 -0
  181. package/static/app/assets/{Brain-tuhI4sOC.js → Brain-DWu1BhFg.js} +2 -2
  182. package/static/app/assets/BrainHome-95Hilr9R.js +2 -0
  183. package/static/app/assets/BrainSignals-QdeqCpAF.js +1 -0
  184. package/static/app/assets/Capture-BHpCxnzb.js +1 -0
  185. package/static/app/assets/Chronicle-B4xYKoed.js +1 -0
  186. package/static/app/assets/CommandPalette-BVXnttSz.js +1 -0
  187. package/static/app/assets/Library-DgYcHome.js +1 -0
  188. package/static/app/assets/{LivingBrain-DBwhto14.js → LivingBrain-CrJLDbf7.js} +1 -1
  189. package/static/app/assets/ProductFlow-DFlScKoJ.js +1 -0
  190. package/static/app/assets/ReviewCard-Cy5f48Pj.js +3 -0
  191. package/static/app/assets/System-NF8IfhTa.js +1 -0
  192. package/static/app/assets/arrow-left-DwkSYrjR.js +1 -0
  193. package/static/app/assets/{bot-Cia42c2h.js → bot-CucuhLhm.js} +1 -1
  194. package/static/app/assets/brain-BBnSryW_.js +1 -0
  195. package/static/app/assets/{button-2j2Ijzgq.js → button-C2GUj2Ai.js} +1 -1
  196. package/static/app/assets/circle-check-CxOVPwYq.js +1 -0
  197. package/static/app/assets/{circle-pause-BEFeWpVW.js → circle-pause-CbkWzBmG.js} +1 -1
  198. package/static/app/assets/{circle-play-ujXMcHxl.js → circle-play-7lEaqHdJ.js} +1 -1
  199. package/static/app/assets/{cpu-k4awryFq.js → cpu-DAlCXlIy.js} +1 -1
  200. package/static/app/assets/{download-DFbLJ_ig.js → download-RNhuuJwh.js} +1 -1
  201. package/static/app/assets/{folder-open-7y_b6xkM.js → folder-open-CLW4odzM.js} +1 -1
  202. package/static/app/assets/{hard-drive-Bidh02Kr.js → hard-drive-NKEiDIAJ.js} +1 -1
  203. package/static/app/assets/{index-DwDl9-8Y.css → index-BLPb5lmE.css} +1 -1
  204. package/static/app/assets/index-DMurvUuR.js +10 -0
  205. package/static/app/assets/input-D2UhPC1X.js +1 -0
  206. package/static/app/assets/link-2-6amKbP_P.js +1 -0
  207. package/static/app/assets/{permissionCopy-Bpb83Hx9.js → permissionCopy-Cu9TZtdR.js} +1 -1
  208. package/static/app/assets/primitives-gPsccucr.js +1 -0
  209. package/static/app/assets/search-Cj_TKk_2.js +1 -0
  210. package/static/app/assets/{share-2-BH1M-WNi.js → share-2-Bau7KkPq.js} +1 -1
  211. package/static/app/assets/{shield-alert-BlKdBXcG.js → shield-alert-BufNYypi.js} +1 -1
  212. package/static/app/assets/{textarea-CCWbUfFB.js → textarea-BQnVWhYs.js} +1 -1
  213. package/static/app/assets/{useFocusTrap-YdHQ7pJ1.js → useFocusTrap-B3_w60si.js} +1 -1
  214. package/static/app/assets/useMutation-BHhCflT6.js +1 -0
  215. package/static/app/assets/{useQuery-CXQiwbVT.js → useQuery-rBWfI-5t.js} +1 -1
  216. package/static/app/assets/utils-V_5-wxr5.js +4 -0
  217. package/static/app/assets/workspace-K1zjYUHj.js +1 -0
  218. package/static/app/index.html +4 -4
  219. package/static/sw.js +1 -1
  220. package/lattice_brain/graph/_kg_common.py +0 -1331
  221. package/lattice_brain/graph/discovery_index.py +0 -1141
  222. package/lattice_brain/graph/retrieval.py +0 -1120
  223. package/lattice_brain/graph/retrieval_vector.py +0 -1293
  224. package/lattice_brain/ingestion.py +0 -1525
  225. package/lattice_brain/multimodal.py +0 -1258
  226. package/latticeai/core/agent.py +0 -1465
  227. package/latticeai/core/embedding_providers.py +0 -1196
  228. package/latticeai/core/file_generation.py +0 -1047
  229. package/latticeai/integrations/telegram_bot.py +0 -1390
  230. package/latticeai/models/router.py +0 -1007
  231. package/latticeai/runtime/build_phases.py +0 -1450
  232. package/latticeai/services/brain_intelligence.py +0 -1083
  233. package/latticeai/services/memory_service.py +0 -1177
  234. package/latticeai/services/model_runtime.py +0 -1281
  235. package/latticeai/setup/wizard.py +0 -1310
  236. package/static/app/assets/Act-AWf0SAKp.js +0 -1
  237. package/static/app/assets/AdminConsole-D0u8Tiyj.js +0 -1
  238. package/static/app/assets/BrainHome-Ts7G_Ila.js +0 -2
  239. package/static/app/assets/BrainSignals-jMYgQ2Ar.js +0 -1
  240. package/static/app/assets/Capture-CqOSzyPr.js +0 -1
  241. package/static/app/assets/CommandPalette-DC0Bzh-I.js +0 -1
  242. package/static/app/assets/Library-CX-bbhmK.js +0 -1
  243. package/static/app/assets/ProductFlow-BHA2cfKI.js +0 -1
  244. package/static/app/assets/ReviewCard-BUhCKRNM.js +0 -3
  245. package/static/app/assets/System-Bu2t5hn1.js +0 -1
  246. package/static/app/assets/arrow-left-Dzwa5zRb.js +0 -1
  247. package/static/app/assets/brain-DJMoqrwx.js +0 -1
  248. package/static/app/assets/index-BpYkzcVm.js +0 -10
  249. package/static/app/assets/input-DSlJJxRs.js +0 -1
  250. package/static/app/assets/primitives-BCx6TvfG.js +0 -1
  251. package/static/app/assets/search-Cgy8cCFJ.js +0 -1
  252. package/static/app/assets/utils-zqPZJxdx.js +0 -4
  253. package/static/app/assets/workspace-DXTihhfU.js +0 -1
@@ -0,0 +1,997 @@
1
+ #!/usr/bin/env python3
2
+ """Build the committed Python↔Rust retrieval parity fixture.
3
+
4
+ ``rust/lattice-retrieval`` is a port, and a port is only worth having if
5
+ something keeps proving it is still one. This script is the Python half of that
6
+ proof: it builds a small, fully deterministic Brain with the **real** write path
7
+ (``KnowledgeGraphStore._upsert_node`` / ``_upsert_chunk`` / ``_upsert_edge``,
8
+ the real hash embedder, the real v2 projection and trigram FTS index), then runs
9
+ the real ``hybrid_search`` / ``search`` / ``vector_search`` over it and writes
10
+ their answers to ``rust/fixtures/golden/``.
11
+
12
+ v11.5.0 widens it past search: the same store also carries the conversation
13
+ corpus, and the same generator drives the KG relationship/traverse reads, the
14
+ service-layer graph and three-channel hybrid, the durable history reads and the
15
+ context assembler (:data:`SUITES`).
16
+
17
+ Two consumers read what it writes:
18
+
19
+ * ``tests/unit/test_rust_parity_contract.py`` re-runs the Python engines against
20
+ the committed database and asserts the goldens still hold — so a change to
21
+ Python semantics fails loudly instead of silently invalidating the contract
22
+ the Rust side is pinned to;
23
+ * ``rust/lattice-retrieval/tests/parity.rs`` runs the Rust port against the same
24
+ database and the same goldens.
25
+
26
+ Determinism is the whole design constraint: every timestamp is written by the
27
+ real code and then **backdated**; the two ``datetime.now()`` calls the ports
28
+ reach are frozen at :data:`FROZEN_NOW` (recorded in the manifest); LLM concept
29
+ extraction is forced off; every environment knob is pinned to its default.
30
+
31
+ Usage::
32
+
33
+ .venv/bin/python scripts/generate_rust_parity_fixtures.py
34
+ """
35
+
36
+ from __future__ import annotations
37
+
38
+ import json
39
+ import logging
40
+ import os
41
+ import shutil
42
+ import sqlite3
43
+ import sys
44
+ from contextlib import contextmanager
45
+ from datetime import datetime
46
+ from pathlib import Path
47
+ from typing import Any, Callable, Dict, Iterator, List, Optional, Tuple
48
+
49
+ REPO_ROOT = Path(__file__).resolve().parents[1]
50
+ if str(REPO_ROOT) not in sys.path:
51
+ sys.path.insert(0, str(REPO_ROOT))
52
+
53
+ FIXTURE_DIR = REPO_ROOT / "rust" / "fixtures"
54
+ GOLDEN_DIR = FIXTURE_DIR / "golden"
55
+ STORE_PATH = FIXTURE_DIR / "parity_store.sqlite"
56
+
57
+ #: The wall clock the ports see: a golden built against a moving clock is none.
58
+ FROZEN_NOW = "2026-08-01T12:00:00"
59
+
60
+ #: Every environment variable the ported path reads, pinned to the configuration
61
+ #: it targets (brute backend, RRF/expansion/rerank off, rewrite on).
62
+ PINNED_ENV: Dict[str, str] = {
63
+ "LATTICEAI_VECTOR_DIM": "384",
64
+ "LATTICEAI_VECTOR_INDEX": "brute",
65
+ "LATTICEAI_VECTOR_MAX_CANDIDATES": "10000",
66
+ "LATTICEAI_KG_READ_V2": "1",
67
+ "LATTICEAI_FUSION_WEIGHTS": "",
68
+ "LATTICEAI_FUSION_STRATEGY": "",
69
+ "LATTICEAI_FUSION_RRF": "",
70
+ "LATTICEAI_GRAPH_EXPANSION": "",
71
+ "LATTICEAI_CROSS_ENCODER_RERANK": "",
72
+ "LATTICEAI_QUERY_REWRITE": "",
73
+ "LATTICEAI_LLM_EXTRACTION": "false",
74
+ }
75
+
76
+ WS_ALPHA = "ws-alpha"
77
+ WS_BETA = "ws-beta"
78
+
79
+ # ── the corpus ───────────────────────────────────────────────────────────────
80
+ # (node_id, type, title, summary, metadata, workspace_id, updated_at)
81
+ #
82
+ # Shaped on purpose: every ``type_boost`` type appears and so do types outside it;
83
+ # titles/summaries are half Korean (the tokenizer, classifier and extractor all
84
+ # branch on script); the ``tie:`` block is five rows sharing one timestamp, one
85
+ # type and nothing to match, pinning the (hits, boost, updated_at) → id ASC
86
+ # tie-break; two workspaces plus NULL-workspace rows cover all three scoping
87
+ # answers (no scoping / empty set / a specific workspace).
88
+ NODES: List[Tuple[str, str, str, str, Dict[str, Any], Optional[str], str]] = [
89
+ ("dec:fusion-alpha", "Decision", "Hybrid retrieval fusion stays alpha weighted",
90
+ "We decided the ranking keeps alpha fusion: lexical rank plus max normalized vector score.",
91
+ {"category": "retrieval", "owner": "jiwon"}, WS_ALPHA, "2026-07-20T09:00:00"),
92
+ ("dec:rust-foundation", "Decision", "Rust 기반 검색 이관 결정",
93
+ "검색 랭킹을 Rust로 이관하기로 결정했습니다. 회의 결과 패리티 증명을 먼저 만듭니다.",
94
+ {"category": "platform"}, WS_ALPHA, "2026-07-18T14:30:00"),
95
+ ("task:parity-harness", "Task", "Build the retrieval parity harness",
96
+ "Generate a golden fixture so the Rust ranking can be compared against the Python ranking.",
97
+ {"status": "open"}, WS_ALPHA, "2026-07-15T11:00:00"),
98
+ ("task:onboarding-checklist", "Task", "온보딩 체크리스트 정리",
99
+ "새로운 사용자가 처음 다섯 걸음을 마칠 수 있도록 온보딩 체크리스트를 정리합니다.",
100
+ {"status": "open"}, WS_BETA, "2026-06-30T08:20:00"),
101
+ ("file:ranking-notes", "File", "ranking-notes.md",
102
+ "Notes on ranking: lexical channel scores one over rank, vector channel is max normalized.",
103
+ {"filename": "ranking-notes.md", "ext": ".md"}, WS_ALPHA, "2026-07-10T16:45:00"),
104
+ ("doc:handbook", "Document", "Lattice AI 사용 안내서",
105
+ "제품 안내서입니다. 검색, 회의 기록, 온보딩 체크리스트를 한 곳에서 설명합니다.",
106
+ {"filename": "handbook.pdf", "ext": ".pdf", "source_node": "doc:handbook"},
107
+ WS_ALPHA, "2026-06-12T10:00:00"),
108
+ ("doc:retrieval-spec", "Document", "Retrieval specification",
109
+ "The retrieval specification describes the lexical channel, the vector channel and their fusion.",
110
+ {"filename": "retrieval-spec.pdf", "ext": ".pdf"}, WS_BETA, "2026-05-28T13:15:00"),
111
+ ("code:hybrid-search", "CodeFile", "hybrid.py",
112
+ "def hybrid_search(query): the ranking pipeline lives here and calls vector_search().",
113
+ {"filename": "hybrid.py", "ext": ".py"}, WS_ALPHA, "2026-07-08T09:30:00"),
114
+ ("code:build-failure", "CodeFile", "build_pipeline.py",
115
+ "빌드 실패 원인을 남겨두는 파일입니다. 컴파일 오류가 나면 여기부터 확인합니다.",
116
+ {"filename": "build_pipeline.py", "ext": ".py"}, WS_BETA, "2026-04-22T18:05:00"),
117
+ ("sheet:metrics", "Spreadsheet", "retrieval-metrics.xlsx",
118
+ "Recall and precision per query class for the ranking experiments.",
119
+ {"filename": "retrieval-metrics.xlsx", "ext": ".xlsx"}, WS_ALPHA, "2026-05-06T09:45:00"),
120
+ ("deck:review", "SlideDeck", "분기 리뷰 발표자료",
121
+ "지난주 분기 리뷰에서 사용한 발표자료입니다. 검색 품질 지표를 담고 있습니다.",
122
+ {"filename": "quarterly-review.pptx", "ext": ".pptx"}, WS_ALPHA, "2026-07-02T15:00:00"),
123
+ ("img:whiteboard", "Image", "whiteboard-retrieval.png",
124
+ "Whiteboard photo from the retrieval design session.",
125
+ {"filename": "whiteboard-retrieval.png", "ext": ".png"}, WS_BETA, "2026-03-19T11:11:00"),
126
+ ("img:ocr-notes", "ImageText", "화이트보드 OCR 텍스트",
127
+ "회의 결정 사항: 랭킹은 alpha 융합을 유지한다.",
128
+ {"filename": "whiteboard-retrieval.png", "ocr": True}, WS_BETA, "2026-03-19T11:12:00"),
129
+ ("audio:standup", "Audio", "standup-2026-07-13.m4a",
130
+ "Standup recording where the parity harness was assigned.",
131
+ {"filename": "standup-2026-07-13.m4a", "ext": ".m4a"}, WS_ALPHA, "2026-07-13T09:05:00"),
132
+ ("page:onboarding", "Page", "온보딩 첫 다섯 걸음",
133
+ "온보딩 페이지입니다. 체크리스트와 안내 문구를 담습니다.", {}, None, "2026-06-01T12:00:00"),
134
+ ("slide:fusion", "Slide", "Fusion slide",
135
+ "One slide explaining alpha fusion of the lexical and vector channels.",
136
+ {}, None, "2026-05-14T12:00:00"),
137
+ ("concept:retrieval", "Concept", "Retrieval",
138
+ "Retrieval is the act of finding the right memory for a question.",
139
+ {}, WS_ALPHA, "2026-04-02T10:00:00"),
140
+ ("concept:ranking", "Concept", "Ranking",
141
+ "Ranking orders candidates so the best answer is first.", {}, WS_BETA, "2026-04-03T10:00:00"),
142
+ ("person:jiwon", "Person", "김지원 님",
143
+ "검색 품질 담당자입니다. 온보딩 개선도 함께 맡고 있습니다.",
144
+ {"role": "owner"}, WS_ALPHA, "2026-04-11T10:00:00"),
145
+ ("person:minseo", "Person", "박민서 님",
146
+ "빌드 파이프라인 담당자입니다.", {"role": "reviewer"}, WS_BETA, "2026-04-12T10:00:00"),
147
+ ("meeting:weekly", "Meeting", "주간 회의 2026-07-14",
148
+ "지난주 주간 회의 기록입니다. 검색 랭킹과 온보딩을 논의했습니다.",
149
+ {}, WS_ALPHA, "2026-07-14T10:00:00"),
150
+ ("meeting:kickoff", "Meeting", "Kickoff meeting",
151
+ "Kickoff meeting for the Rust foundation work.", {}, WS_BETA, "2026-02-09T10:00:00"),
152
+ ("chat:ranking", "Chat", "랭킹 관련 대화",
153
+ "랭킹이 왜 이렇게 나오는지 물어본 대화입니다.",
154
+ {"conversation_id": "conv-1"}, WS_ALPHA, "2026-07-05T21:00:00"),
155
+ ("source:repo", "Source", "github.com/lattice/ai",
156
+ "Source repository for the product.", {"source": "git"}, WS_ALPHA, "2026-02-20T10:00:00"),
157
+ ("repo:lattice", "Repository", "lattice-ai",
158
+ "The monorepo holding the Python worker and the Rust workspace.",
159
+ {}, WS_ALPHA, "2026-02-21T10:00:00"),
160
+ ("org:lattice", "Organization", "Lattice", "The organization.", {}, None, "2026-02-22T10:00:00"),
161
+ ("workflow:nightly", "Workflow", "Nightly reindex",
162
+ "A workflow that reindexes the vector store overnight.", {}, WS_BETA, "2026-03-02T10:00:00"),
163
+ ("agent:librarian", "Agent", "Librarian agent",
164
+ "An agent that files new documents into the graph.", {}, WS_BETA, "2026-03-03T10:00:00"),
165
+ ("error:timeout", "Error", "Search timeout",
166
+ "An error where the ranking took longer than the request budget.",
167
+ {}, WS_ALPHA, "2026-03-04T10:00:00"),
168
+ ("feature:command-palette", "Feature", "Command palette",
169
+ "The palette that opens search from anywhere.", {}, WS_ALPHA, "2026-03-05T10:00:00"),
170
+ ("topic:quality", "Topic", "검색 품질",
171
+ "검색 품질에 대한 주제입니다.", {}, WS_ALPHA, "2026-03-06T10:00:00"),
172
+ # ── the tie block: identical type, identical timestamp, nothing to match ──
173
+ ("tie:a", "Concept", "Tie candidate A", "", {}, WS_ALPHA, "2026-02-01T00:00:00"),
174
+ ("tie:b", "Concept", "Tie candidate B", "", {}, WS_ALPHA, "2026-02-01T00:00:00"),
175
+ ("tie:c", "Concept", "Tie candidate C", "", {}, WS_BETA, "2026-02-01T00:00:00"),
176
+ ("tie:d", "Concept", "Tie candidate D", "", {}, None, "2026-02-01T00:00:00"),
177
+ ("tie:e", "Concept", "Tie candidate E", "", {}, None, "2026-02-01T00:00:00"),
178
+ # ── boosted twins on the same timestamp: type_boost breaks nothing here,
179
+ # id ASC does, and that is the assertion.
180
+ ("twin:doc-a", "Document", "Twin document A", "동일한 시각의 문서 A", {}, WS_ALPHA,
181
+ "2026-02-05T00:00:00"),
182
+ ("twin:doc-b", "Document", "Twin document B", "동일한 시각의 문서 B", {}, WS_ALPHA,
183
+ "2026-02-05T00:00:00"),
184
+ ("twin:concept-a", "Concept", "Twin concept A", "동일한 시각의 개념 A", {}, WS_ALPHA,
185
+ "2026-02-05T00:00:00"),
186
+ ("legacy:global-note", "Document", "Legacy global note",
187
+ "A legacy row with no workspace, visible only with include_legacy_global.",
188
+ {}, None, "2026-06-20T10:00:00"),
189
+ ]
190
+
191
+ # (chunk_id, parent_node_id, index, node_title, text, chunk_fields, workspace, updated_at)
192
+ #
193
+ # The shape mirrors ``KnowledgeGraphIngestMixin``: every chunk is BOTH a ``Chunk``
194
+ # node (lexical lane + workspace scoping) and a ``chunks`` row with its own
195
+ # embedding (vector lane, rolled up to its parent). Getting that duality wrong is
196
+ # the whole reason chunk-heavy queries are in the query set.
197
+ CHUNKS: List[Tuple[str, str, int, str, str, Dict[str, Any], Optional[str], str]] = [
198
+ ("chunk:handbook:1", "doc:handbook", 0, "handbook.pdf chunk 1",
199
+ "온보딩 체크리스트: 첫째, 폴더를 연결합니다. 둘째, 질문을 합니다. 셋째, 근거를 확인합니다.",
200
+ {"heading_path": "안내서 > 온보딩", "page": 3, "page_end": 4, "start_char": 0},
201
+ WS_ALPHA, "2026-06-12T10:00:01"),
202
+ ("chunk:handbook:2", "doc:handbook", 1, "handbook.pdf chunk 2",
203
+ "검색 화면에서는 회의 결정 사항을 한 번에 찾을 수 있습니다.",
204
+ {"heading_path": "안내서 > 검색", "page": 7, "start_char": 1200},
205
+ WS_ALPHA, "2026-06-12T10:00:02"),
206
+ ("chunk:spec:1", "doc:retrieval-spec", 0, "retrieval-spec.pdf chunk 1",
207
+ "The lexical channel scores one over rank. The vector channel is max normalized before fusion.",
208
+ {"heading_path": "Retrieval > Fusion", "page": 2, "start_char": 0},
209
+ WS_BETA, "2026-05-28T13:15:01"),
210
+ ("chunk:spec:2", "doc:retrieval-spec", 1, "retrieval-spec.pdf chunk 2",
211
+ "Ranking ties are broken by node id ascending, which keeps the answer stable across runs.",
212
+ {"start_char": 900}, WS_BETA, "2026-05-28T13:15:02"),
213
+ ]
214
+
215
+ # (from, to, type, weight, created_at)
216
+ #
217
+ # Weights and timestamps are assigned rather than inherited: ``relationship_search``
218
+ # orders by ``weight DESC, created_at DESC, id ASC`` and ``traverse`` caps every BFS
219
+ # round with ``ORDER BY weight DESC, id ASC``, so an edge set sharing one weight and
220
+ # one clock proves nothing. Three shapes are deliberate — a weight tie broken by
221
+ # ``created_at`` (the ``org:lattice`` pair), a weight *and* clock tie broken by edge
222
+ # id (``topic:quality``/``deck:review``), and legacy-global endpoints for scoping.
223
+ # Types are canonicalized by the write door (``relates_to``/``owns`` → ``MENTIONS``),
224
+ # which is why the goldens record uppercase names the fixture never spells.
225
+ EDGES: List[Tuple[str, str, str, float, str]] = [
226
+ ("dec:fusion-alpha", "concept:retrieval", "mentions", 0.9, "2026-07-01T00:00:00"),
227
+ ("dec:fusion-alpha", "concept:ranking", "mentions", 0.8, "2026-07-02T00:00:00"),
228
+ ("task:parity-harness", "dec:rust-foundation", "relates_to", 1.0, "2026-07-03T00:00:00"),
229
+ ("doc:handbook", "page:onboarding", "contains", 0.7, "2026-07-04T00:00:00"),
230
+ ("meeting:weekly", "dec:fusion-alpha", "discusses", 0.95, "2026-07-05T00:00:00"),
231
+ ("person:jiwon", "task:parity-harness", "owns", 0.6, "2026-07-06T00:00:00"),
232
+ ("person:minseo", "code:build-failure", "owns", 0.6, "2026-07-07T00:00:00"),
233
+ ("meeting:weekly", "task:onboarding-checklist", "discusses", 0.55, "2026-07-08T00:00:00"),
234
+ ("meeting:kickoff", "dec:rust-foundation", "discusses", 0.5, "2026-06-01T00:00:00"),
235
+ ("doc:retrieval-spec", "concept:ranking", "mentions", 0.45, "2026-06-02T00:00:00"),
236
+ ("concept:retrieval", "concept:ranking", "relates_to", 0.4, "2026-06-03T00:00:00"),
237
+ ("file:ranking-notes", "dec:fusion-alpha", "relates_to", 0.35, "2026-06-04T00:00:00"),
238
+ ("code:hybrid-search", "file:ranking-notes", "relates_to", 0.3, "2026-06-05T00:00:00"),
239
+ # Same weight, different clock → created_at DESC decides.
240
+ ("org:lattice", "person:jiwon", "contains", 0.25, "2026-06-06T00:00:00"),
241
+ ("org:lattice", "person:minseo", "contains", 0.25, "2026-06-07T00:00:00"),
242
+ ("slide:fusion", "concept:retrieval", "mentions", 0.25, "2026-06-08T00:00:00"),
243
+ # Same weight AND same clock → the id ASC tie-break is the only thing left.
244
+ ("topic:quality", "deck:review", "relates_to", 0.2, "2026-06-09T00:00:00"),
245
+ ("deck:review", "meeting:weekly", "relates_to", 0.2, "2026-06-09T00:00:00"),
246
+ ("task:onboarding-checklist", "page:onboarding", "relates_to", 0.15, "2026-05-01T00:00:00"),
247
+ ("doc:handbook", "doc:retrieval-spec", "relates_to", 0.1, "2026-05-02T00:00:00"),
248
+ ("tie:a", "tie:b", "relates_to", 1.0, "2026-07-03T00:00:00"),
249
+ ]
250
+
251
+ # ── the conversation corpus (episodic memory, same database file) ────────────
252
+ # (conversation_id, role, content, user_email, nickname, source, timestamp,
253
+ # workspace_id, organization_id, extra)
254
+ #
255
+ # Every branch the history reads take: two users × three workspaces, NULL and
256
+ # empty-string workspaces (the legacy rows ``_scope_sql`` admits), rows with no
257
+ # ``conversation_id`` (the ``legacy-previous-history`` bucket), a whitespace-only
258
+ # first message (the ``새 대화`` placeholder and its later upgrade), an
259
+ # assistant-first conversation (no upgrade), an empty timestamp (the ``or ""``
260
+ # fallbacks), extra keys ``metadata_json`` merges flat, and ko/en content.
261
+ MESSAGES: List[Tuple[Any, ...]] = [
262
+ ("conv-a", "user", "온보딩 체크리스트 어떻게 시작해?", "jiwon@lattice.ai", "지원", "web", "2026-07-20T09:00:00", WS_ALPHA, "org-1", {}),
263
+ ("conv-a", "assistant", "먼저 폴더를 연결하세요. 그다음 질문하면 됩니다.", "jiwon@lattice.ai", None, "web", "2026-07-20T09:00:05", WS_ALPHA, "org-1", {}),
264
+ ("conv-a", "user", "고마워", "jiwon@lattice.ai", "지원", "web", "2026-07-20T09:01:00", WS_ALPHA, "org-1", {"trace_id": "t-1"}),
265
+ ("conv-b", "user", "How does hybrid retrieval ranking work?", "minseo@lattice.ai", "Minseo", "telegram", "2026-07-21T10:00:00", WS_BETA, "org-1", {}),
266
+ ("conv-b", "assistant", "The lexical channel scores one over rank.", "minseo@lattice.ai", None, "telegram", "2026-07-21T10:00:07", WS_BETA, "org-1", {"tokens": 42, "cited": ["doc:retrieval-spec"]}),
267
+ ("conv-c", "user", " \n ", None, None, None, "2026-07-22T08:00:00", None, None, {}),
268
+ ("conv-c", "assistant", "무엇을 도와드릴까요?", None, None, None, "2026-07-22T08:00:05", None, None, {}),
269
+ ("conv-c", "user", "지난주 회의 기록 보여줘", None, None, None, "2026-07-22T08:01:00", None, None, {}),
270
+ (None, "user", "이전 대화 기록입니다", "jiwon@lattice.ai", "지원", "web", "2026-06-01T09:00:00", "", None, {}),
271
+ (None, "assistant", "네, 확인했습니다.", "jiwon@lattice.ai", None, "web", "2026-06-01T09:00:03", "", None, {}),
272
+ (None, "user", "legacy english message about ranking", None, None, None, "", None, None, {}),
273
+ ("conv-d", "user", "빌드 실패 원인 알려줘", "minseo@lattice.ai", "Minseo", "vscode", "2026-07-23T11:00:00", WS_ALPHA, "org-1", {}),
274
+ ("conv-d", "assistant", "컴파일 오류 로그를 확인하세요.", "minseo@lattice.ai", None, "vscode", "2026-07-23T11:00:04", WS_ALPHA, "org-1", {}),
275
+ ("conv-e", "user", "Ranking ties are broken by node id", "jiwon@lattice.ai", "지원", "web", "2026-07-24T12:00:00", WS_BETA, "org-2", {}),
276
+ ("conv-e", "assistant", "Yes — id ascending keeps it stable.", "jiwon@lattice.ai", None, "web", "2026-07-24T12:00:05", WS_BETA, "org-2", {}),
277
+ ("conv-f", "user", "검색 품질을 어떻게 측정하나요? 재현율과 정밀도를 모두 보고 싶고 주간 회의에서 공유할 예정입니다.", None, None, "web", "2026-07-25T13:00:00", None, None, {}),
278
+ ("conv-f", "assistant", "재현율/정밀도 지표는 분기 리뷰 발표자료에 있습니다.", None, None, "web", "2026-07-25T13:00:05", None, None, {}),
279
+ ("conv-g", "assistant", "assistant-first conversation", None, None, "web", "2026-07-26T14:00:00", None, None, {}),
280
+ ("conv-g", "user", "follow up question about ranking", None, None, "web", "2026-07-26T14:00:10", None, None, {}),
281
+ ("conv-h", "user", "Empty user and empty workspace row", "", "", "", "2026-07-27T15:00:00", "", "", {}),
282
+ ("conv-h", "assistant", "회의 결정 사항을 정리했습니다.", "", "", "", "2026-07-27T15:00:06", "", "", {}),
283
+ ]
284
+
285
+ # ── the query set ────────────────────────────────────────────────────────────
286
+ # ``allowed`` is None (no scoping), [] (a caller who may read nothing) or a list
287
+ # of workspace ids; ``legacy`` is ``include_legacy_global``.
288
+ QUERIES: List[Dict[str, Any]] = [
289
+ {"key": "en_fact", "query": "hybrid retrieval ranking"},
290
+ {"key": "ko_fact", "query": "회의 결정 사항"},
291
+ {"key": "en_code", "query": "vector_search() returns"},
292
+ {"key": "ko_code", "query": "빌드 실패 원인"},
293
+ {"key": "en_person", "query": "who owns the onboarding checklist"},
294
+ {"key": "ko_person", "query": "담당자 누구"},
295
+ {"key": "en_recency", "query": "recent decisions last week"},
296
+ {"key": "ko_recency", "query": "지난주 회의 기록"},
297
+ {"key": "en_filler", "query": " what is the retrieval specification please "},
298
+ {"key": "ko_filler", "query": "온보딩 체크리스트 좀 알려줘"},
299
+ {"key": "short_query", "query": "ai"},
300
+ {"key": "no_hit", "query": "zzqq wumpus nonsense"},
301
+ {"key": "tie_heavy", "query": "Tie candidate"},
302
+ {"key": "chunk_heavy", "query": "온보딩 체크리스트"},
303
+ {"key": "quoted", "query": 'said "retrieval" twice'},
304
+ {"key": "empty_query", "query": ""},
305
+ {"key": "ws_empty", "query": "hybrid retrieval ranking", "allowed": []},
306
+ {"key": "ws_empty_legacy", "query": "회의 결정 사항", "allowed": [], "legacy": True},
307
+ {"key": "ws_alpha", "query": "hybrid retrieval ranking", "allowed": [WS_ALPHA]},
308
+ {"key": "ws_alpha_legacy", "query": "회의 결정 사항", "allowed": [WS_ALPHA], "legacy": True},
309
+ {"key": "ws_beta", "query": "온보딩 체크리스트", "allowed": [WS_BETA]},
310
+ # A vector floor nothing clears — the only way to reach the stale-embedder
311
+ # probe and the lexical-only fusion label without breaking the store.
312
+ {"key": "min_vector_floor", "query": "hybrid retrieval ranking", "min_vector": 0.95},
313
+ # top_k small enough that the rerank window (top_k * 2) cuts the candidates.
314
+ {"key": "top_k_small", "query": "hybrid retrieval ranking", "top_k": 3},
315
+ # Pinned alpha: no policy, so no class, no rewrite, and no age decay on a
316
+ # query that would otherwise be recency-classed.
317
+ {"key": "alpha_pinned", "query": "지난주 회의 기록", "alpha": 0.2},
318
+ # A limit below the FTS hit count, so `ORDER BY rank LIMIT ?` decides which
319
+ # rows exist at all — the one place bm25 ordering (and therefore a SQLite
320
+ # version difference between the two runtimes) is observable.
321
+ {"key": "fts_rank_cut", "query": "Tie candidate", "limit": 2, "top_k": 2},
322
+ ]
323
+
324
+ #: Every branch of the memories section: a workspace hit, one with no kind and an
325
+ #: empty snippet, a non-workspace row to drop, and one past the limit.
326
+ CONTEXT_MEMORIES: Dict[str, Any] = {
327
+ "results": [
328
+ {"id": "mem-1", "kind": "preference", "snippet": "답변은 한국어로", "score": 0.91, "source": "workspace"},
329
+ {"id": "mem-2", "kind": None, "snippet": "", "score": 0.0, "source": "workspace"},
330
+ {"id": "mem-3", "kind": "decision", "snippet": "ranking keeps alpha fusion", "score": 0.4, "source": "personal"},
331
+ {"id": "mem-4", "kind": "fact", "snippet": "온보딩은 다섯 걸음", "score": 0.3, "source": "workspace"},
332
+ ]
333
+ }
334
+
335
+ #: A ledger with a pathless row, a non-dict row, and more than the ten-row cut.
336
+ CONTEXT_ARTIFACTS: List[Any] = [
337
+ {"path": "notes/ranking.md", "at": "2026-07-20T09:00:00", "run_id": "run-1"},
338
+ {"path": "notes/onboarding.md", "run_id": "run-2"},
339
+ {"path": "", "at": "2026-07-20T09:00:01"},
340
+ "not-a-dict",
341
+ ] + [{"path": f"out/file-{index}.md", "at": None, "run_id": f"r{index}"} for index in range(10)]
342
+
343
+ #: The Phase-2/3 suites: one spec list per ported entry point.
344
+ SUITES: Dict[str, List[Dict[str, Any]]] = {
345
+ "relationship": [
346
+ {"key": "all"},
347
+ {"key": "by_type_mention", "relationship_type": "mention"},
348
+ {"key": "by_type_contains", "relationship_type": "CONTAINS"},
349
+ {"key": "by_type_unknown", "relationship_type": "relates_to"},
350
+ {"key": "by_node", "node_id": "dec:fusion-alpha"},
351
+ {"key": "by_query_ko", "query": "회의"},
352
+ {"key": "by_query_en", "query": "ranking"},
353
+ {"key": "by_query_meta", "query": "lattice"},
354
+ {"key": "combined", "node_id": "dec:fusion-alpha", "relationship_type": "mention", "query": "retrieval"},
355
+ {"key": "limit_one", "limit": 1},
356
+ {"key": "limit_zero", "limit": 0},
357
+ {"key": "limit_over", "limit": 500},
358
+ {"key": "scoped_alpha", "allowed": [WS_ALPHA]},
359
+ {"key": "scoped_alpha_legacy", "allowed": [WS_ALPHA], "legacy": True},
360
+ {"key": "scoped_beta", "allowed": [WS_BETA]},
361
+ {"key": "scoped_empty", "allowed": []},
362
+ {"key": "no_hit", "query": "zzqq wumpus"},
363
+ ],
364
+ "traverse": [
365
+ # 9 clamps to 4 and -1 clamps to 0; both are on purpose.
366
+ *[{"key": f"hub_d{depth}", "node_id": "dec:fusion-alpha", "depth": depth}
367
+ for depth in (0, 1, 2, 3, 9)],
368
+ {"key": "hub_dneg", "node_id": "dec:fusion-alpha", "depth": -1},
369
+ {"key": "leaf_d2", "node_id": "code:hybrid-search", "depth": 2},
370
+ {"key": "isolated", "node_id": "tie:c", "depth": 2},
371
+ {"key": "limit_two", "node_id": "dec:fusion-alpha", "depth": 3, "limit": 2},
372
+ {"key": "limit_five", "node_id": "dec:fusion-alpha", "depth": 3, "limit": 5},
373
+ {"key": "limit_zero", "node_id": "dec:fusion-alpha", "depth": 2, "limit": 0},
374
+ {"key": "limit_over", "node_id": "dec:fusion-alpha", "depth": 2, "limit": 900},
375
+ {"key": "org_hub", "node_id": "org:lattice", "depth": 2},
376
+ {"key": "tie_pair", "node_id": "tie:a", "depth": 2},
377
+ {"key": "scoped_alpha", "node_id": "dec:fusion-alpha", "depth": 2, "allowed": [WS_ALPHA]},
378
+ {"key": "scoped_alpha_legacy", "node_id": "dec:fusion-alpha", "depth": 2, "allowed": [WS_ALPHA], "legacy": True},
379
+ {"key": "scoped_seed_hidden", "node_id": "dec:fusion-alpha", "allowed": [WS_BETA]},
380
+ {"key": "scoped_empty", "node_id": "dec:fusion-alpha", "allowed": []},
381
+ {"key": "empty_id", "node_id": ""},
382
+ {"key": "missing_seed", "node_id": "nope:missing"},
383
+ ],
384
+ "graph_search": [
385
+ {"key": "en_fact", "query": "hybrid retrieval ranking"},
386
+ {"key": "ko_fact", "query": "회의 결정 사항"},
387
+ {"key": "person", "query": "who owns the onboarding checklist"},
388
+ {"key": "code", "query": "빌드 실패 원인"},
389
+ {"key": "expand0", "query": "hybrid retrieval ranking", "expand_depth": 0},
390
+ {"key": "expand3", "query": "hybrid retrieval ranking", "expand_depth": 3},
391
+ {"key": "expand_clamp", "query": "회의 결정 사항", "expand_depth": 9},
392
+ {"key": "limit_small", "query": "hybrid retrieval ranking", "limit": 3},
393
+ {"key": "limit_over", "query": "회의 결정 사항", "limit": 500},
394
+ {"key": "scoped_alpha", "query": "hybrid retrieval ranking", "allowed": [WS_ALPHA]},
395
+ {"key": "scoped_alpha_legacy", "query": "회의 결정 사항", "allowed": [WS_ALPHA], "legacy": True},
396
+ {"key": "scoped_empty", "query": "hybrid retrieval ranking", "allowed": []},
397
+ {"key": "no_hit", "query": "zzqq wumpus nonsense"},
398
+ {"key": "empty_query", "query": ""},
399
+ ],
400
+ "service_hybrid": [
401
+ {"key": "en_fact", "query": "hybrid retrieval ranking"},
402
+ {"key": "ko_recency", "query": "지난주 회의 기록"},
403
+ {"key": "en_code", "query": "vector_search() returns"},
404
+ {"key": "ko_person", "query": "담당자 누구"},
405
+ {"key": "en_filler", "query": " what is the retrieval specification please "},
406
+ # Explicit weights disable BOTH the rewrite and the age decay, on a
407
+ # query that would otherwise get both. That asymmetry is the contract.
408
+ {"key": "pinned_recency", "query": "지난주 회의 기록", "weights": {"keyword": 0.5, "vector": 0.3, "graph": 0.2}},
409
+ {"key": "pinned_partial", "query": "hybrid retrieval ranking", "weights": {"graph": 1.0}},
410
+ {"key": "pinned_zero", "query": "회의 결정 사항", "weights": {"keyword": 0.0, "vector": 0.0, "graph": 0.0}},
411
+ {"key": "limit_small", "query": "hybrid retrieval ranking", "limit": 3},
412
+ {"key": "limit_over", "query": "회의 결정 사항", "limit": 500},
413
+ {"key": "channel_limits", "query": "hybrid retrieval ranking", "keyword_limit": 5, "vector_limit": 5, "graph_limit": 5},
414
+ {"key": "scoped_alpha", "query": "hybrid retrieval ranking", "allowed": [WS_ALPHA]},
415
+ {"key": "scoped_alpha_legacy", "query": "회의 결정 사항", "allowed": [WS_ALPHA], "legacy": True},
416
+ {"key": "scoped_empty", "query": "hybrid retrieval ranking", "allowed": []},
417
+ {"key": "no_hit", "query": "zzqq wumpus nonsense"},
418
+ {"key": "empty_query", "query": ""},
419
+ ],
420
+ "history": [
421
+ {"key": "all"},
422
+ {"key": "limit_two", "limit": 2},
423
+ {"key": "limit_zero", "limit": 0},
424
+ {"key": "conv_a", "conversation_id": "conv-a"},
425
+ {"key": "conv_missing", "conversation_id": "nope"},
426
+ {"key": "conv_null", "conversation_id": ""},
427
+ {"key": "user_jiwon", "user_email": "jiwon@lattice.ai"},
428
+ {"key": "user_jiwon_strict", "user_email": "jiwon@lattice.ai", "legacy": False},
429
+ {"key": "user_unknown", "user_email": "ghost@lattice.ai"},
430
+ {"key": "ws_alpha", "allowed": [WS_ALPHA]},
431
+ {"key": "ws_alpha_strict", "allowed": [WS_ALPHA], "legacy": False},
432
+ {"key": "ws_both_strict", "allowed": [WS_ALPHA, WS_BETA], "legacy": False},
433
+ {"key": "ws_empty_legacy", "allowed": []},
434
+ {"key": "ws_empty_strict", "allowed": [], "legacy": False},
435
+ {"key": "user_and_ws", "user_email": "jiwon@lattice.ai", "allowed": [WS_ALPHA], "legacy": False},
436
+ {"key": "conv_and_user", "conversation_id": "conv-a", "user_email": "jiwon@lattice.ai", "legacy": False},
437
+ {"key": "ws_blank_only", "allowed": [""], "legacy": False},
438
+ ],
439
+ "conversations": [
440
+ {"key": "all"},
441
+ {"key": "user_jiwon", "user_email": "jiwon@lattice.ai"},
442
+ {"key": "user_jiwon_strict", "user_email": "jiwon@lattice.ai", "legacy": False},
443
+ {"key": "ws_alpha_strict", "allowed": [WS_ALPHA], "legacy": False},
444
+ {"key": "ws_beta_strict", "allowed": [WS_BETA], "legacy": False},
445
+ {"key": "ws_empty_strict", "allowed": [], "legacy": False},
446
+ ],
447
+ "conversation_messages": [
448
+ {"key": "conv_a", "conversation_id": "conv-a"},
449
+ {"key": "legacy_bucket", "conversation_id": "legacy-previous-history"},
450
+ {"key": "missing", "conversation_id": "nope"},
451
+ {"key": "scoped_alpha_strict", "conversation_id": "conv-a", "allowed": [WS_ALPHA], "legacy": False},
452
+ {"key": "legacy_bucket_scoped", "conversation_id": "legacy-previous-history", "allowed": [WS_ALPHA], "legacy": False},
453
+ ],
454
+ "history_search": [
455
+ {"key": "ko_hit", "query": "회의"},
456
+ {"key": "ko_partial", "query": "체크리스트"},
457
+ {"key": "en_hit", "query": "ranking"},
458
+ {"key": "case_insensitive", "query": "RANKING"},
459
+ {"key": "blank", "query": " "},
460
+ {"key": "no_hit", "query": "zzqq"},
461
+ {"key": "limit_one", "query": "ranking", "limit": 1},
462
+ {"key": "scoped_strict", "query": "ranking", "allowed": [WS_BETA], "legacy": False},
463
+ ],
464
+ "context_assemble": [
465
+ {"key": "all_seams", "query": "회의 결정 사항", "memories": CONTEXT_MEMORIES, "artifacts": CONTEXT_ARTIFACTS, "notes": "정원 노트: 랭킹은 alpha 융합을 유지한다.", "recent": {"limit": 4}},
466
+ {"key": "knowledge_only", "query": "hybrid retrieval ranking"},
467
+ {"key": "no_seams", "query": "hybrid retrieval ranking", "knowledge": False},
468
+ {"key": "memories_only", "query": "온보딩", "knowledge": False, "memories": CONTEXT_MEMORIES, "memory_limit": 2},
469
+ {"key": "artifacts_only", "query": "온보딩", "knowledge": False, "artifacts": CONTEXT_ARTIFACTS},
470
+ {"key": "notes_blank", "query": "온보딩", "knowledge": False, "notes": " "},
471
+ {"key": "recent_conversation", "query": "온보딩", "knowledge": False, "recent": {"conversation_id": "conv-a", "limit": 10}},
472
+ {"key": "recent_personal_workspace", "query": "온보딩", "knowledge": False, "recent": {"workspace_id": "personal", "limit": 6}},
473
+ {"key": "recent_user_scoped", "query": "온보딩", "knowledge": False, "recent": {"user_email": "jiwon@lattice.ai", "limit": 5}},
474
+ {"key": "budget_tiny", "query": "회의 결정 사항", "memories": CONTEXT_MEMORIES, "artifacts": CONTEXT_ARTIFACTS, "notes": "정원 노트: 랭킹은 alpha 융합을 유지한다.", "recent": {"limit": 4}, "budget": 20},
475
+ {"key": "budget_one", "query": "회의 결정 사항", "memories": CONTEXT_MEMORIES, "artifacts": CONTEXT_ARTIFACTS, "notes": "노트", "recent": {"limit": 4}, "budget": 1},
476
+ {"key": "budget_zero", "query": "회의 결정 사항", "memories": CONTEXT_MEMORIES, "notes": "노트", "budget": 0},
477
+ {"key": "knowledge_limit_one", "query": "hybrid retrieval ranking", "knowledge_limit": 1},
478
+ ],
479
+ }
480
+
481
+ #: Texts whose tokenizer output, hashes and vectors pin the embedding port.
482
+ EMBEDDING_TEXTS: List[str] = [
483
+ "hybrid retrieval ranking",
484
+ "회의 결정 사항",
485
+ "온보딩 체크리스트 v2",
486
+ "vector_search() returns a dict",
487
+ "!!! ??? ...",
488
+ "MixedCase and snake_case and kebab-case",
489
+ "2026-07-20T09:00:00",
490
+ "a",
491
+ ]
492
+
493
+ #: Values separating CPython round-half-even from naive scaling, plus real shapes.
494
+ ROUNDING_VALUES: List[float] = [
495
+ 0.0, 1.0, 5e-07, 1.5e-06, 2.5e-06, 2.6535895, 1 / 3, 1 / 7,
496
+ 0.1234565, 0.1234575, 0.9999999999, 123456.7890625, -0.0000005,
497
+ 0.6 * 0.5 + 0.4 * (1 / 3), 0.35 * 0.25 + 0.65 * (1 / 7), 0.5 + 0.5 * 0.7071067811865476,
498
+ ]
499
+
500
+
501
+ @contextmanager
502
+ def pinned_environment() -> Iterator[None]:
503
+ """Pin every env knob the retrieval stack reads, then restore."""
504
+ previous = {key: os.environ.get(key) for key in PINNED_ENV}
505
+ os.environ.update(PINNED_ENV)
506
+ try:
507
+ yield
508
+ finally:
509
+ for key, value in previous.items():
510
+ if value is None:
511
+ os.environ.pop(key, None)
512
+ else:
513
+ os.environ[key] = value
514
+
515
+
516
+ @contextmanager
517
+ def frozen_clock() -> Iterator[None]:
518
+ """Freeze every ported ``datetime.now()`` at :data:`FROZEN_NOW`.
519
+
520
+ Only the recency-class age decay reads the clock, through the ``datetime``
521
+ name its module imported — so rebinding that name is the whole patch. Two
522
+ modules do it: the graph-layer and the service-layer ``hybrid_search``.
523
+ """
524
+ from lattice_brain.graph.retrieval import hybrid as hybrid_module
525
+ from latticeai.services import search_service as service_module
526
+
527
+ frozen = datetime.fromisoformat(FROZEN_NOW)
528
+
529
+ class _FrozenDatetime(datetime):
530
+ @classmethod
531
+ def now(cls, tz=None): # noqa: ARG003 — mirrors datetime.now's signature
532
+ return frozen
533
+
534
+ modules = (hybrid_module, service_module)
535
+ originals = [module.datetime for module in modules]
536
+ for module in modules:
537
+ module.datetime = _FrozenDatetime
538
+ try:
539
+ yield
540
+ finally:
541
+ for module, original in zip(modules, originals, strict=True):
542
+ module.datetime = original
543
+
544
+
545
+ @contextmanager
546
+ def rules_only_extraction() -> Iterator[None]:
547
+ """Force ``_topic_candidates`` down its rule-based path.
548
+
549
+ The LLM path needs a bound router no fixture run has, but "no router happens
550
+ to be bound" is an accident and this contract cannot rest on one.
551
+ """
552
+ from lattice_brain.graph._kg_common import extraction
553
+
554
+ original = extraction.ENABLE_LLM_EXTRACTION
555
+ extraction.ENABLE_LLM_EXTRACTION = False
556
+ try:
557
+ yield
558
+ finally:
559
+ extraction.ENABLE_LLM_EXTRACTION = original
560
+
561
+
562
+ def open_store(db_path: Path):
563
+ """A ``KnowledgeGraphStore`` over ``db_path`` (blobs beside it)."""
564
+ from lattice_brain.graph import KnowledgeGraphStore
565
+ return KnowledgeGraphStore(Path(db_path), Path(db_path).parent / "blobs")
566
+
567
+
568
+ def open_conversations(db_path: Path):
569
+ """A ``ConversationStore`` over ``db_path`` (the same file as the graph)."""
570
+ from lattice_brain.conversations import ConversationStore
571
+ return ConversationStore(Path(db_path))
572
+
573
+
574
+ def write_conversations(db_path: Path) -> None:
575
+ """Append :data:`MESSAGES` through the real durable-history write path."""
576
+ conversations = open_conversations(db_path)
577
+ for conv_id, role, content, email, nick, source, stamp, workspace, org, extra in MESSAGES:
578
+ conversations.append({
579
+ "conversation_id": conv_id, "role": role, "content": content,
580
+ "user_email": email, "user_nickname": nick, "source": source,
581
+ "timestamp": stamp, "workspace_id": workspace, "organization_id": org,
582
+ **extra,
583
+ })
584
+
585
+
586
+ def _backdate(conn: sqlite3.Connection, node_id: str, stamp: str) -> None:
587
+ conn.execute("UPDATE nodes SET created_at=?, updated_at=? WHERE id=?", (stamp, stamp, node_id))
588
+ conn.execute(
589
+ "UPDATE nodes_v2 SET created_at=?, updated_at=? WHERE id=?", (stamp, stamp, node_id)
590
+ )
591
+
592
+
593
+ def build_store(db_path: Path) -> None:
594
+ """Create the fixture database from scratch with the real write path."""
595
+ for sibling in (db_path, Path(f"{db_path}-wal"), Path(f"{db_path}-shm")):
596
+ if sibling.exists():
597
+ sibling.unlink()
598
+ blob_dir = db_path.parent / "blobs"
599
+ if blob_dir.exists():
600
+ shutil.rmtree(blob_dir)
601
+
602
+ store = open_store(db_path)
603
+ with store._connect() as conn:
604
+ for node_id, node_type, title, summary, metadata, workspace, stamp in NODES:
605
+ store._upsert_node(
606
+ conn, node_id, node_type, title, summary, metadata, workspace_id=workspace
607
+ )
608
+ _backdate(conn, node_id, stamp)
609
+ for chunk_id, parent, index, title, text, fields, workspace, stamp in CHUNKS:
610
+ metadata = {"index": index, "source_node": parent, **fields}
611
+ store._upsert_node(
612
+ conn, chunk_id, "Chunk", title, text[:500], metadata, workspace_id=workspace
613
+ )
614
+ store._upsert_chunk(
615
+ conn, chunk_id=chunk_id, source_node=parent, text=text, metadata=metadata
616
+ )
617
+ _backdate(conn, chunk_id, stamp)
618
+ for from_node, to_node, edge_type, weight, stamp in EDGES:
619
+ # ``_upsert_edge`` stamps ``created_at`` from the wall clock in BOTH
620
+ # tables, and the read path is the ``kgv2_edges`` view over ``edges_v2``:
621
+ # backdating only the legacy table (as this generator first did) left
622
+ # the relationship ordering moving with the clock.
623
+ edge_id = store._upsert_edge(conn, from_node, to_node, edge_type, weight, {})
624
+ conn.execute("UPDATE edges SET created_at=? WHERE id=?", (stamp, edge_id))
625
+ conn.execute("UPDATE edges_v2 SET created_at=? WHERE id=?", (stamp, edge_id))
626
+ # ``indexed_at`` decides the candidate scan order (and, when the cap bites,
627
+ # which candidates exist at all), so it is assigned rather than inherited.
628
+ item_ids = [row["item_id"] for row in conn.execute(
629
+ "SELECT item_id FROM vector_embeddings ORDER BY item_id ASC"
630
+ ).fetchall()]
631
+ for seq, item_id in enumerate(item_ids):
632
+ stamp = f"2026-03-01T{seq // 60:02d}:{seq % 60:02d}:00"
633
+ conn.execute(
634
+ "UPDATE vector_embeddings SET indexed_at=? WHERE item_id=?", (stamp, item_id)
635
+ )
636
+ store.record_embedder_fingerprint()
637
+ write_conversations(db_path)
638
+
639
+ # Leave one self-contained file behind: checkpoint the WAL, drop back to a
640
+ # rollback journal so no sidecar is committed, and compact.
641
+ with sqlite3.connect(str(db_path)) as conn:
642
+ conn.execute("PRAGMA wal_checkpoint(TRUNCATE)")
643
+ conn.execute("PRAGMA journal_mode=DELETE")
644
+ with sqlite3.connect(str(db_path)) as conn:
645
+ conn.execute("VACUUM")
646
+ conn.close()
647
+ for sibling in (Path(f"{db_path}-wal"), Path(f"{db_path}-shm")):
648
+ if sibling.exists():
649
+ sibling.unlink()
650
+ if blob_dir.exists():
651
+ shutil.rmtree(blob_dir)
652
+
653
+
654
+ def _allowed(spec: Dict[str, Any]):
655
+ allowed = spec.get("allowed")
656
+ return None if allowed is None else set(allowed)
657
+
658
+
659
+ ENGINES: Dict[str, Callable[[Any, Dict[str, Any]], Dict[str, Any]]] = {
660
+ "hybrid": lambda store, spec: store.hybrid_search(
661
+ spec["query"], top_k=spec.get("top_k", 20), alpha=spec.get("alpha"),
662
+ allowed_workspaces=_allowed(spec), min_vector_score=spec.get("min_vector", 0.0),
663
+ include_legacy_global=spec.get("legacy", False),
664
+ ),
665
+ "keyword": lambda store, spec: store.search(
666
+ spec["query"], spec.get("limit", 30), allowed_workspaces=_allowed(spec),
667
+ include_legacy_global=spec.get("legacy", False),
668
+ ),
669
+ "vector": lambda store, spec: store.vector_search(
670
+ spec["query"], limit=spec.get("limit", 30), min_score=spec.get("min_score", 0.0)
671
+ ),
672
+ }
673
+
674
+
675
+ def run_engine(store, engine: str, spec: Dict[str, Any]) -> Dict[str, Any]:
676
+ """Run one engine for one query spec under the frozen clock."""
677
+ with frozen_clock(), rules_only_extraction():
678
+ return ENGINES[engine](store, spec)
679
+
680
+
681
+ # ── suites (v11.5.0): KG reads, the service layer, history, context ──────────
682
+ #
683
+ # The Phase-1 engines share one query shape (a query set × an engine set); the
684
+ # Phase-2/3 ports do not, so each gets its own spec list and runner, both carried
685
+ # in the manifest so Rust and Python enumerate exactly the same work.
686
+
687
+
688
+ class Harness:
689
+ """Every Python entry point the v11.5.0 goldens are produced from.
690
+
691
+ ``require_auth=False`` is the loopback-owner configuration the native routes
692
+ reproduce: the history scope is whatever the caller passes explicitly.
693
+ """
694
+
695
+ def __init__(self, db_path: Path):
696
+ from latticeai.runtime.history_runtime import build_history_query_runtime
697
+ from latticeai.services.chat_service import ChatService
698
+ from latticeai.services.search_service import SearchService
699
+ self.store = open_store(db_path)
700
+ self.service = SearchService(graph_store=self.store)
701
+ self.conversations = open_conversations(db_path)
702
+ self.history_runtime = build_history_query_runtime(
703
+ conversations=self.conversations,
704
+ workspace_service=None,
705
+ require_auth=False,
706
+ logging=logging,
707
+ )
708
+ self.chat = ChatService(store=None, get_history=self.history_runtime["get_history"])
709
+
710
+
711
+ def _allowed_list(spec: Dict[str, Any]):
712
+ """``allowed`` as the graph layer wants it: ``None`` or a set."""
713
+ allowed = spec.get("allowed")
714
+ return None if allowed is None else set(allowed)
715
+
716
+
717
+ def _history_scope(spec: Dict[str, Any]) -> Dict[str, Any]:
718
+ """The identity/workspace scope every history read takes.
719
+
720
+ ``include_legacy_global`` defaults to ``True`` — ``ConversationStore``'s own
721
+ default, the opposite of the graph layer's and the kind of asymmetry a port
722
+ gets wrong.
723
+ """
724
+ allowed = spec.get("allowed")
725
+ return {
726
+ "user_email": spec.get("user_email"),
727
+ "allowed_workspaces": None if allowed is None else list(allowed),
728
+ "include_legacy_global": spec.get("legacy", True),
729
+ }
730
+
731
+
732
+ def _run_relationship(h: Harness, spec: Dict[str, Any]) -> Dict[str, Any]:
733
+ return h.store.relationship_search(
734
+ query=spec.get("query", ""), node_id=spec.get("node_id", ""),
735
+ relationship_type=spec.get("relationship_type", ""), limit=spec.get("limit", 30),
736
+ allowed_workspaces=_allowed_list(spec),
737
+ include_legacy_global=spec.get("legacy", False),
738
+ )
739
+
740
+
741
+ def _run_traverse(h: Harness, spec: Dict[str, Any]) -> Dict[str, Any]:
742
+ try:
743
+ return h.store.traverse(
744
+ spec.get("node_id", ""), depth=spec.get("depth", 1),
745
+ limit=spec.get("limit", 100), allowed_workspaces=_allowed_list(spec),
746
+ include_legacy_global=spec.get("legacy", False),
747
+ )
748
+ except ValueError as exc:
749
+ # The two documented refusals (blank id, seed invisible to the caller's
750
+ # scope) are contract, so they are recorded rather than skipped; a payload
751
+ # never carries an "error" key, so the golden stays unambiguous.
752
+ return {"error": str(exc)}
753
+
754
+
755
+ def _run_graph_search(h: Harness, spec: Dict[str, Any]) -> Dict[str, Any]:
756
+ return h.service.graph_search(
757
+ spec["query"], limit=spec.get("limit", 30),
758
+ expand_depth=spec.get("expand_depth", 1), allowed_workspaces=_allowed_list(spec),
759
+ include_legacy_global=spec.get("legacy", False),
760
+ )
761
+
762
+
763
+ def _run_service_hybrid(h: Harness, spec: Dict[str, Any]) -> Dict[str, Any]:
764
+ return h.service.hybrid_search(
765
+ spec["query"], limit=spec.get("limit", 30),
766
+ keyword_limit=spec.get("keyword_limit", 30),
767
+ vector_limit=spec.get("vector_limit", 30),
768
+ graph_limit=spec.get("graph_limit", 30), weights=spec.get("weights"),
769
+ allowed_workspaces=_allowed_list(spec),
770
+ include_legacy_global=spec.get("legacy", False),
771
+ )
772
+
773
+
774
+ def _run_history(h: Harness, spec: Dict[str, Any]) -> List[Dict[str, Any]]:
775
+ return h.conversations.history(
776
+ conversation_id=spec.get("conversation_id"), limit=spec.get("limit"),
777
+ **_history_scope(spec),
778
+ )
779
+
780
+
781
+ def _run_conversations(h: Harness, spec: Dict[str, Any]) -> List[Dict[str, Any]]:
782
+ history = h.history_runtime["get_history"](**_history_scope(spec))
783
+ return h.history_runtime["group_history_conversations"](history)
784
+
785
+
786
+ def _run_conversation_messages(h: Harness, spec: Dict[str, Any]) -> List[Dict[str, Any]]:
787
+ return h.history_runtime["get_conversation_messages"](
788
+ spec["conversation_id"], **_history_scope(spec)
789
+ )
790
+
791
+
792
+ def _run_history_search(h: Harness, spec: Dict[str, Any]) -> List[Dict[str, Any]]:
793
+ return h.chat.search_history(
794
+ spec["query"], scope=_history_scope(spec), limit=spec.get("limit", 30),
795
+ conversation_title=h.history_runtime["conversation_title"],
796
+ )
797
+
798
+
799
+ def _context_seams(h: Harness, spec: Dict[str, Any]) -> Dict[str, Any]:
800
+ """The seam set for one context spec — data seams plus the real engines.
801
+
802
+ ``memories`` / ``artifacts`` / ``notes`` are *data* seams: the payload is the
803
+ spec, so both runtimes feed the assembler the same bytes and what is under
804
+ test is the assembler. ``knowledge`` and ``recent`` are real — the
805
+ service-layer hybrid search and the durable history reader.
806
+ """
807
+ from latticeai.api.chat_helpers import build_recent_chat_context
808
+
809
+ # Signatures matter: the assembler inspects them to decide which context
810
+ # fields a seam may be handed, so each one declares exactly what it accepts.
811
+ seams: Dict[str, Any] = {}
812
+ if spec.get("memories") is not None:
813
+ memories = spec["memories"]
814
+ seams["memory_recall"] = (
815
+ lambda query, *, user_email=None, workspace_id=None, limit=5: memories
816
+ )
817
+ if spec.get("artifacts") is not None:
818
+ artifacts = spec["artifacts"]
819
+ seams["recent_artifacts"] = (
820
+ lambda *, user_email=None, conversation_id=None, workspace_id=None: artifacts
821
+ )
822
+ if spec.get("knowledge", True):
823
+ # Loopback trust: no workspace scoping, exactly as on the native route.
824
+ seams["hybrid_search"] = (
825
+ lambda query, *, limit=5, user_email=None, workspace_id=None:
826
+ h.service.hybrid_search(query, limit=limit)
827
+ )
828
+ if spec.get("notes") is not None:
829
+ notes = spec["notes"]
830
+ seams["notes_context"] = lambda query, *, user_email=None, workspace_id=None: notes
831
+ if spec.get("recent") is not None:
832
+ recent = spec["recent"]
833
+ seams["recent_chat"] = (
834
+ lambda *, user_email=None, conversation_id=None, workspace_id=None:
835
+ build_recent_chat_context(
836
+ get_history=h.history_runtime["get_history"],
837
+ limit=recent.get("limit", 10),
838
+ include_image_missing_replies=recent.get("images", True),
839
+ user_email=recent.get("user_email"),
840
+ conversation_id=recent.get("conversation_id"),
841
+ workspace_id=recent.get("workspace_id"),
842
+ )
843
+ )
844
+ return seams
845
+
846
+
847
+ def _run_context_assemble(h: Harness, spec: Dict[str, Any]) -> Dict[str, Any]:
848
+ from lattice_brain.context import ContextAssembler
849
+
850
+ assembled = ContextAssembler(**_context_seams(h, spec)).assemble(
851
+ spec["query"], user_email=spec.get("user_email"),
852
+ workspace_id=spec.get("workspace_id"),
853
+ conversation_id=spec.get("conversation_id"), budget=spec.get("budget", 2000),
854
+ memory_limit=spec.get("memory_limit", 5),
855
+ knowledge_limit=spec.get("knowledge_limit", 5),
856
+ )
857
+ return {"text": assembled.text, "approx_tokens": assembled.approx_tokens,
858
+ "trace": assembled.trace()}
859
+
860
+
861
+ SUITE_RUNNERS: Dict[str, Callable[[Harness, Dict[str, Any]], Any]] = {
862
+ "relationship": _run_relationship,
863
+ "traverse": _run_traverse,
864
+ "graph_search": _run_graph_search,
865
+ "service_hybrid": _run_service_hybrid,
866
+ "history": _run_history,
867
+ "conversations": _run_conversations,
868
+ "conversation_messages": _run_conversation_messages,
869
+ "history_search": _run_history_search,
870
+ "context_assemble": _run_context_assemble,
871
+ }
872
+
873
+
874
+ def run_suite(harness: Harness, suite: str, spec: Dict[str, Any]) -> Any:
875
+ """Run one suite spec under the frozen clock and the rule-based extractor."""
876
+ with frozen_clock(), rules_only_extraction():
877
+ return SUITE_RUNNERS[suite](harness, spec)
878
+
879
+
880
+ def golden_path(engine: str, key: str) -> Path:
881
+ return GOLDEN_DIR / f"{engine}__{key}.json"
882
+
883
+
884
+ def golden_payload(engine: str, spec: Dict[str, Any], result: Dict[str, Any]) -> Dict[str, Any]:
885
+ return {
886
+ "engine": engine,
887
+ "key": spec["key"],
888
+ "query": spec["query"],
889
+ "params": {
890
+ "top_k": spec.get("top_k", 20), "alpha": spec.get("alpha"),
891
+ "limit": spec.get("limit", 30), "min_score": spec.get("min_score", 0.0),
892
+ "min_vector_score": spec.get("min_vector", 0.0),
893
+ "allowed_workspaces": spec.get("allowed"),
894
+ "include_legacy_global": spec.get("legacy", False),
895
+ },
896
+ "result": result,
897
+ }
898
+
899
+
900
+ def suite_payload(suite: str, spec: Dict[str, Any], result: Any) -> Dict[str, Any]:
901
+ """One suite golden: the spec that produced it, verbatim, and the answer.
902
+
903
+ The spec rides along rather than being flattened into a fixed ``params``
904
+ block: these entry points share no parameter shape.
905
+ """
906
+ return {"suite": suite, "key": spec["key"], "spec": spec, "result": result}
907
+
908
+
909
+ def _dump(path: Path, payload: Any) -> None:
910
+ path.write_text(
911
+ json.dumps(payload, ensure_ascii=False, sort_keys=True, indent=2) + "\n",
912
+ encoding="utf-8",
913
+ )
914
+
915
+
916
+ def embeddings_golden() -> Dict[str, Any]:
917
+ """Tokenizer output, hash pairs and full vectors for the pinned texts."""
918
+ from lattice_brain.embeddings import LocalEmbeddingModel, _hash_to_index, _tokenize
919
+
920
+ model = LocalEmbeddingModel()
921
+ cases = []
922
+ for text in EMBEDDING_TEXTS:
923
+ features = _tokenize(text)
924
+ cases.append(
925
+ {
926
+ "text": text,
927
+ "features": features,
928
+ "hashes": [
929
+ {"feature": feature, "index": index, "sign": sign}
930
+ for feature, (index, sign) in (
931
+ (feature, _hash_to_index(feature, model.dim))
932
+ for feature in features[:8]
933
+ )
934
+ ],
935
+ "vector": model.embed(text),
936
+ "encoded_hex": model.encode(model.embed(text)).hex(),
937
+ }
938
+ )
939
+ return {"model_id": model.model_id, "dim": model.dim, "cases": cases}
940
+
941
+
942
+ def rounding_golden() -> List[Dict[str, float]]:
943
+ """``round(x, 6)`` for values where the tie rule is observable."""
944
+ return [{"input": v, "expected": round(v, 6)} for v in ROUNDING_VALUES]
945
+
946
+
947
+ def manifest() -> Dict[str, Any]:
948
+ from lattice_brain.embeddings import LocalEmbeddingModel
949
+ model = LocalEmbeddingModel()
950
+ return {
951
+ "frozen_now": FROZEN_NOW,
952
+ "store": STORE_PATH.name,
953
+ "embedding_model": model.model_id,
954
+ "embedding_dim": model.dim,
955
+ "engines": sorted(ENGINES),
956
+ "queries": QUERIES,
957
+ "suites": {suite: SUITES[suite] for suite in sorted(SUITES)},
958
+ "pinned_env": PINNED_ENV,
959
+ }
960
+
961
+
962
+ def main() -> int:
963
+ FIXTURE_DIR.mkdir(parents=True, exist_ok=True)
964
+ if GOLDEN_DIR.exists():
965
+ shutil.rmtree(GOLDEN_DIR)
966
+ GOLDEN_DIR.mkdir(parents=True, exist_ok=True)
967
+ with pinned_environment():
968
+ build_store(STORE_PATH)
969
+ store = open_store(STORE_PATH)
970
+ written = 0
971
+ for spec in QUERIES:
972
+ for engine in sorted(ENGINES):
973
+ result = run_engine(store, engine, spec)
974
+ _dump(golden_path(engine, spec["key"]), golden_payload(engine, spec, result))
975
+ written += 1
976
+ harness = Harness(STORE_PATH)
977
+ for suite in sorted(SUITES):
978
+ for spec in SUITES[suite]:
979
+ result = run_suite(harness, suite, spec)
980
+ _dump(golden_path(suite, spec["key"]), suite_payload(suite, spec, result))
981
+ written += 1
982
+ _dump(GOLDEN_DIR / "embeddings_golden.json", embeddings_golden())
983
+ _dump(GOLDEN_DIR / "rounding_golden.json", rounding_golden())
984
+ _dump(GOLDEN_DIR / "manifest.json", manifest())
985
+ # ``KnowledgeGraphStore.__init__`` creates its blob directory eagerly and the
986
+ # fixture has no blobs; an empty directory beside an artefact is just noise.
987
+ blob_dir = STORE_PATH.parent / "blobs"
988
+ if blob_dir.is_dir() and not any(blob_dir.iterdir()):
989
+ blob_dir.rmdir()
990
+ size_kb = STORE_PATH.stat().st_size / 1024
991
+ print(f"store: {STORE_PATH.relative_to(REPO_ROOT)} ({size_kb:.1f} KiB)")
992
+ print(f"golden: {written} engine files + embeddings + rounding + manifest")
993
+ return 0
994
+
995
+
996
+ if __name__ == "__main__":
997
+ raise SystemExit(main())