memorix 1.2.1 → 1.2.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (199) hide show
  1. package/CHANGELOG.md +16 -0
  2. package/README.md +14 -2
  3. package/README.zh-CN.md +14 -2
  4. package/TEAM.md +86 -86
  5. package/dist/cli/index.js +15407 -13779
  6. package/dist/cli/index.js.map +1 -1
  7. package/dist/index.js +1321 -529
  8. package/dist/index.js.map +1 -1
  9. package/dist/maintenance-runner.d.ts +1 -1
  10. package/dist/maintenance-runner.js +8458 -8087
  11. package/dist/maintenance-runner.js.map +1 -1
  12. package/dist/memcode-runtime/CHANGELOG.md +16 -0
  13. package/dist/sdk.d.ts +7 -2
  14. package/dist/sdk.js +1349 -535
  15. package/dist/sdk.js.map +1 -1
  16. package/dist/types.d.ts +49 -1
  17. package/dist/types.js.map +1 -1
  18. package/docs/1.2.2-MEMORY-CONTROL-PLANE.md +434 -0
  19. package/docs/AGENT_OPERATOR_PLAYBOOK.md +4 -0
  20. package/docs/API_REFERENCE.md +24 -4
  21. package/docs/DESIGN_DECISIONS.md +357 -357
  22. package/docs/README.md +1 -1
  23. package/docs/dev-log/progress.txt +91 -11
  24. package/package.json +1 -1
  25. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  26. package/src/audit/index.ts +156 -156
  27. package/src/cli/command-guide.ts +192 -0
  28. package/src/cli/commands/audit-list.ts +89 -89
  29. package/src/cli/commands/audit.ts +9 -4
  30. package/src/cli/commands/background.ts +659 -659
  31. package/src/cli/commands/cleanup.ts +5 -1
  32. package/src/cli/commands/codegraph.ts +15 -5
  33. package/src/cli/commands/context.ts +3 -2
  34. package/src/cli/commands/doctor.ts +4 -2
  35. package/src/cli/commands/explain.ts +9 -3
  36. package/src/cli/commands/formation.ts +48 -48
  37. package/src/cli/commands/git-hook-install.ts +111 -111
  38. package/src/cli/commands/handoff.ts +75 -61
  39. package/src/cli/commands/hooks-status.ts +63 -63
  40. package/src/cli/commands/identity.ts +116 -0
  41. package/src/cli/commands/ingest-commit.ts +153 -153
  42. package/src/cli/commands/ingest-image.ts +71 -69
  43. package/src/cli/commands/ingest-log.ts +180 -180
  44. package/src/cli/commands/ingest.ts +44 -44
  45. package/src/cli/commands/integrate-shared.ts +15 -15
  46. package/src/cli/commands/lock.ts +93 -92
  47. package/src/cli/commands/memory.ts +58 -21
  48. package/src/cli/commands/message.ts +123 -118
  49. package/src/cli/commands/operator-shared.ts +98 -3
  50. package/src/cli/commands/poll.ts +74 -64
  51. package/src/cli/commands/purge-all-memory.ts +85 -85
  52. package/src/cli/commands/purge-project-memory.ts +83 -83
  53. package/src/cli/commands/reasoning.ts +135 -121
  54. package/src/cli/commands/retention.ts +9 -4
  55. package/src/cli/commands/serve-http.ts +8 -2
  56. package/src/cli/commands/serve-shared.ts +118 -118
  57. package/src/cli/commands/session.ts +29 -3
  58. package/src/cli/commands/skills.ts +124 -119
  59. package/src/cli/commands/status.ts +4 -3
  60. package/src/cli/commands/task.ts +193 -184
  61. package/src/cli/commands/team.ts +14 -10
  62. package/src/cli/commands/transfer.ts +108 -55
  63. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  64. package/src/cli/identity.ts +89 -0
  65. package/src/cli/index.ts +96 -19
  66. package/src/cli/invocation.ts +115 -0
  67. package/src/cli/tui/ChatView.tsx +234 -234
  68. package/src/cli/tui/CommandBar.tsx +312 -312
  69. package/src/cli/tui/ContextRail.tsx +118 -118
  70. package/src/cli/tui/HeaderBar.tsx +72 -72
  71. package/src/cli/tui/LogoBanner.tsx +51 -51
  72. package/src/cli/tui/Sidebar.tsx +179 -179
  73. package/src/cli/tui/chat-service.ts +41 -18
  74. package/src/cli/tui/data.ts +23 -44
  75. package/src/cli/tui/index.ts +41 -41
  76. package/src/cli/tui/markdown-render.tsx +371 -371
  77. package/src/cli/tui/operator-context.ts +60 -0
  78. package/src/cli/tui/use-mouse.ts +157 -157
  79. package/src/cli/tui/useNavigation.ts +56 -56
  80. package/src/cli/tui/views/MemoryView.tsx +10 -8
  81. package/src/cli/update-checker.ts +211 -211
  82. package/src/cli/version.ts +7 -7
  83. package/src/cli/workbench.ts +1 -1
  84. package/src/codegraph/auto-context.ts +31 -2
  85. package/src/codegraph/context-pack.ts +1 -0
  86. package/src/codegraph/project-context.ts +2 -0
  87. package/src/compact/engine.ts +26 -10
  88. package/src/compact/index-format.ts +25 -2
  89. package/src/compact/token-budget.ts +74 -74
  90. package/src/dashboard/project-classification.ts +64 -64
  91. package/src/dashboard/server.ts +46 -9
  92. package/src/embedding/fastembed-provider.ts +142 -142
  93. package/src/embedding/transformers-provider.ts +111 -111
  94. package/src/git/extractor.ts +209 -209
  95. package/src/git/hooks-path.ts +85 -85
  96. package/src/hooks/admission.ts +117 -0
  97. package/src/hooks/handler.ts +98 -91
  98. package/src/hooks/pattern-detector.ts +173 -173
  99. package/src/hooks/significance-filter.ts +250 -250
  100. package/src/knowledge/context-assembly.ts +97 -0
  101. package/src/knowledge/workset.ts +179 -10
  102. package/src/llm/memory-manager.ts +328 -328
  103. package/src/llm/provider.ts +885 -885
  104. package/src/llm/quality.ts +248 -248
  105. package/src/memory/admission.ts +57 -0
  106. package/src/memory/attribution-guard.ts +249 -249
  107. package/src/memory/consolidation.ts +13 -2
  108. package/src/memory/disclosure-policy.ts +140 -135
  109. package/src/memory/entity-extractor.ts +197 -197
  110. package/src/memory/export-import.ts +11 -3
  111. package/src/memory/formation/evaluate.ts +217 -217
  112. package/src/memory/formation/extract.ts +361 -361
  113. package/src/memory/formation/index.ts +417 -417
  114. package/src/memory/formation/resolve.ts +344 -344
  115. package/src/memory/formation/types.ts +315 -315
  116. package/src/memory/freshness.ts +122 -122
  117. package/src/memory/graph-context.ts +8 -2
  118. package/src/memory/graph.ts +197 -197
  119. package/src/memory/observations.ts +162 -4
  120. package/src/memory/quality-audit.ts +2 -0
  121. package/src/memory/refs.ts +94 -94
  122. package/src/memory/retention.ts +22 -2
  123. package/src/memory/secret-filter.ts +79 -79
  124. package/src/memory/session.ts +5 -2
  125. package/src/memory/visibility.ts +80 -0
  126. package/src/multimodal/image-loader.ts +143 -143
  127. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  128. package/src/orchestrate/adapters/claude.ts +111 -111
  129. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  130. package/src/orchestrate/adapters/codex.ts +41 -41
  131. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  132. package/src/orchestrate/adapters/gemini.ts +42 -42
  133. package/src/orchestrate/adapters/index.ts +73 -73
  134. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  135. package/src/orchestrate/adapters/opencode.ts +47 -47
  136. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  137. package/src/orchestrate/adapters/types.ts +77 -77
  138. package/src/orchestrate/capability-router.ts +284 -284
  139. package/src/orchestrate/context-compact.ts +188 -188
  140. package/src/orchestrate/cost-tracker.ts +219 -219
  141. package/src/orchestrate/error-recovery.ts +191 -191
  142. package/src/orchestrate/evidence.ts +140 -140
  143. package/src/orchestrate/ledger.ts +110 -110
  144. package/src/orchestrate/memorix-bridge.ts +378 -340
  145. package/src/orchestrate/output-budget.ts +80 -80
  146. package/src/orchestrate/permission.ts +152 -152
  147. package/src/orchestrate/pipeline-trace.ts +131 -131
  148. package/src/orchestrate/prompt-builder.ts +155 -155
  149. package/src/orchestrate/ring-buffer.ts +37 -37
  150. package/src/orchestrate/task-graph.ts +389 -389
  151. package/src/orchestrate/worktree.ts +232 -232
  152. package/src/project/aliases.ts +374 -374
  153. package/src/project/detector.ts +268 -268
  154. package/src/rules/adapters/claude-code.ts +99 -99
  155. package/src/rules/adapters/codex.ts +97 -97
  156. package/src/rules/adapters/copilot.ts +124 -124
  157. package/src/rules/adapters/cursor.ts +114 -114
  158. package/src/rules/adapters/kiro.ts +126 -126
  159. package/src/rules/adapters/trae.ts +56 -56
  160. package/src/rules/adapters/windsurf.ts +83 -83
  161. package/src/rules/syncer.ts +235 -235
  162. package/src/runtime/control-plane-maintenance.ts +1 -0
  163. package/src/runtime/isolated-maintenance.ts +1 -0
  164. package/src/runtime/lifecycle.ts +18 -0
  165. package/src/runtime/maintenance-jobs.ts +1 -0
  166. package/src/runtime/maintenance-runner.ts +2 -0
  167. package/src/runtime/project-maintenance.ts +89 -0
  168. package/src/sdk.ts +334 -304
  169. package/src/search/intent-detector.ts +289 -289
  170. package/src/search/query-expansion.ts +52 -52
  171. package/src/server/formation-timeout.ts +27 -27
  172. package/src/server.ts +260 -81
  173. package/src/skills/mini-skills.ts +386 -386
  174. package/src/store/chat-store.ts +119 -119
  175. package/src/store/graph-store.ts +249 -249
  176. package/src/store/mini-skill-store.ts +349 -349
  177. package/src/store/orama-store.ts +61 -6
  178. package/src/store/persistence-json.ts +212 -212
  179. package/src/store/persistence.ts +291 -291
  180. package/src/store/project-affinity.ts +195 -195
  181. package/src/store/sqlite-db.ts +23 -1
  182. package/src/store/sqlite-store.ts +12 -2
  183. package/src/team/event-bus.ts +76 -76
  184. package/src/team/file-locks.ts +173 -173
  185. package/src/team/handoff.ts +168 -161
  186. package/src/team/messages.ts +203 -203
  187. package/src/team/poll.ts +132 -132
  188. package/src/team/tasks.ts +211 -211
  189. package/src/types.ts +51 -0
  190. package/src/wiki/generator.ts +2 -0
  191. package/src/workspace/mcp-adapters/codex.ts +191 -191
  192. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  193. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  194. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  195. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  196. package/src/workspace/mcp-adapters/trae.ts +134 -134
  197. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  198. package/src/workspace/sanitizer.ts +60 -60
  199. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,142 +1,142 @@
1
- /**
2
- * FastEmbed Provider
3
- *
4
- * Local ONNX-based embedding using fastembed (Qdrant).
5
- * Model: BAAI/bge-small-en-v1.5 (384 dimensions, ~30MB)
6
- *
7
- * This is an optional dependency — if fastembed is not installed,
8
- * the provider module gracefully falls back to fulltext-only search.
9
- *
10
- * Persistent disk cache: embeddings are saved to ~/.memorix/data/.embedding-cache.json
11
- * so server restarts don't need to regenerate them (saves minutes of CPU on 500+ obs).
12
- */
13
-
14
- import { createHash } from 'node:crypto';
15
- import { readFile, writeFile, mkdir } from 'node:fs/promises';
16
- import { join } from 'node:path';
17
- import { homedir } from 'node:os';
18
- import type { EmbeddingProvider } from './provider.js';
19
-
20
- const CACHE_DIR = process.env.MEMORIX_DATA_DIR || join(homedir(), '.memorix', 'data');
21
- const CACHE_FILE = join(CACHE_DIR, '.embedding-cache.json');
22
-
23
- // In-memory cache keyed by text hash → embedding
24
- const cache = new Map<string, number[]>();
25
- const MAX_CACHE_SIZE = 5000;
26
- let diskCacheDirty = false;
27
-
28
- function textHash(text: string): string {
29
- return createHash('sha256').update(text).digest('hex').slice(0, 16);
30
- }
31
-
32
- async function loadDiskCache(): Promise<void> {
33
- try {
34
- const raw = await readFile(CACHE_FILE, 'utf-8');
35
- const entries: [string, number[]][] = JSON.parse(raw);
36
- for (const [k, v] of entries) cache.set(k, v);
37
- console.error(`[memorix] Loaded ${entries.length} cached embeddings from disk`);
38
- } catch {
39
- // No cache file or corrupt — start fresh
40
- }
41
- }
42
-
43
- async function saveDiskCache(): Promise<void> {
44
- if (!diskCacheDirty) return;
45
- try {
46
- await mkdir(CACHE_DIR, { recursive: true });
47
- const entries = Array.from(cache.entries());
48
- await writeFile(CACHE_FILE, JSON.stringify(entries));
49
- diskCacheDirty = false;
50
- } catch {
51
- // Ignore write errors — cache is an optimization, not critical
52
- }
53
- }
54
-
55
- export class FastEmbedProvider implements EmbeddingProvider {
56
- readonly name = 'fastembed-bge-small';
57
- readonly dimensions = 384;
58
-
59
- private model: { embed: (docs: string[], batchSize?: number) => AsyncGenerator<number[][]>; queryEmbed: (query: string) => Promise<number[]> };
60
-
61
- private constructor(model: FastEmbedProvider['model']) {
62
- this.model = model;
63
- }
64
-
65
- /**
66
- * Initialize the FastEmbed provider.
67
- * Downloads model on first use (~30MB), cached locally after.
68
- * Loads persistent embedding cache from disk.
69
- */
70
- static async create(): Promise<FastEmbedProvider> {
71
- // Dynamic import — throws if fastembed is not installed
72
- const { EmbeddingModel, FlagEmbedding } = await import('fastembed');
73
- const model = await FlagEmbedding.init({
74
- model: EmbeddingModel.BGESmallENV15,
75
- });
76
- // Load disk cache before returning — subsequent embedBatch calls will hit cache
77
- await loadDiskCache();
78
- return new FastEmbedProvider(model);
79
- }
80
-
81
- async embed(text: string): Promise<number[]> {
82
- const hash = textHash(text);
83
- const cached = cache.get(hash);
84
- if (cached) return cached;
85
-
86
- const raw = await this.model.queryEmbed(text);
87
- // Ensure plain number[] (fastembed may return Float32Array)
88
- const result = Array.from(raw) as number[];
89
- if (result.length !== this.dimensions) {
90
- throw new Error(`Expected ${this.dimensions}d embedding, got ${result.length}d`);
91
- }
92
- this.cacheSet(hash, result);
93
- return result;
94
- }
95
-
96
- async embedBatch(texts: string[]): Promise<number[][]> {
97
- const results: number[][] = new Array(texts.length);
98
- const uncachedIndices: number[] = [];
99
- const uncachedTexts: string[] = [];
100
-
101
- // Check cache for each text (by hash)
102
- for (let i = 0; i < texts.length; i++) {
103
- const hash = textHash(texts[i]);
104
- const cached = cache.get(hash);
105
- if (cached) {
106
- results[i] = cached;
107
- } else {
108
- uncachedIndices.push(i);
109
- uncachedTexts.push(texts[i]);
110
- }
111
- }
112
-
113
- // Batch embed uncached texts
114
- if (uncachedTexts.length > 0) {
115
- console.error(`[memorix] Embedding ${uncachedTexts.length}/${texts.length} uncached texts (${texts.length - uncachedTexts.length} from cache)`);
116
- let batchIdx = 0;
117
- for await (const batch of this.model.embed(uncachedTexts, 64)) {
118
- for (const vec of batch) {
119
- const originalIdx = uncachedIndices[batchIdx];
120
- const plain = Array.from(vec) as number[];
121
- results[originalIdx] = plain;
122
- this.cacheSet(textHash(uncachedTexts[batchIdx]), plain);
123
- batchIdx++;
124
- }
125
- }
126
- // Persist cache to disk after batch operations
127
- await saveDiskCache();
128
- }
129
-
130
- return results;
131
- }
132
-
133
- private cacheSet(hash: string, value: number[]): void {
134
- // Evict oldest entries if cache is full
135
- if (cache.size >= MAX_CACHE_SIZE) {
136
- const firstKey = cache.keys().next().value;
137
- if (firstKey !== undefined) cache.delete(firstKey);
138
- }
139
- cache.set(hash, value);
140
- diskCacheDirty = true;
141
- }
142
- }
1
+ /**
2
+ * FastEmbed Provider
3
+ *
4
+ * Local ONNX-based embedding using fastembed (Qdrant).
5
+ * Model: BAAI/bge-small-en-v1.5 (384 dimensions, ~30MB)
6
+ *
7
+ * This is an optional dependency — if fastembed is not installed,
8
+ * the provider module gracefully falls back to fulltext-only search.
9
+ *
10
+ * Persistent disk cache: embeddings are saved to ~/.memorix/data/.embedding-cache.json
11
+ * so server restarts don't need to regenerate them (saves minutes of CPU on 500+ obs).
12
+ */
13
+
14
+ import { createHash } from 'node:crypto';
15
+ import { readFile, writeFile, mkdir } from 'node:fs/promises';
16
+ import { join } from 'node:path';
17
+ import { homedir } from 'node:os';
18
+ import type { EmbeddingProvider } from './provider.js';
19
+
20
+ const CACHE_DIR = process.env.MEMORIX_DATA_DIR || join(homedir(), '.memorix', 'data');
21
+ const CACHE_FILE = join(CACHE_DIR, '.embedding-cache.json');
22
+
23
+ // In-memory cache keyed by text hash → embedding
24
+ const cache = new Map<string, number[]>();
25
+ const MAX_CACHE_SIZE = 5000;
26
+ let diskCacheDirty = false;
27
+
28
+ function textHash(text: string): string {
29
+ return createHash('sha256').update(text).digest('hex').slice(0, 16);
30
+ }
31
+
32
+ async function loadDiskCache(): Promise<void> {
33
+ try {
34
+ const raw = await readFile(CACHE_FILE, 'utf-8');
35
+ const entries: [string, number[]][] = JSON.parse(raw);
36
+ for (const [k, v] of entries) cache.set(k, v);
37
+ console.error(`[memorix] Loaded ${entries.length} cached embeddings from disk`);
38
+ } catch {
39
+ // No cache file or corrupt — start fresh
40
+ }
41
+ }
42
+
43
+ async function saveDiskCache(): Promise<void> {
44
+ if (!diskCacheDirty) return;
45
+ try {
46
+ await mkdir(CACHE_DIR, { recursive: true });
47
+ const entries = Array.from(cache.entries());
48
+ await writeFile(CACHE_FILE, JSON.stringify(entries));
49
+ diskCacheDirty = false;
50
+ } catch {
51
+ // Ignore write errors — cache is an optimization, not critical
52
+ }
53
+ }
54
+
55
+ export class FastEmbedProvider implements EmbeddingProvider {
56
+ readonly name = 'fastembed-bge-small';
57
+ readonly dimensions = 384;
58
+
59
+ private model: { embed: (docs: string[], batchSize?: number) => AsyncGenerator<number[][]>; queryEmbed: (query: string) => Promise<number[]> };
60
+
61
+ private constructor(model: FastEmbedProvider['model']) {
62
+ this.model = model;
63
+ }
64
+
65
+ /**
66
+ * Initialize the FastEmbed provider.
67
+ * Downloads model on first use (~30MB), cached locally after.
68
+ * Loads persistent embedding cache from disk.
69
+ */
70
+ static async create(): Promise<FastEmbedProvider> {
71
+ // Dynamic import — throws if fastembed is not installed
72
+ const { EmbeddingModel, FlagEmbedding } = await import('fastembed');
73
+ const model = await FlagEmbedding.init({
74
+ model: EmbeddingModel.BGESmallENV15,
75
+ });
76
+ // Load disk cache before returning — subsequent embedBatch calls will hit cache
77
+ await loadDiskCache();
78
+ return new FastEmbedProvider(model);
79
+ }
80
+
81
+ async embed(text: string): Promise<number[]> {
82
+ const hash = textHash(text);
83
+ const cached = cache.get(hash);
84
+ if (cached) return cached;
85
+
86
+ const raw = await this.model.queryEmbed(text);
87
+ // Ensure plain number[] (fastembed may return Float32Array)
88
+ const result = Array.from(raw) as number[];
89
+ if (result.length !== this.dimensions) {
90
+ throw new Error(`Expected ${this.dimensions}d embedding, got ${result.length}d`);
91
+ }
92
+ this.cacheSet(hash, result);
93
+ return result;
94
+ }
95
+
96
+ async embedBatch(texts: string[]): Promise<number[][]> {
97
+ const results: number[][] = new Array(texts.length);
98
+ const uncachedIndices: number[] = [];
99
+ const uncachedTexts: string[] = [];
100
+
101
+ // Check cache for each text (by hash)
102
+ for (let i = 0; i < texts.length; i++) {
103
+ const hash = textHash(texts[i]);
104
+ const cached = cache.get(hash);
105
+ if (cached) {
106
+ results[i] = cached;
107
+ } else {
108
+ uncachedIndices.push(i);
109
+ uncachedTexts.push(texts[i]);
110
+ }
111
+ }
112
+
113
+ // Batch embed uncached texts
114
+ if (uncachedTexts.length > 0) {
115
+ console.error(`[memorix] Embedding ${uncachedTexts.length}/${texts.length} uncached texts (${texts.length - uncachedTexts.length} from cache)`);
116
+ let batchIdx = 0;
117
+ for await (const batch of this.model.embed(uncachedTexts, 64)) {
118
+ for (const vec of batch) {
119
+ const originalIdx = uncachedIndices[batchIdx];
120
+ const plain = Array.from(vec) as number[];
121
+ results[originalIdx] = plain;
122
+ this.cacheSet(textHash(uncachedTexts[batchIdx]), plain);
123
+ batchIdx++;
124
+ }
125
+ }
126
+ // Persist cache to disk after batch operations
127
+ await saveDiskCache();
128
+ }
129
+
130
+ return results;
131
+ }
132
+
133
+ private cacheSet(hash: string, value: number[]): void {
134
+ // Evict oldest entries if cache is full
135
+ if (cache.size >= MAX_CACHE_SIZE) {
136
+ const firstKey = cache.keys().next().value;
137
+ if (firstKey !== undefined) cache.delete(firstKey);
138
+ }
139
+ cache.set(hash, value);
140
+ diskCacheDirty = true;
141
+ }
142
+ }
@@ -1,111 +1,111 @@
1
- /**
2
- * Transformers.js Provider
3
- *
4
- * Pure JavaScript embedding using @huggingface/transformers (HuggingFace).
5
- * Model: Xenova/all-MiniLM-L6-v2 (384 dimensions, ~22MB quantized)
6
- *
7
- * Key advantages over fastembed:
8
- * - No native ONNX binding required (pure JS / WASM)
9
- * - Works out-of-the-box on Windows, macOS, Linux
10
- * - Supports quantized models (q8, q4) for smaller footprint
11
- *
12
- * This is an optional dependency — if @huggingface/transformers is not
13
- * installed, the provider module gracefully falls back to the next option.
14
- *
15
- * Inspired by Mem0's multi-provider embedding architecture.
16
- */
17
-
18
- import type { EmbeddingProvider } from './provider.js';
19
-
20
- // In-memory LRU cache
21
- const cache = new Map<string, number[]>();
22
- const MAX_CACHE_SIZE = 5000;
23
-
24
- export class TransformersProvider implements EmbeddingProvider {
25
- readonly name = 'transformers-minilm';
26
- readonly dimensions = 384;
27
-
28
- private extractor: any; // Pipeline instance
29
-
30
- private constructor(extractor: any) {
31
- this.extractor = extractor;
32
- }
33
-
34
- /**
35
- * Initialize the Transformers.js provider.
36
- * Downloads model on first use (~22MB quantized), cached locally after.
37
- */
38
- static async create(): Promise<TransformersProvider> {
39
- // Dynamic import — throws if @huggingface/transformers is not installed
40
- const { pipeline } = await import('@huggingface/transformers');
41
- const extractor = await pipeline(
42
- 'feature-extraction',
43
- 'Xenova/all-MiniLM-L6-v2',
44
- { dtype: 'q8' }, // Quantized for small footprint
45
- );
46
- return new TransformersProvider(extractor);
47
- }
48
-
49
- async embed(text: string): Promise<number[]> {
50
- // Check cache first
51
- const cached = cache.get(text);
52
- if (cached) return cached;
53
-
54
- const output = await this.extractor(text, {
55
- pooling: 'mean',
56
- normalize: true,
57
- });
58
-
59
- // output.tolist() returns [[...384 floats]]
60
- const result: number[] = Array.from(output.tolist()[0]);
61
- if (result.length !== this.dimensions) {
62
- throw new Error(`Expected ${this.dimensions}d embedding, got ${result.length}d`);
63
- }
64
-
65
- this.cacheSet(text, result);
66
- return result;
67
- }
68
-
69
- async embedBatch(texts: string[]): Promise<number[][]> {
70
- const results: number[][] = new Array(texts.length);
71
- const uncachedIndices: number[] = [];
72
- const uncachedTexts: string[] = [];
73
-
74
- // Check cache for each text
75
- for (let i = 0; i < texts.length; i++) {
76
- const cached = cache.get(texts[i]);
77
- if (cached) {
78
- results[i] = cached;
79
- } else {
80
- uncachedIndices.push(i);
81
- uncachedTexts.push(texts[i]);
82
- }
83
- }
84
-
85
- // Batch embed uncached texts
86
- if (uncachedTexts.length > 0) {
87
- const output = await this.extractor(uncachedTexts, {
88
- pooling: 'mean',
89
- normalize: true,
90
- });
91
- const allVecs: number[][] = output.tolist();
92
-
93
- for (let i = 0; i < allVecs.length; i++) {
94
- const vec = Array.from(allVecs[i]) as number[];
95
- const originalIdx = uncachedIndices[i];
96
- results[originalIdx] = vec;
97
- this.cacheSet(uncachedTexts[i], vec);
98
- }
99
- }
100
-
101
- return results;
102
- }
103
-
104
- private cacheSet(key: string, value: number[]): void {
105
- if (cache.size >= MAX_CACHE_SIZE) {
106
- const firstKey = cache.keys().next().value;
107
- if (firstKey !== undefined) cache.delete(firstKey);
108
- }
109
- cache.set(key, value);
110
- }
111
- }
1
+ /**
2
+ * Transformers.js Provider
3
+ *
4
+ * Pure JavaScript embedding using @huggingface/transformers (HuggingFace).
5
+ * Model: Xenova/all-MiniLM-L6-v2 (384 dimensions, ~22MB quantized)
6
+ *
7
+ * Key advantages over fastembed:
8
+ * - No native ONNX binding required (pure JS / WASM)
9
+ * - Works out-of-the-box on Windows, macOS, Linux
10
+ * - Supports quantized models (q8, q4) for smaller footprint
11
+ *
12
+ * This is an optional dependency — if @huggingface/transformers is not
13
+ * installed, the provider module gracefully falls back to the next option.
14
+ *
15
+ * Inspired by Mem0's multi-provider embedding architecture.
16
+ */
17
+
18
+ import type { EmbeddingProvider } from './provider.js';
19
+
20
+ // In-memory LRU cache
21
+ const cache = new Map<string, number[]>();
22
+ const MAX_CACHE_SIZE = 5000;
23
+
24
+ export class TransformersProvider implements EmbeddingProvider {
25
+ readonly name = 'transformers-minilm';
26
+ readonly dimensions = 384;
27
+
28
+ private extractor: any; // Pipeline instance
29
+
30
+ private constructor(extractor: any) {
31
+ this.extractor = extractor;
32
+ }
33
+
34
+ /**
35
+ * Initialize the Transformers.js provider.
36
+ * Downloads model on first use (~22MB quantized), cached locally after.
37
+ */
38
+ static async create(): Promise<TransformersProvider> {
39
+ // Dynamic import — throws if @huggingface/transformers is not installed
40
+ const { pipeline } = await import('@huggingface/transformers');
41
+ const extractor = await pipeline(
42
+ 'feature-extraction',
43
+ 'Xenova/all-MiniLM-L6-v2',
44
+ { dtype: 'q8' }, // Quantized for small footprint
45
+ );
46
+ return new TransformersProvider(extractor);
47
+ }
48
+
49
+ async embed(text: string): Promise<number[]> {
50
+ // Check cache first
51
+ const cached = cache.get(text);
52
+ if (cached) return cached;
53
+
54
+ const output = await this.extractor(text, {
55
+ pooling: 'mean',
56
+ normalize: true,
57
+ });
58
+
59
+ // output.tolist() returns [[...384 floats]]
60
+ const result: number[] = Array.from(output.tolist()[0]);
61
+ if (result.length !== this.dimensions) {
62
+ throw new Error(`Expected ${this.dimensions}d embedding, got ${result.length}d`);
63
+ }
64
+
65
+ this.cacheSet(text, result);
66
+ return result;
67
+ }
68
+
69
+ async embedBatch(texts: string[]): Promise<number[][]> {
70
+ const results: number[][] = new Array(texts.length);
71
+ const uncachedIndices: number[] = [];
72
+ const uncachedTexts: string[] = [];
73
+
74
+ // Check cache for each text
75
+ for (let i = 0; i < texts.length; i++) {
76
+ const cached = cache.get(texts[i]);
77
+ if (cached) {
78
+ results[i] = cached;
79
+ } else {
80
+ uncachedIndices.push(i);
81
+ uncachedTexts.push(texts[i]);
82
+ }
83
+ }
84
+
85
+ // Batch embed uncached texts
86
+ if (uncachedTexts.length > 0) {
87
+ const output = await this.extractor(uncachedTexts, {
88
+ pooling: 'mean',
89
+ normalize: true,
90
+ });
91
+ const allVecs: number[][] = output.tolist();
92
+
93
+ for (let i = 0; i < allVecs.length; i++) {
94
+ const vec = Array.from(allVecs[i]) as number[];
95
+ const originalIdx = uncachedIndices[i];
96
+ results[originalIdx] = vec;
97
+ this.cacheSet(uncachedTexts[i], vec);
98
+ }
99
+ }
100
+
101
+ return results;
102
+ }
103
+
104
+ private cacheSet(key: string, value: number[]): void {
105
+ if (cache.size >= MAX_CACHE_SIZE) {
106
+ const firstKey = cache.keys().next().value;
107
+ if (firstKey !== undefined) cache.delete(firstKey);
108
+ }
109
+ cache.set(key, value);
110
+ }
111
+ }