@livekit/agents 0.7.8 → 1.0.0-next.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (627) hide show
  1. package/dist/_exceptions.cjs +109 -0
  2. package/dist/_exceptions.cjs.map +1 -0
  3. package/dist/_exceptions.d.cts +64 -0
  4. package/dist/_exceptions.d.ts +64 -0
  5. package/dist/_exceptions.d.ts.map +1 -0
  6. package/dist/_exceptions.js +80 -0
  7. package/dist/_exceptions.js.map +1 -0
  8. package/dist/audio.cjs +10 -3
  9. package/dist/audio.cjs.map +1 -1
  10. package/dist/audio.d.cts +2 -0
  11. package/dist/audio.d.ts +2 -0
  12. package/dist/audio.d.ts.map +1 -1
  13. package/dist/audio.js +8 -2
  14. package/dist/audio.js.map +1 -1
  15. package/dist/cli.cjs +25 -0
  16. package/dist/cli.cjs.map +1 -1
  17. package/dist/cli.d.ts.map +1 -1
  18. package/dist/cli.js +25 -0
  19. package/dist/cli.js.map +1 -1
  20. package/dist/constants.cjs +6 -0
  21. package/dist/constants.cjs.map +1 -1
  22. package/dist/constants.d.cts +2 -0
  23. package/dist/constants.d.ts +2 -0
  24. package/dist/constants.d.ts.map +1 -1
  25. package/dist/constants.js +4 -0
  26. package/dist/constants.js.map +1 -1
  27. package/dist/http_server.cjs.map +1 -1
  28. package/dist/http_server.d.cts +1 -0
  29. package/dist/http_server.d.ts +1 -0
  30. package/dist/http_server.d.ts.map +1 -1
  31. package/dist/http_server.js.map +1 -1
  32. package/dist/index.cjs +27 -20
  33. package/dist/index.cjs.map +1 -1
  34. package/dist/index.d.cts +13 -10
  35. package/dist/index.d.ts +13 -10
  36. package/dist/index.d.ts.map +1 -1
  37. package/dist/index.js +15 -11
  38. package/dist/index.js.map +1 -1
  39. package/dist/inference_runner.cjs +0 -1
  40. package/dist/inference_runner.cjs.map +1 -1
  41. package/dist/inference_runner.d.cts +2 -3
  42. package/dist/inference_runner.d.ts +2 -3
  43. package/dist/inference_runner.d.ts.map +1 -1
  44. package/dist/inference_runner.js +0 -1
  45. package/dist/inference_runner.js.map +1 -1
  46. package/dist/ipc/inference_proc_executor.cjs +2 -2
  47. package/dist/ipc/inference_proc_executor.cjs.map +1 -1
  48. package/dist/ipc/inference_proc_executor.js +2 -2
  49. package/dist/ipc/inference_proc_executor.js.map +1 -1
  50. package/dist/ipc/job_executor.cjs.map +1 -1
  51. package/dist/ipc/job_executor.js.map +1 -1
  52. package/dist/ipc/job_proc_executor.cjs +1 -0
  53. package/dist/ipc/job_proc_executor.cjs.map +1 -1
  54. package/dist/ipc/job_proc_executor.js +1 -0
  55. package/dist/ipc/job_proc_executor.js.map +1 -1
  56. package/dist/ipc/job_proc_lazy_main.cjs +1 -1
  57. package/dist/ipc/job_proc_lazy_main.cjs.map +1 -1
  58. package/dist/ipc/job_proc_lazy_main.js +1 -1
  59. package/dist/ipc/job_proc_lazy_main.js.map +1 -1
  60. package/dist/ipc/supervised_proc.d.cts +1 -1
  61. package/dist/ipc/supervised_proc.d.ts +1 -1
  62. package/dist/ipc/supervised_proc.d.ts.map +1 -1
  63. package/dist/job.cjs +14 -2
  64. package/dist/job.cjs.map +1 -1
  65. package/dist/job.d.cts +8 -0
  66. package/dist/job.d.ts +8 -0
  67. package/dist/job.d.ts.map +1 -1
  68. package/dist/job.js +12 -1
  69. package/dist/job.js.map +1 -1
  70. package/dist/llm/chat_context.cjs +332 -82
  71. package/dist/llm/chat_context.cjs.map +1 -1
  72. package/dist/llm/chat_context.d.cts +152 -48
  73. package/dist/llm/chat_context.d.ts +152 -48
  74. package/dist/llm/chat_context.d.ts.map +1 -1
  75. package/dist/llm/chat_context.js +327 -81
  76. package/dist/llm/chat_context.js.map +1 -1
  77. package/dist/llm/chat_context.test.cjs +380 -0
  78. package/dist/llm/chat_context.test.cjs.map +1 -0
  79. package/dist/llm/chat_context.test.js +385 -0
  80. package/dist/llm/chat_context.test.js.map +1 -0
  81. package/dist/llm/index.cjs +37 -8
  82. package/dist/llm/index.cjs.map +1 -1
  83. package/dist/llm/index.d.cts +7 -3
  84. package/dist/llm/index.d.ts +7 -3
  85. package/dist/llm/index.d.ts.map +1 -1
  86. package/dist/llm/index.js +39 -9
  87. package/dist/llm/index.js.map +1 -1
  88. package/dist/llm/llm.cjs +98 -33
  89. package/dist/llm/llm.cjs.map +1 -1
  90. package/dist/llm/llm.d.cts +50 -24
  91. package/dist/llm/llm.d.ts +50 -24
  92. package/dist/llm/llm.d.ts.map +1 -1
  93. package/dist/llm/llm.js +99 -33
  94. package/dist/llm/llm.js.map +1 -1
  95. package/dist/llm/provider_format/google.cjs +128 -0
  96. package/dist/llm/provider_format/google.cjs.map +1 -0
  97. package/dist/llm/provider_format/google.d.cts +6 -0
  98. package/dist/llm/provider_format/google.d.ts +6 -0
  99. package/dist/llm/provider_format/google.d.ts.map +1 -0
  100. package/dist/llm/provider_format/google.js +104 -0
  101. package/dist/llm/provider_format/google.js.map +1 -0
  102. package/dist/llm/provider_format/google.test.cjs +676 -0
  103. package/dist/llm/provider_format/google.test.cjs.map +1 -0
  104. package/dist/llm/provider_format/google.test.js +675 -0
  105. package/dist/llm/provider_format/google.test.js.map +1 -0
  106. package/dist/llm/provider_format/index.cjs +40 -0
  107. package/dist/llm/provider_format/index.cjs.map +1 -0
  108. package/dist/llm/provider_format/index.d.cts +4 -0
  109. package/dist/llm/provider_format/index.d.ts +4 -0
  110. package/dist/llm/provider_format/index.d.ts.map +1 -0
  111. package/dist/llm/provider_format/index.js +16 -0
  112. package/dist/llm/provider_format/index.js.map +1 -0
  113. package/dist/llm/provider_format/openai.cjs +116 -0
  114. package/dist/llm/provider_format/openai.cjs.map +1 -0
  115. package/dist/llm/provider_format/openai.d.cts +3 -0
  116. package/dist/llm/provider_format/openai.d.ts +3 -0
  117. package/dist/llm/provider_format/openai.d.ts.map +1 -0
  118. package/dist/llm/provider_format/openai.js +92 -0
  119. package/dist/llm/provider_format/openai.js.map +1 -0
  120. package/dist/llm/provider_format/openai.test.cjs +490 -0
  121. package/dist/llm/provider_format/openai.test.cjs.map +1 -0
  122. package/dist/llm/provider_format/openai.test.js +489 -0
  123. package/dist/llm/provider_format/openai.test.js.map +1 -0
  124. package/dist/llm/provider_format/utils.cjs +146 -0
  125. package/dist/llm/provider_format/utils.cjs.map +1 -0
  126. package/dist/llm/provider_format/utils.d.cts +38 -0
  127. package/dist/llm/provider_format/utils.d.ts +38 -0
  128. package/dist/llm/provider_format/utils.d.ts.map +1 -0
  129. package/dist/llm/provider_format/utils.js +122 -0
  130. package/dist/llm/provider_format/utils.js.map +1 -0
  131. package/dist/llm/realtime.cjs +77 -0
  132. package/dist/llm/realtime.cjs.map +1 -0
  133. package/dist/llm/realtime.d.cts +98 -0
  134. package/dist/llm/realtime.d.ts +98 -0
  135. package/dist/llm/realtime.d.ts.map +1 -0
  136. package/dist/llm/realtime.js +52 -0
  137. package/dist/llm/realtime.js.map +1 -0
  138. package/dist/llm/remote_chat_context.cjs +112 -0
  139. package/dist/llm/remote_chat_context.cjs.map +1 -0
  140. package/dist/llm/remote_chat_context.d.cts +23 -0
  141. package/dist/llm/remote_chat_context.d.ts +23 -0
  142. package/dist/llm/remote_chat_context.d.ts.map +1 -0
  143. package/dist/llm/remote_chat_context.js +88 -0
  144. package/dist/llm/remote_chat_context.js.map +1 -0
  145. package/dist/llm/remote_chat_context.test.cjs +225 -0
  146. package/dist/llm/remote_chat_context.test.cjs.map +1 -0
  147. package/dist/llm/remote_chat_context.test.js +224 -0
  148. package/dist/llm/remote_chat_context.test.js.map +1 -0
  149. package/dist/llm/tool_context.cjs +111 -0
  150. package/dist/llm/tool_context.cjs.map +1 -0
  151. package/dist/llm/tool_context.d.cts +125 -0
  152. package/dist/llm/tool_context.d.ts +125 -0
  153. package/dist/llm/tool_context.d.ts.map +1 -0
  154. package/dist/llm/tool_context.js +80 -0
  155. package/dist/llm/tool_context.js.map +1 -0
  156. package/dist/llm/tool_context.test.cjs +162 -0
  157. package/dist/llm/tool_context.test.cjs.map +1 -0
  158. package/dist/llm/tool_context.test.js +161 -0
  159. package/dist/llm/tool_context.test.js.map +1 -0
  160. package/dist/llm/tool_context.type.test.cjs +92 -0
  161. package/dist/llm/tool_context.type.test.cjs.map +1 -0
  162. package/dist/llm/tool_context.type.test.js +91 -0
  163. package/dist/llm/tool_context.type.test.js.map +1 -0
  164. package/dist/llm/utils.cjs +260 -0
  165. package/dist/llm/utils.cjs.map +1 -0
  166. package/dist/llm/utils.d.cts +42 -0
  167. package/dist/llm/utils.d.ts +42 -0
  168. package/dist/llm/utils.d.ts.map +1 -0
  169. package/dist/llm/utils.js +223 -0
  170. package/dist/llm/utils.js.map +1 -0
  171. package/dist/llm/utils.test.cjs +513 -0
  172. package/dist/llm/utils.test.cjs.map +1 -0
  173. package/dist/llm/utils.test.js +490 -0
  174. package/dist/llm/utils.test.js.map +1 -0
  175. package/dist/metrics/base.cjs +0 -27
  176. package/dist/metrics/base.cjs.map +1 -1
  177. package/dist/metrics/base.d.cts +105 -63
  178. package/dist/metrics/base.d.ts +105 -63
  179. package/dist/metrics/base.d.ts.map +1 -1
  180. package/dist/metrics/base.js +0 -19
  181. package/dist/metrics/base.js.map +1 -1
  182. package/dist/metrics/index.cjs +0 -3
  183. package/dist/metrics/index.cjs.map +1 -1
  184. package/dist/metrics/index.d.cts +2 -3
  185. package/dist/metrics/index.d.ts +2 -3
  186. package/dist/metrics/index.d.ts.map +1 -1
  187. package/dist/metrics/index.js +0 -2
  188. package/dist/metrics/index.js.map +1 -1
  189. package/dist/metrics/usage_collector.cjs +17 -12
  190. package/dist/metrics/usage_collector.cjs.map +1 -1
  191. package/dist/metrics/usage_collector.d.cts +3 -2
  192. package/dist/metrics/usage_collector.d.ts +3 -2
  193. package/dist/metrics/usage_collector.d.ts.map +1 -1
  194. package/dist/metrics/usage_collector.js +17 -12
  195. package/dist/metrics/usage_collector.js.map +1 -1
  196. package/dist/metrics/utils.cjs +22 -59
  197. package/dist/metrics/utils.cjs.map +1 -1
  198. package/dist/metrics/utils.d.cts +1 -8
  199. package/dist/metrics/utils.d.ts +1 -8
  200. package/dist/metrics/utils.d.ts.map +1 -1
  201. package/dist/metrics/utils.js +22 -52
  202. package/dist/metrics/utils.js.map +1 -1
  203. package/dist/multimodal/index.cjs +0 -2
  204. package/dist/multimodal/index.cjs.map +1 -1
  205. package/dist/multimodal/index.d.cts +0 -1
  206. package/dist/multimodal/index.d.ts +0 -1
  207. package/dist/multimodal/index.d.ts.map +1 -1
  208. package/dist/multimodal/index.js +0 -1
  209. package/dist/multimodal/index.js.map +1 -1
  210. package/dist/plugin.cjs +24 -8
  211. package/dist/plugin.cjs.map +1 -1
  212. package/dist/plugin.d.cts +18 -4
  213. package/dist/plugin.d.ts +18 -4
  214. package/dist/plugin.d.ts.map +1 -1
  215. package/dist/plugin.js +22 -7
  216. package/dist/plugin.js.map +1 -1
  217. package/dist/stream/deferred_stream.cjs +98 -0
  218. package/dist/stream/deferred_stream.cjs.map +1 -0
  219. package/dist/stream/deferred_stream.d.cts +27 -0
  220. package/dist/stream/deferred_stream.d.ts +27 -0
  221. package/dist/stream/deferred_stream.d.ts.map +1 -0
  222. package/dist/stream/deferred_stream.js +73 -0
  223. package/dist/stream/deferred_stream.js.map +1 -0
  224. package/dist/stream/deferred_stream.test.cjs +527 -0
  225. package/dist/stream/deferred_stream.test.cjs.map +1 -0
  226. package/dist/stream/deferred_stream.test.js +526 -0
  227. package/dist/stream/deferred_stream.test.js.map +1 -0
  228. package/dist/stream/identity_transform.cjs +42 -0
  229. package/dist/stream/identity_transform.cjs.map +1 -0
  230. package/dist/stream/identity_transform.d.cts +6 -0
  231. package/dist/stream/identity_transform.d.ts +6 -0
  232. package/dist/stream/identity_transform.d.ts.map +1 -0
  233. package/dist/stream/identity_transform.js +18 -0
  234. package/dist/stream/identity_transform.js.map +1 -0
  235. package/dist/stream/identity_transform.test.cjs +125 -0
  236. package/dist/stream/identity_transform.test.cjs.map +1 -0
  237. package/dist/stream/identity_transform.test.js +124 -0
  238. package/dist/stream/identity_transform.test.js.map +1 -0
  239. package/dist/stream/index.cjs +38 -0
  240. package/dist/stream/index.cjs.map +1 -0
  241. package/dist/stream/index.d.cts +5 -0
  242. package/dist/stream/index.d.ts +5 -0
  243. package/dist/stream/index.d.ts.map +1 -0
  244. package/dist/stream/index.js +11 -0
  245. package/dist/stream/index.js.map +1 -0
  246. package/dist/stream/merge_readable_streams.cjs +59 -0
  247. package/dist/stream/merge_readable_streams.cjs.map +1 -0
  248. package/dist/stream/merge_readable_streams.d.cts +4 -0
  249. package/dist/stream/merge_readable_streams.d.ts +4 -0
  250. package/dist/stream/merge_readable_streams.d.ts.map +1 -0
  251. package/dist/stream/merge_readable_streams.js +35 -0
  252. package/dist/stream/merge_readable_streams.js.map +1 -0
  253. package/dist/stream/stream_channel.cjs +47 -0
  254. package/dist/stream/stream_channel.cjs.map +1 -0
  255. package/dist/stream/stream_channel.d.cts +9 -0
  256. package/dist/stream/stream_channel.d.ts +9 -0
  257. package/dist/stream/stream_channel.d.ts.map +1 -0
  258. package/dist/stream/stream_channel.js +23 -0
  259. package/dist/stream/stream_channel.js.map +1 -0
  260. package/dist/stream/stream_channel.test.cjs +97 -0
  261. package/dist/stream/stream_channel.test.cjs.map +1 -0
  262. package/dist/stream/stream_channel.test.js +96 -0
  263. package/dist/stream/stream_channel.test.js.map +1 -0
  264. package/dist/stt/stream_adapter.cjs +3 -4
  265. package/dist/stt/stream_adapter.cjs.map +1 -1
  266. package/dist/stt/stream_adapter.d.cts +1 -0
  267. package/dist/stt/stream_adapter.d.ts +1 -0
  268. package/dist/stt/stream_adapter.d.ts.map +1 -1
  269. package/dist/stt/stream_adapter.js +3 -4
  270. package/dist/stt/stream_adapter.js.map +1 -1
  271. package/dist/stt/stt.cjs +101 -10
  272. package/dist/stt/stt.cjs.map +1 -1
  273. package/dist/stt/stt.d.cts +26 -5
  274. package/dist/stt/stt.d.ts +26 -5
  275. package/dist/stt/stt.d.ts.map +1 -1
  276. package/dist/stt/stt.js +102 -11
  277. package/dist/stt/stt.js.map +1 -1
  278. package/dist/tokenize/basic/basic.cjs +10 -5
  279. package/dist/tokenize/basic/basic.cjs.map +1 -1
  280. package/dist/tokenize/basic/basic.d.cts +7 -1
  281. package/dist/tokenize/basic/basic.d.ts +7 -1
  282. package/dist/tokenize/basic/basic.d.ts.map +1 -1
  283. package/dist/tokenize/basic/basic.js +10 -5
  284. package/dist/tokenize/basic/basic.js.map +1 -1
  285. package/dist/tokenize/basic/sentence.cjs +14 -6
  286. package/dist/tokenize/basic/sentence.cjs.map +1 -1
  287. package/dist/tokenize/basic/sentence.d.cts +1 -1
  288. package/dist/tokenize/basic/sentence.d.ts +1 -1
  289. package/dist/tokenize/basic/sentence.d.ts.map +1 -1
  290. package/dist/tokenize/basic/sentence.js +14 -6
  291. package/dist/tokenize/basic/sentence.js.map +1 -1
  292. package/dist/tokenize/token_stream.cjs +5 -3
  293. package/dist/tokenize/token_stream.cjs.map +1 -1
  294. package/dist/tokenize/token_stream.d.cts +1 -0
  295. package/dist/tokenize/token_stream.d.ts +1 -0
  296. package/dist/tokenize/token_stream.d.ts.map +1 -1
  297. package/dist/tokenize/token_stream.js +6 -4
  298. package/dist/tokenize/token_stream.js.map +1 -1
  299. package/dist/transcription.cjs +1 -2
  300. package/dist/transcription.cjs.map +1 -1
  301. package/dist/transcription.d.ts.map +1 -1
  302. package/dist/transcription.js +2 -3
  303. package/dist/transcription.js.map +1 -1
  304. package/dist/tts/index.cjs +2 -4
  305. package/dist/tts/index.cjs.map +1 -1
  306. package/dist/tts/index.d.cts +1 -1
  307. package/dist/tts/index.d.ts +1 -1
  308. package/dist/tts/index.d.ts.map +1 -1
  309. package/dist/tts/index.js +1 -3
  310. package/dist/tts/index.js.map +1 -1
  311. package/dist/tts/stream_adapter.cjs +26 -13
  312. package/dist/tts/stream_adapter.cjs.map +1 -1
  313. package/dist/tts/stream_adapter.d.cts +1 -1
  314. package/dist/tts/stream_adapter.d.ts +1 -1
  315. package/dist/tts/stream_adapter.d.ts.map +1 -1
  316. package/dist/tts/stream_adapter.js +27 -14
  317. package/dist/tts/stream_adapter.js.map +1 -1
  318. package/dist/tts/tts.cjs +157 -25
  319. package/dist/tts/tts.cjs.map +1 -1
  320. package/dist/tts/tts.d.cts +29 -5
  321. package/dist/tts/tts.d.ts +29 -5
  322. package/dist/tts/tts.d.ts.map +1 -1
  323. package/dist/tts/tts.js +157 -24
  324. package/dist/tts/tts.js.map +1 -1
  325. package/dist/types.cjs +60 -0
  326. package/dist/types.cjs.map +1 -0
  327. package/dist/types.d.cts +13 -0
  328. package/dist/types.d.ts +13 -0
  329. package/dist/types.d.ts.map +1 -0
  330. package/dist/types.js +35 -0
  331. package/dist/types.js.map +1 -0
  332. package/dist/utils.cjs +281 -27
  333. package/dist/utils.cjs.map +1 -1
  334. package/dist/utils.d.cts +134 -9
  335. package/dist/utils.d.ts +134 -9
  336. package/dist/utils.d.ts.map +1 -1
  337. package/dist/utils.js +265 -26
  338. package/dist/utils.js.map +1 -1
  339. package/dist/utils.test.cjs +492 -0
  340. package/dist/utils.test.cjs.map +1 -0
  341. package/dist/utils.test.js +498 -0
  342. package/dist/utils.test.js.map +1 -0
  343. package/dist/vad.cjs +76 -20
  344. package/dist/vad.cjs.map +1 -1
  345. package/dist/vad.d.cts +25 -5
  346. package/dist/vad.d.ts +25 -5
  347. package/dist/vad.d.ts.map +1 -1
  348. package/dist/vad.js +76 -20
  349. package/dist/vad.js.map +1 -1
  350. package/dist/voice/agent.cjs +245 -0
  351. package/dist/voice/agent.cjs.map +1 -0
  352. package/dist/voice/agent.d.cts +78 -0
  353. package/dist/voice/agent.d.ts +78 -0
  354. package/dist/voice/agent.d.ts.map +1 -0
  355. package/dist/voice/agent.js +220 -0
  356. package/dist/voice/agent.js.map +1 -0
  357. package/dist/voice/agent.test.cjs +61 -0
  358. package/dist/voice/agent.test.cjs.map +1 -0
  359. package/dist/voice/agent.test.js +60 -0
  360. package/dist/voice/agent.test.js.map +1 -0
  361. package/dist/voice/agent_activity.cjs +1453 -0
  362. package/dist/voice/agent_activity.cjs.map +1 -0
  363. package/dist/voice/agent_activity.d.cts +94 -0
  364. package/dist/voice/agent_activity.d.ts +94 -0
  365. package/dist/voice/agent_activity.d.ts.map +1 -0
  366. package/dist/voice/agent_activity.js +1449 -0
  367. package/dist/voice/agent_activity.js.map +1 -0
  368. package/dist/voice/agent_session.cjs +312 -0
  369. package/dist/voice/agent_session.cjs.map +1 -0
  370. package/dist/voice/agent_session.d.cts +121 -0
  371. package/dist/voice/agent_session.d.ts +121 -0
  372. package/dist/voice/agent_session.d.ts.map +1 -0
  373. package/dist/voice/agent_session.js +295 -0
  374. package/dist/voice/agent_session.js.map +1 -0
  375. package/dist/voice/audio_recognition.cjs +375 -0
  376. package/dist/voice/audio_recognition.cjs.map +1 -0
  377. package/dist/voice/audio_recognition.d.cts +80 -0
  378. package/dist/voice/audio_recognition.d.ts +80 -0
  379. package/dist/voice/audio_recognition.d.ts.map +1 -0
  380. package/dist/voice/audio_recognition.js +351 -0
  381. package/dist/voice/audio_recognition.js.map +1 -0
  382. package/dist/voice/events.cjs +145 -0
  383. package/dist/voice/events.cjs.map +1 -0
  384. package/dist/voice/events.d.cts +124 -0
  385. package/dist/voice/events.d.ts +124 -0
  386. package/dist/voice/events.d.ts.map +1 -0
  387. package/dist/voice/events.js +110 -0
  388. package/dist/voice/events.js.map +1 -0
  389. package/dist/voice/generation.cjs +700 -0
  390. package/dist/voice/generation.cjs.map +1 -0
  391. package/dist/voice/generation.d.cts +115 -0
  392. package/dist/voice/generation.d.ts +115 -0
  393. package/dist/voice/generation.d.ts.map +1 -0
  394. package/dist/voice/generation.js +672 -0
  395. package/dist/voice/generation.js.map +1 -0
  396. package/dist/voice/index.cjs +40 -0
  397. package/dist/voice/index.cjs.map +1 -0
  398. package/dist/voice/index.d.cts +5 -0
  399. package/dist/voice/index.d.ts +5 -0
  400. package/dist/voice/index.d.ts.map +1 -0
  401. package/dist/voice/index.js +11 -0
  402. package/dist/voice/index.js.map +1 -0
  403. package/dist/voice/io.cjs +245 -0
  404. package/dist/voice/io.cjs.map +1 -0
  405. package/dist/voice/io.d.cts +101 -0
  406. package/dist/voice/io.d.ts +101 -0
  407. package/dist/voice/io.d.ts.map +1 -0
  408. package/dist/voice/io.js +217 -0
  409. package/dist/voice/io.js.map +1 -0
  410. package/dist/voice/room_io/_input.cjs +121 -0
  411. package/dist/voice/room_io/_input.cjs.map +1 -0
  412. package/dist/voice/room_io/_input.d.cts +24 -0
  413. package/dist/voice/room_io/_input.d.ts +24 -0
  414. package/dist/voice/room_io/_input.d.ts.map +1 -0
  415. package/dist/voice/room_io/_input.js +102 -0
  416. package/dist/voice/room_io/_input.js.map +1 -0
  417. package/dist/voice/room_io/_output.cjs +358 -0
  418. package/dist/voice/room_io/_output.cjs.map +1 -0
  419. package/dist/voice/room_io/_output.d.cts +75 -0
  420. package/dist/voice/room_io/_output.d.ts +75 -0
  421. package/dist/voice/room_io/_output.d.ts.map +1 -0
  422. package/dist/voice/room_io/_output.js +342 -0
  423. package/dist/voice/room_io/_output.js.map +1 -0
  424. package/dist/voice/room_io/index.cjs +25 -0
  425. package/dist/voice/room_io/index.cjs.map +1 -0
  426. package/dist/voice/room_io/index.d.cts +3 -0
  427. package/dist/voice/room_io/index.d.ts +3 -0
  428. package/dist/voice/room_io/index.d.ts.map +1 -0
  429. package/dist/voice/room_io/index.js +3 -0
  430. package/dist/voice/room_io/index.js.map +1 -0
  431. package/dist/voice/room_io/room_io.cjs +370 -0
  432. package/dist/voice/room_io/room_io.cjs.map +1 -0
  433. package/dist/voice/room_io/room_io.d.cts +73 -0
  434. package/dist/voice/room_io/room_io.d.ts +73 -0
  435. package/dist/voice/room_io/room_io.d.ts.map +1 -0
  436. package/dist/voice/room_io/room_io.js +361 -0
  437. package/dist/voice/room_io/room_io.js.map +1 -0
  438. package/dist/{pipeline/index.cjs → voice/run_context.cjs} +16 -11
  439. package/dist/voice/run_context.cjs.map +1 -0
  440. package/dist/voice/run_context.d.cts +12 -0
  441. package/dist/voice/run_context.d.ts +12 -0
  442. package/dist/voice/run_context.d.ts.map +1 -0
  443. package/dist/voice/run_context.js +14 -0
  444. package/dist/voice/run_context.js.map +1 -0
  445. package/dist/voice/speech_handle.cjs +105 -0
  446. package/dist/voice/speech_handle.cjs.map +1 -0
  447. package/dist/voice/speech_handle.d.cts +46 -0
  448. package/dist/voice/speech_handle.d.ts +46 -0
  449. package/dist/voice/speech_handle.d.ts.map +1 -0
  450. package/dist/voice/speech_handle.js +81 -0
  451. package/dist/voice/speech_handle.js.map +1 -0
  452. package/dist/voice/transcription/_utils.cjs +45 -0
  453. package/dist/voice/transcription/_utils.cjs.map +1 -0
  454. package/dist/voice/transcription/_utils.d.cts +3 -0
  455. package/dist/voice/transcription/_utils.d.ts +3 -0
  456. package/dist/voice/transcription/_utils.d.ts.map +1 -0
  457. package/dist/voice/transcription/_utils.js +21 -0
  458. package/dist/voice/transcription/_utils.js.map +1 -0
  459. package/dist/voice/transcription/index.cjs +23 -0
  460. package/dist/voice/transcription/index.cjs.map +1 -0
  461. package/dist/voice/transcription/index.d.cts +2 -0
  462. package/dist/voice/transcription/index.d.ts +2 -0
  463. package/dist/voice/transcription/index.d.ts.map +1 -0
  464. package/dist/voice/transcription/index.js +2 -0
  465. package/dist/voice/transcription/index.js.map +1 -0
  466. package/dist/voice/transcription/synchronizer.cjs +380 -0
  467. package/dist/voice/transcription/synchronizer.cjs.map +1 -0
  468. package/dist/voice/transcription/synchronizer.d.cts +86 -0
  469. package/dist/voice/transcription/synchronizer.d.ts +86 -0
  470. package/dist/voice/transcription/synchronizer.d.ts.map +1 -0
  471. package/dist/voice/transcription/synchronizer.js +355 -0
  472. package/dist/voice/transcription/synchronizer.js.map +1 -0
  473. package/dist/worker.cjs +22 -4
  474. package/dist/worker.cjs.map +1 -1
  475. package/dist/worker.d.cts +1 -1
  476. package/dist/worker.d.ts +1 -1
  477. package/dist/worker.d.ts.map +1 -1
  478. package/dist/worker.js +22 -4
  479. package/dist/worker.js.map +1 -1
  480. package/package.json +9 -2
  481. package/src/_exceptions.ts +137 -0
  482. package/src/audio.ts +12 -1
  483. package/src/cli.ts +37 -0
  484. package/src/constants.ts +2 -0
  485. package/src/http_server.ts +1 -0
  486. package/src/index.ts +13 -10
  487. package/src/inference_runner.ts +2 -3
  488. package/src/ipc/inference_proc_executor.ts +2 -2
  489. package/src/ipc/job_executor.ts +1 -1
  490. package/src/ipc/job_proc_executor.ts +1 -1
  491. package/src/ipc/job_proc_lazy_main.ts +1 -1
  492. package/src/job.ts +18 -0
  493. package/src/llm/__snapshots__/chat_context.test.ts.snap +527 -0
  494. package/src/llm/__snapshots__/tool_context.test.ts.snap +177 -0
  495. package/src/llm/__snapshots__/utils.test.ts.snap +65 -0
  496. package/src/llm/chat_context.test.ts +450 -0
  497. package/src/llm/chat_context.ts +501 -103
  498. package/src/llm/index.ts +53 -18
  499. package/src/llm/llm.ts +149 -50
  500. package/src/llm/provider_format/google.test.ts +772 -0
  501. package/src/llm/provider_format/google.ts +130 -0
  502. package/src/llm/provider_format/index.ts +23 -0
  503. package/src/llm/provider_format/openai.test.ts +581 -0
  504. package/src/llm/provider_format/openai.ts +118 -0
  505. package/src/llm/provider_format/utils.ts +183 -0
  506. package/src/llm/realtime.ts +151 -0
  507. package/src/llm/remote_chat_context.test.ts +290 -0
  508. package/src/llm/remote_chat_context.ts +114 -0
  509. package/src/llm/tool_context.test.ts +198 -0
  510. package/src/llm/tool_context.ts +259 -0
  511. package/src/llm/tool_context.type.test.ts +115 -0
  512. package/src/llm/utils.test.ts +670 -0
  513. package/src/llm/utils.ts +324 -0
  514. package/src/metrics/base.ts +110 -78
  515. package/src/metrics/index.ts +3 -9
  516. package/src/metrics/usage_collector.ts +19 -13
  517. package/src/metrics/utils.ts +24 -69
  518. package/src/multimodal/index.ts +0 -1
  519. package/src/plugin.ts +26 -8
  520. package/src/stream/deferred_stream.test.ts +755 -0
  521. package/src/stream/deferred_stream.ts +110 -0
  522. package/src/stream/identity_transform.test.ts +179 -0
  523. package/src/stream/identity_transform.ts +18 -0
  524. package/src/stream/index.ts +7 -0
  525. package/src/stream/merge_readable_streams.ts +40 -0
  526. package/src/stream/stream_channel.test.ts +129 -0
  527. package/src/stream/stream_channel.ts +32 -0
  528. package/src/stt/stream_adapter.ts +3 -5
  529. package/src/stt/stt.ts +135 -17
  530. package/src/tokenize/basic/basic.ts +13 -5
  531. package/src/tokenize/basic/sentence.ts +20 -6
  532. package/src/tokenize/token_stream.ts +7 -4
  533. package/src/transcription.ts +2 -3
  534. package/src/tts/index.ts +0 -1
  535. package/src/tts/stream_adapter.ts +42 -16
  536. package/src/tts/tts.ts +203 -21
  537. package/src/types.ts +42 -0
  538. package/src/utils.test.ts +658 -0
  539. package/src/utils.ts +375 -44
  540. package/src/vad.ts +90 -22
  541. package/src/voice/agent.test.ts +80 -0
  542. package/src/voice/agent.ts +332 -0
  543. package/src/voice/agent_activity.ts +1913 -0
  544. package/src/voice/agent_session.ts +460 -0
  545. package/src/voice/audio_recognition.ts +474 -0
  546. package/src/voice/events.ts +252 -0
  547. package/src/voice/generation.ts +881 -0
  548. package/src/voice/index.ts +7 -0
  549. package/src/voice/io.ts +304 -0
  550. package/src/voice/room_io/_input.ts +144 -0
  551. package/src/voice/room_io/_output.ts +436 -0
  552. package/src/voice/room_io/index.ts +5 -0
  553. package/src/voice/room_io/room_io.ts +495 -0
  554. package/src/voice/run_context.ts +20 -0
  555. package/src/voice/speech_handle.ts +104 -0
  556. package/src/voice/transcription/_utils.ts +25 -0
  557. package/src/voice/transcription/index.ts +4 -0
  558. package/src/voice/transcription/synchronizer.ts +478 -0
  559. package/src/worker.ts +22 -2
  560. package/dist/llm/function_context.cjs +0 -103
  561. package/dist/llm/function_context.cjs.map +0 -1
  562. package/dist/llm/function_context.d.cts +0 -47
  563. package/dist/llm/function_context.d.ts +0 -47
  564. package/dist/llm/function_context.d.ts.map +0 -1
  565. package/dist/llm/function_context.js +0 -78
  566. package/dist/llm/function_context.js.map +0 -1
  567. package/dist/llm/function_context.test.cjs +0 -218
  568. package/dist/llm/function_context.test.cjs.map +0 -1
  569. package/dist/llm/function_context.test.js +0 -217
  570. package/dist/llm/function_context.test.js.map +0 -1
  571. package/dist/multimodal/multimodal_agent.cjs +0 -451
  572. package/dist/multimodal/multimodal_agent.cjs.map +0 -1
  573. package/dist/multimodal/multimodal_agent.d.cts +0 -48
  574. package/dist/multimodal/multimodal_agent.d.ts +0 -48
  575. package/dist/multimodal/multimodal_agent.d.ts.map +0 -1
  576. package/dist/multimodal/multimodal_agent.js +0 -425
  577. package/dist/multimodal/multimodal_agent.js.map +0 -1
  578. package/dist/pipeline/agent_output.cjs +0 -197
  579. package/dist/pipeline/agent_output.cjs.map +0 -1
  580. package/dist/pipeline/agent_output.d.cts +0 -33
  581. package/dist/pipeline/agent_output.d.ts +0 -33
  582. package/dist/pipeline/agent_output.d.ts.map +0 -1
  583. package/dist/pipeline/agent_output.js +0 -172
  584. package/dist/pipeline/agent_output.js.map +0 -1
  585. package/dist/pipeline/agent_playout.cjs +0 -175
  586. package/dist/pipeline/agent_playout.cjs.map +0 -1
  587. package/dist/pipeline/agent_playout.d.cts +0 -40
  588. package/dist/pipeline/agent_playout.d.ts +0 -40
  589. package/dist/pipeline/agent_playout.d.ts.map +0 -1
  590. package/dist/pipeline/agent_playout.js +0 -139
  591. package/dist/pipeline/agent_playout.js.map +0 -1
  592. package/dist/pipeline/human_input.cjs +0 -171
  593. package/dist/pipeline/human_input.cjs.map +0 -1
  594. package/dist/pipeline/human_input.d.cts +0 -30
  595. package/dist/pipeline/human_input.d.ts +0 -30
  596. package/dist/pipeline/human_input.d.ts.map +0 -1
  597. package/dist/pipeline/human_input.js +0 -146
  598. package/dist/pipeline/human_input.js.map +0 -1
  599. package/dist/pipeline/index.cjs.map +0 -1
  600. package/dist/pipeline/index.d.cts +0 -2
  601. package/dist/pipeline/index.d.ts +0 -2
  602. package/dist/pipeline/index.d.ts.map +0 -1
  603. package/dist/pipeline/index.js +0 -11
  604. package/dist/pipeline/index.js.map +0 -1
  605. package/dist/pipeline/pipeline_agent.cjs +0 -849
  606. package/dist/pipeline/pipeline_agent.cjs.map +0 -1
  607. package/dist/pipeline/pipeline_agent.d.cts +0 -150
  608. package/dist/pipeline/pipeline_agent.d.ts +0 -150
  609. package/dist/pipeline/pipeline_agent.d.ts.map +0 -1
  610. package/dist/pipeline/pipeline_agent.js +0 -826
  611. package/dist/pipeline/pipeline_agent.js.map +0 -1
  612. package/dist/pipeline/speech_handle.cjs +0 -176
  613. package/dist/pipeline/speech_handle.cjs.map +0 -1
  614. package/dist/pipeline/speech_handle.d.cts +0 -37
  615. package/dist/pipeline/speech_handle.d.ts +0 -37
  616. package/dist/pipeline/speech_handle.d.ts.map +0 -1
  617. package/dist/pipeline/speech_handle.js +0 -152
  618. package/dist/pipeline/speech_handle.js.map +0 -1
  619. package/src/llm/function_context.test.ts +0 -248
  620. package/src/llm/function_context.ts +0 -142
  621. package/src/multimodal/multimodal_agent.ts +0 -555
  622. package/src/pipeline/agent_output.ts +0 -219
  623. package/src/pipeline/agent_playout.ts +0 -192
  624. package/src/pipeline/human_input.ts +0 -188
  625. package/src/pipeline/index.ts +0 -15
  626. package/src/pipeline/pipeline_agent.ts +0 -1185
  627. package/src/pipeline/speech_handle.ts +0 -201
@@ -0,0 +1,772 @@
1
+ // SPDX-FileCopyrightText: 2025 LiveKit, Inc.
2
+ //
3
+ // SPDX-License-Identifier: Apache-2.0
4
+ import { VideoBufferType, VideoFrame } from '@livekit/rtc-node';
5
+ import { beforeEach, describe, expect, it, vi } from 'vitest';
6
+ import { initializeLogger } from '../../log.js';
7
+ import { ChatContext, FunctionCall, FunctionCallOutput } from '../chat_context.js';
8
+ import { serializeImage } from '../utils.js';
9
+ import { toChatCtx } from './google.js';
10
+
11
+ vi.mock('../utils.js', () => ({
12
+ serializeImage: vi.fn(),
13
+ }));
14
+
15
+ describe('Google Provider Format - toChatCtx', () => {
16
+ const serializeImageMock = vi.mocked(serializeImage);
17
+
18
+ initializeLogger({ level: 'silent', pretty: false });
19
+
20
+ beforeEach(async () => {
21
+ vi.clearAllMocks();
22
+ });
23
+
24
+ it('should convert simple text messages', async () => {
25
+ const ctx = ChatContext.empty();
26
+ ctx.addMessage({ role: 'user', content: 'Hello' });
27
+ ctx.addMessage({ role: 'assistant', content: 'Hi there!' });
28
+
29
+ const [result, formatData] = await toChatCtx(ctx, false);
30
+
31
+ expect(result).toEqual([
32
+ {
33
+ role: 'user',
34
+ parts: [{ text: 'Hello' }],
35
+ },
36
+ {
37
+ role: 'model',
38
+ parts: [{ text: 'Hi there!' }],
39
+ },
40
+ ]);
41
+ expect(formatData.systemMessages).toBeNull();
42
+ });
43
+
44
+ it('should handle system messages separately', async () => {
45
+ const ctx = ChatContext.empty();
46
+ ctx.addMessage({ role: 'system', content: 'You are a helpful assistant' });
47
+ ctx.addMessage({ role: 'user', content: 'Hello' });
48
+
49
+ const [result, formatData] = await toChatCtx(ctx, false);
50
+
51
+ expect(result).toEqual([
52
+ {
53
+ role: 'user',
54
+ parts: [{ text: 'Hello' }],
55
+ },
56
+ ]);
57
+ expect(formatData.systemMessages).toEqual(['You are a helpful assistant']);
58
+ });
59
+
60
+ it('should handle multiple system messages', async () => {
61
+ const ctx = ChatContext.empty();
62
+ ctx.addMessage({ role: 'system', content: 'You are a helpful assistant' });
63
+ ctx.addMessage({ role: 'system', content: 'Be concise in your responses' });
64
+ ctx.addMessage({ role: 'user', content: 'Hello' });
65
+
66
+ const [result, formatData] = await toChatCtx(ctx, false);
67
+
68
+ expect(result).toEqual([
69
+ {
70
+ role: 'user',
71
+ parts: [{ text: 'Hello' }],
72
+ },
73
+ ]);
74
+ expect(formatData.systemMessages).toEqual([
75
+ 'You are a helpful assistant',
76
+ 'Be concise in your responses',
77
+ ]);
78
+ });
79
+
80
+ it('should handle multi-part text content', async () => {
81
+ const ctx = ChatContext.empty();
82
+ ctx.addMessage({ role: 'user', content: ['Line 1', 'Line 2', 'Line 3'] });
83
+
84
+ const [result, formatData] = await toChatCtx(ctx, false);
85
+
86
+ expect(result).toEqual([
87
+ {
88
+ role: 'user',
89
+ parts: [{ text: 'Line 1' }, { text: 'Line 2' }, { text: 'Line 3' }],
90
+ },
91
+ ]);
92
+ expect(formatData.systemMessages).toBeNull();
93
+ });
94
+
95
+ it('should handle messages with external URL images', async () => {
96
+ serializeImageMock.mockResolvedValue({
97
+ inferenceDetail: 'high',
98
+ externalUrl: 'https://example.com/image.jpg',
99
+ mimeType: 'image/jpeg',
100
+ });
101
+
102
+ const ctx = ChatContext.empty();
103
+ ctx.addMessage({
104
+ role: 'user',
105
+ content: [
106
+ 'Check out this image:',
107
+ {
108
+ id: 'img1',
109
+ type: 'image_content',
110
+ image: 'https://example.com/image.jpg',
111
+ inferenceDetail: 'high',
112
+ _cache: {},
113
+ },
114
+ ],
115
+ });
116
+
117
+ const [result, formatData] = await toChatCtx(ctx, false);
118
+
119
+ expect(result).toEqual([
120
+ {
121
+ role: 'user',
122
+ parts: [
123
+ { text: 'Check out this image:' },
124
+ {
125
+ fileData: {
126
+ fileUri: 'https://example.com/image.jpg',
127
+ mimeType: 'image/jpeg',
128
+ },
129
+ },
130
+ ],
131
+ },
132
+ ]);
133
+ expect(formatData.systemMessages).toBeNull();
134
+ });
135
+
136
+ it('should handle messages with base64 images', async () => {
137
+ serializeImageMock.mockResolvedValue({
138
+ inferenceDetail: 'auto',
139
+ mimeType: 'image/png',
140
+ base64Data: 'iVBORw0KGgoAAAANSUhEUgAAAAEAAAAB',
141
+ });
142
+
143
+ const ctx = ChatContext.empty();
144
+ ctx.addMessage({
145
+ role: 'assistant',
146
+ content: [
147
+ {
148
+ id: 'img1',
149
+ type: 'image_content',
150
+ image: 'data:image/png;base64,iVBORw0KGgoAAAANSUhEUgAAAAEAAAAB',
151
+ inferenceDetail: 'auto',
152
+ _cache: {},
153
+ },
154
+ 'Here is the image you requested',
155
+ ],
156
+ });
157
+
158
+ const [result, formatData] = await toChatCtx(ctx, false);
159
+
160
+ expect(result).toEqual([
161
+ {
162
+ role: 'model',
163
+ parts: [
164
+ {
165
+ inlineData: {
166
+ data: 'iVBORw0KGgoAAAANSUhEUgAAAAEAAAAB',
167
+ mimeType: 'image/png',
168
+ },
169
+ },
170
+ { text: 'Here is the image you requested' },
171
+ ],
172
+ },
173
+ ]);
174
+ expect(formatData.systemMessages).toBeNull();
175
+ });
176
+
177
+ it('should handle VideoFrame images', async () => {
178
+ serializeImageMock.mockResolvedValue({
179
+ inferenceDetail: 'low',
180
+ mimeType: 'image/jpeg',
181
+ base64Data: '/9j/4AAQSkZJRg==',
182
+ });
183
+
184
+ const frameData = new Uint8Array(4 * 4 * 4); // 4x4 RGBA
185
+ const videoFrame = new VideoFrame(frameData, 4, 4, VideoBufferType.RGBA);
186
+
187
+ const ctx = ChatContext.empty();
188
+ ctx.addMessage({
189
+ role: 'user',
190
+ content: [
191
+ {
192
+ id: 'frame1',
193
+ type: 'image_content',
194
+ image: videoFrame,
195
+ inferenceDetail: 'low',
196
+ _cache: {},
197
+ },
198
+ ],
199
+ });
200
+
201
+ const [result, formatData] = await toChatCtx(ctx, false);
202
+
203
+ expect(result).toEqual([
204
+ {
205
+ role: 'user',
206
+ parts: [
207
+ {
208
+ inlineData: {
209
+ data: '/9j/4AAQSkZJRg==',
210
+ mimeType: 'image/jpeg',
211
+ },
212
+ },
213
+ ],
214
+ },
215
+ ]);
216
+ expect(formatData.systemMessages).toBeNull();
217
+ });
218
+
219
+ it('should cache serialized images', async () => {
220
+ serializeImageMock.mockResolvedValue({
221
+ inferenceDetail: 'high',
222
+ mimeType: 'image/png',
223
+ base64Data: 'cached-data',
224
+ });
225
+
226
+ const imageContent = {
227
+ id: 'img1',
228
+ type: 'image_content' as const,
229
+ image: 'https://example.com/image.jpg',
230
+ inferenceDetail: 'high' as const,
231
+ _cache: {},
232
+ };
233
+
234
+ const ctx = ChatContext.empty();
235
+ ctx.addMessage({ role: 'user', content: [imageContent] });
236
+
237
+ await toChatCtx(ctx, false);
238
+ await toChatCtx(ctx, false);
239
+
240
+ expect(serializeImageMock).toHaveBeenCalledTimes(1);
241
+ expect(imageContent._cache).toHaveProperty('serialized_image');
242
+ });
243
+
244
+ it('should handle tool calls and outputs', async () => {
245
+ const ctx = ChatContext.empty();
246
+
247
+ const msg = ctx.addMessage({ role: 'assistant', content: 'Let me help you with that.' });
248
+ const toolCall = FunctionCall.create({
249
+ id: msg.id + '/tool_1',
250
+ callId: 'call_123',
251
+ name: 'get_weather',
252
+ args: '{"location": "San Francisco"}',
253
+ });
254
+ const toolOutput = FunctionCallOutput.create({
255
+ callId: 'call_123',
256
+ name: 'get_weather',
257
+ output: '{"temperature": 72, "condition": "sunny"}',
258
+ isError: false,
259
+ });
260
+
261
+ ctx.insert([toolCall, toolOutput]);
262
+
263
+ const [result, formatData] = await toChatCtx(ctx, false);
264
+
265
+ expect(result).toEqual([
266
+ {
267
+ role: 'model',
268
+ parts: [
269
+ { text: 'Let me help you with that.' },
270
+ {
271
+ functionCall: {
272
+ id: 'call_123',
273
+ name: 'get_weather',
274
+ args: { location: 'San Francisco' },
275
+ },
276
+ },
277
+ ],
278
+ },
279
+ {
280
+ role: 'user',
281
+ parts: [
282
+ {
283
+ functionResponse: {
284
+ id: 'call_123',
285
+ name: 'get_weather',
286
+ response: { output: '{"temperature": 72, "condition": "sunny"}' },
287
+ },
288
+ },
289
+ ],
290
+ },
291
+ ]);
292
+ expect(formatData.systemMessages).toBeNull();
293
+ });
294
+
295
+ it('should handle multiple tool calls in one message', async () => {
296
+ const ctx = ChatContext.empty();
297
+
298
+ const msg = ctx.addMessage({ role: 'assistant', content: "I'll check both locations." });
299
+ const toolCall1 = new FunctionCall({
300
+ id: msg.id + '/tool_1',
301
+ callId: 'call_1',
302
+ name: 'get_weather',
303
+ args: '{"location": "NYC"}',
304
+ });
305
+ const toolCall2 = new FunctionCall({
306
+ id: msg.id + '/tool_2',
307
+ callId: 'call_2',
308
+ name: 'get_weather',
309
+ args: '{"location": "LA"}',
310
+ });
311
+ const toolOutput1 = new FunctionCallOutput({
312
+ callId: 'call_1',
313
+ name: 'get_weather',
314
+ output: '{"temperature": 65}',
315
+ isError: false,
316
+ });
317
+ const toolOutput2 = new FunctionCallOutput({
318
+ callId: 'call_2',
319
+ name: 'get_weather',
320
+ output: '{"temperature": 78}',
321
+ isError: false,
322
+ });
323
+
324
+ ctx.insert([toolCall1, toolCall2, toolOutput1, toolOutput2]);
325
+
326
+ const [result, formatData] = await toChatCtx(ctx, false);
327
+
328
+ expect(result).toEqual([
329
+ {
330
+ role: 'model',
331
+ parts: [
332
+ { text: "I'll check both locations." },
333
+ {
334
+ functionCall: {
335
+ id: 'call_1',
336
+ name: 'get_weather',
337
+ args: { location: 'NYC' },
338
+ },
339
+ },
340
+ {
341
+ functionCall: {
342
+ id: 'call_2',
343
+ name: 'get_weather',
344
+ args: { location: 'LA' },
345
+ },
346
+ },
347
+ ],
348
+ },
349
+ {
350
+ role: 'user',
351
+ parts: [
352
+ {
353
+ functionResponse: {
354
+ id: 'call_1',
355
+ name: 'get_weather',
356
+ response: { output: '{"temperature": 65}' },
357
+ },
358
+ },
359
+ {
360
+ functionResponse: {
361
+ id: 'call_2',
362
+ name: 'get_weather',
363
+ response: { output: '{"temperature": 78}' },
364
+ },
365
+ },
366
+ ],
367
+ },
368
+ ]);
369
+ expect(formatData.systemMessages).toBeNull();
370
+ });
371
+
372
+ it('should handle tool calls without accompanying message', async () => {
373
+ const ctx = ChatContext.empty();
374
+
375
+ const toolCall = new FunctionCall({
376
+ id: 'func_123',
377
+ callId: 'call_456',
378
+ name: 'calculate',
379
+ args: '{"a": 5, "b": 3}',
380
+ });
381
+ const toolOutput = new FunctionCallOutput({
382
+ callId: 'call_456',
383
+ name: 'calculate',
384
+ output: '{"result": 8}',
385
+ isError: false,
386
+ });
387
+
388
+ ctx.insert([toolCall, toolOutput]);
389
+
390
+ const [result, formatData] = await toChatCtx(ctx, false);
391
+
392
+ expect(result).toEqual([
393
+ {
394
+ role: 'model',
395
+ parts: [
396
+ {
397
+ functionCall: {
398
+ id: 'call_456',
399
+ name: 'calculate',
400
+ args: { a: 5, b: 3 },
401
+ },
402
+ },
403
+ ],
404
+ },
405
+ {
406
+ role: 'user',
407
+ parts: [
408
+ {
409
+ functionResponse: {
410
+ id: 'call_456',
411
+ name: 'calculate',
412
+ response: { output: '{"result": 8}' },
413
+ },
414
+ },
415
+ ],
416
+ },
417
+ ]);
418
+ expect(formatData.systemMessages).toBeNull();
419
+ });
420
+
421
+ it('should handle tool call errors', async () => {
422
+ const ctx = ChatContext.empty();
423
+
424
+ const toolCall = new FunctionCall({
425
+ id: 'func_error',
426
+ callId: 'call_error',
427
+ name: 'failing_function',
428
+ args: '{}',
429
+ });
430
+ const toolOutput = new FunctionCallOutput({
431
+ callId: 'call_error',
432
+ name: 'failing_function',
433
+ output: 'Function failed to execute',
434
+ isError: true,
435
+ });
436
+
437
+ ctx.insert([toolCall, toolOutput]);
438
+
439
+ const [result, formatData] = await toChatCtx(ctx, false);
440
+
441
+ expect(result).toEqual([
442
+ {
443
+ role: 'model',
444
+ parts: [
445
+ {
446
+ functionCall: {
447
+ id: 'call_error',
448
+ name: 'failing_function',
449
+ args: {},
450
+ },
451
+ },
452
+ ],
453
+ },
454
+ {
455
+ role: 'user',
456
+ parts: [
457
+ {
458
+ functionResponse: {
459
+ id: 'call_error',
460
+ name: 'failing_function',
461
+ response: { error: 'Function failed to execute' },
462
+ },
463
+ },
464
+ ],
465
+ },
466
+ ]);
467
+ expect(formatData.systemMessages).toBeNull();
468
+ });
469
+
470
+ it('should inject dummy user message when last turn is not user', async () => {
471
+ const ctx = ChatContext.empty();
472
+ ctx.addMessage({ role: 'user', content: 'Hello' });
473
+ ctx.addMessage({ role: 'assistant', content: 'Hi there!' });
474
+
475
+ const [result, formatData] = await toChatCtx(ctx, true);
476
+
477
+ expect(result).toEqual([
478
+ {
479
+ role: 'user',
480
+ parts: [{ text: 'Hello' }],
481
+ },
482
+ {
483
+ role: 'model',
484
+ parts: [{ text: 'Hi there!' }],
485
+ },
486
+ {
487
+ role: 'user',
488
+ parts: [{ text: '.' }],
489
+ },
490
+ ]);
491
+ expect(formatData.systemMessages).toBeNull();
492
+ });
493
+
494
+ it('should not inject dummy user message when last turn is already user', async () => {
495
+ const ctx = ChatContext.empty();
496
+ ctx.addMessage({ role: 'assistant', content: 'Hi there!' });
497
+ ctx.addMessage({ role: 'user', content: 'Hello' });
498
+
499
+ const [result, formatData] = await toChatCtx(ctx, true);
500
+
501
+ expect(result).toEqual([
502
+ {
503
+ role: 'model',
504
+ parts: [{ text: 'Hi there!' }],
505
+ },
506
+ {
507
+ role: 'user',
508
+ parts: [{ text: 'Hello' }],
509
+ },
510
+ ]);
511
+ expect(formatData.systemMessages).toBeNull();
512
+ });
513
+
514
+ it('should not inject dummy user message when disabled', async () => {
515
+ const ctx = ChatContext.empty();
516
+ ctx.addMessage({ role: 'user', content: 'Hello' });
517
+ ctx.addMessage({ role: 'assistant', content: 'Hi there!' });
518
+
519
+ const [result, formatData] = await toChatCtx(ctx, false);
520
+
521
+ expect(result).toEqual([
522
+ {
523
+ role: 'user',
524
+ parts: [{ text: 'Hello' }],
525
+ },
526
+ {
527
+ role: 'model',
528
+ parts: [{ text: 'Hi there!' }],
529
+ },
530
+ ]);
531
+ expect(formatData.systemMessages).toBeNull();
532
+ });
533
+
534
+ it('should handle mixed content with text and multiple images', async () => {
535
+ serializeImageMock
536
+ .mockResolvedValueOnce({
537
+ inferenceDetail: 'high',
538
+ externalUrl: 'https://example.com/image1.jpg',
539
+ mimeType: 'image/jpeg',
540
+ })
541
+ .mockResolvedValueOnce({
542
+ inferenceDetail: 'low',
543
+ mimeType: 'image/png',
544
+ base64Data: 'base64data',
545
+ });
546
+
547
+ const ctx = ChatContext.empty();
548
+ ctx.addMessage({
549
+ role: 'user',
550
+ content: [
551
+ 'Here are two images:',
552
+ {
553
+ id: 'img1',
554
+ type: 'image_content',
555
+ image: 'https://example.com/image1.jpg',
556
+ inferenceDetail: 'high',
557
+ _cache: {},
558
+ },
559
+ 'And the second one:',
560
+ {
561
+ id: 'img2',
562
+ type: 'image_content',
563
+ image: 'data:image/png;base64,base64data',
564
+ inferenceDetail: 'low',
565
+ _cache: {},
566
+ },
567
+ 'What do you think?',
568
+ ],
569
+ });
570
+
571
+ const [result, formatData] = await toChatCtx(ctx, false);
572
+
573
+ expect(result).toEqual([
574
+ {
575
+ role: 'user',
576
+ parts: [
577
+ { text: 'Here are two images:' },
578
+ {
579
+ fileData: {
580
+ fileUri: 'https://example.com/image1.jpg',
581
+ mimeType: 'image/jpeg',
582
+ },
583
+ },
584
+ { text: 'And the second one:' },
585
+ {
586
+ inlineData: {
587
+ data: 'base64data',
588
+ mimeType: 'image/png',
589
+ },
590
+ },
591
+ { text: 'What do you think?' },
592
+ ],
593
+ },
594
+ ]);
595
+ expect(formatData.systemMessages).toBeNull();
596
+ });
597
+
598
+ it('should handle content with only images and no text', async () => {
599
+ serializeImageMock.mockResolvedValue({
600
+ inferenceDetail: 'auto',
601
+ externalUrl: 'https://example.com/image.jpg',
602
+ mimeType: 'image/jpeg',
603
+ });
604
+
605
+ const ctx = ChatContext.empty();
606
+ ctx.addMessage({
607
+ role: 'user',
608
+ content: [
609
+ {
610
+ id: 'img1',
611
+ type: 'image_content',
612
+ image: 'https://example.com/image.jpg',
613
+ inferenceDetail: 'auto',
614
+ _cache: {},
615
+ },
616
+ ],
617
+ });
618
+
619
+ const [result, formatData] = await toChatCtx(ctx, false);
620
+
621
+ expect(result).toEqual([
622
+ {
623
+ role: 'user',
624
+ parts: [
625
+ {
626
+ fileData: {
627
+ fileUri: 'https://example.com/image.jpg',
628
+ mimeType: 'image/jpeg',
629
+ },
630
+ },
631
+ ],
632
+ },
633
+ ]);
634
+ expect(formatData.systemMessages).toBeNull();
635
+ });
636
+
637
+ it('should group consecutive messages by role', async () => {
638
+ const ctx = ChatContext.empty();
639
+ ctx.addMessage({ role: 'user', content: 'First user message' });
640
+ ctx.addMessage({ role: 'user', content: 'Second user message' });
641
+ ctx.addMessage({ role: 'assistant', content: 'First assistant response' });
642
+ ctx.addMessage({ role: 'assistant', content: 'Second assistant response' });
643
+
644
+ const [result, formatData] = await toChatCtx(ctx, false);
645
+
646
+ expect(result).toEqual([
647
+ {
648
+ role: 'user',
649
+ parts: [{ text: 'First user message' }, { text: 'Second user message' }],
650
+ },
651
+ {
652
+ role: 'model',
653
+ parts: [{ text: 'First assistant response' }, { text: 'Second assistant response' }],
654
+ },
655
+ ]);
656
+ expect(formatData.systemMessages).toBeNull();
657
+ });
658
+
659
+ it('should handle empty content arrays', async () => {
660
+ const ctx = ChatContext.empty();
661
+ ctx.addMessage({ role: 'user', content: [] });
662
+
663
+ const [result, formatData] = await toChatCtx(ctx, false);
664
+
665
+ expect(result).toEqual([]);
666
+ expect(formatData.systemMessages).toBeNull();
667
+ });
668
+
669
+ it('should handle image with default MIME type when missing', async () => {
670
+ serializeImageMock.mockResolvedValue({
671
+ inferenceDetail: 'high',
672
+ externalUrl: 'https://example.com/image.jpg',
673
+ // No mimeType provided
674
+ });
675
+
676
+ const ctx = ChatContext.empty();
677
+ ctx.addMessage({
678
+ role: 'user',
679
+ content: [
680
+ {
681
+ id: 'img1',
682
+ type: 'image_content',
683
+ image: 'https://example.com/image.jpg',
684
+ inferenceDetail: 'high',
685
+ _cache: {},
686
+ },
687
+ ],
688
+ });
689
+
690
+ const [result, formatData] = await toChatCtx(ctx, false);
691
+
692
+ expect(result).toEqual([
693
+ {
694
+ role: 'user',
695
+ parts: [
696
+ {
697
+ fileData: {
698
+ fileUri: 'https://example.com/image.jpg',
699
+ mimeType: 'image/jpeg', // Should default to image/jpeg
700
+ },
701
+ },
702
+ ],
703
+ },
704
+ ]);
705
+ expect(formatData.systemMessages).toBeNull();
706
+ });
707
+
708
+ it('should skip empty groups', async () => {
709
+ const ctx = ChatContext.empty();
710
+ ctx.addMessage({ role: 'user', content: 'Hello', createdAt: 1000 });
711
+
712
+ const orphanOutput = new FunctionCallOutput({
713
+ callId: 'orphan_call',
714
+ output: 'This should be ignored',
715
+ isError: false,
716
+ createdAt: 2000,
717
+ });
718
+ ctx.insert(orphanOutput);
719
+
720
+ ctx.addMessage({ role: 'assistant', content: 'Hi!', createdAt: 3000 });
721
+
722
+ const [result, formatData] = await toChatCtx(ctx, false);
723
+
724
+ expect(result).toEqual([
725
+ {
726
+ role: 'user',
727
+ parts: [{ text: 'Hello' }],
728
+ },
729
+ {
730
+ role: 'model',
731
+ parts: [{ text: 'Hi!' }],
732
+ },
733
+ ]);
734
+ expect(formatData.systemMessages).toBeNull();
735
+ });
736
+
737
+ it('should filter out standalone function calls without outputs', async () => {
738
+ const ctx = ChatContext.empty();
739
+
740
+ const funcCall = new FunctionCall({
741
+ id: 'func_standalone',
742
+ callId: 'call_999',
743
+ name: 'standalone_function',
744
+ args: '{}',
745
+ });
746
+
747
+ ctx.insert(funcCall);
748
+
749
+ const [result, formatData] = await toChatCtx(ctx, false);
750
+
751
+ expect(result).toEqual([]);
752
+ expect(formatData.systemMessages).toBeNull();
753
+ });
754
+
755
+ it('should handle mixed content types', async () => {
756
+ const ctx = ChatContext.empty();
757
+ ctx.addMessage({
758
+ role: 'user',
759
+ content: ['First part', 'Second part'],
760
+ });
761
+
762
+ const [result, formatData] = await toChatCtx(ctx, false);
763
+
764
+ expect(result).toEqual([
765
+ {
766
+ role: 'user',
767
+ parts: [{ text: 'First part' }, { text: 'Second part' }],
768
+ },
769
+ ]);
770
+ expect(formatData.systemMessages).toBeNull();
771
+ });
772
+ });