@genai-fi/nanogpt 0.10.1 → 0.10.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (206) hide show
  1. package/dist/Generator.js +14 -14
  2. package/dist/{RealDiv-DgA3z9oO.js → RealDiv-zz7FpkKX.js} +17 -17
  3. package/dist/{Reshape-CF6odzV4.js → Reshape-CDVLyVfz.js} +3 -3
  4. package/dist/{Reshape-_kILl6tK.js → Reshape-CHdUjC72.js} +4 -4
  5. package/dist/TeachableLLM.js +8 -8
  6. package/dist/{axis_util-BvHEw88j.js → axis_util-BsIr9ZNu.js} +1 -1
  7. package/dist/backend.js +2 -2
  8. package/dist/{backend_util-D-rUb2ty.js → backend_util-B1XRLuq9.js} +31 -31
  9. package/dist/{backend_webgpu-B0u2ndUn.js → backend_webgpu-CqpfEImu.js} +5 -5
  10. package/dist/{broadcast_to-CwF7XIeu.js → broadcast_to-B0ChcDaz.js} +4 -4
  11. package/dist/checks/appendCache.js +2 -2
  12. package/dist/checks/attentionMask.js +3 -3
  13. package/dist/checks/gelu.js +2 -2
  14. package/dist/checks/matMulGelu.js +5 -5
  15. package/dist/checks/normRMS.js +4 -4
  16. package/dist/checks/normRMSGrad.js +3 -3
  17. package/dist/checks/packUnpack.js +2 -2
  18. package/dist/checks/qkv.js +3 -3
  19. package/dist/checks/rope.js +2 -2
  20. package/dist/{complex-CSlYz-2T.js → complex-BBiRlsVq.js} +3 -3
  21. package/dist/{concat-BHlIJeyT.js → concat-DmBLPVGC.js} +3 -3
  22. package/dist/{concat_util-DcJk7YHS.js → concat_util-iBYIyuQe.js} +1 -1
  23. package/dist/{dataset-0xP8GjwI.js → dataset-D2P7rHAw.js} +5 -5
  24. package/dist/{dropout-C1pM3f11.js → dropout-B1x1kYMa.js} +3 -3
  25. package/dist/{expand_dims-BPG4fwBP.js → expand_dims-ouvfxQ1n.js} +3 -3
  26. package/dist/{exports_initializers-xuidcwI4.js → exports_initializers-CZSUJoVE.js} +1 -1
  27. package/dist/{gather-DykLGqmW.js → gather-CH9sdacz.js} +2 -2
  28. package/dist/{gelu-CNLFZWea.js → gelu-Bmhopi0J.js} +2 -2
  29. package/dist/{gpgpu_math-DDVJCn6-.js → gpgpu_math-DsCcikas.js} +3 -3
  30. package/dist/{index-ZyQhjEPo.js → index-D6Q1lPZO.js} +55 -55
  31. package/dist/{index-CjOj7j-u.js → index-DRyE072i.js} +15 -15
  32. package/dist/{kernel_funcs_utils-Dg_-E44D.js → kernel_funcs_utils-CWfOAPGO.js} +9 -9
  33. package/dist/layers/BaseLayer.js +10 -10
  34. package/dist/layers/CausalSelfAttention.js +6 -6
  35. package/dist/layers/MLP.js +4 -4
  36. package/dist/layers/PositionEmbedding.js +5 -5
  37. package/dist/layers/RMSNorm.js +3 -3
  38. package/dist/layers/RoPECache.js +4 -4
  39. package/dist/layers/TiedEmbedding.js +6 -6
  40. package/dist/layers/TransformerBlock.js +1 -1
  41. package/dist/loader/loadTransformers.js +1 -1
  42. package/dist/loader/oldZipLoad.js +8 -8
  43. package/dist/{log_sum_exp-DWI-76TI.js → log_sum_exp-D3ftBNY5.js} +6 -6
  44. package/dist/main.js +8 -8
  45. package/dist/{matMul16--R5hOwDG.js → matMul16-fEAJ4smh.js} +4 -4
  46. package/dist/{mat_mul-DeAh4uTH.js → mat_mul-C59XWcJd.js} +2 -2
  47. package/dist/{mod-Gt1rMB4n.js → mod-DESSvHIU.js} +2 -2
  48. package/dist/models/NanoGPTV1.js +2 -2
  49. package/dist/models/model.js +8 -8
  50. package/dist/{mulmat_packed_gpu-BMFhLwta.js → mulmat_packed_gpu-Coh6qbJk.js} +1 -1
  51. package/dist/{ones-CAMiP4I2.js → ones-jU9jlQvM.js} +4 -4
  52. package/dist/ops/adamAdjust.js +1 -1
  53. package/dist/ops/adamMoments.js +1 -1
  54. package/dist/ops/add16.js +1 -1
  55. package/dist/ops/appendCache.js +3 -3
  56. package/dist/ops/attentionMask.js +1 -1
  57. package/dist/ops/concat16.js +2 -2
  58. package/dist/ops/cpu/adamAdjust.js +2 -2
  59. package/dist/ops/cpu/adamMoments.js +3 -3
  60. package/dist/ops/cpu/appendCache.js +3 -3
  61. package/dist/ops/cpu/attentionMask.js +6 -6
  62. package/dist/ops/cpu/fusedSoftmax.js +3 -3
  63. package/dist/ops/cpu/gatherSub.js +4 -4
  64. package/dist/ops/cpu/gelu.js +2 -2
  65. package/dist/ops/cpu/matMul16.js +3 -3
  66. package/dist/ops/cpu/matMulGelu.js +4 -4
  67. package/dist/ops/cpu/matMulMul.js +2 -2
  68. package/dist/ops/cpu/mulDropout.js +2 -2
  69. package/dist/ops/cpu/normRMS.js +2 -2
  70. package/dist/ops/cpu/qkv.js +4 -4
  71. package/dist/ops/cpu/rope.js +6 -6
  72. package/dist/ops/cpu/scatterSub.js +7 -7
  73. package/dist/ops/dot16.js +2 -2
  74. package/dist/ops/gatherSub.js +1 -1
  75. package/dist/ops/gelu.js +2 -2
  76. package/dist/ops/grads/add16.js +2 -2
  77. package/dist/ops/grads/attentionMask.js +3 -3
  78. package/dist/ops/grads/gelu.js +3 -3
  79. package/dist/ops/grads/matMul16.js +4 -4
  80. package/dist/ops/grads/matMulGelu.js +2 -2
  81. package/dist/ops/grads/normRMS.js +2 -2
  82. package/dist/ops/grads/pack16.js +4 -4
  83. package/dist/ops/grads/qkv.js +4 -4
  84. package/dist/ops/grads/rope.js +3 -3
  85. package/dist/ops/grads/softmax16.js +2 -2
  86. package/dist/ops/grads/unpack16.js +3 -3
  87. package/dist/ops/matMul16.js +3 -3
  88. package/dist/ops/matMulGelu.js +1 -1
  89. package/dist/ops/matMulMul.js +1 -1
  90. package/dist/ops/mul16.js +1 -1
  91. package/dist/ops/mulDrop.js +1 -1
  92. package/dist/ops/normRMS.js +1 -1
  93. package/dist/ops/pack16.js +2 -2
  94. package/dist/ops/qkv.js +1 -1
  95. package/dist/ops/reshape16.js +3 -3
  96. package/dist/ops/rope.js +5 -5
  97. package/dist/ops/scatterSub.js +1 -1
  98. package/dist/ops/slice16.js +2 -2
  99. package/dist/ops/softmax16.js +1 -1
  100. package/dist/ops/sub16.js +1 -1
  101. package/dist/ops/sum16.js +2 -2
  102. package/dist/ops/transpose16.js +4 -4
  103. package/dist/ops/unpack16.js +2 -2
  104. package/dist/ops/webgl/adamAdjust.js +3 -3
  105. package/dist/ops/webgl/adamMoments.js +2 -2
  106. package/dist/ops/webgl/appendCache.js +2 -2
  107. package/dist/ops/webgl/attentionMask.js +2 -2
  108. package/dist/ops/webgl/fusedSoftmax.js +6 -6
  109. package/dist/ops/webgl/gatherSub.js +2 -2
  110. package/dist/ops/webgl/gelu.js +3 -3
  111. package/dist/ops/webgl/log.js +4 -4
  112. package/dist/ops/webgl/matMul16.js +5 -5
  113. package/dist/ops/webgl/matMulGelu.js +6 -6
  114. package/dist/ops/webgl/matMulMul.js +2 -2
  115. package/dist/ops/webgl/mulDropout.js +2 -2
  116. package/dist/ops/webgl/normRMS.js +3 -3
  117. package/dist/ops/webgl/qkv.js +2 -2
  118. package/dist/ops/webgl/rope.js +2 -2
  119. package/dist/ops/webgl/scatterSub.js +2 -2
  120. package/dist/ops/webgpu/adamAdjust.js +5 -5
  121. package/dist/ops/webgpu/adamMoments.js +5 -5
  122. package/dist/ops/webgpu/add16.js +2 -2
  123. package/dist/ops/webgpu/appendCache.js +5 -5
  124. package/dist/ops/webgpu/attentionMask.js +4 -4
  125. package/dist/ops/webgpu/attentionMask32_program.js +2 -2
  126. package/dist/ops/webgpu/concat16.js +7 -7
  127. package/dist/ops/webgpu/gatherSub.js +5 -5
  128. package/dist/ops/webgpu/gelu.js +4 -4
  129. package/dist/ops/webgpu/matMul16.js +6 -6
  130. package/dist/ops/webgpu/matMul16_program.js +3 -3
  131. package/dist/ops/webgpu/mul16.js +2 -2
  132. package/dist/ops/webgpu/normRMS.js +4 -4
  133. package/dist/ops/webgpu/normRMSGrad.js +6 -6
  134. package/dist/ops/webgpu/pack16.js +2 -2
  135. package/dist/ops/webgpu/pack16_program.js +2 -2
  136. package/dist/ops/webgpu/qkv.js +4 -4
  137. package/dist/ops/webgpu/rope.js +5 -5
  138. package/dist/ops/webgpu/scatterSub.js +5 -5
  139. package/dist/ops/webgpu/slice16.js +6 -6
  140. package/dist/ops/webgpu/softmax16.js +4 -4
  141. package/dist/ops/webgpu/softmax16_program.js +2 -2
  142. package/dist/ops/webgpu/softmax16_subgroup_program.js +2 -2
  143. package/dist/ops/webgpu/softmax16grad.js +2 -2
  144. package/dist/ops/webgpu/sub16.js +2 -2
  145. package/dist/ops/webgpu/sum16.js +5 -5
  146. package/dist/ops/webgpu/transpose16.js +3 -3
  147. package/dist/ops/webgpu/transpose16_program.js +2 -2
  148. package/dist/ops/webgpu/transpose16_shared_program.js +4 -4
  149. package/dist/ops/webgpu/unpack16.js +4 -4
  150. package/dist/ops/webgpu/utils/binary_op.js +4 -4
  151. package/dist/ops/webgpu/utils/reductions.js +5 -5
  152. package/dist/{ops-CNI3TwqM.js → ops-BFDtP6th.js} +24 -24
  153. package/dist/{pack16-CFUqumar.js → pack16-CmVZs6af.js} +3 -3
  154. package/dist/patches/PackedTensor.js +1 -1
  155. package/dist/patches/engine.js +7 -5
  156. package/dist/patches/tape.js +1 -1
  157. package/dist/patches/webgpu_backend.js +5 -5
  158. package/dist/patches/webgpu_base.js +1 -1
  159. package/dist/patches/webgpu_program.js +3 -3
  160. package/dist/{random_width-DY6Kk2Dl.js → random_width-BVV9HveY.js} +31 -31
  161. package/dist/{range-BMS52eQi.js → range-ZZZD60Fx.js} +2 -2
  162. package/dist/{reciprocal-CTmshQ9J.js → reciprocal-CrYlsAGD.js} +2 -2
  163. package/dist/{register_all_kernels-Bwu1PTuU.js → register_all_kernels-nvj2k7OC.js} +41 -41
  164. package/dist/{relu-yZ2-7WxU.js → relu-BYDneVPn.js} +2 -2
  165. package/dist/{reshape-DevtBWtf.js → reshape-CaPQzFvz.js} +2 -2
  166. package/dist/{rope-B5UUMsPi.js → rope-s4W2XO9B.js} +5 -5
  167. package/dist/{scatter_nd_util-5EL-8VAQ.js → scatter_nd_util-C7zXRT_h.js} +1 -1
  168. package/dist/{selu_util-D1w6yyTO.js → selu_util-BGPXmd4B.js} +16 -16
  169. package/dist/{shared-BRksrJb3.js → shared-CHhxz-O5.js} +1 -1
  170. package/dist/{shared-BuAXb4CI.js → shared-D2NP_CpY.js} +8 -8
  171. package/dist/{sin-BGfy2HZo.js → sin-Djs4aQiu.js} +2 -2
  172. package/dist/{slice-D_gkkqZK.js → slice-DvovR5wq.js} +2 -2
  173. package/dist/{slice_util-DtEldBfK.js → slice_util-DyjSAD0u.js} +1 -1
  174. package/dist/{softmax-ZHVebtR1.js → softmax-C9JQEtnO.js} +2 -2
  175. package/dist/{split-DrfihRpZ.js → split-DBck65sX.js} +2 -2
  176. package/dist/{squeeze-DZEpeblb.js → squeeze-C00Ipm_7.js} +3 -3
  177. package/dist/{stack-yOIAalTq.js → stack-ChnHwRpX.js} +3 -3
  178. package/dist/{sum-_fzj5ZTB.js → sum-ywRJj3Zr.js} +2 -2
  179. package/dist/{tensor-f35l8Odg.js → tensor-0r5yOo2R.js} +1 -1
  180. package/dist/{tensor-DdQUJZlz.js → tensor-CzmOBsdf.js} +21 -21
  181. package/dist/{tensor1d-CeZuc-Rv.js → tensor1d-BlUT89BP.js} +2 -2
  182. package/dist/{tensor2d-G4Ys2GxX.js → tensor2d-CSB4KOb0.js} +2 -2
  183. package/dist/{tensor4d-B8roDgtc.js → tensor4d-D7bLqGqz.js} +2 -2
  184. package/dist/{tensor_util-DV-FP5Q3.js → tensor_util-DfwaWayG.js} +12 -12
  185. package/dist/{tfjs_backend-kNyO5L2d.js → tfjs_backend-CNkSTL0c.js} +38 -38
  186. package/dist/{tile-BzyEiF-F.js → tile-CR074jmp.js} +3 -3
  187. package/dist/training/Adam.js +2 -2
  188. package/dist/training/AdamExt.js +1 -1
  189. package/dist/training/DatasetBuilder.js +2 -2
  190. package/dist/training/FullTrainer.js +1 -1
  191. package/dist/training/Trainer.js +2 -2
  192. package/dist/training/sparseCrossEntropy.js +3 -3
  193. package/dist/{transpose-DKELTqhe.js → transpose-DH4gmHvu.js} +4 -4
  194. package/dist/utilities/dummy.js +3 -3
  195. package/dist/utilities/multinomialCPU.js +2 -2
  196. package/dist/utilities/packed.js +338 -304
  197. package/dist/utilities/performance.js +1 -1
  198. package/dist/utilities/profile.js +1 -1
  199. package/dist/utilities/safetensors.js +2 -2
  200. package/dist/utilities/sentences.js +5 -5
  201. package/dist/utilities/weights.js +2 -2
  202. package/dist/{variable-Bhn5bHYv.js → variable-DzfrwYuP.js} +1 -1
  203. package/dist/{webgpu_program-Cigz-7RF.js → webgpu_program-DzaQiqel.js} +2 -2
  204. package/dist/{webgpu_util-BBCnKm2X.js → webgpu_util-0_ubCEHJ.js} +2 -2
  205. package/dist/{zeros-2gldETuK.js → zeros-DBFVbpv5.js} +3 -3
  206. package/package.json +1 -1
@@ -1,4 +1,4 @@
1
- import { t as s } from "../index-ZyQhjEPo.js";
1
+ import { t as s } from "../index-D6Q1lPZO.js";
2
2
  async function f(e, o = 10, r = !1) {
3
3
  for (let t = 0; t < 100; t++) {
4
4
  const a = r ? await e() : s(e);
@@ -1,4 +1,4 @@
1
- import { a } from "../index-ZyQhjEPo.js";
1
+ import { a } from "../index-D6Q1lPZO.js";
2
2
  const s = 1024 * 1024;
3
3
  class l {
4
4
  log = /* @__PURE__ */ new Map();
@@ -1,5 +1,5 @@
1
- import "../index-ZyQhjEPo.js";
2
- import { t as y } from "../tensor-f35l8Odg.js";
1
+ import "../index-D6Q1lPZO.js";
2
+ import { t as y } from "../tensor-0r5yOo2R.js";
3
3
  function l(t) {
4
4
  if (t === "float32") return "F32";
5
5
  if (t === "int32") return "I32";
@@ -1,8 +1,8 @@
1
- import { m as w } from "../index-ZyQhjEPo.js";
2
- import { t as g } from "../tensor2d-G4Ys2GxX.js";
3
- import { e as y } from "../expand_dims-BPG4fwBP.js";
4
- import { s as h } from "../sum-_fzj5ZTB.js";
5
- import { c as T } from "../concat-BHlIJeyT.js";
1
+ import { m as w } from "../index-D6Q1lPZO.js";
2
+ import { t as g } from "../tensor2d-CSB4KOb0.js";
3
+ import { e as y } from "../expand_dims-ouvfxQ1n.js";
4
+ import { s as h } from "../sum-ywRJj3Zr.js";
5
+ import { c as T } from "../concat-DmBLPVGC.js";
6
6
  const p = 16;
7
7
  function A(o, t) {
8
8
  if (!t)
@@ -1,5 +1,5 @@
1
- import "../index-ZyQhjEPo.js";
2
- import { t as p } from "../tensor-f35l8Odg.js";
1
+ import "../index-D6Q1lPZO.js";
2
+ import { t as p } from "../tensor-0r5yOo2R.js";
3
3
  function h(n) {
4
4
  const e = n.reduce((s, o) => s + o.length, 0), a = new Float32Array(e);
5
5
  let t = 0;
@@ -1,4 +1,4 @@
1
- import { E as i } from "./index-ZyQhjEPo.js";
1
+ import { E as i } from "./index-D6Q1lPZO.js";
2
2
  function m(r, a = !0, e, t) {
3
3
  return i.makeVariable(r, a, e, t);
4
4
  }
@@ -1,5 +1,5 @@
1
- import { h as k } from "./index-ZyQhjEPo.js";
2
- import { b as z, e as E, i as j, a as A } from "./tensor-DdQUJZlz.js";
1
+ import { h as k } from "./index-D6Q1lPZO.js";
2
+ import { b as z, e as E, i as j, a as A } from "./tensor-CzmOBsdf.js";
3
3
  function L(e, s) {
4
4
  if (Math.max(...e) > 5)
5
5
  throw new Error("Cannot symbolically compute strides for rank > 6 tensor.");
@@ -1,5 +1,5 @@
1
- import "./index-ZyQhjEPo.js";
2
- import { a as u } from "./tensor-DdQUJZlz.js";
1
+ import "./index-D6Q1lPZO.js";
2
+ import { a as u } from "./tensor-CzmOBsdf.js";
3
3
  const e = (r) => {
4
4
  let n = 1;
5
5
  for (let t = 0; t < r.length; t++)
@@ -1,6 +1,6 @@
1
- import { E as m } from "./index-ZyQhjEPo.js";
2
- import { c as n } from "./complex-CSlYz-2T.js";
3
- import { d as i, f, s as c } from "./tensor-DdQUJZlz.js";
1
+ import { E as m } from "./index-D6Q1lPZO.js";
2
+ import { c as n } from "./complex-BBiRlsVq.js";
3
+ import { d as i, f, s as c } from "./tensor-CzmOBsdf.js";
4
4
  function e(o, r = "float32") {
5
5
  if (i(o), r === "complex64") {
6
6
  const a = e(o, "float32"), t = e(o, "float32");
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@genai-fi/nanogpt",
3
- "version": "0.10.1",
3
+ "version": "0.10.2",
4
4
  "type": "module",
5
5
  "main": "dist/main.js",
6
6
  "types": "dist/main.d.ts",