cui-llama.rn 1.4.3 → 1.4.6
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +93 -114
- package/android/src/main/CMakeLists.txt +5 -0
- package/android/src/main/java/com/rnllama/LlamaContext.java +91 -17
- package/android/src/main/java/com/rnllama/RNLlama.java +37 -4
- package/android/src/main/jni-utils.h +6 -0
- package/android/src/main/jni.cpp +289 -31
- package/android/src/main/jniLibs/arm64-v8a/librnllama.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2_dotprod.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2_dotprod_i8mm.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2_i8mm.so +0 -0
- package/android/src/main/jniLibs/x86_64/librnllama.so +0 -0
- package/android/src/main/jniLibs/x86_64/librnllama_x86_64.so +0 -0
- package/android/src/newarch/java/com/rnllama/RNLlamaModule.java +7 -2
- package/android/src/oldarch/java/com/rnllama/RNLlamaModule.java +7 -2
- package/cpp/chat-template.hpp +529 -0
- package/cpp/chat.cpp +1779 -0
- package/cpp/chat.h +135 -0
- package/cpp/common.cpp +2064 -1873
- package/cpp/common.h +700 -699
- package/cpp/ggml-alloc.c +1039 -1042
- package/cpp/ggml-alloc.h +1 -1
- package/cpp/ggml-backend-impl.h +255 -255
- package/cpp/ggml-backend-reg.cpp +586 -582
- package/cpp/ggml-backend.cpp +2004 -2002
- package/cpp/ggml-backend.h +354 -354
- package/cpp/ggml-common.h +1851 -1853
- package/cpp/ggml-cpp.h +39 -39
- package/cpp/ggml-cpu-aarch64.cpp +4248 -4247
- package/cpp/ggml-cpu-aarch64.h +8 -8
- package/cpp/ggml-cpu-impl.h +531 -386
- package/cpp/ggml-cpu-quants.c +12527 -10920
- package/cpp/ggml-cpu-traits.cpp +36 -36
- package/cpp/ggml-cpu-traits.h +38 -38
- package/cpp/ggml-cpu.c +15766 -14391
- package/cpp/ggml-cpu.cpp +655 -635
- package/cpp/ggml-cpu.h +138 -135
- package/cpp/ggml-impl.h +567 -567
- package/cpp/ggml-metal-impl.h +235 -0
- package/cpp/ggml-metal.h +1 -1
- package/cpp/ggml-metal.m +5146 -4884
- package/cpp/ggml-opt.cpp +854 -854
- package/cpp/ggml-opt.h +216 -216
- package/cpp/ggml-quants.c +5238 -5238
- package/cpp/ggml-threading.h +14 -14
- package/cpp/ggml.c +6529 -6514
- package/cpp/ggml.h +2198 -2194
- package/cpp/gguf.cpp +1329 -1329
- package/cpp/gguf.h +202 -202
- package/cpp/json-schema-to-grammar.cpp +1024 -1045
- package/cpp/json-schema-to-grammar.h +21 -8
- package/cpp/json.hpp +24766 -24766
- package/cpp/llama-adapter.cpp +347 -347
- package/cpp/llama-adapter.h +74 -74
- package/cpp/llama-arch.cpp +1513 -1487
- package/cpp/llama-arch.h +403 -400
- package/cpp/llama-batch.cpp +368 -368
- package/cpp/llama-batch.h +88 -88
- package/cpp/llama-chat.cpp +588 -578
- package/cpp/llama-chat.h +53 -52
- package/cpp/llama-context.cpp +1775 -1775
- package/cpp/llama-context.h +128 -128
- package/cpp/llama-cparams.cpp +1 -1
- package/cpp/llama-cparams.h +37 -37
- package/cpp/llama-cpp.h +30 -30
- package/cpp/llama-grammar.cpp +1219 -1139
- package/cpp/llama-grammar.h +173 -143
- package/cpp/llama-hparams.cpp +71 -71
- package/cpp/llama-hparams.h +139 -139
- package/cpp/llama-impl.cpp +167 -167
- package/cpp/llama-impl.h +61 -61
- package/cpp/llama-kv-cache.cpp +718 -718
- package/cpp/llama-kv-cache.h +219 -218
- package/cpp/llama-mmap.cpp +600 -590
- package/cpp/llama-mmap.h +68 -67
- package/cpp/llama-model-loader.cpp +1124 -1124
- package/cpp/llama-model-loader.h +167 -167
- package/cpp/llama-model.cpp +4087 -3997
- package/cpp/llama-model.h +370 -370
- package/cpp/llama-sampling.cpp +2558 -2408
- package/cpp/llama-sampling.h +32 -32
- package/cpp/llama-vocab.cpp +3264 -3247
- package/cpp/llama-vocab.h +125 -125
- package/cpp/llama.cpp +10284 -10077
- package/cpp/llama.h +1354 -1323
- package/cpp/log.cpp +393 -401
- package/cpp/log.h +132 -121
- package/cpp/minja/chat-template.hpp +529 -0
- package/cpp/minja/minja.hpp +2915 -0
- package/cpp/minja.hpp +2915 -0
- package/cpp/rn-llama.cpp +66 -6
- package/cpp/rn-llama.h +26 -1
- package/cpp/sampling.cpp +570 -505
- package/cpp/sampling.h +3 -0
- package/cpp/sgemm.cpp +2598 -2597
- package/cpp/sgemm.h +14 -14
- package/cpp/speculative.cpp +278 -277
- package/cpp/speculative.h +28 -28
- package/cpp/unicode.cpp +9 -2
- package/ios/CMakeLists.txt +6 -0
- package/ios/RNLlama.h +0 -8
- package/ios/RNLlama.mm +27 -3
- package/ios/RNLlamaContext.h +10 -1
- package/ios/RNLlamaContext.mm +269 -57
- package/jest/mock.js +21 -2
- package/lib/commonjs/NativeRNLlama.js.map +1 -1
- package/lib/commonjs/grammar.js +3 -0
- package/lib/commonjs/grammar.js.map +1 -1
- package/lib/commonjs/index.js +87 -13
- package/lib/commonjs/index.js.map +1 -1
- package/lib/module/NativeRNLlama.js.map +1 -1
- package/lib/module/grammar.js +3 -0
- package/lib/module/grammar.js.map +1 -1
- package/lib/module/index.js +86 -13
- package/lib/module/index.js.map +1 -1
- package/lib/typescript/NativeRNLlama.d.ts +107 -2
- package/lib/typescript/NativeRNLlama.d.ts.map +1 -1
- package/lib/typescript/grammar.d.ts.map +1 -1
- package/lib/typescript/index.d.ts +32 -7
- package/lib/typescript/index.d.ts.map +1 -1
- package/llama-rn.podspec +1 -1
- package/package.json +3 -2
- package/src/NativeRNLlama.ts +115 -3
- package/src/grammar.ts +3 -0
- package/src/index.ts +138 -21
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CMakeCCompiler.cmake +0 -81
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CMakeSystem.cmake +0 -15
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdC/CMakeCCompilerId.c +0 -904
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdC/CMakeCCompilerId.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdCXX/CMakeCXXCompilerId.cpp +0 -919
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdCXX/CMakeCXXCompilerId.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/CMakeConfigureLog.yaml +0 -55
- package/cpp/rn-llama.hpp +0 -913
package/cpp/llama-chat.h
CHANGED
@@ -1,52 +1,53 @@
|
|
1
|
-
#pragma once
|
2
|
-
|
3
|
-
#include <string>
|
4
|
-
#include <vector>
|
5
|
-
#include <cstdint>
|
6
|
-
|
7
|
-
enum llm_chat_template {
|
8
|
-
LLM_CHAT_TEMPLATE_CHATML,
|
9
|
-
LLM_CHAT_TEMPLATE_LLAMA_2,
|
10
|
-
LLM_CHAT_TEMPLATE_LLAMA_2_SYS,
|
11
|
-
LLM_CHAT_TEMPLATE_LLAMA_2_SYS_BOS,
|
12
|
-
LLM_CHAT_TEMPLATE_LLAMA_2_SYS_STRIP,
|
13
|
-
LLM_CHAT_TEMPLATE_MISTRAL_V1,
|
14
|
-
LLM_CHAT_TEMPLATE_MISTRAL_V3,
|
15
|
-
LLM_CHAT_TEMPLATE_MISTRAL_V3_TEKKEN,
|
16
|
-
LLM_CHAT_TEMPLATE_MISTRAL_V7,
|
17
|
-
LLM_CHAT_TEMPLATE_PHI_3,
|
18
|
-
LLM_CHAT_TEMPLATE_PHI_4,
|
19
|
-
LLM_CHAT_TEMPLATE_FALCON_3,
|
20
|
-
LLM_CHAT_TEMPLATE_ZEPHYR,
|
21
|
-
LLM_CHAT_TEMPLATE_MONARCH,
|
22
|
-
LLM_CHAT_TEMPLATE_GEMMA,
|
23
|
-
LLM_CHAT_TEMPLATE_ORION,
|
24
|
-
LLM_CHAT_TEMPLATE_OPENCHAT,
|
25
|
-
LLM_CHAT_TEMPLATE_VICUNA,
|
26
|
-
LLM_CHAT_TEMPLATE_VICUNA_ORCA,
|
27
|
-
LLM_CHAT_TEMPLATE_DEEPSEEK,
|
28
|
-
LLM_CHAT_TEMPLATE_DEEPSEEK_2,
|
29
|
-
LLM_CHAT_TEMPLATE_DEEPSEEK_3,
|
30
|
-
LLM_CHAT_TEMPLATE_COMMAND_R,
|
31
|
-
LLM_CHAT_TEMPLATE_LLAMA_3,
|
32
|
-
LLM_CHAT_TEMPLATE_CHATGML_3,
|
33
|
-
LLM_CHAT_TEMPLATE_CHATGML_4,
|
34
|
-
|
35
|
-
|
36
|
-
|
37
|
-
|
38
|
-
|
39
|
-
|
40
|
-
|
41
|
-
|
42
|
-
|
43
|
-
|
44
|
-
|
45
|
-
|
46
|
-
|
47
|
-
|
48
|
-
|
49
|
-
|
50
|
-
|
51
|
-
|
52
|
-
std::
|
1
|
+
#pragma once
|
2
|
+
|
3
|
+
#include <string>
|
4
|
+
#include <vector>
|
5
|
+
#include <cstdint>
|
6
|
+
|
7
|
+
enum llm_chat_template {
|
8
|
+
LLM_CHAT_TEMPLATE_CHATML,
|
9
|
+
LLM_CHAT_TEMPLATE_LLAMA_2,
|
10
|
+
LLM_CHAT_TEMPLATE_LLAMA_2_SYS,
|
11
|
+
LLM_CHAT_TEMPLATE_LLAMA_2_SYS_BOS,
|
12
|
+
LLM_CHAT_TEMPLATE_LLAMA_2_SYS_STRIP,
|
13
|
+
LLM_CHAT_TEMPLATE_MISTRAL_V1,
|
14
|
+
LLM_CHAT_TEMPLATE_MISTRAL_V3,
|
15
|
+
LLM_CHAT_TEMPLATE_MISTRAL_V3_TEKKEN,
|
16
|
+
LLM_CHAT_TEMPLATE_MISTRAL_V7,
|
17
|
+
LLM_CHAT_TEMPLATE_PHI_3,
|
18
|
+
LLM_CHAT_TEMPLATE_PHI_4,
|
19
|
+
LLM_CHAT_TEMPLATE_FALCON_3,
|
20
|
+
LLM_CHAT_TEMPLATE_ZEPHYR,
|
21
|
+
LLM_CHAT_TEMPLATE_MONARCH,
|
22
|
+
LLM_CHAT_TEMPLATE_GEMMA,
|
23
|
+
LLM_CHAT_TEMPLATE_ORION,
|
24
|
+
LLM_CHAT_TEMPLATE_OPENCHAT,
|
25
|
+
LLM_CHAT_TEMPLATE_VICUNA,
|
26
|
+
LLM_CHAT_TEMPLATE_VICUNA_ORCA,
|
27
|
+
LLM_CHAT_TEMPLATE_DEEPSEEK,
|
28
|
+
LLM_CHAT_TEMPLATE_DEEPSEEK_2,
|
29
|
+
LLM_CHAT_TEMPLATE_DEEPSEEK_3,
|
30
|
+
LLM_CHAT_TEMPLATE_COMMAND_R,
|
31
|
+
LLM_CHAT_TEMPLATE_LLAMA_3,
|
32
|
+
LLM_CHAT_TEMPLATE_CHATGML_3,
|
33
|
+
LLM_CHAT_TEMPLATE_CHATGML_4,
|
34
|
+
LLM_CHAT_TEMPLATE_GLMEDGE,
|
35
|
+
LLM_CHAT_TEMPLATE_MINICPM,
|
36
|
+
LLM_CHAT_TEMPLATE_EXAONE_3,
|
37
|
+
LLM_CHAT_TEMPLATE_RWKV_WORLD,
|
38
|
+
LLM_CHAT_TEMPLATE_GRANITE,
|
39
|
+
LLM_CHAT_TEMPLATE_GIGACHAT,
|
40
|
+
LLM_CHAT_TEMPLATE_MEGREZ,
|
41
|
+
LLM_CHAT_TEMPLATE_UNKNOWN,
|
42
|
+
};
|
43
|
+
|
44
|
+
struct llama_chat_message;
|
45
|
+
|
46
|
+
llm_chat_template llm_chat_template_from_str(const std::string & name);
|
47
|
+
|
48
|
+
llm_chat_template llm_chat_detect_template(const std::string & tmpl);
|
49
|
+
|
50
|
+
int32_t llm_chat_apply_template(
|
51
|
+
llm_chat_template tmpl,
|
52
|
+
const std::vector<const llama_chat_message *> & chat,
|
53
|
+
std::string & dest, bool add_ass);
|