cui-llama.rn 1.4.4 → 1.4.6
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/android/src/main/CMakeLists.txt +2 -2
- package/android/src/main/jni.cpp +12 -10
- package/android/src/main/jniLibs/arm64-v8a/librnllama.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2_dotprod.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2_dotprod_i8mm.so +0 -0
- package/android/src/main/jniLibs/arm64-v8a/librnllama_v8_2_i8mm.so +0 -0
- package/android/src/main/jniLibs/x86_64/librnllama.so +0 -0
- package/android/src/main/jniLibs/x86_64/librnllama_x86_64.so +0 -0
- package/cpp/chat-template.hpp +529 -529
- package/cpp/chat.cpp +959 -265
- package/cpp/chat.h +135 -0
- package/cpp/common.cpp +2064 -1996
- package/cpp/common.h +700 -744
- package/cpp/ggml-alloc.c +1039 -1030
- package/cpp/ggml-alloc.h +1 -1
- package/cpp/ggml-backend-impl.h +255 -255
- package/cpp/ggml-backend-reg.cpp +586 -582
- package/cpp/ggml-backend.cpp +2004 -2002
- package/cpp/ggml-backend.h +354 -354
- package/cpp/ggml-common.h +1851 -1851
- package/cpp/ggml-cpp.h +39 -39
- package/cpp/ggml-cpu-aarch64.cpp +4248 -4247
- package/cpp/ggml-cpu-aarch64.h +8 -8
- package/cpp/ggml-cpu-impl.h +531 -380
- package/cpp/ggml-cpu-quants.c +12527 -11517
- package/cpp/ggml-cpu-traits.cpp +36 -36
- package/cpp/ggml-cpu-traits.h +38 -38
- package/cpp/ggml-cpu.c +15766 -14485
- package/cpp/ggml-cpu.cpp +655 -633
- package/cpp/ggml-cpu.h +138 -135
- package/cpp/ggml-impl.h +567 -567
- package/cpp/ggml-metal-impl.h +235 -0
- package/cpp/ggml-metal.h +66 -66
- package/cpp/ggml-metal.m +5146 -5002
- package/cpp/ggml-opt.cpp +854 -854
- package/cpp/ggml-opt.h +216 -216
- package/cpp/ggml-quants.c +5238 -5238
- package/cpp/ggml-threading.h +14 -14
- package/cpp/ggml.c +6529 -6524
- package/cpp/ggml.h +2198 -2194
- package/cpp/gguf.cpp +1329 -1329
- package/cpp/gguf.h +202 -202
- package/cpp/json-schema-to-grammar.cpp +1024 -1025
- package/cpp/json-schema-to-grammar.h +21 -22
- package/cpp/json.hpp +24766 -24766
- package/cpp/llama-adapter.cpp +347 -347
- package/cpp/llama-adapter.h +74 -74
- package/cpp/llama-arch.cpp +1513 -1492
- package/cpp/llama-arch.h +403 -402
- package/cpp/llama-batch.cpp +368 -368
- package/cpp/llama-batch.h +88 -88
- package/cpp/llama-chat.cpp +588 -587
- package/cpp/llama-chat.h +53 -53
- package/cpp/llama-context.cpp +1775 -1775
- package/cpp/llama-context.h +128 -128
- package/cpp/llama-cparams.cpp +1 -1
- package/cpp/llama-cparams.h +37 -37
- package/cpp/llama-cpp.h +30 -30
- package/cpp/llama-grammar.cpp +1219 -1219
- package/cpp/llama-grammar.h +173 -164
- package/cpp/llama-hparams.cpp +71 -71
- package/cpp/llama-hparams.h +139 -139
- package/cpp/llama-impl.cpp +167 -167
- package/cpp/llama-impl.h +61 -61
- package/cpp/llama-kv-cache.cpp +718 -718
- package/cpp/llama-kv-cache.h +219 -218
- package/cpp/llama-mmap.cpp +600 -590
- package/cpp/llama-mmap.h +68 -68
- package/cpp/llama-model-loader.cpp +1124 -1124
- package/cpp/llama-model-loader.h +167 -167
- package/cpp/llama-model.cpp +4087 -4023
- package/cpp/llama-model.h +370 -370
- package/cpp/llama-sampling.cpp +2558 -2525
- package/cpp/llama-sampling.h +32 -32
- package/cpp/llama-vocab.cpp +3264 -3252
- package/cpp/llama-vocab.h +125 -125
- package/cpp/llama.cpp +10284 -10137
- package/cpp/llama.h +1354 -1340
- package/cpp/log.cpp +393 -423
- package/cpp/log.h +132 -132
- package/cpp/minja/chat-template.hpp +529 -0
- package/cpp/minja/minja.hpp +2915 -0
- package/cpp/minja.hpp +2915 -2883
- package/cpp/rn-llama.cpp +20 -37
- package/cpp/rn-llama.h +12 -2
- package/cpp/sampling.cpp +570 -532
- package/cpp/sgemm.cpp +2598 -2598
- package/cpp/sgemm.h +14 -14
- package/cpp/speculative.cpp +278 -277
- package/cpp/speculative.h +28 -28
- package/package.json +1 -1
- package/android/src/main/build-arm64/CMakeCache.txt +0 -429
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CMakeCCompiler.cmake +0 -81
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CMakeCXXCompiler.cmake +0 -101
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CMakeDetermineCompilerABI_C.bin +0 -0
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CMakeDetermineCompilerABI_CXX.bin +0 -0
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CMakeSystem.cmake +0 -15
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdC/CMakeCCompilerId.c +0 -904
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdC/CMakeCCompilerId.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdCXX/CMakeCXXCompilerId.cpp +0 -919
- package/android/src/main/build-arm64/CMakeFiles/3.31.4/CompilerIdCXX/CMakeCXXCompilerId.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/CMakeConfigureLog.yaml +0 -431
- package/android/src/main/build-arm64/CMakeFiles/CMakeDirectoryInformation.cmake +0 -16
- package/android/src/main/build-arm64/CMakeFiles/Makefile.cmake +0 -165
- package/android/src/main/build-arm64/CMakeFiles/Makefile2 +0 -297
- package/android/src/main/build-arm64/CMakeFiles/Progress/1 +0 -1
- package/android/src/main/build-arm64/CMakeFiles/Progress/2 +0 -1
- package/android/src/main/build-arm64/CMakeFiles/Progress/3 +0 -1
- package/android/src/main/build-arm64/CMakeFiles/Progress/4 +0 -1
- package/android/src/main/build-arm64/CMakeFiles/Progress/5 +0 -1
- package/android/src/main/build-arm64/CMakeFiles/Progress/6 +0 -1
- package/android/src/main/build-arm64/CMakeFiles/Progress/count.txt +0 -1
- package/android/src/main/build-arm64/CMakeFiles/TargetDirectories.txt +0 -8
- package/android/src/main/build-arm64/CMakeFiles/cmake.check_cache +0 -1
- package/android/src/main/build-arm64/CMakeFiles/progress.marks +0 -1
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-alloc.c.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-alloc.c.o.d +0 -58
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-backend-reg.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-backend-reg.cpp.o.d +0 -756
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-backend.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-backend.cpp.o.d +0 -709
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu-aarch64.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu-aarch64.cpp.o.d +0 -714
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu-quants.c.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu-quants.c.o.d +0 -62
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu-traits.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu-traits.cpp.o.d +0 -708
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu.c.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu.c.o.d +0 -113
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-cpu.cpp.o.d +0 -713
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-opt.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-opt.cpp.o.d +0 -763
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-quants.c.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-quants.c.o.d +0 -61
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-threading.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml-threading.cpp.o.d +0 -707
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml.c.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/ggml.c.o.d +0 -104
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/gguf.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/gguf.cpp.o.d +0 -714
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/log.cpp.o +0 -0
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/D_/dev/react-native/cui-llama.rn/cpp/log.cpp.o.d +0 -723
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/DependInfo.cmake +0 -62
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/build.make +0 -722
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/cmake_clean.cmake +0 -89
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/compiler_depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/compiler_depend.ts +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/flags.make +0 -17
- package/android/src/main/build-arm64/CMakeFiles/rnllama.dir/progress.make +0 -41
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/DependInfo.cmake +0 -62
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/build.make +0 -722
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/cmake_clean.cmake +0 -89
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/compiler_depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/compiler_depend.ts +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/flags.make +0 -17
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8.dir/progress.make +0 -41
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/DependInfo.cmake +0 -62
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/build.make +0 -722
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/cmake_clean.cmake +0 -89
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/compiler_depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/compiler_depend.ts +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/flags.make +0 -17
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2.dir/progress.make +0 -41
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/DependInfo.cmake +0 -62
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/build.make +0 -722
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/cmake_clean.cmake +0 -89
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/compiler_depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/compiler_depend.ts +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/flags.make +0 -17
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod.dir/progress.make +0 -41
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/DependInfo.cmake +0 -62
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/build.make +0 -722
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/cmake_clean.cmake +0 -89
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/compiler_depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/compiler_depend.ts +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/flags.make +0 -17
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_dotprod_i8mm.dir/progress.make +0 -41
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/DependInfo.cmake +0 -62
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/build.make +0 -722
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/cmake_clean.cmake +0 -89
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/compiler_depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/compiler_depend.ts +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/depend.make +0 -2
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/flags.make +0 -17
- package/android/src/main/build-arm64/CMakeFiles/rnllama_v8_2_i8mm.dir/progress.make +0 -41
- package/android/src/main/build-arm64/Makefile +0 -1862
- package/android/src/main/build-arm64/cmake_install.cmake +0 -66
- package/cpp/chat.hpp +0 -55
- package/cpp/rn-llama.hpp +0 -913
package/cpp/llama-vocab.h
CHANGED
@@ -1,125 +1,125 @@
|
|
1
|
-
#pragma once
|
2
|
-
|
3
|
-
#include "llama.h"
|
4
|
-
|
5
|
-
#include <string>
|
6
|
-
#include <vector>
|
7
|
-
#include <memory>
|
8
|
-
|
9
|
-
struct LLM_KV;
|
10
|
-
struct llama_model_loader;
|
11
|
-
|
12
|
-
struct llama_vocab {
|
13
|
-
struct token_data {
|
14
|
-
std::string text;
|
15
|
-
float score;
|
16
|
-
llama_token_attr attr;
|
17
|
-
};
|
18
|
-
|
19
|
-
llama_vocab();
|
20
|
-
~llama_vocab();
|
21
|
-
|
22
|
-
void load(llama_model_loader & ml, const LLM_KV & kv);
|
23
|
-
|
24
|
-
enum llama_vocab_type get_type() const;
|
25
|
-
enum llama_vocab_pre_type get_pre_type() const;
|
26
|
-
|
27
|
-
uint32_t n_tokens() const;
|
28
|
-
uint32_t n_token_types() const;
|
29
|
-
|
30
|
-
std::string type_name() const;
|
31
|
-
|
32
|
-
bool is_normal (llama_token id) const;
|
33
|
-
bool is_unknown (llama_token id) const;
|
34
|
-
bool is_control (llama_token id) const;
|
35
|
-
bool is_byte (llama_token id) const;
|
36
|
-
bool is_user_defined(llama_token id) const;
|
37
|
-
bool is_unused (llama_token id) const;
|
38
|
-
bool is_eog (llama_token id) const;
|
39
|
-
|
40
|
-
uint8_t token_to_byte(llama_token id) const;
|
41
|
-
llama_token byte_to_token(uint8_t ch) const;
|
42
|
-
|
43
|
-
llama_token text_to_token(const std::string & text) const;
|
44
|
-
|
45
|
-
const token_data & get_token_data(llama_token id) const;
|
46
|
-
|
47
|
-
const char * token_get_text (llama_token id) const;
|
48
|
-
float token_get_score(llama_token id) const;
|
49
|
-
llama_token_attr token_get_attr (llama_token id) const;
|
50
|
-
|
51
|
-
llama_token token_bos() const;
|
52
|
-
llama_token token_eos() const;
|
53
|
-
llama_token token_eot() const;
|
54
|
-
llama_token token_eom() const;
|
55
|
-
llama_token token_unk() const;
|
56
|
-
llama_token token_sep() const;
|
57
|
-
llama_token token_nl () const;
|
58
|
-
llama_token token_pad() const;
|
59
|
-
|
60
|
-
llama_token token_prefix() const;
|
61
|
-
llama_token token_middle() const;
|
62
|
-
llama_token token_suffix() const;
|
63
|
-
|
64
|
-
llama_token token_fim_pre() const;
|
65
|
-
llama_token token_fim_suf() const;
|
66
|
-
llama_token token_fim_mid() const;
|
67
|
-
llama_token token_fim_pad() const;
|
68
|
-
llama_token token_fim_rep() const;
|
69
|
-
llama_token token_fim_sep() const;
|
70
|
-
|
71
|
-
bool get_add_space_prefix () const;
|
72
|
-
bool get_add_bos () const;
|
73
|
-
bool get_add_eos () const;
|
74
|
-
bool get_ignore_merges () const;
|
75
|
-
bool get_clean_spaces () const;
|
76
|
-
bool get_remove_extra_whitespaces () const;
|
77
|
-
bool get_escape_whitespaces () const;
|
78
|
-
bool get_treat_whitespace_as_suffix() const;
|
79
|
-
|
80
|
-
int max_token_len() const;
|
81
|
-
|
82
|
-
int find_bpe_rank(const std::string & token_left, const std::string & token_right) const;
|
83
|
-
|
84
|
-
int32_t tokenize(
|
85
|
-
const char * text,
|
86
|
-
int32_t text_len,
|
87
|
-
llama_token * tokens,
|
88
|
-
int32_t n_tokens_max,
|
89
|
-
bool add_special,
|
90
|
-
bool parse_special) const;
|
91
|
-
|
92
|
-
std::vector<llama_token> tokenize(
|
93
|
-
const std::string & raw_text,
|
94
|
-
bool add_special,
|
95
|
-
bool parse_special = false) const;
|
96
|
-
|
97
|
-
// does not write null-terminator to buf
|
98
|
-
int32_t token_to_piece(
|
99
|
-
llama_token token,
|
100
|
-
char * buf,
|
101
|
-
int32_t length,
|
102
|
-
int32_t lstrip,
|
103
|
-
bool special) const;
|
104
|
-
|
105
|
-
// use cached data
|
106
|
-
const std::string & token_to_piece(llama_token token) const;
|
107
|
-
|
108
|
-
int32_t detokenize(
|
109
|
-
const llama_token * tokens,
|
110
|
-
int32_t n_tokens,
|
111
|
-
char * text,
|
112
|
-
int32_t text_len_max,
|
113
|
-
bool remove_special,
|
114
|
-
bool unparse_special) const;
|
115
|
-
|
116
|
-
std::string detokenize(
|
117
|
-
const std::vector<llama_token> & tokens,
|
118
|
-
bool special) const;
|
119
|
-
|
120
|
-
void print_info() const;
|
121
|
-
|
122
|
-
private:
|
123
|
-
struct impl;
|
124
|
-
std::unique_ptr<impl> pimpl;
|
125
|
-
};
|
1
|
+
#pragma once
|
2
|
+
|
3
|
+
#include "llama.h"
|
4
|
+
|
5
|
+
#include <string>
|
6
|
+
#include <vector>
|
7
|
+
#include <memory>
|
8
|
+
|
9
|
+
struct LLM_KV;
|
10
|
+
struct llama_model_loader;
|
11
|
+
|
12
|
+
struct llama_vocab {
|
13
|
+
struct token_data {
|
14
|
+
std::string text;
|
15
|
+
float score;
|
16
|
+
llama_token_attr attr;
|
17
|
+
};
|
18
|
+
|
19
|
+
llama_vocab();
|
20
|
+
~llama_vocab();
|
21
|
+
|
22
|
+
void load(llama_model_loader & ml, const LLM_KV & kv);
|
23
|
+
|
24
|
+
enum llama_vocab_type get_type() const;
|
25
|
+
enum llama_vocab_pre_type get_pre_type() const;
|
26
|
+
|
27
|
+
uint32_t n_tokens() const;
|
28
|
+
uint32_t n_token_types() const;
|
29
|
+
|
30
|
+
std::string type_name() const;
|
31
|
+
|
32
|
+
bool is_normal (llama_token id) const;
|
33
|
+
bool is_unknown (llama_token id) const;
|
34
|
+
bool is_control (llama_token id) const;
|
35
|
+
bool is_byte (llama_token id) const;
|
36
|
+
bool is_user_defined(llama_token id) const;
|
37
|
+
bool is_unused (llama_token id) const;
|
38
|
+
bool is_eog (llama_token id) const;
|
39
|
+
|
40
|
+
uint8_t token_to_byte(llama_token id) const;
|
41
|
+
llama_token byte_to_token(uint8_t ch) const;
|
42
|
+
|
43
|
+
llama_token text_to_token(const std::string & text) const;
|
44
|
+
|
45
|
+
const token_data & get_token_data(llama_token id) const;
|
46
|
+
|
47
|
+
const char * token_get_text (llama_token id) const;
|
48
|
+
float token_get_score(llama_token id) const;
|
49
|
+
llama_token_attr token_get_attr (llama_token id) const;
|
50
|
+
|
51
|
+
llama_token token_bos() const;
|
52
|
+
llama_token token_eos() const;
|
53
|
+
llama_token token_eot() const;
|
54
|
+
llama_token token_eom() const;
|
55
|
+
llama_token token_unk() const;
|
56
|
+
llama_token token_sep() const;
|
57
|
+
llama_token token_nl () const;
|
58
|
+
llama_token token_pad() const;
|
59
|
+
|
60
|
+
llama_token token_prefix() const;
|
61
|
+
llama_token token_middle() const;
|
62
|
+
llama_token token_suffix() const;
|
63
|
+
|
64
|
+
llama_token token_fim_pre() const;
|
65
|
+
llama_token token_fim_suf() const;
|
66
|
+
llama_token token_fim_mid() const;
|
67
|
+
llama_token token_fim_pad() const;
|
68
|
+
llama_token token_fim_rep() const;
|
69
|
+
llama_token token_fim_sep() const;
|
70
|
+
|
71
|
+
bool get_add_space_prefix () const;
|
72
|
+
bool get_add_bos () const;
|
73
|
+
bool get_add_eos () const;
|
74
|
+
bool get_ignore_merges () const;
|
75
|
+
bool get_clean_spaces () const;
|
76
|
+
bool get_remove_extra_whitespaces () const;
|
77
|
+
bool get_escape_whitespaces () const;
|
78
|
+
bool get_treat_whitespace_as_suffix() const;
|
79
|
+
|
80
|
+
int max_token_len() const;
|
81
|
+
|
82
|
+
int find_bpe_rank(const std::string & token_left, const std::string & token_right) const;
|
83
|
+
|
84
|
+
int32_t tokenize(
|
85
|
+
const char * text,
|
86
|
+
int32_t text_len,
|
87
|
+
llama_token * tokens,
|
88
|
+
int32_t n_tokens_max,
|
89
|
+
bool add_special,
|
90
|
+
bool parse_special) const;
|
91
|
+
|
92
|
+
std::vector<llama_token> tokenize(
|
93
|
+
const std::string & raw_text,
|
94
|
+
bool add_special,
|
95
|
+
bool parse_special = false) const;
|
96
|
+
|
97
|
+
// does not write null-terminator to buf
|
98
|
+
int32_t token_to_piece(
|
99
|
+
llama_token token,
|
100
|
+
char * buf,
|
101
|
+
int32_t length,
|
102
|
+
int32_t lstrip,
|
103
|
+
bool special) const;
|
104
|
+
|
105
|
+
// use cached data
|
106
|
+
const std::string & token_to_piece(llama_token token) const;
|
107
|
+
|
108
|
+
int32_t detokenize(
|
109
|
+
const llama_token * tokens,
|
110
|
+
int32_t n_tokens,
|
111
|
+
char * text,
|
112
|
+
int32_t text_len_max,
|
113
|
+
bool remove_special,
|
114
|
+
bool unparse_special) const;
|
115
|
+
|
116
|
+
std::string detokenize(
|
117
|
+
const std::vector<llama_token> & tokens,
|
118
|
+
bool special) const;
|
119
|
+
|
120
|
+
void print_info() const;
|
121
|
+
|
122
|
+
private:
|
123
|
+
struct impl;
|
124
|
+
std::unique_ptr<impl> pimpl;
|
125
|
+
};
|