modelparams 0.0.13__tar.gz → 0.0.15__tar.gz

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (33) hide show
  1. {modelparams-0.0.13 → modelparams-0.0.15}/PKG-INFO +1 -1
  2. {modelparams-0.0.13 → modelparams-0.0.15}/pyproject.toml +1 -1
  3. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/catalog.json +195 -0
  4. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/model_ids.py +6 -0
  5. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/registry.py +3 -0
  6. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/alibaba.py +43 -0
  7. {modelparams-0.0.13 → modelparams-0.0.15}/.gitignore +0 -0
  8. {modelparams-0.0.13 → modelparams-0.0.15}/LICENSE +0 -0
  9. {modelparams-0.0.13 → modelparams-0.0.15}/README.md +0 -0
  10. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/__init__.py +0 -0
  11. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/_generated/__init__.py +0 -0
  12. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/catalog.py +0 -0
  13. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/models.py +0 -0
  14. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/py.typed +0 -0
  15. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/__init__.py +0 -0
  16. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/anthropic.py +0 -0
  17. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/cerebras.py +0 -0
  18. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/cohere.py +0 -0
  19. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/deepseek.py +0 -0
  20. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/google.py +0 -0
  21. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/groq.py +0 -0
  22. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/meta.py +0 -0
  23. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/minimax.py +0 -0
  24. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/mistral.py +0 -0
  25. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/moonshot.py +0 -0
  26. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/nvidia.py +0 -0
  27. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/openai.py +0 -0
  28. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/perplexity.py +0 -0
  29. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/thinking_machines.py +0 -0
  30. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/xai.py +0 -0
  31. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/xiaomi.py +0 -0
  32. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/types/z_ai.py +0 -0
  33. {modelparams-0.0.13 → modelparams-0.0.15}/src/modelparams/validation.py +0 -0
@@ -1,6 +1,6 @@
1
1
  Metadata-Version: 2.5
2
2
  Name: modelparams
3
- Version: 0.0.13
3
+ Version: 0.0.15
4
4
  Summary: Typed model parameters for Python, generated from the modelparams.dev catalog.
5
5
  Project-URL: Homepage, https://modelparams.dev
6
6
  Project-URL: Repository, https://github.com/mnfst/modelparameters.dev
@@ -4,7 +4,7 @@ build-backend = "hatchling.build"
4
4
 
5
5
  [project]
6
6
  name = "modelparams"
7
- version = "0.0.13"
7
+ version = "0.0.15"
8
8
  description = "Typed model parameters for Python, generated from the modelparams.dev catalog."
9
9
  readme = "README.md"
10
10
  requires-python = ">=3.10"
@@ -59,6 +59,66 @@
59
59
  }
60
60
  ]
61
61
  },
62
+ {
63
+ "provider": "alibaba",
64
+ "authType": "api_key",
65
+ "model": "qwen-max",
66
+ "params": [
67
+ {
68
+ "path": "max_tokens",
69
+ "label": "Max tokens",
70
+ "description": "Maximum number of output tokens the model may generate.",
71
+ "group": "generation_length",
72
+ "type": "integer",
73
+ "range": {
74
+ "min": 1
75
+ }
76
+ },
77
+ {
78
+ "path": "temperature",
79
+ "label": "Temperature",
80
+ "description": "Controls randomness. Lower values make outputs more focused; higher values make them more varied.",
81
+ "group": "sampling",
82
+ "type": "number",
83
+ "range": {
84
+ "min": 0,
85
+ "max": 1.9,
86
+ "step": 0.1
87
+ }
88
+ },
89
+ {
90
+ "path": "top_p",
91
+ "label": "Top P",
92
+ "description": "Controls nucleus sampling by limiting generation to tokens within the selected cumulative probability.",
93
+ "group": "sampling",
94
+ "type": "number",
95
+ "range": {
96
+ "min": 0,
97
+ "max": 1,
98
+ "step": 0.01
99
+ }
100
+ },
101
+ {
102
+ "path": "extra_body.top_k",
103
+ "label": "Top K",
104
+ "description": "Limits generation to the selected number of highest-probability tokens.",
105
+ "group": "sampling",
106
+ "type": "integer",
107
+ "default": 20,
108
+ "range": {
109
+ "min": 1
110
+ }
111
+ },
112
+ {
113
+ "path": "extra_body.chat_template_kwargs.enable_thinking",
114
+ "label": "Enable thinking",
115
+ "description": "Controls Qwen3 thinking mode when using OpenAI-compatible clients that pass provider-specific extra body fields.",
116
+ "group": "reasoning",
117
+ "type": "boolean",
118
+ "default": true
119
+ }
120
+ ]
121
+ },
62
122
  {
63
123
  "provider": "alibaba",
64
124
  "authType": "api_key",
@@ -119,6 +179,66 @@
119
179
  }
120
180
  ]
121
181
  },
182
+ {
183
+ "provider": "alibaba",
184
+ "authType": "api_key",
185
+ "model": "qwen-turbo",
186
+ "params": [
187
+ {
188
+ "path": "max_tokens",
189
+ "label": "Max tokens",
190
+ "description": "Maximum number of output tokens the model may generate.",
191
+ "group": "generation_length",
192
+ "type": "integer",
193
+ "range": {
194
+ "min": 1
195
+ }
196
+ },
197
+ {
198
+ "path": "temperature",
199
+ "label": "Temperature",
200
+ "description": "Controls randomness. Lower values make outputs more focused; higher values make them more varied.",
201
+ "group": "sampling",
202
+ "type": "number",
203
+ "range": {
204
+ "min": 0,
205
+ "max": 1.9,
206
+ "step": 0.1
207
+ }
208
+ },
209
+ {
210
+ "path": "top_p",
211
+ "label": "Top P",
212
+ "description": "Controls nucleus sampling by limiting generation to tokens within the selected cumulative probability.",
213
+ "group": "sampling",
214
+ "type": "number",
215
+ "range": {
216
+ "min": 0,
217
+ "max": 1,
218
+ "step": 0.01
219
+ }
220
+ },
221
+ {
222
+ "path": "extra_body.top_k",
223
+ "label": "Top K",
224
+ "description": "Limits generation to the selected number of highest-probability tokens.",
225
+ "group": "sampling",
226
+ "type": "integer",
227
+ "default": 20,
228
+ "range": {
229
+ "min": 1
230
+ }
231
+ },
232
+ {
233
+ "path": "extra_body.chat_template_kwargs.enable_thinking",
234
+ "label": "Enable thinking",
235
+ "description": "Controls Qwen3 thinking mode when using OpenAI-compatible clients that pass provider-specific extra body fields.",
236
+ "group": "reasoning",
237
+ "type": "boolean",
238
+ "default": true
239
+ }
240
+ ]
241
+ },
122
242
  {
123
243
  "provider": "alibaba",
124
244
  "authType": "api_key",
@@ -628,6 +748,81 @@
628
748
  }
629
749
  ]
630
750
  },
751
+ {
752
+ "provider": "alibaba",
753
+ "authType": "api_key",
754
+ "model": "qwen3.8-2.4t-a95b",
755
+ "params": [
756
+ {
757
+ "path": "max_completion_tokens",
758
+ "label": "Max tokens",
759
+ "description": "Maximum number of tokens to generate, including both reasoning and the final answer.",
760
+ "group": "generation_length",
761
+ "type": "integer",
762
+ "range": {
763
+ "min": 1
764
+ }
765
+ },
766
+ {
767
+ "path": "temperature",
768
+ "label": "Temperature",
769
+ "description": "Controls randomness. Lower values make outputs more focused; higher values make them more varied.",
770
+ "group": "sampling",
771
+ "type": "number",
772
+ "range": {
773
+ "min": 0,
774
+ "max": 2,
775
+ "step": 0.1
776
+ }
777
+ },
778
+ {
779
+ "path": "top_p",
780
+ "label": "Top P",
781
+ "description": "Controls nucleus sampling by limiting generation to tokens within the selected cumulative probability.",
782
+ "group": "sampling",
783
+ "type": "number",
784
+ "range": {
785
+ "min": 0,
786
+ "max": 1,
787
+ "step": 0.01
788
+ }
789
+ },
790
+ {
791
+ "path": "extra_body.top_k",
792
+ "label": "Top K",
793
+ "description": "Limits generation to the selected number of highest-probability tokens. Values above 100 disable top-k sampling.",
794
+ "group": "sampling",
795
+ "type": "integer",
796
+ "default": 20,
797
+ "range": {
798
+ "min": 0
799
+ }
800
+ },
801
+ {
802
+ "path": "extra_body.enable_thinking",
803
+ "label": "Enable thinking",
804
+ "description": "Toggles the model's hybrid thinking mode, sent as a provider-specific extra body field on OpenAI-compatible clients.",
805
+ "group": "reasoning",
806
+ "type": "boolean",
807
+ "default": true
808
+ },
809
+ {
810
+ "path": "extra_body.thinking_budget",
811
+ "label": "Thinking budget",
812
+ "description": "Maximum number of tokens the model may spend on reasoning before it starts the final answer; defaults to the model's maximum reasoning length.",
813
+ "group": "reasoning",
814
+ "applicability": {
815
+ "only": {
816
+ "extra_body.enable_thinking": true
817
+ }
818
+ },
819
+ "type": "integer",
820
+ "range": {
821
+ "min": 1
822
+ }
823
+ }
824
+ ]
825
+ },
631
826
  {
632
827
  "provider": "alibaba",
633
828
  "authType": "api_key",
@@ -5,7 +5,9 @@ from typing import Literal
5
5
 
6
6
  ModelId = Literal[
7
7
  "alibaba/qwen-flash",
8
+ "alibaba/qwen-max",
8
9
  "alibaba/qwen-plus",
10
+ "alibaba/qwen-turbo",
9
11
  "alibaba/qwen3-coder-flash",
10
12
  "alibaba/qwen3-coder-plus",
11
13
  "alibaba/qwen3-max",
@@ -14,6 +16,7 @@ ModelId = Literal[
14
16
  "alibaba/qwen3.6-flash",
15
17
  "alibaba/qwen3.7-max",
16
18
  "alibaba/qwen3.7-plus",
19
+ "alibaba/qwen3.8-2.4t-a95b",
17
20
  "alibaba/qwen3.8-max",
18
21
  "alibaba/qwq-plus",
19
22
  "anthropic/claude-3-5-haiku-20241022",
@@ -263,7 +266,9 @@ ModelId = Literal[
263
266
 
264
267
  MODEL_IDS: tuple[ModelId, ...] = (
265
268
  "alibaba/qwen-flash",
269
+ "alibaba/qwen-max",
266
270
  "alibaba/qwen-plus",
271
+ "alibaba/qwen-turbo",
267
272
  "alibaba/qwen3-coder-flash",
268
273
  "alibaba/qwen3-coder-plus",
269
274
  "alibaba/qwen3-max",
@@ -272,6 +277,7 @@ MODEL_IDS: tuple[ModelId, ...] = (
272
277
  "alibaba/qwen3.6-flash",
273
278
  "alibaba/qwen3.7-max",
274
279
  "alibaba/qwen3.7-plus",
280
+ "alibaba/qwen3.8-2.4t-a95b",
275
281
  "alibaba/qwen3.8-max",
276
282
  "alibaba/qwq-plus",
277
283
  "anthropic/claude-3-5-haiku-20241022",
@@ -26,7 +26,9 @@ from .model_ids import ModelId
26
26
 
27
27
  PARAM_TYPES: dict[ModelId, Any] = {
28
28
  "alibaba/qwen-flash": alibaba.Qwen_FlashParams,
29
+ "alibaba/qwen-max": alibaba.Qwen_MaxParams,
29
30
  "alibaba/qwen-plus": alibaba.Qwen_PlusParams,
31
+ "alibaba/qwen-turbo": alibaba.Qwen_TurboParams,
30
32
  "alibaba/qwen3-coder-flash": alibaba.Qwen3_Coder_FlashParams,
31
33
  "alibaba/qwen3-coder-plus": alibaba.Qwen3_Coder_PlusParams,
32
34
  "alibaba/qwen3-max": alibaba.Qwen3_MaxParams,
@@ -35,6 +37,7 @@ PARAM_TYPES: dict[ModelId, Any] = {
35
37
  "alibaba/qwen3.6-flash": alibaba.Qwen3_6_FlashParams,
36
38
  "alibaba/qwen3.7-max": alibaba.Qwen3_7_MaxParams,
37
39
  "alibaba/qwen3.7-plus": alibaba.Qwen3_7_PlusParams,
40
+ "alibaba/qwen3.8-2.4t-a95b": alibaba.Qwen3_8_2_4t_A95bParams,
38
41
  "alibaba/qwen3.8-max": alibaba.Qwen3_8_MaxParams,
39
42
  "alibaba/qwq-plus": alibaba.Qwq_PlusParams,
40
43
  "anthropic/claude-3-5-haiku-20241022": anthropic.Claude_3_5_Haiku_20241022Params,
@@ -23,6 +23,19 @@ Qwen_FlashParams = TypedDict(
23
23
  )
24
24
  setattr(Qwen_FlashParams, "__pydantic_config__", _PARAMS_CONFIG)
25
25
 
26
+ Qwen_MaxParams = TypedDict(
27
+ "Qwen_MaxParams",
28
+ {
29
+ "max_tokens": Annotated[int, Field(ge=1)],
30
+ "temperature": Annotated[float, Field(ge=0, le=1.9)],
31
+ "top_p": Annotated[float, Field(ge=0, le=1)],
32
+ "extra_body.top_k": Annotated[int, Field(ge=1)],
33
+ "extra_body.chat_template_kwargs.enable_thinking": bool,
34
+ },
35
+ total=False,
36
+ )
37
+ setattr(Qwen_MaxParams, "__pydantic_config__", _PARAMS_CONFIG)
38
+
26
39
  Qwen_PlusParams = TypedDict(
27
40
  "Qwen_PlusParams",
28
41
  {
@@ -36,6 +49,19 @@ Qwen_PlusParams = TypedDict(
36
49
  )
37
50
  setattr(Qwen_PlusParams, "__pydantic_config__", _PARAMS_CONFIG)
38
51
 
52
+ Qwen_TurboParams = TypedDict(
53
+ "Qwen_TurboParams",
54
+ {
55
+ "max_tokens": Annotated[int, Field(ge=1)],
56
+ "temperature": Annotated[float, Field(ge=0, le=1.9)],
57
+ "top_p": Annotated[float, Field(ge=0, le=1)],
58
+ "extra_body.top_k": Annotated[int, Field(ge=1)],
59
+ "extra_body.chat_template_kwargs.enable_thinking": bool,
60
+ },
61
+ total=False,
62
+ )
63
+ setattr(Qwen_TurboParams, "__pydantic_config__", _PARAMS_CONFIG)
64
+
39
65
  Qwen3_Coder_FlashParams = TypedDict(
40
66
  "Qwen3_Coder_FlashParams",
41
67
  {
@@ -141,6 +167,20 @@ Qwen3_7_PlusParams = TypedDict(
141
167
  )
142
168
  setattr(Qwen3_7_PlusParams, "__pydantic_config__", _PARAMS_CONFIG)
143
169
 
170
+ Qwen3_8_2_4t_A95bParams = TypedDict(
171
+ "Qwen3_8_2_4t_A95bParams",
172
+ {
173
+ "max_completion_tokens": Annotated[int, Field(ge=1)],
174
+ "temperature": Annotated[float, Field(ge=0, le=2)],
175
+ "top_p": Annotated[float, Field(ge=0, le=1)],
176
+ "extra_body.top_k": Annotated[int, Field(ge=0)],
177
+ "extra_body.enable_thinking": bool,
178
+ "extra_body.thinking_budget": Annotated[int, Field(ge=1)],
179
+ },
180
+ total=False,
181
+ )
182
+ setattr(Qwen3_8_2_4t_A95bParams, "__pydantic_config__", _PARAMS_CONFIG)
183
+
144
184
  Qwen3_8_MaxParams = TypedDict(
145
185
  "Qwen3_8_MaxParams",
146
186
  {
@@ -169,7 +209,9 @@ setattr(Qwq_PlusParams, "__pydantic_config__", _PARAMS_CONFIG)
169
209
 
170
210
  __all__ = [
171
211
  "Qwen_FlashParams",
212
+ "Qwen_MaxParams",
172
213
  "Qwen_PlusParams",
214
+ "Qwen_TurboParams",
173
215
  "Qwen3_Coder_FlashParams",
174
216
  "Qwen3_Coder_PlusParams",
175
217
  "Qwen3_MaxParams",
@@ -178,6 +220,7 @@ __all__ = [
178
220
  "Qwen3_6_FlashParams",
179
221
  "Qwen3_7_MaxParams",
180
222
  "Qwen3_7_PlusParams",
223
+ "Qwen3_8_2_4t_A95bParams",
181
224
  "Qwen3_8_MaxParams",
182
225
  "Qwq_PlusParams",
183
226
  ]
File without changes
File without changes
File without changes