numba-cuda 0.22.0__cp312-cp312-manylinux_2_24_aarch64.manylinux_2_28_aarch64.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.

Potentially problematic release.


This version of numba-cuda might be problematic. Click here for more details.

Files changed (487) hide show
  1. _numba_cuda_redirector.pth +4 -0
  2. _numba_cuda_redirector.py +89 -0
  3. numba_cuda/VERSION +1 -0
  4. numba_cuda/__init__.py +6 -0
  5. numba_cuda/_version.py +11 -0
  6. numba_cuda/numba/cuda/__init__.py +70 -0
  7. numba_cuda/numba/cuda/_internal/cuda_bf16.py +16394 -0
  8. numba_cuda/numba/cuda/_internal/cuda_fp16.py +8112 -0
  9. numba_cuda/numba/cuda/api.py +580 -0
  10. numba_cuda/numba/cuda/api_util.py +76 -0
  11. numba_cuda/numba/cuda/args.py +72 -0
  12. numba_cuda/numba/cuda/bf16.py +397 -0
  13. numba_cuda/numba/cuda/cache_hints.py +287 -0
  14. numba_cuda/numba/cuda/cext/__init__.py +2 -0
  15. numba_cuda/numba/cuda/cext/_devicearray.cpp +159 -0
  16. numba_cuda/numba/cuda/cext/_devicearray.cpython-312-aarch64-linux-gnu.so +0 -0
  17. numba_cuda/numba/cuda/cext/_devicearray.h +29 -0
  18. numba_cuda/numba/cuda/cext/_dispatcher.cpp +1098 -0
  19. numba_cuda/numba/cuda/cext/_dispatcher.cpython-312-aarch64-linux-gnu.so +0 -0
  20. numba_cuda/numba/cuda/cext/_hashtable.cpp +532 -0
  21. numba_cuda/numba/cuda/cext/_hashtable.h +135 -0
  22. numba_cuda/numba/cuda/cext/_helperlib.c +71 -0
  23. numba_cuda/numba/cuda/cext/_helperlib.cpython-312-aarch64-linux-gnu.so +0 -0
  24. numba_cuda/numba/cuda/cext/_helpermod.c +82 -0
  25. numba_cuda/numba/cuda/cext/_pymodule.h +38 -0
  26. numba_cuda/numba/cuda/cext/_typeconv.cpp +206 -0
  27. numba_cuda/numba/cuda/cext/_typeconv.cpython-312-aarch64-linux-gnu.so +0 -0
  28. numba_cuda/numba/cuda/cext/_typeof.cpp +1159 -0
  29. numba_cuda/numba/cuda/cext/_typeof.h +19 -0
  30. numba_cuda/numba/cuda/cext/capsulethunk.h +111 -0
  31. numba_cuda/numba/cuda/cext/mviewbuf.c +385 -0
  32. numba_cuda/numba/cuda/cext/mviewbuf.cpython-312-aarch64-linux-gnu.so +0 -0
  33. numba_cuda/numba/cuda/cext/typeconv.cpp +212 -0
  34. numba_cuda/numba/cuda/cext/typeconv.hpp +101 -0
  35. numba_cuda/numba/cuda/cg.py +67 -0
  36. numba_cuda/numba/cuda/cgutils.py +1294 -0
  37. numba_cuda/numba/cuda/cloudpickle/__init__.py +21 -0
  38. numba_cuda/numba/cuda/cloudpickle/cloudpickle.py +1598 -0
  39. numba_cuda/numba/cuda/cloudpickle/cloudpickle_fast.py +17 -0
  40. numba_cuda/numba/cuda/codegen.py +541 -0
  41. numba_cuda/numba/cuda/compiler.py +1396 -0
  42. numba_cuda/numba/cuda/core/analysis.py +758 -0
  43. numba_cuda/numba/cuda/core/annotations/__init__.py +0 -0
  44. numba_cuda/numba/cuda/core/annotations/pretty_annotate.py +288 -0
  45. numba_cuda/numba/cuda/core/annotations/type_annotations.py +305 -0
  46. numba_cuda/numba/cuda/core/base.py +1332 -0
  47. numba_cuda/numba/cuda/core/boxing.py +1411 -0
  48. numba_cuda/numba/cuda/core/bytecode.py +728 -0
  49. numba_cuda/numba/cuda/core/byteflow.py +2346 -0
  50. numba_cuda/numba/cuda/core/caching.py +744 -0
  51. numba_cuda/numba/cuda/core/callconv.py +392 -0
  52. numba_cuda/numba/cuda/core/codegen.py +171 -0
  53. numba_cuda/numba/cuda/core/compiler.py +199 -0
  54. numba_cuda/numba/cuda/core/compiler_lock.py +85 -0
  55. numba_cuda/numba/cuda/core/compiler_machinery.py +497 -0
  56. numba_cuda/numba/cuda/core/config.py +650 -0
  57. numba_cuda/numba/cuda/core/consts.py +124 -0
  58. numba_cuda/numba/cuda/core/controlflow.py +989 -0
  59. numba_cuda/numba/cuda/core/entrypoints.py +57 -0
  60. numba_cuda/numba/cuda/core/environment.py +66 -0
  61. numba_cuda/numba/cuda/core/errors.py +917 -0
  62. numba_cuda/numba/cuda/core/event.py +511 -0
  63. numba_cuda/numba/cuda/core/funcdesc.py +330 -0
  64. numba_cuda/numba/cuda/core/generators.py +387 -0
  65. numba_cuda/numba/cuda/core/imputils.py +509 -0
  66. numba_cuda/numba/cuda/core/inline_closurecall.py +1787 -0
  67. numba_cuda/numba/cuda/core/interpreter.py +3617 -0
  68. numba_cuda/numba/cuda/core/ir.py +1812 -0
  69. numba_cuda/numba/cuda/core/ir_utils.py +2638 -0
  70. numba_cuda/numba/cuda/core/optional.py +129 -0
  71. numba_cuda/numba/cuda/core/options.py +262 -0
  72. numba_cuda/numba/cuda/core/postproc.py +249 -0
  73. numba_cuda/numba/cuda/core/pythonapi.py +1859 -0
  74. numba_cuda/numba/cuda/core/registry.py +46 -0
  75. numba_cuda/numba/cuda/core/removerefctpass.py +123 -0
  76. numba_cuda/numba/cuda/core/rewrites/__init__.py +26 -0
  77. numba_cuda/numba/cuda/core/rewrites/ir_print.py +91 -0
  78. numba_cuda/numba/cuda/core/rewrites/registry.py +104 -0
  79. numba_cuda/numba/cuda/core/rewrites/static_binop.py +41 -0
  80. numba_cuda/numba/cuda/core/rewrites/static_getitem.py +189 -0
  81. numba_cuda/numba/cuda/core/rewrites/static_raise.py +100 -0
  82. numba_cuda/numba/cuda/core/sigutils.py +68 -0
  83. numba_cuda/numba/cuda/core/ssa.py +498 -0
  84. numba_cuda/numba/cuda/core/targetconfig.py +330 -0
  85. numba_cuda/numba/cuda/core/tracing.py +231 -0
  86. numba_cuda/numba/cuda/core/transforms.py +956 -0
  87. numba_cuda/numba/cuda/core/typed_passes.py +867 -0
  88. numba_cuda/numba/cuda/core/typeinfer.py +1950 -0
  89. numba_cuda/numba/cuda/core/unsafe/__init__.py +0 -0
  90. numba_cuda/numba/cuda/core/unsafe/bytes.py +67 -0
  91. numba_cuda/numba/cuda/core/unsafe/eh.py +67 -0
  92. numba_cuda/numba/cuda/core/unsafe/refcount.py +98 -0
  93. numba_cuda/numba/cuda/core/untyped_passes.py +1979 -0
  94. numba_cuda/numba/cuda/cpython/builtins.py +1153 -0
  95. numba_cuda/numba/cuda/cpython/charseq.py +1218 -0
  96. numba_cuda/numba/cuda/cpython/cmathimpl.py +560 -0
  97. numba_cuda/numba/cuda/cpython/enumimpl.py +103 -0
  98. numba_cuda/numba/cuda/cpython/iterators.py +167 -0
  99. numba_cuda/numba/cuda/cpython/listobj.py +1326 -0
  100. numba_cuda/numba/cuda/cpython/mathimpl.py +499 -0
  101. numba_cuda/numba/cuda/cpython/numbers.py +1475 -0
  102. numba_cuda/numba/cuda/cpython/rangeobj.py +289 -0
  103. numba_cuda/numba/cuda/cpython/slicing.py +322 -0
  104. numba_cuda/numba/cuda/cpython/tupleobj.py +456 -0
  105. numba_cuda/numba/cuda/cpython/unicode.py +2865 -0
  106. numba_cuda/numba/cuda/cpython/unicode_support.py +1597 -0
  107. numba_cuda/numba/cuda/cpython/unsafe/__init__.py +0 -0
  108. numba_cuda/numba/cuda/cpython/unsafe/numbers.py +64 -0
  109. numba_cuda/numba/cuda/cpython/unsafe/tuple.py +92 -0
  110. numba_cuda/numba/cuda/cuda_paths.py +691 -0
  111. numba_cuda/numba/cuda/cudadecl.py +543 -0
  112. numba_cuda/numba/cuda/cudadrv/__init__.py +14 -0
  113. numba_cuda/numba/cuda/cudadrv/devicearray.py +954 -0
  114. numba_cuda/numba/cuda/cudadrv/devices.py +249 -0
  115. numba_cuda/numba/cuda/cudadrv/driver.py +3238 -0
  116. numba_cuda/numba/cuda/cudadrv/drvapi.py +435 -0
  117. numba_cuda/numba/cuda/cudadrv/dummyarray.py +562 -0
  118. numba_cuda/numba/cuda/cudadrv/enums.py +613 -0
  119. numba_cuda/numba/cuda/cudadrv/error.py +48 -0
  120. numba_cuda/numba/cuda/cudadrv/libs.py +220 -0
  121. numba_cuda/numba/cuda/cudadrv/linkable_code.py +184 -0
  122. numba_cuda/numba/cuda/cudadrv/mappings.py +14 -0
  123. numba_cuda/numba/cuda/cudadrv/ndarray.py +26 -0
  124. numba_cuda/numba/cuda/cudadrv/nvrtc.py +193 -0
  125. numba_cuda/numba/cuda/cudadrv/nvvm.py +756 -0
  126. numba_cuda/numba/cuda/cudadrv/rtapi.py +13 -0
  127. numba_cuda/numba/cuda/cudadrv/runtime.py +34 -0
  128. numba_cuda/numba/cuda/cudaimpl.py +983 -0
  129. numba_cuda/numba/cuda/cudamath.py +149 -0
  130. numba_cuda/numba/cuda/datamodel/__init__.py +7 -0
  131. numba_cuda/numba/cuda/datamodel/cuda_manager.py +66 -0
  132. numba_cuda/numba/cuda/datamodel/cuda_models.py +1446 -0
  133. numba_cuda/numba/cuda/datamodel/cuda_packer.py +224 -0
  134. numba_cuda/numba/cuda/datamodel/cuda_registry.py +22 -0
  135. numba_cuda/numba/cuda/datamodel/cuda_testing.py +153 -0
  136. numba_cuda/numba/cuda/datamodel/manager.py +11 -0
  137. numba_cuda/numba/cuda/datamodel/models.py +9 -0
  138. numba_cuda/numba/cuda/datamodel/packer.py +9 -0
  139. numba_cuda/numba/cuda/datamodel/registry.py +11 -0
  140. numba_cuda/numba/cuda/datamodel/testing.py +11 -0
  141. numba_cuda/numba/cuda/debuginfo.py +997 -0
  142. numba_cuda/numba/cuda/decorators.py +294 -0
  143. numba_cuda/numba/cuda/descriptor.py +35 -0
  144. numba_cuda/numba/cuda/device_init.py +155 -0
  145. numba_cuda/numba/cuda/deviceufunc.py +1021 -0
  146. numba_cuda/numba/cuda/dispatcher.py +2463 -0
  147. numba_cuda/numba/cuda/errors.py +72 -0
  148. numba_cuda/numba/cuda/extending.py +697 -0
  149. numba_cuda/numba/cuda/flags.py +178 -0
  150. numba_cuda/numba/cuda/fp16.py +357 -0
  151. numba_cuda/numba/cuda/include/12/cuda_bf16.h +5118 -0
  152. numba_cuda/numba/cuda/include/12/cuda_bf16.hpp +3865 -0
  153. numba_cuda/numba/cuda/include/12/cuda_fp16.h +5363 -0
  154. numba_cuda/numba/cuda/include/12/cuda_fp16.hpp +3483 -0
  155. numba_cuda/numba/cuda/include/13/cuda_bf16.h +5118 -0
  156. numba_cuda/numba/cuda/include/13/cuda_bf16.hpp +3865 -0
  157. numba_cuda/numba/cuda/include/13/cuda_fp16.h +5363 -0
  158. numba_cuda/numba/cuda/include/13/cuda_fp16.hpp +3483 -0
  159. numba_cuda/numba/cuda/initialize.py +24 -0
  160. numba_cuda/numba/cuda/intrinsics.py +531 -0
  161. numba_cuda/numba/cuda/itanium_mangler.py +214 -0
  162. numba_cuda/numba/cuda/kernels/__init__.py +2 -0
  163. numba_cuda/numba/cuda/kernels/reduction.py +265 -0
  164. numba_cuda/numba/cuda/kernels/transpose.py +65 -0
  165. numba_cuda/numba/cuda/libdevice.py +3386 -0
  166. numba_cuda/numba/cuda/libdevicedecl.py +20 -0
  167. numba_cuda/numba/cuda/libdevicefuncs.py +1060 -0
  168. numba_cuda/numba/cuda/libdeviceimpl.py +88 -0
  169. numba_cuda/numba/cuda/locks.py +19 -0
  170. numba_cuda/numba/cuda/lowering.py +1980 -0
  171. numba_cuda/numba/cuda/mathimpl.py +374 -0
  172. numba_cuda/numba/cuda/memory_management/__init__.py +4 -0
  173. numba_cuda/numba/cuda/memory_management/memsys.cu +99 -0
  174. numba_cuda/numba/cuda/memory_management/memsys.cuh +22 -0
  175. numba_cuda/numba/cuda/memory_management/nrt.cu +212 -0
  176. numba_cuda/numba/cuda/memory_management/nrt.cuh +48 -0
  177. numba_cuda/numba/cuda/memory_management/nrt.py +390 -0
  178. numba_cuda/numba/cuda/memory_management/nrt_context.py +438 -0
  179. numba_cuda/numba/cuda/misc/appdirs.py +594 -0
  180. numba_cuda/numba/cuda/misc/cffiimpl.py +24 -0
  181. numba_cuda/numba/cuda/misc/coverage_support.py +43 -0
  182. numba_cuda/numba/cuda/misc/dump_style.py +41 -0
  183. numba_cuda/numba/cuda/misc/findlib.py +75 -0
  184. numba_cuda/numba/cuda/misc/firstlinefinder.py +96 -0
  185. numba_cuda/numba/cuda/misc/gdb_hook.py +240 -0
  186. numba_cuda/numba/cuda/misc/literal.py +28 -0
  187. numba_cuda/numba/cuda/misc/llvm_pass_timings.py +412 -0
  188. numba_cuda/numba/cuda/misc/special.py +94 -0
  189. numba_cuda/numba/cuda/models.py +56 -0
  190. numba_cuda/numba/cuda/np/arraymath.py +5130 -0
  191. numba_cuda/numba/cuda/np/arrayobj.py +7635 -0
  192. numba_cuda/numba/cuda/np/extensions.py +11 -0
  193. numba_cuda/numba/cuda/np/linalg.py +3087 -0
  194. numba_cuda/numba/cuda/np/math/__init__.py +0 -0
  195. numba_cuda/numba/cuda/np/math/cmathimpl.py +558 -0
  196. numba_cuda/numba/cuda/np/math/mathimpl.py +487 -0
  197. numba_cuda/numba/cuda/np/math/numbers.py +1461 -0
  198. numba_cuda/numba/cuda/np/npdatetime.py +969 -0
  199. numba_cuda/numba/cuda/np/npdatetime_helpers.py +217 -0
  200. numba_cuda/numba/cuda/np/npyfuncs.py +1808 -0
  201. numba_cuda/numba/cuda/np/npyimpl.py +1027 -0
  202. numba_cuda/numba/cuda/np/numpy_support.py +798 -0
  203. numba_cuda/numba/cuda/np/polynomial/__init__.py +4 -0
  204. numba_cuda/numba/cuda/np/polynomial/polynomial_core.py +242 -0
  205. numba_cuda/numba/cuda/np/polynomial/polynomial_functions.py +380 -0
  206. numba_cuda/numba/cuda/np/ufunc/__init__.py +4 -0
  207. numba_cuda/numba/cuda/np/ufunc/decorators.py +203 -0
  208. numba_cuda/numba/cuda/np/ufunc/sigparse.py +68 -0
  209. numba_cuda/numba/cuda/np/ufunc/ufuncbuilder.py +65 -0
  210. numba_cuda/numba/cuda/np/ufunc_db.py +1282 -0
  211. numba_cuda/numba/cuda/np/unsafe/__init__.py +0 -0
  212. numba_cuda/numba/cuda/np/unsafe/ndarray.py +84 -0
  213. numba_cuda/numba/cuda/nvvmutils.py +254 -0
  214. numba_cuda/numba/cuda/printimpl.py +126 -0
  215. numba_cuda/numba/cuda/random.py +308 -0
  216. numba_cuda/numba/cuda/reshape_funcs.cu +156 -0
  217. numba_cuda/numba/cuda/serialize.py +267 -0
  218. numba_cuda/numba/cuda/simulator/__init__.py +63 -0
  219. numba_cuda/numba/cuda/simulator/_internal/__init__.py +4 -0
  220. numba_cuda/numba/cuda/simulator/_internal/cuda_bf16.py +2 -0
  221. numba_cuda/numba/cuda/simulator/api.py +179 -0
  222. numba_cuda/numba/cuda/simulator/bf16.py +4 -0
  223. numba_cuda/numba/cuda/simulator/compiler.py +38 -0
  224. numba_cuda/numba/cuda/simulator/cudadrv/__init__.py +11 -0
  225. numba_cuda/numba/cuda/simulator/cudadrv/devicearray.py +462 -0
  226. numba_cuda/numba/cuda/simulator/cudadrv/devices.py +122 -0
  227. numba_cuda/numba/cuda/simulator/cudadrv/driver.py +66 -0
  228. numba_cuda/numba/cuda/simulator/cudadrv/drvapi.py +7 -0
  229. numba_cuda/numba/cuda/simulator/cudadrv/dummyarray.py +7 -0
  230. numba_cuda/numba/cuda/simulator/cudadrv/error.py +10 -0
  231. numba_cuda/numba/cuda/simulator/cudadrv/libs.py +10 -0
  232. numba_cuda/numba/cuda/simulator/cudadrv/linkable_code.py +61 -0
  233. numba_cuda/numba/cuda/simulator/cudadrv/nvrtc.py +11 -0
  234. numba_cuda/numba/cuda/simulator/cudadrv/nvvm.py +32 -0
  235. numba_cuda/numba/cuda/simulator/cudadrv/runtime.py +22 -0
  236. numba_cuda/numba/cuda/simulator/dispatcher.py +11 -0
  237. numba_cuda/numba/cuda/simulator/kernel.py +320 -0
  238. numba_cuda/numba/cuda/simulator/kernelapi.py +509 -0
  239. numba_cuda/numba/cuda/simulator/memory_management/__init__.py +4 -0
  240. numba_cuda/numba/cuda/simulator/memory_management/nrt.py +21 -0
  241. numba_cuda/numba/cuda/simulator/reduction.py +19 -0
  242. numba_cuda/numba/cuda/simulator/tests/support.py +4 -0
  243. numba_cuda/numba/cuda/simulator/vector_types.py +65 -0
  244. numba_cuda/numba/cuda/simulator_init.py +18 -0
  245. numba_cuda/numba/cuda/stubs.py +624 -0
  246. numba_cuda/numba/cuda/target.py +505 -0
  247. numba_cuda/numba/cuda/testing.py +347 -0
  248. numba_cuda/numba/cuda/tests/__init__.py +62 -0
  249. numba_cuda/numba/cuda/tests/benchmarks/__init__.py +0 -0
  250. numba_cuda/numba/cuda/tests/benchmarks/test_kernel_launch.py +119 -0
  251. numba_cuda/numba/cuda/tests/cloudpickle_main_class.py +9 -0
  252. numba_cuda/numba/cuda/tests/core/serialize_usecases.py +113 -0
  253. numba_cuda/numba/cuda/tests/core/test_itanium_mangler.py +83 -0
  254. numba_cuda/numba/cuda/tests/core/test_serialize.py +371 -0
  255. numba_cuda/numba/cuda/tests/cudadrv/__init__.py +9 -0
  256. numba_cuda/numba/cuda/tests/cudadrv/test_array_attr.py +147 -0
  257. numba_cuda/numba/cuda/tests/cudadrv/test_context_stack.py +161 -0
  258. numba_cuda/numba/cuda/tests/cudadrv/test_cuda_array_slicing.py +397 -0
  259. numba_cuda/numba/cuda/tests/cudadrv/test_cuda_auto_context.py +24 -0
  260. numba_cuda/numba/cuda/tests/cudadrv/test_cuda_devicerecord.py +180 -0
  261. numba_cuda/numba/cuda/tests/cudadrv/test_cuda_driver.py +313 -0
  262. numba_cuda/numba/cuda/tests/cudadrv/test_cuda_memory.py +191 -0
  263. numba_cuda/numba/cuda/tests/cudadrv/test_cuda_ndarray.py +621 -0
  264. numba_cuda/numba/cuda/tests/cudadrv/test_deallocations.py +247 -0
  265. numba_cuda/numba/cuda/tests/cudadrv/test_detect.py +100 -0
  266. numba_cuda/numba/cuda/tests/cudadrv/test_emm_plugins.py +200 -0
  267. numba_cuda/numba/cuda/tests/cudadrv/test_events.py +53 -0
  268. numba_cuda/numba/cuda/tests/cudadrv/test_host_alloc.py +72 -0
  269. numba_cuda/numba/cuda/tests/cudadrv/test_init.py +138 -0
  270. numba_cuda/numba/cuda/tests/cudadrv/test_inline_ptx.py +43 -0
  271. numba_cuda/numba/cuda/tests/cudadrv/test_is_fp16.py +15 -0
  272. numba_cuda/numba/cuda/tests/cudadrv/test_linkable_code.py +58 -0
  273. numba_cuda/numba/cuda/tests/cudadrv/test_linker.py +348 -0
  274. numba_cuda/numba/cuda/tests/cudadrv/test_managed_alloc.py +128 -0
  275. numba_cuda/numba/cuda/tests/cudadrv/test_module_callbacks.py +301 -0
  276. numba_cuda/numba/cuda/tests/cudadrv/test_nvjitlink.py +174 -0
  277. numba_cuda/numba/cuda/tests/cudadrv/test_nvrtc.py +28 -0
  278. numba_cuda/numba/cuda/tests/cudadrv/test_nvvm_driver.py +185 -0
  279. numba_cuda/numba/cuda/tests/cudadrv/test_pinned.py +39 -0
  280. numba_cuda/numba/cuda/tests/cudadrv/test_profiler.py +23 -0
  281. numba_cuda/numba/cuda/tests/cudadrv/test_reset_device.py +38 -0
  282. numba_cuda/numba/cuda/tests/cudadrv/test_runtime.py +48 -0
  283. numba_cuda/numba/cuda/tests/cudadrv/test_select_device.py +44 -0
  284. numba_cuda/numba/cuda/tests/cudadrv/test_streams.py +127 -0
  285. numba_cuda/numba/cuda/tests/cudapy/__init__.py +9 -0
  286. numba_cuda/numba/cuda/tests/cudapy/cache_usecases.py +231 -0
  287. numba_cuda/numba/cuda/tests/cudapy/cache_with_cpu_usecases.py +50 -0
  288. numba_cuda/numba/cuda/tests/cudapy/cg_cache_usecases.py +36 -0
  289. numba_cuda/numba/cuda/tests/cudapy/complex_usecases.py +116 -0
  290. numba_cuda/numba/cuda/tests/cudapy/enum_usecases.py +59 -0
  291. numba_cuda/numba/cuda/tests/cudapy/extensions_usecases.py +62 -0
  292. numba_cuda/numba/cuda/tests/cudapy/jitlink.ptx +28 -0
  293. numba_cuda/numba/cuda/tests/cudapy/overload_usecases.py +33 -0
  294. numba_cuda/numba/cuda/tests/cudapy/recursion_usecases.py +104 -0
  295. numba_cuda/numba/cuda/tests/cudapy/test_alignment.py +47 -0
  296. numba_cuda/numba/cuda/tests/cudapy/test_analysis.py +1122 -0
  297. numba_cuda/numba/cuda/tests/cudapy/test_array.py +344 -0
  298. numba_cuda/numba/cuda/tests/cudapy/test_array_alignment.py +268 -0
  299. numba_cuda/numba/cuda/tests/cudapy/test_array_args.py +203 -0
  300. numba_cuda/numba/cuda/tests/cudapy/test_array_methods.py +63 -0
  301. numba_cuda/numba/cuda/tests/cudapy/test_array_reductions.py +360 -0
  302. numba_cuda/numba/cuda/tests/cudapy/test_atomics.py +1815 -0
  303. numba_cuda/numba/cuda/tests/cudapy/test_bfloat16.py +599 -0
  304. numba_cuda/numba/cuda/tests/cudapy/test_bfloat16_bindings.py +377 -0
  305. numba_cuda/numba/cuda/tests/cudapy/test_blackscholes.py +160 -0
  306. numba_cuda/numba/cuda/tests/cudapy/test_boolean.py +27 -0
  307. numba_cuda/numba/cuda/tests/cudapy/test_byteflow.py +98 -0
  308. numba_cuda/numba/cuda/tests/cudapy/test_cache_hints.py +210 -0
  309. numba_cuda/numba/cuda/tests/cudapy/test_caching.py +683 -0
  310. numba_cuda/numba/cuda/tests/cudapy/test_casting.py +265 -0
  311. numba_cuda/numba/cuda/tests/cudapy/test_cffi.py +42 -0
  312. numba_cuda/numba/cuda/tests/cudapy/test_compiler.py +718 -0
  313. numba_cuda/numba/cuda/tests/cudapy/test_complex.py +370 -0
  314. numba_cuda/numba/cuda/tests/cudapy/test_complex_kernel.py +23 -0
  315. numba_cuda/numba/cuda/tests/cudapy/test_const_string.py +142 -0
  316. numba_cuda/numba/cuda/tests/cudapy/test_constmem.py +178 -0
  317. numba_cuda/numba/cuda/tests/cudapy/test_cooperative_groups.py +193 -0
  318. numba_cuda/numba/cuda/tests/cudapy/test_copy_propagate.py +131 -0
  319. numba_cuda/numba/cuda/tests/cudapy/test_cuda_array_interface.py +438 -0
  320. numba_cuda/numba/cuda/tests/cudapy/test_cuda_jit_no_types.py +94 -0
  321. numba_cuda/numba/cuda/tests/cudapy/test_datetime.py +101 -0
  322. numba_cuda/numba/cuda/tests/cudapy/test_debug.py +105 -0
  323. numba_cuda/numba/cuda/tests/cudapy/test_debuginfo.py +978 -0
  324. numba_cuda/numba/cuda/tests/cudapy/test_debuginfo_types.py +476 -0
  325. numba_cuda/numba/cuda/tests/cudapy/test_device_func.py +500 -0
  326. numba_cuda/numba/cuda/tests/cudapy/test_dispatcher.py +820 -0
  327. numba_cuda/numba/cuda/tests/cudapy/test_enums.py +152 -0
  328. numba_cuda/numba/cuda/tests/cudapy/test_errors.py +111 -0
  329. numba_cuda/numba/cuda/tests/cudapy/test_exception.py +170 -0
  330. numba_cuda/numba/cuda/tests/cudapy/test_extending.py +1088 -0
  331. numba_cuda/numba/cuda/tests/cudapy/test_extending_types.py +71 -0
  332. numba_cuda/numba/cuda/tests/cudapy/test_fastmath.py +265 -0
  333. numba_cuda/numba/cuda/tests/cudapy/test_flow_control.py +1433 -0
  334. numba_cuda/numba/cuda/tests/cudapy/test_forall.py +57 -0
  335. numba_cuda/numba/cuda/tests/cudapy/test_freevar.py +34 -0
  336. numba_cuda/numba/cuda/tests/cudapy/test_frexp_ldexp.py +69 -0
  337. numba_cuda/numba/cuda/tests/cudapy/test_globals.py +62 -0
  338. numba_cuda/numba/cuda/tests/cudapy/test_gufunc.py +474 -0
  339. numba_cuda/numba/cuda/tests/cudapy/test_gufunc_scalar.py +167 -0
  340. numba_cuda/numba/cuda/tests/cudapy/test_gufunc_scheduling.py +92 -0
  341. numba_cuda/numba/cuda/tests/cudapy/test_idiv.py +39 -0
  342. numba_cuda/numba/cuda/tests/cudapy/test_inline.py +170 -0
  343. numba_cuda/numba/cuda/tests/cudapy/test_inspect.py +255 -0
  344. numba_cuda/numba/cuda/tests/cudapy/test_intrinsics.py +1219 -0
  345. numba_cuda/numba/cuda/tests/cudapy/test_ipc.py +263 -0
  346. numba_cuda/numba/cuda/tests/cudapy/test_ir.py +598 -0
  347. numba_cuda/numba/cuda/tests/cudapy/test_ir_utils.py +276 -0
  348. numba_cuda/numba/cuda/tests/cudapy/test_iterators.py +101 -0
  349. numba_cuda/numba/cuda/tests/cudapy/test_lang.py +68 -0
  350. numba_cuda/numba/cuda/tests/cudapy/test_laplace.py +123 -0
  351. numba_cuda/numba/cuda/tests/cudapy/test_libdevice.py +194 -0
  352. numba_cuda/numba/cuda/tests/cudapy/test_lineinfo.py +220 -0
  353. numba_cuda/numba/cuda/tests/cudapy/test_localmem.py +173 -0
  354. numba_cuda/numba/cuda/tests/cudapy/test_make_function_to_jit_function.py +364 -0
  355. numba_cuda/numba/cuda/tests/cudapy/test_mandel.py +47 -0
  356. numba_cuda/numba/cuda/tests/cudapy/test_math.py +842 -0
  357. numba_cuda/numba/cuda/tests/cudapy/test_matmul.py +76 -0
  358. numba_cuda/numba/cuda/tests/cudapy/test_minmax.py +78 -0
  359. numba_cuda/numba/cuda/tests/cudapy/test_montecarlo.py +25 -0
  360. numba_cuda/numba/cuda/tests/cudapy/test_multigpu.py +145 -0
  361. numba_cuda/numba/cuda/tests/cudapy/test_multiprocessing.py +39 -0
  362. numba_cuda/numba/cuda/tests/cudapy/test_multithreads.py +82 -0
  363. numba_cuda/numba/cuda/tests/cudapy/test_nondet.py +53 -0
  364. numba_cuda/numba/cuda/tests/cudapy/test_operator.py +504 -0
  365. numba_cuda/numba/cuda/tests/cudapy/test_optimization.py +93 -0
  366. numba_cuda/numba/cuda/tests/cudapy/test_overload.py +402 -0
  367. numba_cuda/numba/cuda/tests/cudapy/test_powi.py +128 -0
  368. numba_cuda/numba/cuda/tests/cudapy/test_print.py +193 -0
  369. numba_cuda/numba/cuda/tests/cudapy/test_py2_div_issue.py +37 -0
  370. numba_cuda/numba/cuda/tests/cudapy/test_random.py +117 -0
  371. numba_cuda/numba/cuda/tests/cudapy/test_record_dtype.py +614 -0
  372. numba_cuda/numba/cuda/tests/cudapy/test_recursion.py +130 -0
  373. numba_cuda/numba/cuda/tests/cudapy/test_reduction.py +94 -0
  374. numba_cuda/numba/cuda/tests/cudapy/test_retrieve_autoconverted_arrays.py +83 -0
  375. numba_cuda/numba/cuda/tests/cudapy/test_serialize.py +86 -0
  376. numba_cuda/numba/cuda/tests/cudapy/test_slicing.py +40 -0
  377. numba_cuda/numba/cuda/tests/cudapy/test_sm.py +457 -0
  378. numba_cuda/numba/cuda/tests/cudapy/test_sm_creation.py +233 -0
  379. numba_cuda/numba/cuda/tests/cudapy/test_ssa.py +454 -0
  380. numba_cuda/numba/cuda/tests/cudapy/test_stream_api.py +56 -0
  381. numba_cuda/numba/cuda/tests/cudapy/test_sync.py +277 -0
  382. numba_cuda/numba/cuda/tests/cudapy/test_tracing.py +200 -0
  383. numba_cuda/numba/cuda/tests/cudapy/test_transpose.py +90 -0
  384. numba_cuda/numba/cuda/tests/cudapy/test_typeconv.py +333 -0
  385. numba_cuda/numba/cuda/tests/cudapy/test_typeinfer.py +538 -0
  386. numba_cuda/numba/cuda/tests/cudapy/test_ufuncs.py +585 -0
  387. numba_cuda/numba/cuda/tests/cudapy/test_userexc.py +42 -0
  388. numba_cuda/numba/cuda/tests/cudapy/test_vector_type.py +485 -0
  389. numba_cuda/numba/cuda/tests/cudapy/test_vectorize.py +312 -0
  390. numba_cuda/numba/cuda/tests/cudapy/test_vectorize_complex.py +23 -0
  391. numba_cuda/numba/cuda/tests/cudapy/test_vectorize_decor.py +183 -0
  392. numba_cuda/numba/cuda/tests/cudapy/test_vectorize_device.py +40 -0
  393. numba_cuda/numba/cuda/tests/cudapy/test_vectorize_scalar_arg.py +40 -0
  394. numba_cuda/numba/cuda/tests/cudapy/test_warning.py +206 -0
  395. numba_cuda/numba/cuda/tests/cudapy/test_warp_ops.py +446 -0
  396. numba_cuda/numba/cuda/tests/cudasim/__init__.py +9 -0
  397. numba_cuda/numba/cuda/tests/cudasim/support.py +9 -0
  398. numba_cuda/numba/cuda/tests/cudasim/test_cudasim_issues.py +111 -0
  399. numba_cuda/numba/cuda/tests/data/__init__.py +2 -0
  400. numba_cuda/numba/cuda/tests/data/cta_barrier.cu +28 -0
  401. numba_cuda/numba/cuda/tests/data/cuda_include.cu +10 -0
  402. numba_cuda/numba/cuda/tests/data/error.cu +12 -0
  403. numba_cuda/numba/cuda/tests/data/include/add.cuh +8 -0
  404. numba_cuda/numba/cuda/tests/data/jitlink.cu +28 -0
  405. numba_cuda/numba/cuda/tests/data/jitlink.ptx +49 -0
  406. numba_cuda/numba/cuda/tests/data/warn.cu +12 -0
  407. numba_cuda/numba/cuda/tests/doc_examples/__init__.py +9 -0
  408. numba_cuda/numba/cuda/tests/doc_examples/ffi/__init__.py +2 -0
  409. numba_cuda/numba/cuda/tests/doc_examples/ffi/functions.cu +54 -0
  410. numba_cuda/numba/cuda/tests/doc_examples/ffi/include/mul.cuh +8 -0
  411. numba_cuda/numba/cuda/tests/doc_examples/ffi/saxpy.cu +14 -0
  412. numba_cuda/numba/cuda/tests/doc_examples/test_cg.py +86 -0
  413. numba_cuda/numba/cuda/tests/doc_examples/test_cpointer.py +68 -0
  414. numba_cuda/numba/cuda/tests/doc_examples/test_cpu_gpu_compat.py +81 -0
  415. numba_cuda/numba/cuda/tests/doc_examples/test_ffi.py +141 -0
  416. numba_cuda/numba/cuda/tests/doc_examples/test_laplace.py +160 -0
  417. numba_cuda/numba/cuda/tests/doc_examples/test_matmul.py +180 -0
  418. numba_cuda/numba/cuda/tests/doc_examples/test_montecarlo.py +119 -0
  419. numba_cuda/numba/cuda/tests/doc_examples/test_random.py +66 -0
  420. numba_cuda/numba/cuda/tests/doc_examples/test_reduction.py +80 -0
  421. numba_cuda/numba/cuda/tests/doc_examples/test_sessionize.py +206 -0
  422. numba_cuda/numba/cuda/tests/doc_examples/test_ufunc.py +53 -0
  423. numba_cuda/numba/cuda/tests/doc_examples/test_vecadd.py +76 -0
  424. numba_cuda/numba/cuda/tests/nocuda/__init__.py +9 -0
  425. numba_cuda/numba/cuda/tests/nocuda/test_dummyarray.py +452 -0
  426. numba_cuda/numba/cuda/tests/nocuda/test_function_resolution.py +48 -0
  427. numba_cuda/numba/cuda/tests/nocuda/test_import.py +63 -0
  428. numba_cuda/numba/cuda/tests/nocuda/test_library_lookup.py +252 -0
  429. numba_cuda/numba/cuda/tests/nocuda/test_nvvm.py +59 -0
  430. numba_cuda/numba/cuda/tests/nrt/__init__.py +9 -0
  431. numba_cuda/numba/cuda/tests/nrt/test_nrt.py +387 -0
  432. numba_cuda/numba/cuda/tests/nrt/test_nrt_refct.py +124 -0
  433. numba_cuda/numba/cuda/tests/support.py +900 -0
  434. numba_cuda/numba/cuda/typeconv/__init__.py +4 -0
  435. numba_cuda/numba/cuda/typeconv/castgraph.py +137 -0
  436. numba_cuda/numba/cuda/typeconv/rules.py +63 -0
  437. numba_cuda/numba/cuda/typeconv/typeconv.py +121 -0
  438. numba_cuda/numba/cuda/types/__init__.py +233 -0
  439. numba_cuda/numba/cuda/types/__init__.pyi +167 -0
  440. numba_cuda/numba/cuda/types/abstract.py +9 -0
  441. numba_cuda/numba/cuda/types/common.py +9 -0
  442. numba_cuda/numba/cuda/types/containers.py +9 -0
  443. numba_cuda/numba/cuda/types/cuda_abstract.py +533 -0
  444. numba_cuda/numba/cuda/types/cuda_common.py +110 -0
  445. numba_cuda/numba/cuda/types/cuda_containers.py +971 -0
  446. numba_cuda/numba/cuda/types/cuda_function_type.py +230 -0
  447. numba_cuda/numba/cuda/types/cuda_functions.py +798 -0
  448. numba_cuda/numba/cuda/types/cuda_iterators.py +120 -0
  449. numba_cuda/numba/cuda/types/cuda_misc.py +569 -0
  450. numba_cuda/numba/cuda/types/cuda_npytypes.py +690 -0
  451. numba_cuda/numba/cuda/types/cuda_scalars.py +280 -0
  452. numba_cuda/numba/cuda/types/ext_types.py +101 -0
  453. numba_cuda/numba/cuda/types/function_type.py +11 -0
  454. numba_cuda/numba/cuda/types/functions.py +9 -0
  455. numba_cuda/numba/cuda/types/iterators.py +9 -0
  456. numba_cuda/numba/cuda/types/misc.py +9 -0
  457. numba_cuda/numba/cuda/types/npytypes.py +9 -0
  458. numba_cuda/numba/cuda/types/scalars.py +9 -0
  459. numba_cuda/numba/cuda/typing/__init__.py +19 -0
  460. numba_cuda/numba/cuda/typing/arraydecl.py +939 -0
  461. numba_cuda/numba/cuda/typing/asnumbatype.py +130 -0
  462. numba_cuda/numba/cuda/typing/bufproto.py +70 -0
  463. numba_cuda/numba/cuda/typing/builtins.py +1209 -0
  464. numba_cuda/numba/cuda/typing/cffi_utils.py +219 -0
  465. numba_cuda/numba/cuda/typing/cmathdecl.py +47 -0
  466. numba_cuda/numba/cuda/typing/collections.py +138 -0
  467. numba_cuda/numba/cuda/typing/context.py +782 -0
  468. numba_cuda/numba/cuda/typing/ctypes_utils.py +125 -0
  469. numba_cuda/numba/cuda/typing/dictdecl.py +63 -0
  470. numba_cuda/numba/cuda/typing/enumdecl.py +74 -0
  471. numba_cuda/numba/cuda/typing/listdecl.py +147 -0
  472. numba_cuda/numba/cuda/typing/mathdecl.py +158 -0
  473. numba_cuda/numba/cuda/typing/npdatetime.py +322 -0
  474. numba_cuda/numba/cuda/typing/npydecl.py +749 -0
  475. numba_cuda/numba/cuda/typing/setdecl.py +115 -0
  476. numba_cuda/numba/cuda/typing/templates.py +1446 -0
  477. numba_cuda/numba/cuda/typing/typeof.py +301 -0
  478. numba_cuda/numba/cuda/ufuncs.py +746 -0
  479. numba_cuda/numba/cuda/utils.py +724 -0
  480. numba_cuda/numba/cuda/vector_types.py +214 -0
  481. numba_cuda/numba/cuda/vectorizers.py +260 -0
  482. numba_cuda-0.22.0.dist-info/METADATA +109 -0
  483. numba_cuda-0.22.0.dist-info/RECORD +487 -0
  484. numba_cuda-0.22.0.dist-info/WHEEL +6 -0
  485. numba_cuda-0.22.0.dist-info/licenses/LICENSE +26 -0
  486. numba_cuda-0.22.0.dist-info/licenses/LICENSE.numba +24 -0
  487. numba_cuda-0.22.0.dist-info/top_level.txt +1 -0
@@ -0,0 +1,203 @@
1
+ # SPDX-FileCopyrightText: Copyright (c) 2025 NVIDIA CORPORATION & AFFILIATES. All rights reserved.
2
+ # SPDX-License-Identifier: BSD-2-Clause
3
+
4
+ import numpy as np
5
+ from collections import namedtuple
6
+
7
+ from numba import cuda
8
+ from numba.cuda.testing import unittest, CUDATestCase
9
+
10
+
11
+ class TestCudaArrayArg(CUDATestCase):
12
+ def test_array_ary(self):
13
+ @cuda.jit("double(double[:],int64)", device=True, inline="always")
14
+ def device_function(a, c):
15
+ return a[c]
16
+
17
+ @cuda.jit("void(double[:],double[:])")
18
+ def kernel(x, y):
19
+ i = cuda.grid(1)
20
+ y[i] = device_function(x, i)
21
+
22
+ x = np.arange(10, dtype=np.double)
23
+ y = np.zeros_like(x)
24
+ kernel[10, 1](x, y)
25
+ self.assertTrue(np.all(x == y))
26
+
27
+ def test_unituple(self):
28
+ @cuda.jit
29
+ def f(r, x):
30
+ r[0] = x[0]
31
+ r[1] = x[1]
32
+ r[2] = x[2]
33
+
34
+ x = (1, 2, 3)
35
+ r = np.zeros(len(x), dtype=np.int64)
36
+ f[1, 1](r, x)
37
+
38
+ for i in range(len(x)):
39
+ self.assertEqual(r[i], x[i])
40
+
41
+ def test_tuple(self):
42
+ @cuda.jit
43
+ def f(r1, r2, x):
44
+ r1[0] = x[0]
45
+ r1[1] = x[1]
46
+ r1[2] = x[2]
47
+ r2[0] = x[3]
48
+ r2[1] = x[4]
49
+ r2[2] = x[5]
50
+
51
+ x = (1, 2, 3, 4.5, 5.5, 6.5)
52
+ r1 = np.zeros(len(x) // 2, dtype=np.int64)
53
+ r2 = np.zeros(len(x) // 2, dtype=np.float64)
54
+ f[1, 1](r1, r2, x)
55
+
56
+ for i in range(len(r1)):
57
+ self.assertEqual(r1[i], x[i])
58
+
59
+ for i in range(len(r2)):
60
+ self.assertEqual(r2[i], x[i + len(r1)])
61
+
62
+ def test_namedunituple(self):
63
+ @cuda.jit
64
+ def f(r, x):
65
+ r[0] = x.x
66
+ r[1] = x.y
67
+
68
+ Point = namedtuple("Point", ("x", "y"))
69
+ x = Point(1, 2)
70
+ r = np.zeros(len(x), dtype=np.int64)
71
+ f[1, 1](r, x)
72
+
73
+ self.assertEqual(r[0], x.x)
74
+ self.assertEqual(r[1], x.y)
75
+
76
+ def test_namedtuple(self):
77
+ @cuda.jit
78
+ def f(r1, r2, x):
79
+ r1[0] = x.x
80
+ r1[1] = x.y
81
+ r2[0] = x.r
82
+
83
+ Point = namedtuple("Point", ("x", "y", "r"))
84
+ x = Point(1, 2, 2.236)
85
+ r1 = np.zeros(2, dtype=np.int64)
86
+ r2 = np.zeros(1, dtype=np.float64)
87
+ f[1, 1](r1, r2, x)
88
+
89
+ self.assertEqual(r1[0], x.x)
90
+ self.assertEqual(r1[1], x.y)
91
+ self.assertEqual(r2[0], x.r)
92
+
93
+ def test_empty_tuple(self):
94
+ @cuda.jit
95
+ def f(r, x):
96
+ r[0] = len(x)
97
+
98
+ x = tuple()
99
+ r = np.ones(1, dtype=np.int64)
100
+ f[1, 1](r, x)
101
+
102
+ self.assertEqual(r[0], 0)
103
+
104
+ def test_tuple_of_empty_tuples(self):
105
+ @cuda.jit
106
+ def f(r, x):
107
+ r[0] = len(x)
108
+ r[1] = len(x[0])
109
+
110
+ x = ((), (), ())
111
+ r = np.ones(2, dtype=np.int64)
112
+ f[1, 1](r, x)
113
+
114
+ self.assertEqual(r[0], 3)
115
+ self.assertEqual(r[1], 0)
116
+
117
+ def test_tuple_of_tuples(self):
118
+ @cuda.jit
119
+ def f(r, x):
120
+ r[0] = len(x)
121
+ r[1] = len(x[0])
122
+ r[2] = len(x[1])
123
+ r[3] = len(x[2])
124
+ r[4] = x[1][0]
125
+ r[5] = x[1][1]
126
+ r[6] = x[2][0]
127
+ r[7] = x[2][1]
128
+ r[8] = x[2][2]
129
+
130
+ x = ((), (5, 6), (8, 9, 10))
131
+ r = np.ones(9, dtype=np.int64)
132
+ f[1, 1](r, x)
133
+
134
+ self.assertEqual(r[0], 3)
135
+ self.assertEqual(r[1], 0)
136
+ self.assertEqual(r[2], 2)
137
+ self.assertEqual(r[3], 3)
138
+ self.assertEqual(r[4], 5)
139
+ self.assertEqual(r[5], 6)
140
+ self.assertEqual(r[6], 8)
141
+ self.assertEqual(r[7], 9)
142
+ self.assertEqual(r[8], 10)
143
+
144
+ def test_tuple_of_tuples_and_scalars(self):
145
+ @cuda.jit
146
+ def f(r, x):
147
+ r[0] = len(x)
148
+ r[1] = len(x[0])
149
+ r[2] = x[0][0]
150
+ r[3] = x[0][1]
151
+ r[4] = x[0][2]
152
+ r[5] = x[1]
153
+
154
+ x = ((6, 5, 4), 7)
155
+ r = np.ones(9, dtype=np.int64)
156
+ f[1, 1](r, x)
157
+
158
+ self.assertEqual(r[0], 2)
159
+ self.assertEqual(r[1], 3)
160
+ self.assertEqual(r[2], 6)
161
+ self.assertEqual(r[3], 5)
162
+ self.assertEqual(r[4], 4)
163
+ self.assertEqual(r[5], 7)
164
+
165
+ def test_tuple_of_arrays(self):
166
+ @cuda.jit
167
+ def f(x):
168
+ i = cuda.grid(1)
169
+ if i < len(x[0]):
170
+ x[0][i] = x[1][i] + x[2][i]
171
+
172
+ N = 10
173
+ x0 = np.zeros(N)
174
+ x1 = np.ones_like(x0)
175
+ x2 = x1 * 3
176
+ x = (x0, x1, x2)
177
+ f[1, N](x)
178
+
179
+ np.testing.assert_equal(x0, x1 + x2)
180
+
181
+ def test_tuple_of_array_scalar_tuple(self):
182
+ @cuda.jit
183
+ def f(r, x):
184
+ r[0] = x[0][0]
185
+ r[1] = x[0][1]
186
+ r[2] = x[1]
187
+ r[3] = x[2][0]
188
+ r[4] = x[2][1]
189
+
190
+ z = np.arange(2, dtype=np.int64)
191
+ x = (2 * z, 10, (4, 3))
192
+ r = np.zeros(5, dtype=np.int64)
193
+ f[1, 1](r, x)
194
+
195
+ self.assertEqual(r[0], 0)
196
+ self.assertEqual(r[1], 2)
197
+ self.assertEqual(r[2], 10)
198
+ self.assertEqual(r[3], 4)
199
+ self.assertEqual(r[4], 3)
200
+
201
+
202
+ if __name__ == "__main__":
203
+ unittest.main()
@@ -0,0 +1,63 @@
1
+ # SPDX-FileCopyrightText: Copyright (c) 2025 NVIDIA CORPORATION & AFFILIATES. All rights reserved.
2
+ # SPDX-License-Identifier: BSD-2-Clause
3
+
4
+ import numpy as np
5
+ from numba import cuda
6
+ from numba.cuda.testing import CUDATestCase
7
+ import unittest
8
+ from numba.cuda import config
9
+
10
+
11
+ def reinterpret_array_type(byte_arr, start, stop, output):
12
+ # Tested with just one thread
13
+ val = byte_arr[start:stop].view(np.int32)[0]
14
+ output[0] = val
15
+
16
+
17
+ class TestCudaArrayMethods(CUDATestCase):
18
+ def setUp(self):
19
+ self.old_nrt_setting = config.CUDA_ENABLE_NRT
20
+ config.CUDA_ENABLE_NRT = True
21
+ super(TestCudaArrayMethods, self).setUp()
22
+
23
+ def tearDown(self):
24
+ config.CUDA_ENABLE_NRT = self.old_nrt_setting
25
+ super(TestCudaArrayMethods, self).tearDown()
26
+
27
+ def test_reinterpret_array_type(self):
28
+ """
29
+ Reinterpret byte array as int32 in the GPU.
30
+ """
31
+ pyfunc = reinterpret_array_type
32
+ kernel = cuda.jit(pyfunc)
33
+
34
+ byte_arr = np.arange(256, dtype=np.uint8)
35
+ itemsize = np.dtype(np.int32).itemsize
36
+ for start in range(0, 256, itemsize):
37
+ stop = start + itemsize
38
+ expect = byte_arr[start:stop].view(np.int32)[0]
39
+
40
+ output = np.zeros(1, dtype=np.int32)
41
+ kernel[1, 1](byte_arr, start, stop, output)
42
+
43
+ got = output[0]
44
+ self.assertEqual(expect, got)
45
+
46
+ def test_array_copy(self):
47
+ val = np.array([1, 2, 3])[::-1]
48
+
49
+ @cuda.jit
50
+ def kernel(out):
51
+ q = val.copy()
52
+ for i in range(len(out)):
53
+ out[i] = q[i]
54
+
55
+ out = cuda.to_device(np.zeros(len(val), dtype="float64"))
56
+
57
+ kernel[1, 1](out)
58
+ for i, j in zip(out.copy_to_host(), val):
59
+ self.assertEqual(i, j)
60
+
61
+
62
+ if __name__ == "__main__":
63
+ unittest.main()
@@ -0,0 +1,360 @@
1
+ # SPDX-FileCopyrightText: Copyright (c) 2025 NVIDIA CORPORATION & AFFILIATES. All rights reserved.
2
+ # SPDX-License-Identifier: BSD-2-Clause
3
+ import numpy as np
4
+
5
+ from numba.cuda.tests.support import TestCase, MemoryLeakMixin
6
+ from numba import cuda
7
+ from numba.cuda.testing import skip_on_cudasim
8
+ from numba.cuda.misc.special import literal_unroll
9
+ from numba.cuda import config
10
+
11
+
12
+ @skip_on_cudasim("doesn't work in the simulator")
13
+ class TestArrayReductions(MemoryLeakMixin, TestCase):
14
+ """
15
+ Test array reduction methods and functions such as .sum(), .max(), etc.
16
+ """
17
+
18
+ def setUp(self):
19
+ super(TestArrayReductions, self).setUp()
20
+ np.random.seed(42)
21
+ self.old_nrt_setting = config.CUDA_ENABLE_NRT
22
+ self.old_perf_warnings_setting = config.DISABLE_PERFORMANCE_WARNINGS
23
+ config.CUDA_ENABLE_NRT = True
24
+ config.DISABLE_PERFORMANCE_WARNINGS = 1
25
+
26
+ def tearDown(self):
27
+ config.CUDA_ENABLE_NRT = self.old_nrt_setting
28
+ config.DISABLE_PERFORMANCE_WARNINGS = self.old_perf_warnings_setting
29
+ super(TestArrayReductions, self).tearDown()
30
+
31
+ def test_all_basic(self):
32
+ cases = (
33
+ np.float64([1.0, 0.0, float("inf"), float("nan")]),
34
+ np.float64([1.0, -0.0, float("inf"), float("nan")]),
35
+ np.float64([1.0, 1.5, float("inf"), float("nan")]),
36
+ np.float64([[1.0, 1.5], [float("inf"), float("nan")]]),
37
+ np.float64([[1.0, 1.5], [1.5, 1.0]]),
38
+ )
39
+
40
+ @cuda.jit
41
+ def kernel(out):
42
+ i = 0
43
+ for case in literal_unroll(cases):
44
+ out[i] = np.all(case)
45
+ i += 1
46
+
47
+ expected = np.array([np.all(a) for a in cases], dtype=np.bool_)
48
+ out = cuda.to_device(np.zeros(len(cases), dtype=np.bool_))
49
+ kernel[1, 1](out)
50
+ got = out.copy_to_host()
51
+ self.assertPreciseEqual(expected, got)
52
+
53
+ def test_any_basic(self):
54
+ cases = (
55
+ np.float64([0.0, -0.0, 0.0, 0.0]),
56
+ np.float64([0.0, -0.0, np.nan, 0.0]),
57
+ np.float64([0.0, -0.0, float("inf"), 0.0]),
58
+ np.float64([0.0, -0.0, 1.5, 0.0]),
59
+ np.float64([[0.0, -0.0], [1.5, 0.0]]),
60
+ np.float64([[0.0, -0.0], [1.5, 0.0]])[::-1],
61
+ )
62
+
63
+ @cuda.jit
64
+ def kernel(out):
65
+ i = 0
66
+ for arr in literal_unroll(cases):
67
+ out[i] = np.any(arr)
68
+ i += 1
69
+
70
+ expected = np.array([np.any(a) for a in cases], dtype=np.bool_)
71
+ out = cuda.to_device(np.zeros(len(cases), dtype=np.bool_))
72
+ kernel[1, 1](out)
73
+ self.assertPreciseEqual(expected, out.copy_to_host())
74
+
75
+ def test_sum_basic(self):
76
+ arrays = (
77
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
78
+ np.float64([-0.0, -1.5]),
79
+ np.float64([-1.5, 2.5, float("inf")]),
80
+ np.float64([-1.5, 2.5, -float("inf")]),
81
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
82
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
83
+ np.float64(
84
+ [np.nan, -1.5, 2.5, np.nan, float("inf"), -float("inf"), 3.0]
85
+ ),
86
+ np.float64([5.0, np.nan, -1.5, np.nan]),
87
+ np.float64([np.nan, np.nan]),
88
+ )
89
+
90
+ @cuda.jit
91
+ def kernel(out):
92
+ i = 0
93
+ for arr in literal_unroll(arrays):
94
+ out[i] = np.sum(arr)
95
+ i += 1
96
+
97
+ expected = np.array([np.sum(a) for a in arrays], dtype=np.float64)
98
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
99
+ kernel[1, 1](out)
100
+ self.assertPreciseEqual(expected, out.copy_to_host())
101
+
102
+ def test_mean_basic(self):
103
+ arrays = (
104
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
105
+ np.float64([-0.0, -1.5]),
106
+ np.float64([-1.5, 2.5, float("inf")]),
107
+ np.float64([-1.5, 2.5, -float("inf")]),
108
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
109
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
110
+ np.float64(
111
+ [np.nan, -1.5, 2.5, np.nan, float("inf"), -float("inf"), 3.0]
112
+ ),
113
+ np.float64([5.0, np.nan, -1.5, np.nan]),
114
+ np.float64([np.nan, np.nan]),
115
+ )
116
+
117
+ @cuda.jit
118
+ def kernel(out):
119
+ i = 0
120
+ for arr in literal_unroll(arrays):
121
+ out[i] = np.mean(arr)
122
+ i += 1
123
+
124
+ expected = np.array([np.mean(a) for a in arrays], dtype=np.float64)
125
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
126
+ kernel[1, 1](out)
127
+ self.assertPreciseEqual(expected, out.copy_to_host())
128
+
129
+ def test_var_basic(self):
130
+ arrays = (
131
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
132
+ np.float64([-0.0, -1.5]),
133
+ np.float64([-1.5, 2.5, float("inf")]),
134
+ np.float64([-1.5, 2.5, -float("inf")]),
135
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
136
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
137
+ np.float64(
138
+ [np.nan, -1.5, 2.5, np.nan, float("inf"), -float("inf"), 3.0]
139
+ ),
140
+ np.float64([5.0, np.nan, -1.5, np.nan]),
141
+ np.float64([np.nan, np.nan]),
142
+ )
143
+
144
+ @cuda.jit
145
+ def kernel(out):
146
+ i = 0
147
+ for arr in literal_unroll(arrays):
148
+ out[i] = np.var(arr)
149
+ i += 1
150
+
151
+ expected = np.array([np.var(a) for a in arrays], dtype=np.float64)
152
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
153
+ kernel[1, 1](out)
154
+ self.assertPreciseEqual(expected, out.copy_to_host(), prec="double")
155
+
156
+ def test_std_basic(self):
157
+ arrays = (
158
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
159
+ np.float64([-0.0, -1.5]),
160
+ np.float64([-1.5, 2.5, float("inf")]),
161
+ np.float64([-1.5, 2.5, -float("inf")]),
162
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
163
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
164
+ np.float64(
165
+ [np.nan, -1.5, 2.5, np.nan, float("inf"), -float("inf"), 3.0]
166
+ ),
167
+ np.float64([5.0, np.nan, -1.5, np.nan]),
168
+ np.float64([np.nan, np.nan]),
169
+ )
170
+
171
+ @cuda.jit
172
+ def kernel(out):
173
+ i = 0
174
+ for arr in literal_unroll(arrays):
175
+ out[i] = np.std(arr)
176
+ i += 1
177
+
178
+ expected = np.array([np.std(a) for a in arrays], dtype=np.float64)
179
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
180
+ kernel[1, 1](out)
181
+ self.assertPreciseEqual(expected, out.copy_to_host())
182
+
183
+ def test_min_basic(self):
184
+ arrays = (
185
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
186
+ np.float64([-0.0, -1.5]),
187
+ np.float64([-1.5, 2.5, float("inf")]),
188
+ np.float64([-1.5, 2.5, -float("inf")]),
189
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
190
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
191
+ np.float64(
192
+ [np.nan, -1.5, 2.5, np.nan, float("inf"), -float("inf"), 3.0]
193
+ ),
194
+ np.float64([5.0, np.nan, -1.5, np.nan]),
195
+ np.float64([np.nan, np.nan]),
196
+ )
197
+
198
+ @cuda.jit
199
+ def kernel(out):
200
+ i = 0
201
+ for arr in literal_unroll(arrays):
202
+ out[i] = np.min(arr)
203
+ i += 1
204
+
205
+ expected = np.array([np.min(a) for a in arrays], dtype=np.float64)
206
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
207
+ kernel[1, 1](out)
208
+ self.assertPreciseEqual(expected, out.copy_to_host())
209
+
210
+ def test_max_basic(self):
211
+ arrays = (
212
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
213
+ np.float64([-0.0, -1.5]),
214
+ np.float64([-1.5, 2.5, float("inf")]),
215
+ np.float64([-1.5, 2.5, -float("inf")]),
216
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
217
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
218
+ np.float64(
219
+ [np.nan, -1.5, 2.5, np.nan, float("inf"), -float("inf"), 3.0]
220
+ ),
221
+ np.float64([5.0, np.nan, -1.5, np.nan]),
222
+ np.float64([np.nan, np.nan]),
223
+ )
224
+
225
+ @cuda.jit
226
+ def kernel(out):
227
+ i = 0
228
+ for arr in literal_unroll(arrays):
229
+ out[i] = np.max(arr)
230
+ i += 1
231
+
232
+ expected = np.array([np.max(a) for a in arrays], dtype=np.float64)
233
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
234
+ kernel[1, 1](out)
235
+ self.assertPreciseEqual(expected, out.copy_to_host())
236
+
237
+ def test_nanmin_basic(self):
238
+ arrays = (
239
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
240
+ np.float64([-0.0, -1.5]),
241
+ np.float64([-1.5, 2.5, np.nan]),
242
+ np.float64([-1.5, 2.5, float("inf")]),
243
+ np.float64([-1.5, 2.5, -float("inf")]),
244
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
245
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
246
+ np.float64([5.0, np.nan, -1.5, np.nan]),
247
+ np.float64([np.nan, np.nan]),
248
+ )
249
+
250
+ @cuda.jit
251
+ def kernel(out):
252
+ i = 0
253
+ for arr in literal_unroll(arrays):
254
+ out[i] = np.nanmin(arr)
255
+ i += 1
256
+
257
+ expected = np.array([np.nanmin(a) for a in arrays], dtype=np.float64)
258
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
259
+ kernel[1, 1](out)
260
+ self.assertPreciseEqual(expected, out.copy_to_host())
261
+
262
+ def test_nanmax_basic(self):
263
+ arrays = (
264
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
265
+ np.float64([-0.0, -1.5]),
266
+ np.float64([-1.5, 2.5, np.nan]),
267
+ np.float64([-1.5, 2.5, float("inf")]),
268
+ np.float64([-1.5, 2.5, -float("inf")]),
269
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
270
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
271
+ np.float64([5.0, np.nan, -1.5, np.nan]),
272
+ np.float64([np.nan, np.nan]),
273
+ )
274
+
275
+ @cuda.jit
276
+ def kernel(out):
277
+ i = 0
278
+ for arr in literal_unroll(arrays):
279
+ out[i] = np.nanmax(arr)
280
+ i += 1
281
+
282
+ expected = np.array([np.nanmax(a) for a in arrays], dtype=np.float64)
283
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
284
+ kernel[1, 1](out)
285
+ self.assertPreciseEqual(expected, out.copy_to_host())
286
+
287
+ def test_nanmean_basic(self):
288
+ arrays = (
289
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
290
+ np.float64([-0.0, -1.5]),
291
+ np.float64([-1.5, 2.5, np.nan]),
292
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
293
+ np.float64(
294
+ [np.nan, -1.5, 2.5, np.nan, float("inf"), -float("inf"), 3.0]
295
+ ),
296
+ np.float64([5.0, np.nan, -1.5, np.nan]),
297
+ np.float64([np.nan, np.nan]),
298
+ )
299
+
300
+ @cuda.jit
301
+ def kernel(out):
302
+ i = 0
303
+ for arr in literal_unroll(arrays):
304
+ out[i] = np.nanmean(arr)
305
+ i += 1
306
+
307
+ expected = np.array([np.nanmean(a) for a in arrays], dtype=np.float64)
308
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
309
+ kernel[1, 1](out)
310
+ self.assertPreciseEqual(expected, out.copy_to_host())
311
+
312
+ def test_nansum_basic(self):
313
+ arrays = (
314
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
315
+ np.float64([-0.0, -1.5]),
316
+ np.float64([-1.5, 2.5, np.nan]),
317
+ np.float64([-1.5, 2.5, float("inf")]),
318
+ np.float64([-1.5, 2.5, -float("inf")]),
319
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
320
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
321
+ np.float64([5.0, np.nan, -1.5, np.nan]),
322
+ np.float64([np.nan, np.nan]),
323
+ )
324
+
325
+ @cuda.jit
326
+ def kernel(out):
327
+ i = 0
328
+ for arr in literal_unroll(arrays):
329
+ out[i] = np.nansum(arr)
330
+ i += 1
331
+
332
+ expected = np.array([np.nansum(a) for a in arrays], dtype=np.float64)
333
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
334
+ kernel[1, 1](out)
335
+ self.assertPreciseEqual(expected, out.copy_to_host())
336
+
337
+ def test_nanprod_basic(self):
338
+ arrays = (
339
+ np.float64([1.0, 2.0, 0.0, -0.0, 1.0, -1.5]),
340
+ np.float64([-0.0, -1.5]),
341
+ np.float64([-1.5, 2.5, np.nan]),
342
+ np.float64([-1.5, 2.5, float("inf")]),
343
+ np.float64([-1.5, 2.5, -float("inf")]),
344
+ np.float64([-1.5, 2.5, float("inf"), -float("inf")]),
345
+ np.float64([np.nan, -1.5, 2.5, np.nan, 3.0]),
346
+ np.float64([5.0, np.nan, -1.5, np.nan]),
347
+ np.float64([np.nan, np.nan]),
348
+ )
349
+
350
+ @cuda.jit
351
+ def kernel(out):
352
+ i = 0
353
+ for arr in literal_unroll(arrays):
354
+ out[i] = np.nanprod(arr)
355
+ i += 1
356
+
357
+ expected = np.array([np.nanprod(a) for a in arrays], dtype=np.float64)
358
+ out = cuda.to_device(np.zeros(len(arrays), dtype=np.float64))
359
+ kernel[1, 1](out)
360
+ self.assertPreciseEqual(expected, out.copy_to_host())