mindstudio-probe 1.0.3__py3-none-any.whl → 1.1.0__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (278) hide show
  1. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/LICENSE +201 -201
  2. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/METADATA +36 -34
  3. mindstudio_probe-1.1.0.dist-info/RECORD +287 -0
  4. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/WHEEL +1 -1
  5. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/entry_points.txt +1 -0
  6. msprobe/README.md +131 -237
  7. msprobe/__init__.py +16 -1
  8. msprobe/{config/config.json → config.json} +47 -49
  9. msprobe/core/advisor/advisor.py +124 -124
  10. msprobe/core/advisor/advisor_const.py +58 -59
  11. msprobe/core/advisor/advisor_result.py +58 -58
  12. msprobe/core/common/const.py +402 -318
  13. msprobe/core/common/exceptions.py +99 -99
  14. msprobe/core/common/{file_check.py → file_utils.py} +523 -283
  15. msprobe/core/common/inplace_op_checker.py +38 -0
  16. msprobe/core/common/inplace_ops.yaml +251 -0
  17. msprobe/core/common/log.py +86 -69
  18. msprobe/core/common/utils.py +371 -616
  19. msprobe/core/common_config.py +78 -71
  20. msprobe/core/compare/acc_compare.py +472 -298
  21. msprobe/core/compare/check.py +180 -95
  22. msprobe/core/compare/compare_cli.py +69 -49
  23. msprobe/core/compare/highlight.py +259 -222
  24. msprobe/core/compare/multiprocessing_compute.py +174 -149
  25. msprobe/core/compare/npy_compare.py +310 -295
  26. msprobe/core/compare/utils.py +464 -429
  27. msprobe/core/data_dump/data_collector.py +153 -144
  28. msprobe/core/data_dump/data_processor/base.py +337 -293
  29. msprobe/core/data_dump/data_processor/factory.py +76 -59
  30. msprobe/core/data_dump/data_processor/mindspore_processor.py +192 -198
  31. msprobe/core/data_dump/data_processor/pytorch_processor.py +383 -389
  32. msprobe/core/data_dump/json_writer.py +117 -116
  33. msprobe/core/data_dump/scope.py +194 -178
  34. msprobe/core/grad_probe/constant.py +74 -70
  35. msprobe/core/grad_probe/grad_compare.py +170 -175
  36. msprobe/core/grad_probe/utils.py +77 -52
  37. msprobe/docs/01.installation.md +99 -0
  38. msprobe/docs/02.config_introduction.md +137 -0
  39. msprobe/docs/03.config_examples.md +237 -0
  40. msprobe/docs/04.acl_config_examples.md +78 -0
  41. msprobe/docs/05.data_dump_PyTorch.md +326 -0
  42. msprobe/docs/06.data_dump_MindSpore.md +285 -0
  43. msprobe/docs/07.accuracy_checker_PyTorch.md +297 -0
  44. msprobe/docs/08.accuracy_checker_online_PyTorch.md +238 -0
  45. msprobe/docs/09.accuracy_checker_MindSpore.md +68 -0
  46. msprobe/docs/10.accuracy_compare_PyTorch.md +327 -0
  47. msprobe/docs/11.accuracy_compare_MindSpore.md +333 -0
  48. msprobe/docs/12.overflow_check_PyTorch.md +79 -0
  49. msprobe/docs/13.overflow_check_MindSpore.md +31 -0
  50. msprobe/{pytorch/doc/parse_tool.md → docs/14.data_parse_PyTorch.md} +283 -286
  51. msprobe/docs/15.free_benchmarking_PyTorch.md +170 -0
  52. msprobe/docs/16.free_benchmarking_MindSpore.md +140 -0
  53. msprobe/{doc/grad_probe/grad_probe.md → docs/17.grad_probe.md} +205 -207
  54. msprobe/{pytorch/doc//321/205/320/254/320/270/321/207/342/225/221/342/224/220/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/206/320/277/320/244/321/205/320/277/342/225/243.md → docs/18.online_dispatch.md} +89 -90
  55. msprobe/docs/FAQ.md +189 -0
  56. msprobe/docs/S02.report_free_benchmarking_validation_performance_baseline.md +146 -0
  57. msprobe/docs/img/free_benchmark_framework.png +0 -0
  58. msprobe/docs/img/ms_dump.png +0 -0
  59. msprobe/docs/img/ms_layer.png +0 -0
  60. msprobe/docs/img/pt_dump.png +0 -0
  61. msprobe/mindspore/__init__.py +2 -1
  62. msprobe/mindspore/api_accuracy_checker/api_accuracy_checker.py +278 -245
  63. msprobe/mindspore/api_accuracy_checker/api_info.py +76 -69
  64. msprobe/mindspore/api_accuracy_checker/api_runner.py +155 -151
  65. msprobe/mindspore/api_accuracy_checker/base_compare_algorithm.py +196 -196
  66. msprobe/mindspore/api_accuracy_checker/cmd_parser.py +6 -0
  67. msprobe/mindspore/api_accuracy_checker/compute_element.py +238 -223
  68. msprobe/mindspore/api_accuracy_checker/main.py +8 -15
  69. msprobe/mindspore/api_accuracy_checker/type_mapping.py +113 -113
  70. msprobe/mindspore/api_accuracy_checker/utils.py +79 -62
  71. msprobe/mindspore/cell_processor.py +58 -34
  72. msprobe/mindspore/common/const.py +108 -87
  73. msprobe/mindspore/common/log.py +37 -37
  74. msprobe/mindspore/common/utils.py +97 -57
  75. msprobe/mindspore/compare/distributed_compare.py +62 -75
  76. msprobe/mindspore/compare/layer_mapping.py +146 -0
  77. msprobe/mindspore/compare/modify_mapping.py +107 -0
  78. msprobe/mindspore/compare/ms_compare.py +357 -117
  79. msprobe/mindspore/compare/ms_graph_compare.py +364 -317
  80. msprobe/mindspore/compare/ms_to_pt_api.yaml +399 -399
  81. msprobe/mindspore/debugger/debugger_config.py +69 -74
  82. msprobe/mindspore/debugger/precision_debugger.py +150 -107
  83. msprobe/mindspore/dump/dump_tool_factory.py +50 -35
  84. msprobe/mindspore/dump/hook_cell/api_registry.py +128 -104
  85. msprobe/mindspore/dump/hook_cell/hook_cell.py +55 -53
  86. msprobe/mindspore/dump/hook_cell/primitive_hooks.py +206 -0
  87. msprobe/mindspore/dump/hook_cell/support_wrap_ops.yaml +994 -925
  88. msprobe/mindspore/dump/hook_cell/wrap_api.py +121 -0
  89. msprobe/mindspore/dump/jit_dump.py +96 -56
  90. msprobe/mindspore/dump/kernel_graph_dump.py +75 -60
  91. msprobe/mindspore/dump/kernel_kbyk_dump.py +79 -65
  92. msprobe/mindspore/free_benchmark/api_pynative_self_check.py +131 -116
  93. msprobe/mindspore/free_benchmark/common/config.py +27 -12
  94. msprobe/mindspore/free_benchmark/common/handler_params.py +32 -17
  95. msprobe/mindspore/free_benchmark/common/utils.py +85 -71
  96. msprobe/mindspore/free_benchmark/data/support_wrap_ops.yaml +842 -842
  97. msprobe/mindspore/free_benchmark/decorator/dec_forward.py +57 -42
  98. msprobe/mindspore/free_benchmark/decorator/decorator_factory.py +122 -107
  99. msprobe/mindspore/free_benchmark/handler/base_handler.py +105 -90
  100. msprobe/mindspore/free_benchmark/handler/check_handler.py +56 -41
  101. msprobe/mindspore/free_benchmark/handler/fix_handler.py +51 -36
  102. msprobe/mindspore/free_benchmark/handler/handler_factory.py +36 -21
  103. msprobe/mindspore/free_benchmark/perturbation/add_noise.py +82 -67
  104. msprobe/mindspore/free_benchmark/perturbation/base_perturbation.py +36 -21
  105. msprobe/mindspore/free_benchmark/perturbation/bit_noise.py +78 -63
  106. msprobe/mindspore/free_benchmark/perturbation/exchange_value.py +77 -0
  107. msprobe/mindspore/free_benchmark/perturbation/improve_precision.py +49 -34
  108. msprobe/mindspore/free_benchmark/perturbation/no_change.py +27 -12
  109. msprobe/mindspore/free_benchmark/perturbation/perturbation_factory.py +44 -27
  110. msprobe/mindspore/free_benchmark/self_check_tool_factory.py +48 -33
  111. msprobe/mindspore/grad_probe/global_context.py +100 -91
  112. msprobe/mindspore/grad_probe/grad_analyzer.py +231 -231
  113. msprobe/mindspore/grad_probe/grad_monitor.py +27 -27
  114. msprobe/mindspore/grad_probe/grad_stat_csv.py +131 -131
  115. msprobe/mindspore/grad_probe/hook.py +94 -92
  116. msprobe/mindspore/grad_probe/utils.py +29 -28
  117. msprobe/mindspore/ms_config.py +128 -126
  118. msprobe/mindspore/overflow_check/kernel_graph_overflow_check.py +60 -45
  119. msprobe/mindspore/overflow_check/overflow_check_tool_factory.py +49 -34
  120. msprobe/mindspore/runtime.py +4 -4
  121. msprobe/mindspore/service.py +297 -354
  122. msprobe/mindspore/task_handler_factory.py +24 -24
  123. msprobe/msprobe.py +105 -107
  124. msprobe/pytorch/__init__.py +23 -4
  125. msprobe/pytorch/api_accuracy_checker/common/config.py +70 -55
  126. msprobe/pytorch/api_accuracy_checker/common/utils.py +246 -165
  127. msprobe/pytorch/api_accuracy_checker/compare/algorithm.py +230 -213
  128. msprobe/pytorch/api_accuracy_checker/compare/api_precision_compare.py +632 -581
  129. msprobe/pytorch/api_accuracy_checker/compare/api_precision_standard.yaml +132 -132
  130. msprobe/pytorch/api_accuracy_checker/compare/api_precision_threshold.yaml +390 -390
  131. msprobe/pytorch/api_accuracy_checker/compare/compare.py +416 -381
  132. msprobe/pytorch/api_accuracy_checker/compare/compare_column.py +90 -73
  133. msprobe/pytorch/api_accuracy_checker/compare/compare_utils.py +265 -244
  134. msprobe/pytorch/api_accuracy_checker/config.yaml +10 -10
  135. msprobe/pytorch/api_accuracy_checker/run_ut/data_generate.py +370 -332
  136. msprobe/pytorch/api_accuracy_checker/run_ut/multi_run_ut.py +221 -199
  137. msprobe/pytorch/api_accuracy_checker/run_ut/run_overflow_check.py +150 -134
  138. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut.py +518 -581
  139. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut_utils.py +213 -74
  140. msprobe/pytorch/api_accuracy_checker/run_ut/torch_ut_setting.json +7 -4
  141. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/attl.py +218 -202
  142. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/client.py +370 -324
  143. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/device_dispatch.py +227 -204
  144. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/dump_dispatch.py +110 -0
  145. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/server.py +244 -218
  146. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/torch_ops_config.yaml +63 -0
  147. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/utils.py +44 -0
  148. msprobe/pytorch/bench_functions/__init__.py +30 -15
  149. msprobe/pytorch/bench_functions/apply_adam_w.py +43 -28
  150. msprobe/pytorch/bench_functions/confusion_transpose.py +34 -19
  151. msprobe/pytorch/bench_functions/fast_gelu.py +70 -55
  152. msprobe/pytorch/bench_functions/layer_norm_eval.py +21 -6
  153. msprobe/pytorch/bench_functions/linear.py +27 -12
  154. msprobe/pytorch/bench_functions/matmul_backward.py +63 -48
  155. msprobe/pytorch/bench_functions/npu_fusion_attention.py +538 -421
  156. msprobe/pytorch/bench_functions/rms_norm.py +30 -15
  157. msprobe/pytorch/bench_functions/rotary_mul.py +71 -52
  158. msprobe/pytorch/bench_functions/scaled_mask_softmax.py +41 -26
  159. msprobe/pytorch/bench_functions/swiglu.py +70 -55
  160. msprobe/pytorch/common/__init__.py +17 -2
  161. msprobe/pytorch/common/compare_script.template +14 -14
  162. msprobe/pytorch/common/log.py +33 -32
  163. msprobe/pytorch/common/parse_json.py +54 -39
  164. msprobe/pytorch/common/utils.py +310 -300
  165. msprobe/pytorch/compare/distributed_compare.py +66 -66
  166. msprobe/pytorch/compare/mapping.yaml +607 -607
  167. msprobe/pytorch/compare/match.py +49 -33
  168. msprobe/pytorch/compare/pt_compare.py +82 -40
  169. msprobe/pytorch/debugger/debugger_config.py +108 -95
  170. msprobe/pytorch/debugger/precision_debugger.py +173 -125
  171. msprobe/pytorch/free_benchmark/__init__.py +23 -8
  172. msprobe/pytorch/free_benchmark/common/constant.py +70 -70
  173. msprobe/pytorch/free_benchmark/common/counter.py +71 -71
  174. msprobe/pytorch/free_benchmark/common/enums.py +65 -37
  175. msprobe/pytorch/free_benchmark/common/params.py +144 -129
  176. msprobe/pytorch/free_benchmark/common/utils.py +118 -102
  177. msprobe/pytorch/free_benchmark/compare/grad_saver.py +200 -179
  178. msprobe/pytorch/free_benchmark/compare/single_benchmark.py +119 -104
  179. msprobe/pytorch/free_benchmark/main.py +120 -105
  180. msprobe/pytorch/free_benchmark/perturbed_layers/base_layer.py +28 -13
  181. msprobe/pytorch/free_benchmark/perturbed_layers/layer_factory.py +56 -41
  182. msprobe/pytorch/free_benchmark/perturbed_layers/npu/add_noise.py +105 -90
  183. msprobe/pytorch/free_benchmark/perturbed_layers/npu/bit_noise.py +119 -104
  184. msprobe/pytorch/free_benchmark/perturbed_layers/npu/change_value.py +87 -63
  185. msprobe/pytorch/free_benchmark/perturbed_layers/npu/improve_precision.py +83 -68
  186. msprobe/pytorch/free_benchmark/perturbed_layers/npu/no_change.py +43 -28
  187. msprobe/pytorch/free_benchmark/perturbed_layers/npu/npu_base_layser.py +60 -45
  188. msprobe/pytorch/free_benchmark/perturbed_layers/run_cpu.py +34 -19
  189. msprobe/pytorch/free_benchmark/result_handlers/base_handler.py +256 -217
  190. msprobe/pytorch/free_benchmark/result_handlers/check_handler.py +54 -39
  191. msprobe/pytorch/free_benchmark/result_handlers/fix_handler.py +38 -23
  192. msprobe/pytorch/free_benchmark/result_handlers/handler_factory.py +45 -30
  193. msprobe/pytorch/free_benchmark/result_handlers/preheat_handler.py +185 -170
  194. msprobe/pytorch/function_factory.py +91 -75
  195. msprobe/pytorch/functional/module_dump.py +84 -0
  196. msprobe/pytorch/grad_probe/grad_monitor.py +91 -90
  197. msprobe/pytorch/grad_probe/grad_stat_csv.py +128 -128
  198. msprobe/pytorch/hook_module/__init__.py +16 -1
  199. msprobe/pytorch/hook_module/api_registry.py +166 -161
  200. msprobe/pytorch/hook_module/hook_module.py +118 -120
  201. msprobe/pytorch/hook_module/support_wrap_ops.yaml +1879 -1877
  202. msprobe/pytorch/hook_module/utils.py +28 -29
  203. msprobe/pytorch/hook_module/wrap_aten.py +111 -110
  204. msprobe/pytorch/hook_module/wrap_distributed.py +77 -78
  205. msprobe/pytorch/hook_module/wrap_functional.py +104 -105
  206. msprobe/pytorch/hook_module/wrap_npu_custom.py +85 -84
  207. msprobe/pytorch/hook_module/wrap_tensor.py +69 -71
  208. msprobe/pytorch/hook_module/wrap_torch.py +84 -86
  209. msprobe/pytorch/hook_module/wrap_vf.py +60 -62
  210. msprobe/pytorch/module_processer.py +153 -138
  211. msprobe/pytorch/online_dispatch/__init__.py +20 -20
  212. msprobe/pytorch/online_dispatch/compare.py +235 -236
  213. msprobe/pytorch/online_dispatch/dispatch.py +271 -271
  214. msprobe/pytorch/online_dispatch/dump_compare.py +155 -156
  215. msprobe/pytorch/online_dispatch/single_compare.py +391 -391
  216. msprobe/pytorch/online_dispatch/torch_ops_config.yaml +57 -49
  217. msprobe/pytorch/online_dispatch/utils.py +127 -146
  218. msprobe/pytorch/parse.py +19 -4
  219. msprobe/pytorch/parse_tool/cli.py +31 -32
  220. msprobe/pytorch/parse_tool/lib/compare.py +259 -271
  221. msprobe/pytorch/parse_tool/lib/config.py +52 -52
  222. msprobe/pytorch/parse_tool/lib/file_desc.py +31 -31
  223. msprobe/pytorch/parse_tool/lib/interactive_cli.py +102 -102
  224. msprobe/pytorch/parse_tool/lib/parse_exception.py +54 -54
  225. msprobe/pytorch/parse_tool/lib/parse_tool.py +161 -158
  226. msprobe/pytorch/parse_tool/lib/utils.py +320 -321
  227. msprobe/pytorch/parse_tool/lib/visualization.py +85 -91
  228. msprobe/pytorch/pt_config.py +317 -187
  229. msprobe/pytorch/service.py +311 -252
  230. mindstudio_probe-1.0.3.dist-info/RECORD +0 -272
  231. msprobe/config/README.md +0 -539
  232. msprobe/mindspore/doc/compare.md +0 -58
  233. msprobe/mindspore/doc/dump.md +0 -217
  234. msprobe/mindspore/dump/hook_cell/wrap_functional.py +0 -91
  235. msprobe/mindspore/dump/hook_cell/wrap_tensor.py +0 -63
  236. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/ssl_config.py +0 -10
  237. msprobe/pytorch/doc/FAQ.md +0 -193
  238. msprobe/pytorch/doc/api_accuracy_checker.md +0 -313
  239. msprobe/pytorch/doc/api_accuracy_checker_online.md +0 -187
  240. msprobe/pytorch/doc/dump.md +0 -260
  241. msprobe/pytorch/doc/msprobe/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/206/320/245/342/226/221/321/206/320/235/320/276dump/321/206/320/260/320/227/321/205/320/227/320/226/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -182
  242. msprobe/pytorch/doc/ptdbg_ascend_compare.md +0 -240
  243. msprobe/pytorch/doc/ptdbg_ascend_overview.md +0 -68
  244. msprobe/pytorch/doc/ptdbg_ascend_quickstart.md +0 -381
  245. msprobe/pytorch/doc/run_overflow_check.md +0 -25
  246. msprobe/pytorch/doc//321/206/320/247/320/260/321/206/320/260/320/227/321/206/320/255/320/226/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/205/320/254/342/225/221/321/206/320/251/320/277/321/211/320/272/320/234/321/210/320/277/320/221/321/205/320/242/320/234/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -151
  247. msprobe/pytorch/functional/data_processor.py +0 -0
  248. msprobe/pytorch/functional/dump_module.py +0 -39
  249. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/top_level.txt +0 -0
  250. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_1.png +0 -0
  251. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_2.png +0 -0
  252. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_3.png +0 -0
  253. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_4.png +0 -0
  254. /msprobe/{pytorch/doc → docs}/img/GPT-3_1.png +0 -0
  255. /msprobe/{pytorch/doc → docs}/img/GPT-3_2.png +0 -0
  256. /msprobe/{pytorch/doc → docs}/img/GPT-3_3.png +0 -0
  257. /msprobe/{pytorch/doc → docs}/img/GPT-3_4.png +0 -0
  258. /msprobe/{pytorch/doc → docs}/img/GPT-3_5.png +0 -0
  259. /msprobe/{pytorch/doc → docs}/img/GPT-3_6.png +0 -0
  260. /msprobe/{pytorch/doc → docs}/img/GPT-3_7.png +0 -0
  261. /msprobe/{pytorch/doc → docs}/img/GPT-3_8.png +0 -0
  262. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_1.png +0 -0
  263. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_2.png +0 -0
  264. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_details.png +0 -0
  265. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_result.png +0 -0
  266. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_details.png +0 -0
  267. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_result.png +0 -0
  268. /msprobe/{pytorch/doc → docs}/img/auto_analyze_log.png +0 -0
  269. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl.png +0 -0
  270. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl_md5.png.png +0 -0
  271. /msprobe/{pytorch/doc → docs}/img/cpu_info.png +0 -0
  272. /msprobe/{config → docs}/img/free_benchmark.png +0 -0
  273. /msprobe/{doc/grad_probe/img/image-1.png → docs/img/grad_probe_image-1.png} +0 -0
  274. /msprobe/{doc/grad_probe/img/image-2.png → docs/img/grad_probe_image-2.png} +0 -0
  275. /msprobe/{doc/grad_probe/img/image-3.png → docs/img/grad_probe_image-3.png} +0 -0
  276. /msprobe/{doc/grad_probe/img/image-4.png → docs/img/grad_probe_image-4.png} +0 -0
  277. /msprobe/{doc/grad_probe/img/image.png → docs/img/grad_probe_image.png} +0 -0
  278. /msprobe/{pytorch/doc → docs}/img/module_compare.png +0 -0
@@ -1,59 +1,76 @@
1
- from msprobe.core.common.const import Const
2
-
3
-
4
- class DataProcessorFactory:
5
- _data_processor = {}
6
- _module_processor = {}
7
-
8
- @classmethod
9
- def register_processor(cls, framework, task, processor_class):
10
- key = (framework, task)
11
- cls._data_processor[key] = processor_class
12
-
13
- @classmethod
14
- def register_module_processor(cls, framework, processor_class):
15
- cls._module_processor[framework] = processor_class
16
-
17
- @classmethod
18
- def get_module_processor(cls, framework):
19
- processor_class = cls._module_processor.get(framework)
20
- if not processor_class:
21
- raise ValueError(f"ModuleProcesser not found for framework: {framework}")
22
- return processor_class
23
-
24
- @classmethod
25
- def create_processor(cls, config, data_writer):
26
- cls.register_processors(config.framework)
27
- task = Const.KERNEL_DUMP if config.level == "L2" else config.task
28
- key = (config.framework, task)
29
- processor_class = cls._data_processor.get(key)
30
- if not processor_class:
31
- raise ValueError(f"Processor not found for framework: {config.framework}, task: {config.task}")
32
- return processor_class(config, data_writer)
33
-
34
- @classmethod
35
- def register_processors(cls, framework):
36
- if framework == Const.PT_FRAMEWORK:
37
- from .pytorch_processor import (
38
- StatisticsDataProcessor as PytorchStatisticsDataProcessor,
39
- TensorDataProcessor as PytorchTensorDataProcessor,
40
- OverflowCheckDataProcessor as PytorchOverflowCheckDataProcessor,
41
- FreeBenchmarkDataProcessor as PytorchFreeBenchmarkDataProcessor,
42
- KernelDumpDataProcessor as PytorchKernelDumpDataProcessor
43
- )
44
- from ....pytorch.module_processer import ModuleProcesser
45
- cls.register_processor(Const.PT_FRAMEWORK, Const.STATISTICS, PytorchStatisticsDataProcessor)
46
- cls.register_processor(Const.PT_FRAMEWORK, Const.TENSOR, PytorchTensorDataProcessor)
47
- cls.register_processor(Const.PT_FRAMEWORK, Const.OVERFLOW_CHECK, PytorchOverflowCheckDataProcessor)
48
- cls.register_processor(Const.PT_FRAMEWORK, Const.FREE_BENCHMARK, PytorchFreeBenchmarkDataProcessor)
49
- cls.register_processor(Const.PT_FRAMEWORK, Const.KERNEL_DUMP, PytorchKernelDumpDataProcessor)
50
- cls.register_module_processor(Const.PT_FRAMEWORK, ModuleProcesser)
51
- elif framework == Const.MS_FRAMEWORK:
52
- from .mindspore_processor import (
53
- StatisticsDataProcessor as MindsporeStatisticsDataProcessor,
54
- TensorDataProcessor as MindsporeTensorDataProcessor,
55
- OverflowCheckDataProcessor as MindsporeOverflowCheckDataProcessor
56
- )
57
- cls.register_processor(Const.MS_FRAMEWORK, Const.STATISTICS, MindsporeStatisticsDataProcessor)
58
- cls.register_processor(Const.MS_FRAMEWORK, Const.TENSOR, MindsporeTensorDataProcessor)
59
- cls.register_processor(Const.MS_FRAMEWORK, Const.OVERFLOW_CHECK, MindsporeOverflowCheckDataProcessor)
1
+ # Copyright (c) 2024-2024, Huawei Technologies Co., Ltd.
2
+ # All rights reserved.
3
+ #
4
+ # Licensed under the Apache License, Version 2.0 (the "License");
5
+ # you may not use this file except in compliance with the License.
6
+ # You may obtain a copy of the License at
7
+ #
8
+ # http://www.apache.org/licenses/LICENSE-2.0
9
+ #
10
+ # Unless required by applicable law or agreed to in writing, software
11
+ # distributed under the License is distributed on an "AS IS" BASIS,
12
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
13
+ # See the License for the specific language governing permissions and
14
+ # limitations under the License.
15
+
16
+ from msprobe.core.common.const import Const
17
+
18
+
19
+ class DataProcessorFactory:
20
+ _data_processor = {}
21
+ _module_processor = {}
22
+
23
+ @classmethod
24
+ def register_processor(cls, framework, task, processor_class):
25
+ key = (framework, task)
26
+ cls._data_processor[key] = processor_class
27
+
28
+ @classmethod
29
+ def register_module_processor(cls, framework, processor_class):
30
+ cls._module_processor[framework] = processor_class
31
+
32
+ @classmethod
33
+ def get_module_processor(cls, framework):
34
+ processor_class = cls._module_processor.get(framework)
35
+ if not processor_class:
36
+ raise ValueError(f"ModuleProcesser not found for framework: {framework}")
37
+ return processor_class
38
+
39
+ @classmethod
40
+ def create_processor(cls, config, data_writer):
41
+ cls.register_processors(config.framework)
42
+ task = Const.KERNEL_DUMP if config.level == "L2" else config.task
43
+ key = (config.framework, task)
44
+ processor_class = cls._data_processor.get(key)
45
+ if not processor_class:
46
+ raise ValueError(f"Processor not found for framework: {config.framework}, task: {config.task}")
47
+ return processor_class(config, data_writer)
48
+
49
+ @classmethod
50
+ def register_processors(cls, framework):
51
+ if framework == Const.PT_FRAMEWORK:
52
+ from msprobe.core.data_dump.data_processor.pytorch_processor import (
53
+ StatisticsDataProcessor as PytorchStatisticsDataProcessor,
54
+ TensorDataProcessor as PytorchTensorDataProcessor,
55
+ OverflowCheckDataProcessor as PytorchOverflowCheckDataProcessor,
56
+ FreeBenchmarkDataProcessor as PytorchFreeBenchmarkDataProcessor,
57
+ KernelDumpDataProcessor as PytorchKernelDumpDataProcessor
58
+ )
59
+ from msprobe.pytorch.module_processer import ModuleProcesser
60
+ cls.register_processor(Const.PT_FRAMEWORK, Const.STATISTICS, PytorchStatisticsDataProcessor)
61
+ cls.register_processor(Const.PT_FRAMEWORK, Const.TENSOR, PytorchTensorDataProcessor)
62
+ cls.register_processor(Const.PT_FRAMEWORK, Const.OVERFLOW_CHECK, PytorchOverflowCheckDataProcessor)
63
+ cls.register_processor(Const.PT_FRAMEWORK, Const.FREE_BENCHMARK, PytorchFreeBenchmarkDataProcessor)
64
+ cls.register_processor(Const.PT_FRAMEWORK, Const.KERNEL_DUMP, PytorchKernelDumpDataProcessor)
65
+ cls.register_module_processor(Const.PT_FRAMEWORK, ModuleProcesser)
66
+ elif framework == Const.MS_FRAMEWORK:
67
+ from msprobe.core.data_dump.data_processor.mindspore_processor import (
68
+ StatisticsDataProcessor as MindsporeStatisticsDataProcessor,
69
+ TensorDataProcessor as MindsporeTensorDataProcessor,
70
+ OverflowCheckDataProcessor as MindsporeOverflowCheckDataProcessor
71
+ )
72
+ from msprobe.mindspore.cell_processor import CellProcessor
73
+ cls.register_processor(Const.MS_FRAMEWORK, Const.STATISTICS, MindsporeStatisticsDataProcessor)
74
+ cls.register_processor(Const.MS_FRAMEWORK, Const.TENSOR, MindsporeTensorDataProcessor)
75
+ cls.register_processor(Const.MS_FRAMEWORK, Const.OVERFLOW_CHECK, MindsporeOverflowCheckDataProcessor)
76
+ cls.register_module_processor(Const.MS_FRAMEWORK, CellProcessor)
@@ -1,198 +1,192 @@
1
- # Copyright 2024 Huawei Technologies Co., Ltd
2
- #
3
- # Licensed under the Apache License, Version 2.0 (the "License");
4
- # you may not use this file except in compliance with the License.
5
- # You may obtain a copy of the License at
6
- #
7
- # http://www.apache.org/licenses/LICENSE-2.0
8
- #
9
- # Unless required by applicable law or agreed to in writing, software
10
- # distributed under the License is distributed on an "AS IS" BASIS,
11
- # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
12
- # See the License for the specific language governing permissions and
13
- # limitations under the License.
14
- # ============================================================================
15
-
16
- import zlib
17
-
18
- import mindspore as ms
19
- from mindspore import ops
20
- import numpy as np
21
-
22
- from msprobe.core.common.const import Const
23
- from msprobe.core.data_dump.data_processor.base import (BaseDataProcessor, TensorStatInfo,
24
- ModuleForwardInputsOutputs, ModuleBackwardInputsOutputs)
25
- from msprobe.core.common.file_check import path_len_exceeds_limit
26
- from msprobe.mindspore.dump.hook_cell.wrap_functional import load_ops_functions
27
- from msprobe.mindspore.common.utils import convert_bf16_to_fp32, save_tensor_as_npy
28
- from msprobe.mindspore.common.log import logger
29
- from msprobe.mindspore.dump.hook_cell.api_registry import api_register
30
-
31
-
32
- class MindsporeDataProcessor(BaseDataProcessor):
33
- mindspore_special_type = tuple([ms.Tensor])
34
- ops_func, mint_ops_func, _ = load_ops_functions()
35
-
36
- def __init__(self, config, data_writer):
37
- super().__init__(config, data_writer)
38
- self.mindspore_object_key = {
39
- "dtype": self.analyze_dtype_in_kwargs
40
- }
41
-
42
- @staticmethod
43
- def get_md5_for_tensor(x):
44
- x = convert_bf16_to_fp32(x)
45
- tensor_bytes = x.asnumpy().tobytes()
46
- crc32_hash = zlib.crc32(tensor_bytes)
47
- return f"{crc32_hash:08x}"
48
-
49
- @staticmethod
50
- def analyze_dtype_in_kwargs(element):
51
- return {"type": "mindspore.dtype", "value": str(element)}
52
-
53
- @staticmethod
54
- def _analyze_builtin(arg):
55
- single_arg = {}
56
- if isinstance(arg, slice):
57
- single_arg.update({"type": "slice"})
58
- # slice参数中可能存在tensor类型,json序列化,需要转换为python数值类型
59
- values = [
60
- value if not isinstance(value, ms.Tensor) else value.item()
61
- for value in [arg.start, arg.stop, arg.step]
62
- ]
63
- single_arg.update({"value": values})
64
- else:
65
- single_arg.update({"type": type(arg).__name__})
66
- single_arg.update({"value": arg})
67
- return single_arg
68
-
69
- @classmethod
70
- def get_special_types(cls):
71
- return super().get_special_types() + cls.mindspore_special_type
72
-
73
- def get_stat_info(self, data):
74
- tensor_stat = TensorStatInfo()
75
- if data.numel() == 0:
76
- return tensor_stat
77
- elif data.dtype == ms.bool_:
78
- data_np = data.asnumpy()
79
- tensor_stat.max = np.max(data_np).item()
80
- tensor_stat.min = np.min(data_np).item()
81
- elif not data.shape:
82
- tensor_stat.max = tensor_stat.min = tensor_stat.mean = tensor_stat.norm = data.item()
83
- elif data.dtype == ms.complex64 or data.dtype == ms.complex128:
84
- data_abs = np.abs(data.asnumpy())
85
- tensor_stat.max = np.max(data_abs).item()
86
- tensor_stat.min = np.min(data_abs).item()
87
- tensor_stat.mean = np.mean(data_abs).item()
88
- tensor_stat.norm = np.linalg.norm(data_abs).item()
89
- else:
90
- if data.dtype == ms.bfloat16 or not ops.is_floating_point(data):
91
- data = data.to(ms.float32)
92
- api_register.norm_inner_op_set_ori_func()
93
- tensor_stat.max = self.mint_ops_func["max"](data).item()
94
- tensor_stat.min = self.mint_ops_func["min"](data).item()
95
- tensor_stat.mean = self.mint_ops_func["mean"](data).item()
96
- tensor_stat.norm = self.ops_func["norm"](data).item()
97
- api_register.norm_inner_op_set_hook_func()
98
- return tensor_stat
99
-
100
- def analyze_single_element(self, element, suffix_stack):
101
- if suffix_stack and suffix_stack[-1] in self.mindspore_object_key:
102
- return self.mindspore_object_key[suffix_stack[-1]](element)
103
-
104
- converted_numpy, numpy_type = self._convert_numpy_to_builtin(element)
105
- if converted_numpy is not element:
106
- return self._analyze_numpy(converted_numpy, numpy_type)
107
- if isinstance(element, ms.Tensor):
108
- return self._analyze_tensor(element, Const.SEP.join(suffix_stack))
109
-
110
- if isinstance(element, (bool, int, float, str, slice)):
111
- return self._analyze_builtin(element)
112
- return {}
113
-
114
- def _analyze_tensor(self, tensor, suffix):
115
- tensor_stat = self.get_stat_info(tensor)
116
- tensor_json = {
117
- 'type': 'mindspore.Tensor',
118
- 'dtype': str(tensor.dtype),
119
- 'shape': tensor.shape,
120
- 'Max': self.transfer_type(tensor_stat.max),
121
- 'Min': self.transfer_type(tensor_stat.min),
122
- 'Mean': self.transfer_type(tensor_stat.mean),
123
- 'Norm': self.transfer_type(tensor_stat.norm),
124
- }
125
- if self.config.summary_mode == Const.MD5:
126
- tensor_md5 = self.get_md5_for_tensor(tensor)
127
- tensor_json.update({Const.MD5: tensor_md5})
128
- return tensor_json
129
-
130
-
131
- class StatisticsDataProcessor(MindsporeDataProcessor):
132
- pass
133
-
134
-
135
- class TensorDataProcessor(MindsporeDataProcessor):
136
- def _analyze_tensor(self, tensor, suffix):
137
- dump_data_name, file_path = self.get_save_file_path(suffix)
138
- single_arg = super()._analyze_tensor(tensor, suffix)
139
- single_arg.update({"data_name": dump_data_name})
140
- save_tensor_as_npy(tensor, file_path)
141
- return single_arg
142
-
143
-
144
- class OverflowCheckDataProcessor(MindsporeDataProcessor):
145
- __slots__ = ["cached_tensors_and_file_paths"]
146
-
147
- def __init__(self, config, data_writer):
148
- super().__init__(config, data_writer)
149
- self.cached_tensors_and_file_paths = {}
150
- self.real_overflow_nums = 0
151
- self.overflow_nums = config.overflow_nums
152
-
153
- @property
154
- def is_terminated(self):
155
- if self.overflow_nums == -1:
156
- return False
157
- if self.real_overflow_nums >= self.overflow_nums:
158
- logger.info(f"[msprobe] 超过预设溢出次数 当前溢出次数: {self.real_overflow_nums}")
159
- return True
160
- return False
161
-
162
- def analyze_forward(self, name, module, module_input_output: ModuleForwardInputsOutputs):
163
- self.has_overflow = False
164
- api_info_struct = super().analyze_forward(name, module, module_input_output)
165
- self.maybe_save_overflow_data()
166
- return api_info_struct if self.has_overflow else None
167
-
168
- def analyze_backward(self, name, module, module_input_output: ModuleBackwardInputsOutputs):
169
- self.has_overflow = False
170
- api_info_struct = super().analyze_backward(name, module, module_input_output)
171
- self.maybe_save_overflow_data()
172
- return api_info_struct if self.has_overflow else None
173
-
174
- def maybe_save_overflow_data(self):
175
- if self.has_overflow:
176
- for file_path, tensor in self.cached_tensors_and_file_paths.items():
177
- save_tensor_as_npy(tensor, file_path)
178
- self.real_overflow_nums += 1
179
- self.cached_tensors_and_file_paths = {}
180
-
181
- def _analyze_maybe_overflow_tensor(self, tensor_json):
182
- if tensor_json['Max'] is None:
183
- return
184
- if np.isinf(tensor_json['Max']) or np.isnan(tensor_json['Max']):
185
- self.has_overflow = True
186
- if np.isinf(tensor_json['Min']) or np.isnan(tensor_json['Min']):
187
- self.has_overflow = True
188
-
189
- def _analyze_tensor(self, tensor, suffix):
190
- dump_data_name, file_path = self.get_save_file_path(suffix)
191
- if not path_len_exceeds_limit(file_path):
192
- self.cached_tensors_and_file_paths.update({file_path: tensor})
193
- else:
194
- logger.warning(f'The file path {file_path} length exceeds limit.')
195
- single_arg = super()._analyze_tensor(tensor, suffix)
196
- self._analyze_maybe_overflow_tensor(single_arg)
197
- single_arg.update({"data_name": dump_data_name})
198
- return single_arg
1
+ # Copyright 2024 Huawei Technologies Co., Ltd
2
+ #
3
+ # Licensed under the Apache License, Version 2.0 (the "License");
4
+ # you may not use this file except in compliance with the License.
5
+ # You may obtain a copy of the License at
6
+ #
7
+ # http://www.apache.org/licenses/LICENSE-2.0
8
+ #
9
+ # Unless required by applicable law or agreed to in writing, software
10
+ # distributed under the License is distributed on an "AS IS" BASIS,
11
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
12
+ # See the License for the specific language governing permissions and
13
+ # limitations under the License.
14
+ # ============================================================================
15
+
16
+ import zlib
17
+
18
+ import mindspore as ms
19
+ from mindspore import mint, ops
20
+ from mindspore._c_expression.typing import Number
21
+ import numpy as np
22
+
23
+ from msprobe.core.common.const import Const
24
+ from msprobe.core.data_dump.data_processor.base import (BaseDataProcessor, TensorStatInfo,
25
+ ModuleForwardInputsOutputs, ModuleBackwardInputsOutputs)
26
+ from msprobe.core.common.file_utils import path_len_exceeds_limit
27
+ from msprobe.mindspore.common.utils import convert_bf16_to_fp32, save_tensor_as_npy
28
+ from msprobe.mindspore.common.log import logger
29
+ from msprobe.mindspore.dump.hook_cell.api_registry import api_register
30
+
31
+
32
+ class MindsporeDataProcessor(BaseDataProcessor):
33
+ mindspore_special_type = tuple([ms.Tensor, Number])
34
+
35
+ def __init__(self, config, data_writer):
36
+ super().__init__(config, data_writer)
37
+ self.mindspore_object_key = {
38
+ "dtype": self.analyze_dtype_in_kwargs
39
+ }
40
+
41
+ @staticmethod
42
+ def get_md5_for_tensor(x):
43
+ x = convert_bf16_to_fp32(x)
44
+ tensor_bytes = x.asnumpy().tobytes()
45
+ crc32_hash = zlib.crc32(tensor_bytes)
46
+ return f"{crc32_hash:08x}"
47
+
48
+ @staticmethod
49
+ def analyze_dtype_in_kwargs(element):
50
+ return {"type": "mindspore.dtype", "value": str(element)}
51
+
52
+ @classmethod
53
+ def get_special_types(cls):
54
+ return super().get_special_types() + cls.mindspore_special_type
55
+
56
+ def get_stat_info(self, data):
57
+ tensor_stat = TensorStatInfo()
58
+ if data.numel() == 0:
59
+ return tensor_stat
60
+ elif data.dtype == ms.bool_:
61
+ data_np = data.asnumpy()
62
+ tensor_stat.max = np.max(data_np).item()
63
+ tensor_stat.min = np.min(data_np).item()
64
+ elif not data.shape:
65
+ tensor_stat.max = tensor_stat.min = tensor_stat.mean = tensor_stat.norm = data.item()
66
+ elif data.dtype == ms.complex64 or data.dtype == ms.complex128:
67
+ data_abs = np.abs(data.asnumpy())
68
+ tensor_stat.max = np.max(data_abs).item()
69
+ tensor_stat.min = np.min(data_abs).item()
70
+ tensor_stat.mean = np.mean(data_abs).item()
71
+ tensor_stat.norm = np.linalg.norm(data_abs).item()
72
+ else:
73
+ if not ops.is_floating_point(data):
74
+ data = data.to(ms.float32)
75
+ api_register.norm_inner_op_set_ori_func()
76
+ get_max_value = api_register.mint_ops_ori_attr.get("max", mint.max)
77
+ get_min_value = api_register.mint_ops_ori_attr.get("min", mint.min)
78
+ get_mean_value = api_register.mint_ops_ori_attr.get("mean", mint.mean)
79
+ if hasattr(mint, "norm"):
80
+ get_norm_value = api_register.mint_ops_ori_attr.get("norm", mint.norm)
81
+ else:
82
+ get_norm_value = api_register.functional_ori_attr.get("norm", ops.norm)
83
+ tensor_stat.max = get_max_value(data).item()
84
+ tensor_stat.min = get_min_value(data).item()
85
+ tensor_stat.mean = get_mean_value(data).item()
86
+ tensor_stat.norm = get_norm_value(data).item()
87
+ api_register.norm_inner_op_set_hook_func()
88
+ return tensor_stat
89
+
90
+ def analyze_single_element(self, element, suffix_stack):
91
+ if suffix_stack and suffix_stack[-1] in self.mindspore_object_key:
92
+ return self.mindspore_object_key[suffix_stack[-1]](element)
93
+
94
+ converted_numpy, numpy_type = self._convert_numpy_to_builtin(element)
95
+ if converted_numpy is not element:
96
+ return self._analyze_numpy(converted_numpy, numpy_type)
97
+ if isinstance(element, Number):
98
+ return self.analyze_dtype_in_kwargs(element)
99
+ if isinstance(element, ms.Tensor):
100
+ return self._analyze_tensor(element, Const.SEP.join(suffix_stack))
101
+ if isinstance(element, (bool, int, float, str, slice, type(Ellipsis))):
102
+ return self._analyze_builtin(element)
103
+ return {}
104
+
105
+ def _analyze_tensor(self, tensor, suffix):
106
+ tensor_stat = self.get_stat_info(tensor)
107
+ tensor_json = {
108
+ 'type': 'mindspore.Tensor',
109
+ 'dtype': str(tensor.dtype),
110
+ 'shape': tensor.shape,
111
+ 'Max': self.transfer_type(tensor_stat.max),
112
+ 'Min': self.transfer_type(tensor_stat.min),
113
+ 'Mean': self.transfer_type(tensor_stat.mean),
114
+ 'Norm': self.transfer_type(tensor_stat.norm),
115
+ }
116
+ if self.config.summary_mode == Const.MD5:
117
+ tensor_md5 = self.get_md5_for_tensor(tensor)
118
+ tensor_json.update({Const.MD5: tensor_md5})
119
+ return tensor_json
120
+
121
+
122
+ class StatisticsDataProcessor(MindsporeDataProcessor):
123
+ pass
124
+
125
+
126
+ class TensorDataProcessor(MindsporeDataProcessor):
127
+ def _analyze_tensor(self, tensor, suffix):
128
+ dump_data_name, file_path = self.get_save_file_path(suffix)
129
+ single_arg = super()._analyze_tensor(tensor, suffix)
130
+ single_arg.update({"data_name": dump_data_name})
131
+ save_tensor_as_npy(tensor, file_path)
132
+ return single_arg
133
+
134
+
135
+ class OverflowCheckDataProcessor(MindsporeDataProcessor):
136
+ __slots__ = ["cached_tensors_and_file_paths"]
137
+
138
+ def __init__(self, config, data_writer):
139
+ super().__init__(config, data_writer)
140
+ self.has_overflow = False
141
+ self.cached_tensors_and_file_paths = {}
142
+ self.real_overflow_nums = 0
143
+ self.overflow_nums = config.overflow_nums
144
+
145
+ @property
146
+ def is_terminated(self):
147
+ if self.overflow_nums == -1:
148
+ return False
149
+ if self.real_overflow_nums >= self.overflow_nums:
150
+ return True
151
+ return False
152
+
153
+ def analyze_forward(self, name, module, module_input_output: ModuleForwardInputsOutputs):
154
+ self.has_overflow = False
155
+ api_info_struct = super().analyze_forward(name, module, module_input_output)
156
+ self.maybe_save_overflow_data()
157
+ return api_info_struct if self.has_overflow else None
158
+
159
+ def analyze_backward(self, name, module, module_input_output: ModuleBackwardInputsOutputs):
160
+ self.has_overflow = False
161
+ api_info_struct = super().analyze_backward(name, module, module_input_output)
162
+ self.maybe_save_overflow_data()
163
+ return api_info_struct if self.has_overflow else None
164
+
165
+ def maybe_save_overflow_data(self):
166
+ if self.has_overflow:
167
+ for file_path, tensor in self.cached_tensors_and_file_paths.items():
168
+ save_tensor_as_npy(tensor, file_path)
169
+ self.real_overflow_nums += 1
170
+ if self.overflow_nums != -1 and self.real_overflow_nums >= self.overflow_nums:
171
+ logger.info(f"[{Const.TOOL_NAME}] Reached the preset overflow times, "
172
+ f"current overflow times: {self.real_overflow_nums}.")
173
+ self.cached_tensors_and_file_paths = {}
174
+
175
+ def _analyze_maybe_overflow_tensor(self, tensor_json):
176
+ if tensor_json['Max'] is None:
177
+ return
178
+ if np.isinf(tensor_json['Max']) or np.isnan(tensor_json['Max']):
179
+ self.has_overflow = True
180
+ if np.isinf(tensor_json['Min']) or np.isnan(tensor_json['Min']):
181
+ self.has_overflow = True
182
+
183
+ def _analyze_tensor(self, tensor, suffix):
184
+ dump_data_name, file_path = self.get_save_file_path(suffix)
185
+ if not path_len_exceeds_limit(file_path):
186
+ self.cached_tensors_and_file_paths.update({file_path: tensor})
187
+ else:
188
+ logger.warning(f'The file path {file_path} length exceeds limit.')
189
+ single_arg = super()._analyze_tensor(tensor, suffix)
190
+ self._analyze_maybe_overflow_tensor(single_arg)
191
+ single_arg.update({"data_name": dump_data_name})
192
+ return single_arg