mindstudio-probe 1.0.3__py3-none-any.whl → 1.1.0__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (278) hide show
  1. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/LICENSE +201 -201
  2. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/METADATA +36 -34
  3. mindstudio_probe-1.1.0.dist-info/RECORD +287 -0
  4. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/WHEEL +1 -1
  5. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/entry_points.txt +1 -0
  6. msprobe/README.md +131 -237
  7. msprobe/__init__.py +16 -1
  8. msprobe/{config/config.json → config.json} +47 -49
  9. msprobe/core/advisor/advisor.py +124 -124
  10. msprobe/core/advisor/advisor_const.py +58 -59
  11. msprobe/core/advisor/advisor_result.py +58 -58
  12. msprobe/core/common/const.py +402 -318
  13. msprobe/core/common/exceptions.py +99 -99
  14. msprobe/core/common/{file_check.py → file_utils.py} +523 -283
  15. msprobe/core/common/inplace_op_checker.py +38 -0
  16. msprobe/core/common/inplace_ops.yaml +251 -0
  17. msprobe/core/common/log.py +86 -69
  18. msprobe/core/common/utils.py +371 -616
  19. msprobe/core/common_config.py +78 -71
  20. msprobe/core/compare/acc_compare.py +472 -298
  21. msprobe/core/compare/check.py +180 -95
  22. msprobe/core/compare/compare_cli.py +69 -49
  23. msprobe/core/compare/highlight.py +259 -222
  24. msprobe/core/compare/multiprocessing_compute.py +174 -149
  25. msprobe/core/compare/npy_compare.py +310 -295
  26. msprobe/core/compare/utils.py +464 -429
  27. msprobe/core/data_dump/data_collector.py +153 -144
  28. msprobe/core/data_dump/data_processor/base.py +337 -293
  29. msprobe/core/data_dump/data_processor/factory.py +76 -59
  30. msprobe/core/data_dump/data_processor/mindspore_processor.py +192 -198
  31. msprobe/core/data_dump/data_processor/pytorch_processor.py +383 -389
  32. msprobe/core/data_dump/json_writer.py +117 -116
  33. msprobe/core/data_dump/scope.py +194 -178
  34. msprobe/core/grad_probe/constant.py +74 -70
  35. msprobe/core/grad_probe/grad_compare.py +170 -175
  36. msprobe/core/grad_probe/utils.py +77 -52
  37. msprobe/docs/01.installation.md +99 -0
  38. msprobe/docs/02.config_introduction.md +137 -0
  39. msprobe/docs/03.config_examples.md +237 -0
  40. msprobe/docs/04.acl_config_examples.md +78 -0
  41. msprobe/docs/05.data_dump_PyTorch.md +326 -0
  42. msprobe/docs/06.data_dump_MindSpore.md +285 -0
  43. msprobe/docs/07.accuracy_checker_PyTorch.md +297 -0
  44. msprobe/docs/08.accuracy_checker_online_PyTorch.md +238 -0
  45. msprobe/docs/09.accuracy_checker_MindSpore.md +68 -0
  46. msprobe/docs/10.accuracy_compare_PyTorch.md +327 -0
  47. msprobe/docs/11.accuracy_compare_MindSpore.md +333 -0
  48. msprobe/docs/12.overflow_check_PyTorch.md +79 -0
  49. msprobe/docs/13.overflow_check_MindSpore.md +31 -0
  50. msprobe/{pytorch/doc/parse_tool.md → docs/14.data_parse_PyTorch.md} +283 -286
  51. msprobe/docs/15.free_benchmarking_PyTorch.md +170 -0
  52. msprobe/docs/16.free_benchmarking_MindSpore.md +140 -0
  53. msprobe/{doc/grad_probe/grad_probe.md → docs/17.grad_probe.md} +205 -207
  54. msprobe/{pytorch/doc//321/205/320/254/320/270/321/207/342/225/221/342/224/220/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/206/320/277/320/244/321/205/320/277/342/225/243.md → docs/18.online_dispatch.md} +89 -90
  55. msprobe/docs/FAQ.md +189 -0
  56. msprobe/docs/S02.report_free_benchmarking_validation_performance_baseline.md +146 -0
  57. msprobe/docs/img/free_benchmark_framework.png +0 -0
  58. msprobe/docs/img/ms_dump.png +0 -0
  59. msprobe/docs/img/ms_layer.png +0 -0
  60. msprobe/docs/img/pt_dump.png +0 -0
  61. msprobe/mindspore/__init__.py +2 -1
  62. msprobe/mindspore/api_accuracy_checker/api_accuracy_checker.py +278 -245
  63. msprobe/mindspore/api_accuracy_checker/api_info.py +76 -69
  64. msprobe/mindspore/api_accuracy_checker/api_runner.py +155 -151
  65. msprobe/mindspore/api_accuracy_checker/base_compare_algorithm.py +196 -196
  66. msprobe/mindspore/api_accuracy_checker/cmd_parser.py +6 -0
  67. msprobe/mindspore/api_accuracy_checker/compute_element.py +238 -223
  68. msprobe/mindspore/api_accuracy_checker/main.py +8 -15
  69. msprobe/mindspore/api_accuracy_checker/type_mapping.py +113 -113
  70. msprobe/mindspore/api_accuracy_checker/utils.py +79 -62
  71. msprobe/mindspore/cell_processor.py +58 -34
  72. msprobe/mindspore/common/const.py +108 -87
  73. msprobe/mindspore/common/log.py +37 -37
  74. msprobe/mindspore/common/utils.py +97 -57
  75. msprobe/mindspore/compare/distributed_compare.py +62 -75
  76. msprobe/mindspore/compare/layer_mapping.py +146 -0
  77. msprobe/mindspore/compare/modify_mapping.py +107 -0
  78. msprobe/mindspore/compare/ms_compare.py +357 -117
  79. msprobe/mindspore/compare/ms_graph_compare.py +364 -317
  80. msprobe/mindspore/compare/ms_to_pt_api.yaml +399 -399
  81. msprobe/mindspore/debugger/debugger_config.py +69 -74
  82. msprobe/mindspore/debugger/precision_debugger.py +150 -107
  83. msprobe/mindspore/dump/dump_tool_factory.py +50 -35
  84. msprobe/mindspore/dump/hook_cell/api_registry.py +128 -104
  85. msprobe/mindspore/dump/hook_cell/hook_cell.py +55 -53
  86. msprobe/mindspore/dump/hook_cell/primitive_hooks.py +206 -0
  87. msprobe/mindspore/dump/hook_cell/support_wrap_ops.yaml +994 -925
  88. msprobe/mindspore/dump/hook_cell/wrap_api.py +121 -0
  89. msprobe/mindspore/dump/jit_dump.py +96 -56
  90. msprobe/mindspore/dump/kernel_graph_dump.py +75 -60
  91. msprobe/mindspore/dump/kernel_kbyk_dump.py +79 -65
  92. msprobe/mindspore/free_benchmark/api_pynative_self_check.py +131 -116
  93. msprobe/mindspore/free_benchmark/common/config.py +27 -12
  94. msprobe/mindspore/free_benchmark/common/handler_params.py +32 -17
  95. msprobe/mindspore/free_benchmark/common/utils.py +85 -71
  96. msprobe/mindspore/free_benchmark/data/support_wrap_ops.yaml +842 -842
  97. msprobe/mindspore/free_benchmark/decorator/dec_forward.py +57 -42
  98. msprobe/mindspore/free_benchmark/decorator/decorator_factory.py +122 -107
  99. msprobe/mindspore/free_benchmark/handler/base_handler.py +105 -90
  100. msprobe/mindspore/free_benchmark/handler/check_handler.py +56 -41
  101. msprobe/mindspore/free_benchmark/handler/fix_handler.py +51 -36
  102. msprobe/mindspore/free_benchmark/handler/handler_factory.py +36 -21
  103. msprobe/mindspore/free_benchmark/perturbation/add_noise.py +82 -67
  104. msprobe/mindspore/free_benchmark/perturbation/base_perturbation.py +36 -21
  105. msprobe/mindspore/free_benchmark/perturbation/bit_noise.py +78 -63
  106. msprobe/mindspore/free_benchmark/perturbation/exchange_value.py +77 -0
  107. msprobe/mindspore/free_benchmark/perturbation/improve_precision.py +49 -34
  108. msprobe/mindspore/free_benchmark/perturbation/no_change.py +27 -12
  109. msprobe/mindspore/free_benchmark/perturbation/perturbation_factory.py +44 -27
  110. msprobe/mindspore/free_benchmark/self_check_tool_factory.py +48 -33
  111. msprobe/mindspore/grad_probe/global_context.py +100 -91
  112. msprobe/mindspore/grad_probe/grad_analyzer.py +231 -231
  113. msprobe/mindspore/grad_probe/grad_monitor.py +27 -27
  114. msprobe/mindspore/grad_probe/grad_stat_csv.py +131 -131
  115. msprobe/mindspore/grad_probe/hook.py +94 -92
  116. msprobe/mindspore/grad_probe/utils.py +29 -28
  117. msprobe/mindspore/ms_config.py +128 -126
  118. msprobe/mindspore/overflow_check/kernel_graph_overflow_check.py +60 -45
  119. msprobe/mindspore/overflow_check/overflow_check_tool_factory.py +49 -34
  120. msprobe/mindspore/runtime.py +4 -4
  121. msprobe/mindspore/service.py +297 -354
  122. msprobe/mindspore/task_handler_factory.py +24 -24
  123. msprobe/msprobe.py +105 -107
  124. msprobe/pytorch/__init__.py +23 -4
  125. msprobe/pytorch/api_accuracy_checker/common/config.py +70 -55
  126. msprobe/pytorch/api_accuracy_checker/common/utils.py +246 -165
  127. msprobe/pytorch/api_accuracy_checker/compare/algorithm.py +230 -213
  128. msprobe/pytorch/api_accuracy_checker/compare/api_precision_compare.py +632 -581
  129. msprobe/pytorch/api_accuracy_checker/compare/api_precision_standard.yaml +132 -132
  130. msprobe/pytorch/api_accuracy_checker/compare/api_precision_threshold.yaml +390 -390
  131. msprobe/pytorch/api_accuracy_checker/compare/compare.py +416 -381
  132. msprobe/pytorch/api_accuracy_checker/compare/compare_column.py +90 -73
  133. msprobe/pytorch/api_accuracy_checker/compare/compare_utils.py +265 -244
  134. msprobe/pytorch/api_accuracy_checker/config.yaml +10 -10
  135. msprobe/pytorch/api_accuracy_checker/run_ut/data_generate.py +370 -332
  136. msprobe/pytorch/api_accuracy_checker/run_ut/multi_run_ut.py +221 -199
  137. msprobe/pytorch/api_accuracy_checker/run_ut/run_overflow_check.py +150 -134
  138. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut.py +518 -581
  139. msprobe/pytorch/api_accuracy_checker/run_ut/run_ut_utils.py +213 -74
  140. msprobe/pytorch/api_accuracy_checker/run_ut/torch_ut_setting.json +7 -4
  141. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/attl.py +218 -202
  142. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/client.py +370 -324
  143. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/device_dispatch.py +227 -204
  144. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/dump_dispatch.py +110 -0
  145. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/server.py +244 -218
  146. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/torch_ops_config.yaml +63 -0
  147. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/utils.py +44 -0
  148. msprobe/pytorch/bench_functions/__init__.py +30 -15
  149. msprobe/pytorch/bench_functions/apply_adam_w.py +43 -28
  150. msprobe/pytorch/bench_functions/confusion_transpose.py +34 -19
  151. msprobe/pytorch/bench_functions/fast_gelu.py +70 -55
  152. msprobe/pytorch/bench_functions/layer_norm_eval.py +21 -6
  153. msprobe/pytorch/bench_functions/linear.py +27 -12
  154. msprobe/pytorch/bench_functions/matmul_backward.py +63 -48
  155. msprobe/pytorch/bench_functions/npu_fusion_attention.py +538 -421
  156. msprobe/pytorch/bench_functions/rms_norm.py +30 -15
  157. msprobe/pytorch/bench_functions/rotary_mul.py +71 -52
  158. msprobe/pytorch/bench_functions/scaled_mask_softmax.py +41 -26
  159. msprobe/pytorch/bench_functions/swiglu.py +70 -55
  160. msprobe/pytorch/common/__init__.py +17 -2
  161. msprobe/pytorch/common/compare_script.template +14 -14
  162. msprobe/pytorch/common/log.py +33 -32
  163. msprobe/pytorch/common/parse_json.py +54 -39
  164. msprobe/pytorch/common/utils.py +310 -300
  165. msprobe/pytorch/compare/distributed_compare.py +66 -66
  166. msprobe/pytorch/compare/mapping.yaml +607 -607
  167. msprobe/pytorch/compare/match.py +49 -33
  168. msprobe/pytorch/compare/pt_compare.py +82 -40
  169. msprobe/pytorch/debugger/debugger_config.py +108 -95
  170. msprobe/pytorch/debugger/precision_debugger.py +173 -125
  171. msprobe/pytorch/free_benchmark/__init__.py +23 -8
  172. msprobe/pytorch/free_benchmark/common/constant.py +70 -70
  173. msprobe/pytorch/free_benchmark/common/counter.py +71 -71
  174. msprobe/pytorch/free_benchmark/common/enums.py +65 -37
  175. msprobe/pytorch/free_benchmark/common/params.py +144 -129
  176. msprobe/pytorch/free_benchmark/common/utils.py +118 -102
  177. msprobe/pytorch/free_benchmark/compare/grad_saver.py +200 -179
  178. msprobe/pytorch/free_benchmark/compare/single_benchmark.py +119 -104
  179. msprobe/pytorch/free_benchmark/main.py +120 -105
  180. msprobe/pytorch/free_benchmark/perturbed_layers/base_layer.py +28 -13
  181. msprobe/pytorch/free_benchmark/perturbed_layers/layer_factory.py +56 -41
  182. msprobe/pytorch/free_benchmark/perturbed_layers/npu/add_noise.py +105 -90
  183. msprobe/pytorch/free_benchmark/perturbed_layers/npu/bit_noise.py +119 -104
  184. msprobe/pytorch/free_benchmark/perturbed_layers/npu/change_value.py +87 -63
  185. msprobe/pytorch/free_benchmark/perturbed_layers/npu/improve_precision.py +83 -68
  186. msprobe/pytorch/free_benchmark/perturbed_layers/npu/no_change.py +43 -28
  187. msprobe/pytorch/free_benchmark/perturbed_layers/npu/npu_base_layser.py +60 -45
  188. msprobe/pytorch/free_benchmark/perturbed_layers/run_cpu.py +34 -19
  189. msprobe/pytorch/free_benchmark/result_handlers/base_handler.py +256 -217
  190. msprobe/pytorch/free_benchmark/result_handlers/check_handler.py +54 -39
  191. msprobe/pytorch/free_benchmark/result_handlers/fix_handler.py +38 -23
  192. msprobe/pytorch/free_benchmark/result_handlers/handler_factory.py +45 -30
  193. msprobe/pytorch/free_benchmark/result_handlers/preheat_handler.py +185 -170
  194. msprobe/pytorch/function_factory.py +91 -75
  195. msprobe/pytorch/functional/module_dump.py +84 -0
  196. msprobe/pytorch/grad_probe/grad_monitor.py +91 -90
  197. msprobe/pytorch/grad_probe/grad_stat_csv.py +128 -128
  198. msprobe/pytorch/hook_module/__init__.py +16 -1
  199. msprobe/pytorch/hook_module/api_registry.py +166 -161
  200. msprobe/pytorch/hook_module/hook_module.py +118 -120
  201. msprobe/pytorch/hook_module/support_wrap_ops.yaml +1879 -1877
  202. msprobe/pytorch/hook_module/utils.py +28 -29
  203. msprobe/pytorch/hook_module/wrap_aten.py +111 -110
  204. msprobe/pytorch/hook_module/wrap_distributed.py +77 -78
  205. msprobe/pytorch/hook_module/wrap_functional.py +104 -105
  206. msprobe/pytorch/hook_module/wrap_npu_custom.py +85 -84
  207. msprobe/pytorch/hook_module/wrap_tensor.py +69 -71
  208. msprobe/pytorch/hook_module/wrap_torch.py +84 -86
  209. msprobe/pytorch/hook_module/wrap_vf.py +60 -62
  210. msprobe/pytorch/module_processer.py +153 -138
  211. msprobe/pytorch/online_dispatch/__init__.py +20 -20
  212. msprobe/pytorch/online_dispatch/compare.py +235 -236
  213. msprobe/pytorch/online_dispatch/dispatch.py +271 -271
  214. msprobe/pytorch/online_dispatch/dump_compare.py +155 -156
  215. msprobe/pytorch/online_dispatch/single_compare.py +391 -391
  216. msprobe/pytorch/online_dispatch/torch_ops_config.yaml +57 -49
  217. msprobe/pytorch/online_dispatch/utils.py +127 -146
  218. msprobe/pytorch/parse.py +19 -4
  219. msprobe/pytorch/parse_tool/cli.py +31 -32
  220. msprobe/pytorch/parse_tool/lib/compare.py +259 -271
  221. msprobe/pytorch/parse_tool/lib/config.py +52 -52
  222. msprobe/pytorch/parse_tool/lib/file_desc.py +31 -31
  223. msprobe/pytorch/parse_tool/lib/interactive_cli.py +102 -102
  224. msprobe/pytorch/parse_tool/lib/parse_exception.py +54 -54
  225. msprobe/pytorch/parse_tool/lib/parse_tool.py +161 -158
  226. msprobe/pytorch/parse_tool/lib/utils.py +320 -321
  227. msprobe/pytorch/parse_tool/lib/visualization.py +85 -91
  228. msprobe/pytorch/pt_config.py +317 -187
  229. msprobe/pytorch/service.py +311 -252
  230. mindstudio_probe-1.0.3.dist-info/RECORD +0 -272
  231. msprobe/config/README.md +0 -539
  232. msprobe/mindspore/doc/compare.md +0 -58
  233. msprobe/mindspore/doc/dump.md +0 -217
  234. msprobe/mindspore/dump/hook_cell/wrap_functional.py +0 -91
  235. msprobe/mindspore/dump/hook_cell/wrap_tensor.py +0 -63
  236. msprobe/pytorch/api_accuracy_checker/tensor_transport_layer/ssl_config.py +0 -10
  237. msprobe/pytorch/doc/FAQ.md +0 -193
  238. msprobe/pytorch/doc/api_accuracy_checker.md +0 -313
  239. msprobe/pytorch/doc/api_accuracy_checker_online.md +0 -187
  240. msprobe/pytorch/doc/dump.md +0 -260
  241. msprobe/pytorch/doc/msprobe/321/207/342/226/223/342/225/233/321/205/342/225/221/320/266/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/206/320/245/342/226/221/321/206/320/235/320/276dump/321/206/320/260/320/227/321/205/320/227/320/226/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -182
  242. msprobe/pytorch/doc/ptdbg_ascend_compare.md +0 -240
  243. msprobe/pytorch/doc/ptdbg_ascend_overview.md +0 -68
  244. msprobe/pytorch/doc/ptdbg_ascend_quickstart.md +0 -381
  245. msprobe/pytorch/doc/run_overflow_check.md +0 -25
  246. msprobe/pytorch/doc//321/206/320/247/320/260/321/206/320/260/320/227/321/206/320/255/320/226/321/205/342/225/226/320/265/321/205/320/225/342/225/226/321/205/320/254/342/225/221/321/206/320/251/320/277/321/211/320/272/320/234/321/210/320/277/320/221/321/205/320/242/320/234/321/206/320/220/320/267/321/210/320/223/342/225/234/321/205/320/257/342/225/221/321/207/342/225/221/342/224/220/321/206/320/232/320/265/321/205/320/241/320/232.md +0 -151
  247. msprobe/pytorch/functional/data_processor.py +0 -0
  248. msprobe/pytorch/functional/dump_module.py +0 -39
  249. {mindstudio_probe-1.0.3.dist-info → mindstudio_probe-1.1.0.dist-info}/top_level.txt +0 -0
  250. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_1.png +0 -0
  251. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_2.png +0 -0
  252. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_3.png +0 -0
  253. /msprobe/{pytorch/doc → docs}/img/BLOOM-7B_4.png +0 -0
  254. /msprobe/{pytorch/doc → docs}/img/GPT-3_1.png +0 -0
  255. /msprobe/{pytorch/doc → docs}/img/GPT-3_2.png +0 -0
  256. /msprobe/{pytorch/doc → docs}/img/GPT-3_3.png +0 -0
  257. /msprobe/{pytorch/doc → docs}/img/GPT-3_4.png +0 -0
  258. /msprobe/{pytorch/doc → docs}/img/GPT-3_5.png +0 -0
  259. /msprobe/{pytorch/doc → docs}/img/GPT-3_6.png +0 -0
  260. /msprobe/{pytorch/doc → docs}/img/GPT-3_7.png +0 -0
  261. /msprobe/{pytorch/doc → docs}/img/GPT-3_8.png +0 -0
  262. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_1.png +0 -0
  263. /msprobe/{pytorch/doc → docs}/img/YOLOV5S_2.png +0 -0
  264. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_details.png +0 -0
  265. /msprobe/{pytorch/doc → docs}/img/accuracy_checking_result.png +0 -0
  266. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_details.png +0 -0
  267. /msprobe/{pytorch/doc → docs}/img/api_precision_compare_result.png +0 -0
  268. /msprobe/{pytorch/doc → docs}/img/auto_analyze_log.png +0 -0
  269. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl.png +0 -0
  270. /msprobe/{pytorch/doc → docs}/img/compare_result_pkl_md5.png.png +0 -0
  271. /msprobe/{pytorch/doc → docs}/img/cpu_info.png +0 -0
  272. /msprobe/{config → docs}/img/free_benchmark.png +0 -0
  273. /msprobe/{doc/grad_probe/img/image-1.png → docs/img/grad_probe_image-1.png} +0 -0
  274. /msprobe/{doc/grad_probe/img/image-2.png → docs/img/grad_probe_image-2.png} +0 -0
  275. /msprobe/{doc/grad_probe/img/image-3.png → docs/img/grad_probe_image-3.png} +0 -0
  276. /msprobe/{doc/grad_probe/img/image-4.png → docs/img/grad_probe_image-4.png} +0 -0
  277. /msprobe/{doc/grad_probe/img/image.png → docs/img/grad_probe_image.png} +0 -0
  278. /msprobe/{pytorch/doc → docs}/img/module_compare.png +0 -0
@@ -1,105 +1,120 @@
1
- from abc import ABC
2
-
3
- import torch
4
- from msprobe.core.common.const import Const
5
- from msprobe.pytorch.free_benchmark import logger
6
- from msprobe.pytorch.free_benchmark.common.constant import CommonField
7
- from msprobe.pytorch.free_benchmark.common.enums import (
8
- DeviceType,
9
- FuzzLevel,
10
- HandlerType,
11
- PerturbationMode,
12
- )
13
- from msprobe.pytorch.free_benchmark.common.params import (
14
- data_pre_deal,
15
- make_handler_params,
16
- )
17
- from msprobe.pytorch.free_benchmark.compare.grad_saver import GradSaver
18
- from msprobe.pytorch.free_benchmark.perturbed_layers.layer_factory import LayerFactory
19
- from msprobe.pytorch.free_benchmark.result_handlers.handler_factory import (
20
- FuzzHandlerFactory,
21
- )
22
-
23
-
24
- class FreeBenchmarkCheck(ABC):
25
-
26
- def __init__(self, config) -> None:
27
- super().__init__()
28
- self.config = config
29
- if self.config.pert_mode is None:
30
- self.config.pert_mode = PerturbationMode.IMPROVE_PRECISION
31
- if self.config.fuzz_level is None:
32
- self.config.fuzz_level = FuzzLevel.BASE_LEVEL
33
- if self.config.fuzz_device is None:
34
- self.config.fuzz_device = DeviceType.NPU
35
- self.current_iter = 0
36
-
37
- def update_iter(self, update_iter):
38
- self.current_iter = update_iter
39
-
40
- def if_fix(self):
41
- if self.config.handler_type==HandlerType.FIX:
42
- return True
43
- return False
44
-
45
- def pre_forward(self, name, module, data_processor, args, kwargs):
46
- if not self.config.fuzz_stage == Const.BACKWARD:
47
- return
48
- origin_func = (
49
- module._slow_forward if torch._C._get_tracing_state() else module.forward
50
- )
51
- handler_params = make_handler_params(name, self.config, self.current_iter)
52
- grad_saver = GradSaver(origin_func, handler_params)
53
- grad_saver.kwargs = kwargs
54
- grad_saver.register_compare_func_for_inputs(args, data_processor)
55
- grad_saver.cache_backward_input(args)
56
- setattr(module, CommonField.GRADSAVER, grad_saver)
57
-
58
- def forward(self, name, module, args, kwargs, output):
59
- if not self.config.fuzz_stage == Const.FORWARD:
60
- return output, []
61
- origin_func = (
62
- module._slow_forward if torch._C._get_tracing_state() else module.forward
63
- )
64
- data_params = data_pre_deal(name, origin_func, args, kwargs)
65
- if data_params.valid_input_index == -1:
66
- return output, []
67
- data_params.original_result = output
68
- data_params.fuzz_stage = self.config.fuzz_stage
69
-
70
- layer = LayerFactory.create(
71
- name, self.config.fuzz_device, self.config.pert_mode
72
- )
73
- layer.handle(data_params)
74
- handler_params = make_handler_params(name, self.config, self.current_iter)
75
- handler = FuzzHandlerFactory.create(handler_params)
76
- perturbed_output = handler.handle(data_params)
77
- return perturbed_output, handler.get_unequal_rows()
78
-
79
- def backward(self, name, module, grad_output):
80
-
81
- if not self.config.fuzz_stage == Const.BACKWARD:
82
- return
83
- try:
84
- grad_saver = getattr(module, CommonField.GRADSAVER)
85
- except AttributeError:
86
- logger.warning_on_rank_0(
87
- f"[msprobe] Free benchmark: get grad saver failed. api_name:{name}"
88
- )
89
- return
90
-
91
- _new_grad_output = grad_output
92
- try:
93
- need_grad_tensors, _inner_args = grad_saver.get_vjp_input()
94
- origin_grad_input = grad_saver.get_grad_input_from_vjp(
95
- tuple(need_grad_tensors), _new_grad_output, _inner_args
96
- )
97
- grad_saver.origin_grad_input = tuple([x.cpu() for x in origin_grad_input])
98
- grad_saver.calculate_perturbed_grad_input(
99
- _new_grad_output, need_grad_tensors, _inner_args
100
- )
101
- except Exception as e:
102
- logger.warning_on_rank_0(
103
- f"[msprobe] Free benchmark: grad vjp calculate failed. api_name:{name} error: {e}"
104
- )
105
- return
1
+ # Copyright (c) 2024-2024, Huawei Technologies Co., Ltd.
2
+ # All rights reserved.
3
+ #
4
+ # Licensed under the Apache License, Version 2.0 (the "License");
5
+ # you may not use this file except in compliance with the License.
6
+ # You may obtain a copy of the License at
7
+ #
8
+ # http://www.apache.org/licenses/LICENSE-2.0
9
+ #
10
+ # Unless required by applicable law or agreed to in writing, software
11
+ # distributed under the License is distributed on an "AS IS" BASIS,
12
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
13
+ # See the License for the specific language governing permissions and
14
+ # limitations under the License.
15
+
16
+ from abc import ABC
17
+
18
+ import torch
19
+ from msprobe.core.common.const import Const
20
+ from msprobe.pytorch.free_benchmark import logger
21
+ from msprobe.pytorch.free_benchmark.common.constant import CommonField
22
+ from msprobe.pytorch.free_benchmark.common.enums import (
23
+ DeviceType,
24
+ FuzzLevel,
25
+ HandlerType,
26
+ PerturbationMode,
27
+ )
28
+ from msprobe.pytorch.free_benchmark.common.params import (
29
+ data_pre_deal,
30
+ make_handler_params,
31
+ )
32
+ from msprobe.pytorch.free_benchmark.compare.grad_saver import GradSaver
33
+ from msprobe.pytorch.free_benchmark.perturbed_layers.layer_factory import LayerFactory
34
+ from msprobe.pytorch.free_benchmark.result_handlers.handler_factory import (
35
+ FuzzHandlerFactory,
36
+ )
37
+
38
+
39
+ class FreeBenchmarkCheck(ABC):
40
+
41
+ def __init__(self, config) -> None:
42
+ super().__init__()
43
+ self.config = config
44
+ if self.config.pert_mode is None:
45
+ self.config.pert_mode = PerturbationMode.IMPROVE_PRECISION
46
+ if self.config.fuzz_level is None:
47
+ self.config.fuzz_level = FuzzLevel.BASE_LEVEL
48
+ if self.config.fuzz_device is None:
49
+ self.config.fuzz_device = DeviceType.NPU
50
+ self.current_iter = 0
51
+
52
+ def update_iter(self, update_iter):
53
+ self.current_iter = update_iter
54
+
55
+ def if_fix(self):
56
+ if self.config.handler_type == HandlerType.FIX:
57
+ return True
58
+ return False
59
+
60
+ def pre_forward(self, name, module, data_processor, args, kwargs):
61
+ if not self.config.fuzz_stage == Const.BACKWARD:
62
+ return
63
+ origin_func = (
64
+ module._slow_forward if torch._C._get_tracing_state() else module.forward
65
+ )
66
+ handler_params = make_handler_params(name, self.config, self.current_iter)
67
+ grad_saver = GradSaver(origin_func, handler_params)
68
+ grad_saver.kwargs = kwargs
69
+ grad_saver.register_compare_func_for_inputs(args, data_processor)
70
+ grad_saver.cache_backward_input(args)
71
+ setattr(module, CommonField.GRADSAVER, grad_saver)
72
+
73
+ def forward(self, name, module, args, kwargs, output):
74
+ if not self.config.fuzz_stage == Const.FORWARD:
75
+ return output, []
76
+ origin_func = (
77
+ module._slow_forward if torch._C._get_tracing_state() else module.forward
78
+ )
79
+ data_params = data_pre_deal(name, origin_func, args, kwargs)
80
+ if data_params.valid_input_index == -1:
81
+ return output, []
82
+ data_params.original_result = output
83
+ data_params.fuzz_stage = self.config.fuzz_stage
84
+
85
+ layer = LayerFactory.create(
86
+ name, self.config.fuzz_device, self.config.pert_mode
87
+ )
88
+ layer.handle(data_params)
89
+ handler_params = make_handler_params(name, self.config, self.current_iter)
90
+ handler = FuzzHandlerFactory.create(handler_params)
91
+ perturbed_output = handler.handle(data_params)
92
+ return perturbed_output, handler.get_unequal_rows()
93
+
94
+ def backward(self, name, module, grad_output):
95
+
96
+ if not self.config.fuzz_stage == Const.BACKWARD:
97
+ return
98
+ try:
99
+ grad_saver = getattr(module, CommonField.GRADSAVER)
100
+ except AttributeError:
101
+ logger.warning_on_rank_0(
102
+ f"[msprobe] Free benchmark: get grad saver failed. api_name:{name}"
103
+ )
104
+ return
105
+
106
+ _new_grad_output = grad_output
107
+ try:
108
+ need_grad_tensors, _inner_args = grad_saver.get_vjp_input()
109
+ origin_grad_input = grad_saver.get_grad_input_from_vjp(
110
+ tuple(need_grad_tensors), _new_grad_output, _inner_args
111
+ )
112
+ grad_saver.origin_grad_input = tuple([x.cpu() for x in origin_grad_input])
113
+ grad_saver.calculate_perturbed_grad_input(
114
+ _new_grad_output, need_grad_tensors, _inner_args
115
+ )
116
+ except Exception as e:
117
+ logger.warning_on_rank_0(
118
+ f"[msprobe] Free benchmark: grad vjp calculate failed. api_name:{name} error: {e}"
119
+ )
120
+ return
@@ -1,13 +1,28 @@
1
- from abc import ABC, abstractmethod
2
- from typing import Any
3
-
4
- from msprobe.pytorch.free_benchmark.common.params import DataParams
5
-
6
-
7
- class BaseLayer(ABC):
8
- def __init__(self, api_name: str) -> None:
9
- self.api_name = api_name
10
-
11
- @abstractmethod
12
- def handle(self, params: DataParams) -> Any:
13
- pass
1
+ # Copyright (c) 2024-2024, Huawei Technologies Co., Ltd.
2
+ # All rights reserved.
3
+ #
4
+ # Licensed under the Apache License, Version 2.0 (the "License");
5
+ # you may not use this file except in compliance with the License.
6
+ # You may obtain a copy of the License at
7
+ #
8
+ # http://www.apache.org/licenses/LICENSE-2.0
9
+ #
10
+ # Unless required by applicable law or agreed to in writing, software
11
+ # distributed under the License is distributed on an "AS IS" BASIS,
12
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
13
+ # See the License for the specific language governing permissions and
14
+ # limitations under the License.
15
+
16
+ from abc import ABC, abstractmethod
17
+ from typing import Any
18
+
19
+ from msprobe.pytorch.free_benchmark.common.params import DataParams
20
+
21
+
22
+ class BaseLayer(ABC):
23
+ def __init__(self, api_name: str) -> None:
24
+ self.api_name = api_name
25
+
26
+ @abstractmethod
27
+ def handle(self, params: DataParams) -> Any:
28
+ pass
@@ -1,41 +1,56 @@
1
- from msprobe.pytorch.free_benchmark import FreeBenchmarkException
2
- from msprobe.pytorch.free_benchmark.common.enums import DeviceType, PerturbationMode
3
- from msprobe.pytorch.free_benchmark.perturbed_layers.npu.improve_precision import (
4
- ImprovePrecisionLayer,
5
- )
6
- from msprobe.pytorch.free_benchmark.perturbed_layers.npu.add_noise import AddNoiseLayer
7
- from msprobe.pytorch.free_benchmark.perturbed_layers.npu.bit_noise import BitNoiseLayer
8
- from msprobe.pytorch.free_benchmark.perturbed_layers.npu.no_change import NoChangeLayer
9
- from msprobe.pytorch.free_benchmark.perturbed_layers.npu.change_value import (
10
- ChangeValueLayer,
11
- )
12
- from msprobe.pytorch.free_benchmark.perturbed_layers.run_cpu import CpuLayer
13
-
14
-
15
- class LayerFactory:
16
- layers = {
17
- DeviceType.NPU: {
18
- PerturbationMode.ADD_NOISE: AddNoiseLayer,
19
- PerturbationMode.CHANGE_VALUE: ChangeValueLayer,
20
- PerturbationMode.NO_CHANGE: NoChangeLayer,
21
- PerturbationMode.BIT_NOISE: BitNoiseLayer,
22
- PerturbationMode.IMPROVE_PRECISION: ImprovePrecisionLayer,
23
- },
24
- DeviceType.CPU: {PerturbationMode.TO_CPU: CpuLayer},
25
- }
26
-
27
- @staticmethod
28
- def create(api_name: str, device_type: str, mode: str):
29
- layer = LayerFactory.layers.get(device_type)
30
- if not layer:
31
- raise FreeBenchmarkException(
32
- FreeBenchmarkException.UnsupportedType,
33
- f"无标杆工具不支持当前设备 {device_type}",
34
- )
35
- layer = layer.get(mode)
36
- if not layer:
37
- raise FreeBenchmarkException(
38
- FreeBenchmarkException.UnsupportedType,
39
- f"无标杆工具无法识别该扰动因子 {mode} on {device_type}",
40
- )
41
- return layer(api_name)
1
+ # Copyright (c) 2024-2024, Huawei Technologies Co., Ltd.
2
+ # All rights reserved.
3
+ #
4
+ # Licensed under the Apache License, Version 2.0 (the "License");
5
+ # you may not use this file except in compliance with the License.
6
+ # You may obtain a copy of the License at
7
+ #
8
+ # http://www.apache.org/licenses/LICENSE-2.0
9
+ #
10
+ # Unless required by applicable law or agreed to in writing, software
11
+ # distributed under the License is distributed on an "AS IS" BASIS,
12
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
13
+ # See the License for the specific language governing permissions and
14
+ # limitations under the License.
15
+
16
+ from msprobe.pytorch.free_benchmark import FreeBenchmarkException
17
+ from msprobe.pytorch.free_benchmark.common.enums import DeviceType, PerturbationMode
18
+ from msprobe.pytorch.free_benchmark.perturbed_layers.npu.add_noise import AddNoiseLayer
19
+ from msprobe.pytorch.free_benchmark.perturbed_layers.npu.bit_noise import BitNoiseLayer
20
+ from msprobe.pytorch.free_benchmark.perturbed_layers.npu.change_value import (
21
+ ChangeValueLayer,
22
+ )
23
+ from msprobe.pytorch.free_benchmark.perturbed_layers.npu.improve_precision import (
24
+ ImprovePrecisionLayer,
25
+ )
26
+ from msprobe.pytorch.free_benchmark.perturbed_layers.npu.no_change import NoChangeLayer
27
+ from msprobe.pytorch.free_benchmark.perturbed_layers.run_cpu import CpuLayer
28
+
29
+
30
+ class LayerFactory:
31
+ layers = {
32
+ DeviceType.NPU: {
33
+ PerturbationMode.ADD_NOISE: AddNoiseLayer,
34
+ PerturbationMode.CHANGE_VALUE: ChangeValueLayer,
35
+ PerturbationMode.NO_CHANGE: NoChangeLayer,
36
+ PerturbationMode.BIT_NOISE: BitNoiseLayer,
37
+ PerturbationMode.IMPROVE_PRECISION: ImprovePrecisionLayer,
38
+ },
39
+ DeviceType.CPU: {PerturbationMode.TO_CPU: CpuLayer},
40
+ }
41
+
42
+ @staticmethod
43
+ def create(api_name: str, device_type: str, mode: str):
44
+ layer = LayerFactory.layers.get(device_type)
45
+ if not layer:
46
+ raise FreeBenchmarkException(
47
+ FreeBenchmarkException.UnsupportedType,
48
+ f"无标杆工具不支持当前设备 {device_type}",
49
+ )
50
+ layer = layer.get(mode)
51
+ if not layer:
52
+ raise FreeBenchmarkException(
53
+ FreeBenchmarkException.UnsupportedType,
54
+ f"无标杆工具无法识别该扰动因子 {mode} on {device_type}",
55
+ )
56
+ return layer(api_name)
@@ -1,90 +1,105 @@
1
- import torch
2
- from msprobe.pytorch.free_benchmark import logger
3
- from msprobe.pytorch.free_benchmark.common.constant import ThresholdConfig
4
- from msprobe.pytorch.free_benchmark.common.enums import PerturbationMode
5
- from msprobe.pytorch.free_benchmark.common.params import DataParams
6
- from msprobe.pytorch.free_benchmark.common.utils import TorchC
7
- from msprobe.pytorch.free_benchmark.perturbed_layers.npu.npu_base_layser import (
8
- NpuBaseLayer,
9
- )
10
-
11
-
12
- class AddNoiseLayer(NpuBaseLayer):
13
-
14
- def add_noise(self, tensor_obj):
15
- if isinstance(tensor_obj, torch.Tensor):
16
- self.perturbed_value = ThresholdConfig.PERTURBATION_VALUE_DICT.get(
17
- tensor_obj.dtype
18
- )
19
- if not self.pre_check(tensor_obj):
20
- return tensor_obj
21
- noise = self._get_noise(tensor_obj)
22
- result = TorchC.where(
23
- TorchC.gt(TorchC.abs(tensor_obj), self.perturbed_value ** 0.5),
24
- TorchC.add(noise, tensor_obj),
25
- tensor_obj,
26
- ).to(tensor_obj.dtype)
27
- self.is_added = True
28
- return result
29
- if isinstance(tensor_obj, dict):
30
- return {key: self.add_noise(value) for key, value in tensor_obj.items()}
31
- if isinstance(tensor_obj, (tuple, list)):
32
- return type(tensor_obj)([self.add_noise(value) for value in tensor_obj])
33
- return tensor_obj
34
-
35
- def handle(self, params: DataParams):
36
- """
37
- 对输入添加扰动并返回
38
- """
39
- logger.info_on_rank_0(
40
- f"[msprobe] Free benchmark: Perturbation is "
41
- f"{PerturbationMode.ADD_NOISE} of {self.api_name}."
42
- )
43
- params.perturbed_value = self.add_noise(params.args[params.valid_input_index])
44
- return self.perturbed_result(params)
45
-
46
- def _get_noise(self, tensor_obj):
47
- dtype = tensor_obj.dtype
48
- device = str(tensor_obj.device)
49
- noise = TorchC.full(
50
- tensor_obj.shape,
51
- self.perturbed_value,
52
- device=device,
53
- dtype=dtype,
54
- )
55
- return noise
56
-
57
- def _check_details(self, tensor_obj):
58
- """
59
- 判断是否需要添加扰动
60
- """
61
- if not self.perturbed_value:
62
- logger.warning_on_rank_0(
63
- f"[msprobe] Free Benchmark: For {self.api_name}, "
64
- f"dtype unsupported. Cancel perturbation."
65
- )
66
- return False
67
- if tensor_obj.numel() == 0:
68
- logger.warning_on_rank_0(
69
- f"[msprobe] Free benchmark: For {self.api_name}, tensor shape must > 0."
70
- f" Cancel adding noise."
71
- )
72
- return False
73
- abs_tol = ThresholdConfig.ABS_TOL_VALUE_DICT.get(
74
- tensor_obj.dtype, ThresholdConfig.NOISE_INPUT_LOWER_BOUND
75
- )
76
- try:
77
- max_val = TorchC.max(TorchC.abs(tensor_obj)).item()
78
- except Exception:
79
- logger.warning_on_rank_0(
80
- f"[msprobe] Free Benchmark: For {self.api_name}, "
81
- f"when calculate maximun value, tensor is changed to float32."
82
- )
83
- max_val = TorchC.max(TorchC.abs(tensor_obj.to(torch.float32))).item()
84
- if max_val < abs_tol:
85
- logger.warning_on_rank_0(
86
- f"[msprobe] Free Benchmark: For {self.api_name}, "
87
- f"Maximun value is less than the minimun threshold. Cancel add noise."
88
- )
89
- return False
90
- return True
1
+ # Copyright (c) 2024-2024, Huawei Technologies Co., Ltd.
2
+ # All rights reserved.
3
+ #
4
+ # Licensed under the Apache License, Version 2.0 (the "License");
5
+ # you may not use this file except in compliance with the License.
6
+ # You may obtain a copy of the License at
7
+ #
8
+ # http://www.apache.org/licenses/LICENSE-2.0
9
+ #
10
+ # Unless required by applicable law or agreed to in writing, software
11
+ # distributed under the License is distributed on an "AS IS" BASIS,
12
+ # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
13
+ # See the License for the specific language governing permissions and
14
+ # limitations under the License.
15
+
16
+ import torch
17
+ from msprobe.pytorch.free_benchmark import logger
18
+ from msprobe.pytorch.free_benchmark.common.constant import ThresholdConfig
19
+ from msprobe.pytorch.free_benchmark.common.enums import PerturbationMode
20
+ from msprobe.pytorch.free_benchmark.common.params import DataParams
21
+ from msprobe.pytorch.free_benchmark.common.utils import TorchC
22
+ from msprobe.pytorch.free_benchmark.perturbed_layers.npu.npu_base_layser import (
23
+ NpuBaseLayer,
24
+ )
25
+
26
+
27
+ class AddNoiseLayer(NpuBaseLayer):
28
+
29
+ def add_noise(self, tensor_obj):
30
+ if isinstance(tensor_obj, torch.Tensor):
31
+ self.perturbed_value = ThresholdConfig.PERTURBATION_VALUE_DICT.get(
32
+ tensor_obj.dtype
33
+ )
34
+ if not self.pre_check(tensor_obj):
35
+ return tensor_obj
36
+ noise = self._get_noise(tensor_obj)
37
+ result = TorchC.where(
38
+ TorchC.gt(TorchC.abs(tensor_obj), self.perturbed_value ** 0.5),
39
+ TorchC.add(noise, tensor_obj),
40
+ tensor_obj,
41
+ ).to(tensor_obj.dtype)
42
+ self.is_added = True
43
+ return result
44
+ if isinstance(tensor_obj, dict):
45
+ return {key: self.add_noise(value) for key, value in tensor_obj.items()}
46
+ if isinstance(tensor_obj, (tuple, list)):
47
+ return type(tensor_obj)([self.add_noise(value) for value in tensor_obj])
48
+ return tensor_obj
49
+
50
+ def handle(self, params: DataParams):
51
+ """
52
+ 对输入添加扰动并返回
53
+ """
54
+ logger.info_on_rank_0(
55
+ f"[msprobe] Free benchmark: Perturbation is "
56
+ f"{PerturbationMode.ADD_NOISE} of {self.api_name}."
57
+ )
58
+ params.perturbed_value = self.add_noise(params.args[params.valid_input_index])
59
+ return self.perturbed_result(params)
60
+
61
+ def _get_noise(self, tensor_obj):
62
+ dtype = tensor_obj.dtype
63
+ device = str(tensor_obj.device)
64
+ noise = TorchC.full(
65
+ tensor_obj.shape,
66
+ self.perturbed_value,
67
+ device=device,
68
+ dtype=dtype,
69
+ )
70
+ return noise
71
+
72
+ def _check_details(self, tensor_obj):
73
+ """
74
+ 判断是否需要添加扰动
75
+ """
76
+ if not self.perturbed_value:
77
+ logger.warning_on_rank_0(
78
+ f"[msprobe] Free Benchmark: For {self.api_name}, "
79
+ f"dtype unsupported. Cancel perturbation."
80
+ )
81
+ return False
82
+ if tensor_obj.numel() == 0:
83
+ logger.warning_on_rank_0(
84
+ f"[msprobe] Free benchmark: For {self.api_name}, tensor shape must > 0."
85
+ f" Cancel adding noise."
86
+ )
87
+ return False
88
+ abs_tol = ThresholdConfig.ABS_TOL_VALUE_DICT.get(
89
+ tensor_obj.dtype, ThresholdConfig.NOISE_INPUT_LOWER_BOUND
90
+ )
91
+ try:
92
+ max_val = TorchC.max(TorchC.abs(tensor_obj)).item()
93
+ except Exception:
94
+ logger.warning_on_rank_0(
95
+ f"[msprobe] Free Benchmark: For {self.api_name}, "
96
+ f"when calculate maximun value, tensor is changed to float32."
97
+ )
98
+ max_val = TorchC.max(TorchC.abs(tensor_obj.to(torch.float32))).item()
99
+ if max_val < abs_tol:
100
+ logger.warning_on_rank_0(
101
+ f"[msprobe] Free Benchmark: For {self.api_name}, "
102
+ f"Maximun value is less than the minimun threshold. Cancel add noise."
103
+ )
104
+ return False
105
+ return True