@smake/eigen 1.0.2 → 1.1.1

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (435) hide show
  1. package/README.md +1 -1
  2. package/eigen/Eigen/AccelerateSupport +52 -0
  3. package/eigen/Eigen/Cholesky +18 -21
  4. package/eigen/Eigen/CholmodSupport +28 -28
  5. package/eigen/Eigen/Core +235 -326
  6. package/eigen/Eigen/Eigenvalues +16 -14
  7. package/eigen/Eigen/Geometry +21 -24
  8. package/eigen/Eigen/Householder +9 -8
  9. package/eigen/Eigen/IterativeLinearSolvers +8 -4
  10. package/eigen/Eigen/Jacobi +14 -14
  11. package/eigen/Eigen/KLUSupport +43 -0
  12. package/eigen/Eigen/LU +16 -20
  13. package/eigen/Eigen/MetisSupport +12 -12
  14. package/eigen/Eigen/OrderingMethods +54 -54
  15. package/eigen/Eigen/PaStiXSupport +23 -20
  16. package/eigen/Eigen/PardisoSupport +17 -14
  17. package/eigen/Eigen/QR +18 -21
  18. package/eigen/Eigen/QtAlignedMalloc +5 -13
  19. package/eigen/Eigen/SPQRSupport +21 -14
  20. package/eigen/Eigen/SVD +23 -18
  21. package/eigen/Eigen/Sparse +1 -4
  22. package/eigen/Eigen/SparseCholesky +18 -23
  23. package/eigen/Eigen/SparseCore +18 -17
  24. package/eigen/Eigen/SparseLU +12 -8
  25. package/eigen/Eigen/SparseQR +16 -14
  26. package/eigen/Eigen/StdDeque +5 -2
  27. package/eigen/Eigen/StdList +5 -2
  28. package/eigen/Eigen/StdVector +5 -2
  29. package/eigen/Eigen/SuperLUSupport +30 -24
  30. package/eigen/Eigen/ThreadPool +80 -0
  31. package/eigen/Eigen/UmfPackSupport +19 -17
  32. package/eigen/Eigen/Version +14 -0
  33. package/eigen/Eigen/src/AccelerateSupport/AccelerateSupport.h +423 -0
  34. package/eigen/Eigen/src/AccelerateSupport/InternalHeaderCheck.h +3 -0
  35. package/eigen/Eigen/src/Cholesky/InternalHeaderCheck.h +3 -0
  36. package/eigen/Eigen/src/Cholesky/LDLT.h +377 -401
  37. package/eigen/Eigen/src/Cholesky/LLT.h +332 -360
  38. package/eigen/Eigen/src/Cholesky/LLT_LAPACKE.h +81 -56
  39. package/eigen/Eigen/src/CholmodSupport/CholmodSupport.h +620 -521
  40. package/eigen/Eigen/src/CholmodSupport/InternalHeaderCheck.h +3 -0
  41. package/eigen/Eigen/src/Core/ArithmeticSequence.h +239 -0
  42. package/eigen/Eigen/src/Core/Array.h +341 -294
  43. package/eigen/Eigen/src/Core/ArrayBase.h +190 -203
  44. package/eigen/Eigen/src/Core/ArrayWrapper.h +127 -171
  45. package/eigen/Eigen/src/Core/Assign.h +30 -40
  46. package/eigen/Eigen/src/Core/AssignEvaluator.h +711 -589
  47. package/eigen/Eigen/src/Core/Assign_MKL.h +130 -125
  48. package/eigen/Eigen/src/Core/BandMatrix.h +268 -283
  49. package/eigen/Eigen/src/Core/Block.h +375 -398
  50. package/eigen/Eigen/src/Core/CommaInitializer.h +86 -97
  51. package/eigen/Eigen/src/Core/ConditionEstimator.h +51 -53
  52. package/eigen/Eigen/src/Core/CoreEvaluators.h +1356 -1026
  53. package/eigen/Eigen/src/Core/CoreIterators.h +73 -59
  54. package/eigen/Eigen/src/Core/CwiseBinaryOp.h +114 -132
  55. package/eigen/Eigen/src/Core/CwiseNullaryOp.h +726 -617
  56. package/eigen/Eigen/src/Core/CwiseTernaryOp.h +77 -103
  57. package/eigen/Eigen/src/Core/CwiseUnaryOp.h +56 -68
  58. package/eigen/Eigen/src/Core/CwiseUnaryView.h +132 -95
  59. package/eigen/Eigen/src/Core/DenseBase.h +632 -571
  60. package/eigen/Eigen/src/Core/DenseCoeffsBase.h +511 -624
  61. package/eigen/Eigen/src/Core/DenseStorage.h +512 -509
  62. package/eigen/Eigen/src/Core/DeviceWrapper.h +153 -0
  63. package/eigen/Eigen/src/Core/Diagonal.h +169 -210
  64. package/eigen/Eigen/src/Core/DiagonalMatrix.h +351 -274
  65. package/eigen/Eigen/src/Core/DiagonalProduct.h +12 -10
  66. package/eigen/Eigen/src/Core/Dot.h +172 -222
  67. package/eigen/Eigen/src/Core/EigenBase.h +75 -85
  68. package/eigen/Eigen/src/Core/Fill.h +138 -0
  69. package/eigen/Eigen/src/Core/FindCoeff.h +464 -0
  70. package/eigen/Eigen/src/Core/ForceAlignedAccess.h +90 -109
  71. package/eigen/Eigen/src/Core/Fuzzy.h +82 -105
  72. package/eigen/Eigen/src/Core/GeneralProduct.h +327 -263
  73. package/eigen/Eigen/src/Core/GenericPacketMath.h +1472 -360
  74. package/eigen/Eigen/src/Core/GlobalFunctions.h +194 -151
  75. package/eigen/Eigen/src/Core/IO.h +147 -139
  76. package/eigen/Eigen/src/Core/IndexedView.h +321 -0
  77. package/eigen/Eigen/src/Core/InnerProduct.h +260 -0
  78. package/eigen/Eigen/src/Core/InternalHeaderCheck.h +3 -0
  79. package/eigen/Eigen/src/Core/Inverse.h +56 -66
  80. package/eigen/Eigen/src/Core/Map.h +124 -142
  81. package/eigen/Eigen/src/Core/MapBase.h +256 -281
  82. package/eigen/Eigen/src/Core/MathFunctions.h +1620 -938
  83. package/eigen/Eigen/src/Core/MathFunctionsImpl.h +233 -71
  84. package/eigen/Eigen/src/Core/Matrix.h +491 -416
  85. package/eigen/Eigen/src/Core/MatrixBase.h +468 -453
  86. package/eigen/Eigen/src/Core/NestByValue.h +66 -85
  87. package/eigen/Eigen/src/Core/NoAlias.h +79 -85
  88. package/eigen/Eigen/src/Core/NumTraits.h +235 -148
  89. package/eigen/Eigen/src/Core/PartialReduxEvaluator.h +253 -0
  90. package/eigen/Eigen/src/Core/PermutationMatrix.h +461 -511
  91. package/eigen/Eigen/src/Core/PlainObjectBase.h +871 -894
  92. package/eigen/Eigen/src/Core/Product.h +260 -139
  93. package/eigen/Eigen/src/Core/ProductEvaluators.h +863 -714
  94. package/eigen/Eigen/src/Core/Random.h +161 -136
  95. package/eigen/Eigen/src/Core/RandomImpl.h +262 -0
  96. package/eigen/Eigen/src/Core/RealView.h +250 -0
  97. package/eigen/Eigen/src/Core/Redux.h +366 -336
  98. package/eigen/Eigen/src/Core/Ref.h +308 -209
  99. package/eigen/Eigen/src/Core/Replicate.h +94 -106
  100. package/eigen/Eigen/src/Core/Reshaped.h +398 -0
  101. package/eigen/Eigen/src/Core/ReturnByValue.h +49 -55
  102. package/eigen/Eigen/src/Core/Reverse.h +136 -145
  103. package/eigen/Eigen/src/Core/Select.h +70 -140
  104. package/eigen/Eigen/src/Core/SelfAdjointView.h +262 -285
  105. package/eigen/Eigen/src/Core/SelfCwiseBinaryOp.h +23 -20
  106. package/eigen/Eigen/src/Core/SkewSymmetricMatrix3.h +382 -0
  107. package/eigen/Eigen/src/Core/Solve.h +97 -111
  108. package/eigen/Eigen/src/Core/SolveTriangular.h +131 -129
  109. package/eigen/Eigen/src/Core/SolverBase.h +138 -101
  110. package/eigen/Eigen/src/Core/StableNorm.h +156 -160
  111. package/eigen/Eigen/src/Core/StlIterators.h +619 -0
  112. package/eigen/Eigen/src/Core/Stride.h +91 -88
  113. package/eigen/Eigen/src/Core/Swap.h +70 -38
  114. package/eigen/Eigen/src/Core/Transpose.h +295 -273
  115. package/eigen/Eigen/src/Core/Transpositions.h +272 -317
  116. package/eigen/Eigen/src/Core/TriangularMatrix.h +670 -755
  117. package/eigen/Eigen/src/Core/VectorBlock.h +59 -72
  118. package/eigen/Eigen/src/Core/VectorwiseOp.h +668 -630
  119. package/eigen/Eigen/src/Core/Visitor.h +480 -216
  120. package/eigen/Eigen/src/Core/arch/AVX/Complex.h +407 -293
  121. package/eigen/Eigen/src/Core/arch/AVX/MathFunctions.h +79 -388
  122. package/eigen/Eigen/src/Core/arch/AVX/PacketMath.h +2935 -491
  123. package/eigen/Eigen/src/Core/arch/AVX/Reductions.h +353 -0
  124. package/eigen/Eigen/src/Core/arch/AVX/TypeCasting.h +279 -22
  125. package/eigen/Eigen/src/Core/arch/AVX512/Complex.h +472 -0
  126. package/eigen/Eigen/src/Core/arch/AVX512/GemmKernel.h +1245 -0
  127. package/eigen/Eigen/src/Core/arch/AVX512/MathFunctions.h +85 -333
  128. package/eigen/Eigen/src/Core/arch/AVX512/MathFunctionsFP16.h +75 -0
  129. package/eigen/Eigen/src/Core/arch/AVX512/PacketMath.h +2490 -649
  130. package/eigen/Eigen/src/Core/arch/AVX512/PacketMathFP16.h +1413 -0
  131. package/eigen/Eigen/src/Core/arch/AVX512/Reductions.h +297 -0
  132. package/eigen/Eigen/src/Core/arch/AVX512/TrsmKernel.h +1167 -0
  133. package/eigen/Eigen/src/Core/arch/AVX512/TrsmUnrolls.inc +1219 -0
  134. package/eigen/Eigen/src/Core/arch/AVX512/TypeCasting.h +277 -0
  135. package/eigen/Eigen/src/Core/arch/AVX512/TypeCastingFP16.h +130 -0
  136. package/eigen/Eigen/src/Core/arch/AltiVec/Complex.h +521 -298
  137. package/eigen/Eigen/src/Core/arch/AltiVec/MathFunctions.h +39 -280
  138. package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProduct.h +3686 -0
  139. package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProductCommon.h +205 -0
  140. package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProductMMA.h +901 -0
  141. package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProductMMAbfloat16.h +742 -0
  142. package/eigen/Eigen/src/Core/arch/AltiVec/MatrixVectorProduct.inc +2818 -0
  143. package/eigen/Eigen/src/Core/arch/AltiVec/PacketMath.h +3391 -723
  144. package/eigen/Eigen/src/Core/arch/AltiVec/TypeCasting.h +153 -0
  145. package/eigen/Eigen/src/Core/arch/Default/BFloat16.h +866 -0
  146. package/eigen/Eigen/src/Core/arch/Default/ConjHelper.h +113 -14
  147. package/eigen/Eigen/src/Core/arch/Default/GenericPacketMathFunctions.h +2634 -0
  148. package/eigen/Eigen/src/Core/arch/Default/GenericPacketMathFunctionsFwd.h +227 -0
  149. package/eigen/Eigen/src/Core/arch/Default/Half.h +1091 -0
  150. package/eigen/Eigen/src/Core/arch/Default/Settings.h +11 -13
  151. package/eigen/Eigen/src/Core/arch/GPU/Complex.h +244 -0
  152. package/eigen/Eigen/src/Core/arch/GPU/MathFunctions.h +104 -0
  153. package/eigen/Eigen/src/Core/arch/GPU/PacketMath.h +1712 -0
  154. package/eigen/Eigen/src/Core/arch/GPU/Tuple.h +268 -0
  155. package/eigen/Eigen/src/Core/arch/GPU/TypeCasting.h +77 -0
  156. package/eigen/Eigen/src/Core/arch/HIP/hcc/math_constants.h +23 -0
  157. package/eigen/Eigen/src/Core/arch/HVX/PacketMath.h +1088 -0
  158. package/eigen/Eigen/src/Core/arch/LSX/Complex.h +520 -0
  159. package/eigen/Eigen/src/Core/arch/LSX/GeneralBlockPanelKernel.h +23 -0
  160. package/eigen/Eigen/src/Core/arch/LSX/MathFunctions.h +43 -0
  161. package/eigen/Eigen/src/Core/arch/LSX/PacketMath.h +2866 -0
  162. package/eigen/Eigen/src/Core/arch/LSX/TypeCasting.h +526 -0
  163. package/eigen/Eigen/src/Core/arch/MSA/Complex.h +620 -0
  164. package/eigen/Eigen/src/Core/arch/MSA/MathFunctions.h +379 -0
  165. package/eigen/Eigen/src/Core/arch/MSA/PacketMath.h +1237 -0
  166. package/eigen/Eigen/src/Core/arch/NEON/Complex.h +531 -289
  167. package/eigen/Eigen/src/Core/arch/NEON/GeneralBlockPanelKernel.h +243 -0
  168. package/eigen/Eigen/src/Core/arch/NEON/MathFunctions.h +50 -73
  169. package/eigen/Eigen/src/Core/arch/NEON/PacketMath.h +5915 -579
  170. package/eigen/Eigen/src/Core/arch/NEON/TypeCasting.h +1642 -0
  171. package/eigen/Eigen/src/Core/arch/NEON/UnaryFunctors.h +57 -0
  172. package/eigen/Eigen/src/Core/arch/SSE/Complex.h +366 -334
  173. package/eigen/Eigen/src/Core/arch/SSE/MathFunctions.h +40 -514
  174. package/eigen/Eigen/src/Core/arch/SSE/PacketMath.h +2164 -675
  175. package/eigen/Eigen/src/Core/arch/SSE/Reductions.h +324 -0
  176. package/eigen/Eigen/src/Core/arch/SSE/TypeCasting.h +188 -35
  177. package/eigen/Eigen/src/Core/arch/SVE/MathFunctions.h +48 -0
  178. package/eigen/Eigen/src/Core/arch/SVE/PacketMath.h +674 -0
  179. package/eigen/Eigen/src/Core/arch/SVE/TypeCasting.h +52 -0
  180. package/eigen/Eigen/src/Core/arch/SYCL/InteropHeaders.h +227 -0
  181. package/eigen/Eigen/src/Core/arch/SYCL/MathFunctions.h +303 -0
  182. package/eigen/Eigen/src/Core/arch/SYCL/PacketMath.h +576 -0
  183. package/eigen/Eigen/src/Core/arch/SYCL/TypeCasting.h +83 -0
  184. package/eigen/Eigen/src/Core/arch/ZVector/Complex.h +434 -261
  185. package/eigen/Eigen/src/Core/arch/ZVector/MathFunctions.h +160 -53
  186. package/eigen/Eigen/src/Core/arch/ZVector/PacketMath.h +1073 -605
  187. package/eigen/Eigen/src/Core/functors/AssignmentFunctors.h +123 -117
  188. package/eigen/Eigen/src/Core/functors/BinaryFunctors.h +594 -322
  189. package/eigen/Eigen/src/Core/functors/NullaryFunctors.h +204 -118
  190. package/eigen/Eigen/src/Core/functors/StlFunctors.h +110 -97
  191. package/eigen/Eigen/src/Core/functors/TernaryFunctors.h +34 -7
  192. package/eigen/Eigen/src/Core/functors/UnaryFunctors.h +1158 -530
  193. package/eigen/Eigen/src/Core/products/GeneralBlockPanelKernel.h +2329 -1333
  194. package/eigen/Eigen/src/Core/products/GeneralMatrixMatrix.h +328 -364
  195. package/eigen/Eigen/src/Core/products/GeneralMatrixMatrixTriangular.h +191 -178
  196. package/eigen/Eigen/src/Core/products/GeneralMatrixMatrixTriangular_BLAS.h +85 -82
  197. package/eigen/Eigen/src/Core/products/GeneralMatrixMatrix_BLAS.h +154 -73
  198. package/eigen/Eigen/src/Core/products/GeneralMatrixVector.h +396 -542
  199. package/eigen/Eigen/src/Core/products/GeneralMatrixVector_BLAS.h +80 -77
  200. package/eigen/Eigen/src/Core/products/Parallelizer.h +208 -92
  201. package/eigen/Eigen/src/Core/products/SelfadjointMatrixMatrix.h +331 -375
  202. package/eigen/Eigen/src/Core/products/SelfadjointMatrixMatrix_BLAS.h +206 -224
  203. package/eigen/Eigen/src/Core/products/SelfadjointMatrixVector.h +139 -146
  204. package/eigen/Eigen/src/Core/products/SelfadjointMatrixVector_BLAS.h +58 -61
  205. package/eigen/Eigen/src/Core/products/SelfadjointProduct.h +71 -71
  206. package/eigen/Eigen/src/Core/products/SelfadjointRank2Update.h +48 -46
  207. package/eigen/Eigen/src/Core/products/TriangularMatrixMatrix.h +294 -369
  208. package/eigen/Eigen/src/Core/products/TriangularMatrixMatrix_BLAS.h +246 -238
  209. package/eigen/Eigen/src/Core/products/TriangularMatrixVector.h +244 -247
  210. package/eigen/Eigen/src/Core/products/TriangularMatrixVector_BLAS.h +212 -192
  211. package/eigen/Eigen/src/Core/products/TriangularSolverMatrix.h +328 -275
  212. package/eigen/Eigen/src/Core/products/TriangularSolverMatrix_BLAS.h +108 -109
  213. package/eigen/Eigen/src/Core/products/TriangularSolverVector.h +70 -93
  214. package/eigen/Eigen/src/Core/util/Assert.h +158 -0
  215. package/eigen/Eigen/src/Core/util/BlasUtil.h +413 -290
  216. package/eigen/Eigen/src/Core/util/ConfigureVectorization.h +543 -0
  217. package/eigen/Eigen/src/Core/util/Constants.h +314 -263
  218. package/eigen/Eigen/src/Core/util/DisableStupidWarnings.h +130 -78
  219. package/eigen/Eigen/src/Core/util/EmulateArray.h +270 -0
  220. package/eigen/Eigen/src/Core/util/ForwardDeclarations.h +450 -224
  221. package/eigen/Eigen/src/Core/util/GpuHipCudaDefines.inc +101 -0
  222. package/eigen/Eigen/src/Core/util/GpuHipCudaUndefines.inc +45 -0
  223. package/eigen/Eigen/src/Core/util/IndexedViewHelper.h +487 -0
  224. package/eigen/Eigen/src/Core/util/IntegralConstant.h +279 -0
  225. package/eigen/Eigen/src/Core/util/MKL_support.h +39 -30
  226. package/eigen/Eigen/src/Core/util/Macros.h +939 -646
  227. package/eigen/Eigen/src/Core/util/MaxSizeVector.h +139 -0
  228. package/eigen/Eigen/src/Core/util/Memory.h +1042 -650
  229. package/eigen/Eigen/src/Core/util/Meta.h +618 -426
  230. package/eigen/Eigen/src/Core/util/MoreMeta.h +638 -0
  231. package/eigen/Eigen/src/Core/util/ReenableStupidWarnings.h +32 -19
  232. package/eigen/Eigen/src/Core/util/ReshapedHelper.h +51 -0
  233. package/eigen/Eigen/src/Core/util/Serializer.h +209 -0
  234. package/eigen/Eigen/src/Core/util/StaticAssert.h +51 -164
  235. package/eigen/Eigen/src/Core/util/SymbolicIndex.h +445 -0
  236. package/eigen/Eigen/src/Core/util/XprHelper.h +793 -538
  237. package/eigen/Eigen/src/Eigenvalues/ComplexEigenSolver.h +246 -277
  238. package/eigen/Eigen/src/Eigenvalues/ComplexSchur.h +299 -319
  239. package/eigen/Eigen/src/Eigenvalues/ComplexSchur_LAPACKE.h +52 -48
  240. package/eigen/Eigen/src/Eigenvalues/EigenSolver.h +413 -456
  241. package/eigen/Eigen/src/Eigenvalues/GeneralizedEigenSolver.h +309 -325
  242. package/eigen/Eigen/src/Eigenvalues/GeneralizedSelfAdjointEigenSolver.h +157 -171
  243. package/eigen/Eigen/src/Eigenvalues/HessenbergDecomposition.h +292 -310
  244. package/eigen/Eigen/src/Eigenvalues/InternalHeaderCheck.h +3 -0
  245. package/eigen/Eigen/src/Eigenvalues/MatrixBaseEigenvalues.h +91 -107
  246. package/eigen/Eigen/src/Eigenvalues/RealQZ.h +539 -606
  247. package/eigen/Eigen/src/Eigenvalues/RealSchur.h +348 -382
  248. package/eigen/Eigen/src/Eigenvalues/RealSchur_LAPACKE.h +41 -35
  249. package/eigen/Eigen/src/Eigenvalues/SelfAdjointEigenSolver.h +579 -600
  250. package/eigen/Eigen/src/Eigenvalues/SelfAdjointEigenSolver_LAPACKE.h +47 -44
  251. package/eigen/Eigen/src/Eigenvalues/Tridiagonalization.h +434 -461
  252. package/eigen/Eigen/src/Geometry/AlignedBox.h +307 -214
  253. package/eigen/Eigen/src/Geometry/AngleAxis.h +135 -137
  254. package/eigen/Eigen/src/Geometry/EulerAngles.h +163 -74
  255. package/eigen/Eigen/src/Geometry/Homogeneous.h +289 -333
  256. package/eigen/Eigen/src/Geometry/Hyperplane.h +152 -161
  257. package/eigen/Eigen/src/Geometry/InternalHeaderCheck.h +3 -0
  258. package/eigen/Eigen/src/Geometry/OrthoMethods.h +168 -145
  259. package/eigen/Eigen/src/Geometry/ParametrizedLine.h +141 -104
  260. package/eigen/Eigen/src/Geometry/Quaternion.h +595 -497
  261. package/eigen/Eigen/src/Geometry/Rotation2D.h +110 -108
  262. package/eigen/Eigen/src/Geometry/RotationBase.h +148 -145
  263. package/eigen/Eigen/src/Geometry/Scaling.h +115 -90
  264. package/eigen/Eigen/src/Geometry/Transform.h +896 -953
  265. package/eigen/Eigen/src/Geometry/Translation.h +100 -98
  266. package/eigen/Eigen/src/Geometry/Umeyama.h +79 -84
  267. package/eigen/Eigen/src/Geometry/arch/Geometry_SIMD.h +154 -0
  268. package/eigen/Eigen/src/Householder/BlockHouseholder.h +54 -42
  269. package/eigen/Eigen/src/Householder/Householder.h +104 -122
  270. package/eigen/Eigen/src/Householder/HouseholderSequence.h +416 -382
  271. package/eigen/Eigen/src/Householder/InternalHeaderCheck.h +3 -0
  272. package/eigen/Eigen/src/IterativeLinearSolvers/BasicPreconditioners.h +153 -166
  273. package/eigen/Eigen/src/IterativeLinearSolvers/BiCGSTAB.h +127 -138
  274. package/eigen/Eigen/src/IterativeLinearSolvers/ConjugateGradient.h +95 -124
  275. package/eigen/Eigen/src/IterativeLinearSolvers/IncompleteCholesky.h +269 -267
  276. package/eigen/Eigen/src/IterativeLinearSolvers/IncompleteLUT.h +246 -259
  277. package/eigen/Eigen/src/IterativeLinearSolvers/InternalHeaderCheck.h +3 -0
  278. package/eigen/Eigen/src/IterativeLinearSolvers/IterativeSolverBase.h +218 -217
  279. package/eigen/Eigen/src/IterativeLinearSolvers/LeastSquareConjugateGradient.h +80 -103
  280. package/eigen/Eigen/src/IterativeLinearSolvers/SolveWithGuess.h +59 -63
  281. package/eigen/Eigen/src/Jacobi/InternalHeaderCheck.h +3 -0
  282. package/eigen/Eigen/src/Jacobi/Jacobi.h +256 -291
  283. package/eigen/Eigen/src/KLUSupport/InternalHeaderCheck.h +3 -0
  284. package/eigen/Eigen/src/KLUSupport/KLUSupport.h +339 -0
  285. package/eigen/Eigen/src/LU/Determinant.h +60 -63
  286. package/eigen/Eigen/src/LU/FullPivLU.h +561 -626
  287. package/eigen/Eigen/src/LU/InternalHeaderCheck.h +3 -0
  288. package/eigen/Eigen/src/LU/InverseImpl.h +213 -275
  289. package/eigen/Eigen/src/LU/PartialPivLU.h +407 -435
  290. package/eigen/Eigen/src/LU/PartialPivLU_LAPACKE.h +54 -40
  291. package/eigen/Eigen/src/LU/arch/InverseSize4.h +353 -0
  292. package/eigen/Eigen/src/MetisSupport/InternalHeaderCheck.h +3 -0
  293. package/eigen/Eigen/src/MetisSupport/MetisSupport.h +81 -93
  294. package/eigen/Eigen/src/OrderingMethods/Amd.h +250 -282
  295. package/eigen/Eigen/src/OrderingMethods/Eigen_Colamd.h +950 -1103
  296. package/eigen/Eigen/src/OrderingMethods/InternalHeaderCheck.h +3 -0
  297. package/eigen/Eigen/src/OrderingMethods/Ordering.h +111 -122
  298. package/eigen/Eigen/src/PaStiXSupport/InternalHeaderCheck.h +3 -0
  299. package/eigen/Eigen/src/PaStiXSupport/PaStiXSupport.h +524 -570
  300. package/eigen/Eigen/src/PardisoSupport/InternalHeaderCheck.h +3 -0
  301. package/eigen/Eigen/src/PardisoSupport/PardisoSupport.h +385 -429
  302. package/eigen/Eigen/src/QR/ColPivHouseholderQR.h +494 -473
  303. package/eigen/Eigen/src/QR/ColPivHouseholderQR_LAPACKE.h +120 -56
  304. package/eigen/Eigen/src/QR/CompleteOrthogonalDecomposition.h +223 -137
  305. package/eigen/Eigen/src/QR/FullPivHouseholderQR.h +517 -460
  306. package/eigen/Eigen/src/QR/HouseholderQR.h +412 -278
  307. package/eigen/Eigen/src/QR/HouseholderQR_LAPACKE.h +32 -23
  308. package/eigen/Eigen/src/QR/InternalHeaderCheck.h +3 -0
  309. package/eigen/Eigen/src/SPQRSupport/InternalHeaderCheck.h +3 -0
  310. package/eigen/Eigen/src/SPQRSupport/SuiteSparseQRSupport.h +263 -261
  311. package/eigen/Eigen/src/SVD/BDCSVD.h +872 -679
  312. package/eigen/Eigen/src/SVD/BDCSVD_LAPACKE.h +174 -0
  313. package/eigen/Eigen/src/SVD/InternalHeaderCheck.h +3 -0
  314. package/eigen/Eigen/src/SVD/JacobiSVD.h +585 -543
  315. package/eigen/Eigen/src/SVD/JacobiSVD_LAPACKE.h +85 -49
  316. package/eigen/Eigen/src/SVD/SVDBase.h +281 -160
  317. package/eigen/Eigen/src/SVD/UpperBidiagonalization.h +202 -237
  318. package/eigen/Eigen/src/SparseCholesky/InternalHeaderCheck.h +3 -0
  319. package/eigen/Eigen/src/SparseCholesky/SimplicialCholesky.h +769 -590
  320. package/eigen/Eigen/src/SparseCholesky/SimplicialCholesky_impl.h +318 -129
  321. package/eigen/Eigen/src/SparseCore/AmbiVector.h +202 -251
  322. package/eigen/Eigen/src/SparseCore/CompressedStorage.h +184 -236
  323. package/eigen/Eigen/src/SparseCore/ConservativeSparseSparseProduct.h +140 -184
  324. package/eigen/Eigen/src/SparseCore/InternalHeaderCheck.h +3 -0
  325. package/eigen/Eigen/src/SparseCore/SparseAssign.h +174 -111
  326. package/eigen/Eigen/src/SparseCore/SparseBlock.h +408 -477
  327. package/eigen/Eigen/src/SparseCore/SparseColEtree.h +100 -112
  328. package/eigen/Eigen/src/SparseCore/SparseCompressedBase.h +531 -280
  329. package/eigen/Eigen/src/SparseCore/SparseCwiseBinaryOp.h +559 -347
  330. package/eigen/Eigen/src/SparseCore/SparseCwiseUnaryOp.h +100 -108
  331. package/eigen/Eigen/src/SparseCore/SparseDenseProduct.h +185 -191
  332. package/eigen/Eigen/src/SparseCore/SparseDiagonalProduct.h +71 -71
  333. package/eigen/Eigen/src/SparseCore/SparseDot.h +49 -47
  334. package/eigen/Eigen/src/SparseCore/SparseFuzzy.h +13 -11
  335. package/eigen/Eigen/src/SparseCore/SparseMap.h +243 -253
  336. package/eigen/Eigen/src/SparseCore/SparseMatrix.h +1614 -1142
  337. package/eigen/Eigen/src/SparseCore/SparseMatrixBase.h +403 -357
  338. package/eigen/Eigen/src/SparseCore/SparsePermutation.h +186 -115
  339. package/eigen/Eigen/src/SparseCore/SparseProduct.h +100 -91
  340. package/eigen/Eigen/src/SparseCore/SparseRedux.h +22 -24
  341. package/eigen/Eigen/src/SparseCore/SparseRef.h +268 -295
  342. package/eigen/Eigen/src/SparseCore/SparseSelfAdjointView.h +371 -414
  343. package/eigen/Eigen/src/SparseCore/SparseSolverBase.h +78 -87
  344. package/eigen/Eigen/src/SparseCore/SparseSparseProductWithPruning.h +81 -95
  345. package/eigen/Eigen/src/SparseCore/SparseTranspose.h +62 -71
  346. package/eigen/Eigen/src/SparseCore/SparseTriangularView.h +132 -144
  347. package/eigen/Eigen/src/SparseCore/SparseUtil.h +146 -115
  348. package/eigen/Eigen/src/SparseCore/SparseVector.h +426 -372
  349. package/eigen/Eigen/src/SparseCore/SparseView.h +164 -193
  350. package/eigen/Eigen/src/SparseCore/TriangularSolver.h +129 -170
  351. package/eigen/Eigen/src/SparseLU/InternalHeaderCheck.h +3 -0
  352. package/eigen/Eigen/src/SparseLU/SparseLU.h +814 -618
  353. package/eigen/Eigen/src/SparseLU/SparseLUImpl.h +61 -48
  354. package/eigen/Eigen/src/SparseLU/SparseLU_Memory.h +102 -118
  355. package/eigen/Eigen/src/SparseLU/SparseLU_Structs.h +38 -35
  356. package/eigen/Eigen/src/SparseLU/SparseLU_SupernodalMatrix.h +273 -255
  357. package/eigen/Eigen/src/SparseLU/SparseLU_Utils.h +44 -49
  358. package/eigen/Eigen/src/SparseLU/SparseLU_column_bmod.h +104 -108
  359. package/eigen/Eigen/src/SparseLU/SparseLU_column_dfs.h +90 -101
  360. package/eigen/Eigen/src/SparseLU/SparseLU_copy_to_ucol.h +57 -58
  361. package/eigen/Eigen/src/SparseLU/SparseLU_heap_relax_snode.h +43 -55
  362. package/eigen/Eigen/src/SparseLU/SparseLU_kernel_bmod.h +74 -71
  363. package/eigen/Eigen/src/SparseLU/SparseLU_panel_bmod.h +125 -133
  364. package/eigen/Eigen/src/SparseLU/SparseLU_panel_dfs.h +136 -159
  365. package/eigen/Eigen/src/SparseLU/SparseLU_pivotL.h +51 -52
  366. package/eigen/Eigen/src/SparseLU/SparseLU_pruneL.h +67 -73
  367. package/eigen/Eigen/src/SparseLU/SparseLU_relax_snode.h +24 -26
  368. package/eigen/Eigen/src/SparseQR/InternalHeaderCheck.h +3 -0
  369. package/eigen/Eigen/src/SparseQR/SparseQR.h +451 -490
  370. package/eigen/Eigen/src/StlSupport/StdDeque.h +28 -105
  371. package/eigen/Eigen/src/StlSupport/StdList.h +28 -84
  372. package/eigen/Eigen/src/StlSupport/StdVector.h +28 -108
  373. package/eigen/Eigen/src/StlSupport/details.h +48 -50
  374. package/eigen/Eigen/src/SuperLUSupport/InternalHeaderCheck.h +3 -0
  375. package/eigen/Eigen/src/SuperLUSupport/SuperLUSupport.h +634 -732
  376. package/eigen/Eigen/src/ThreadPool/Barrier.h +70 -0
  377. package/eigen/Eigen/src/ThreadPool/CoreThreadPoolDevice.h +336 -0
  378. package/eigen/Eigen/src/ThreadPool/EventCount.h +241 -0
  379. package/eigen/Eigen/src/ThreadPool/ForkJoin.h +140 -0
  380. package/eigen/Eigen/src/ThreadPool/InternalHeaderCheck.h +4 -0
  381. package/eigen/Eigen/src/ThreadPool/NonBlockingThreadPool.h +587 -0
  382. package/eigen/Eigen/src/ThreadPool/RunQueue.h +230 -0
  383. package/eigen/Eigen/src/ThreadPool/ThreadCancel.h +21 -0
  384. package/eigen/Eigen/src/ThreadPool/ThreadEnvironment.h +43 -0
  385. package/eigen/Eigen/src/ThreadPool/ThreadLocal.h +289 -0
  386. package/eigen/Eigen/src/ThreadPool/ThreadPoolInterface.h +50 -0
  387. package/eigen/Eigen/src/ThreadPool/ThreadYield.h +16 -0
  388. package/eigen/Eigen/src/UmfPackSupport/InternalHeaderCheck.h +3 -0
  389. package/eigen/Eigen/src/UmfPackSupport/UmfPackSupport.h +480 -380
  390. package/eigen/Eigen/src/misc/Image.h +41 -43
  391. package/eigen/Eigen/src/misc/InternalHeaderCheck.h +3 -0
  392. package/eigen/Eigen/src/misc/Kernel.h +39 -41
  393. package/eigen/Eigen/src/misc/RealSvd2x2.h +19 -21
  394. package/eigen/Eigen/src/misc/blas.h +83 -426
  395. package/eigen/Eigen/src/misc/lapacke.h +9976 -16182
  396. package/eigen/Eigen/src/misc/lapacke_helpers.h +163 -0
  397. package/eigen/Eigen/src/misc/lapacke_mangling.h +4 -5
  398. package/eigen/Eigen/src/plugins/ArrayCwiseBinaryOps.inc +344 -0
  399. package/eigen/Eigen/src/plugins/ArrayCwiseUnaryOps.inc +544 -0
  400. package/eigen/Eigen/src/plugins/BlockMethods.inc +1370 -0
  401. package/eigen/Eigen/src/plugins/CommonCwiseBinaryOps.inc +116 -0
  402. package/eigen/Eigen/src/plugins/CommonCwiseUnaryOps.inc +167 -0
  403. package/eigen/Eigen/src/plugins/IndexedViewMethods.inc +192 -0
  404. package/eigen/Eigen/src/plugins/InternalHeaderCheck.inc +3 -0
  405. package/eigen/Eigen/src/plugins/MatrixCwiseBinaryOps.inc +331 -0
  406. package/eigen/Eigen/src/plugins/MatrixCwiseUnaryOps.inc +118 -0
  407. package/eigen/Eigen/src/plugins/ReshapedMethods.inc +133 -0
  408. package/lib/LibEigen.d.ts +4 -0
  409. package/lib/LibEigen.js +14 -0
  410. package/lib/index.d.ts +1 -1
  411. package/lib/index.js +7 -3
  412. package/package.json +2 -10
  413. package/eigen/Eigen/CMakeLists.txt +0 -19
  414. package/eigen/Eigen/src/Core/BooleanRedux.h +0 -164
  415. package/eigen/Eigen/src/Core/arch/CUDA/Complex.h +0 -103
  416. package/eigen/Eigen/src/Core/arch/CUDA/Half.h +0 -675
  417. package/eigen/Eigen/src/Core/arch/CUDA/MathFunctions.h +0 -91
  418. package/eigen/Eigen/src/Core/arch/CUDA/PacketMath.h +0 -333
  419. package/eigen/Eigen/src/Core/arch/CUDA/PacketMathHalf.h +0 -1124
  420. package/eigen/Eigen/src/Core/arch/CUDA/TypeCasting.h +0 -212
  421. package/eigen/Eigen/src/Core/util/NonMPL2.h +0 -3
  422. package/eigen/Eigen/src/Geometry/arch/Geometry_SSE.h +0 -161
  423. package/eigen/Eigen/src/LU/arch/Inverse_SSE.h +0 -338
  424. package/eigen/Eigen/src/SparseCore/MappedSparseMatrix.h +0 -67
  425. package/eigen/Eigen/src/SparseLU/SparseLU_gemm_kernel.h +0 -280
  426. package/eigen/Eigen/src/misc/lapack.h +0 -152
  427. package/eigen/Eigen/src/plugins/ArrayCwiseBinaryOps.h +0 -332
  428. package/eigen/Eigen/src/plugins/ArrayCwiseUnaryOps.h +0 -552
  429. package/eigen/Eigen/src/plugins/BlockMethods.h +0 -1058
  430. package/eigen/Eigen/src/plugins/CommonCwiseBinaryOps.h +0 -115
  431. package/eigen/Eigen/src/plugins/CommonCwiseUnaryOps.h +0 -163
  432. package/eigen/Eigen/src/plugins/MatrixCwiseBinaryOps.h +0 -152
  433. package/eigen/Eigen/src/plugins/MatrixCwiseUnaryOps.h +0 -85
  434. package/lib/eigen.d.ts +0 -2
  435. package/lib/eigen.js +0 -15
@@ -11,162 +11,320 @@
11
11
  #ifndef EIGEN_COMPLEX32_ALTIVEC_H
12
12
  #define EIGEN_COMPLEX32_ALTIVEC_H
13
13
 
14
+ // IWYU pragma: private
15
+ #include "../../InternalHeaderCheck.h"
16
+
14
17
  namespace Eigen {
15
18
 
16
19
  namespace internal {
17
20
 
18
- static Packet4ui p4ui_CONJ_XOR = vec_mergeh((Packet4ui)p4i_ZERO, (Packet4ui)p4f_MZERO);//{ 0x00000000, 0x80000000, 0x00000000, 0x80000000 };
19
- #ifdef __VSX__
21
+ inline Packet4ui p4ui_CONJ_XOR() {
22
+ return vec_mergeh((Packet4ui)p4i_ZERO, (Packet4ui)p4f_MZERO); //{ 0x00000000, 0x80000000, 0x00000000, 0x80000000 };
23
+ }
24
+ #ifdef EIGEN_VECTORIZE_VSX
20
25
  #if defined(_BIG_ENDIAN)
21
- static Packet2ul p2ul_CONJ_XOR1 = (Packet2ul) vec_sld((Packet4ui) p2d_MZERO, (Packet4ui) p2l_ZERO, 8);//{ 0x8000000000000000, 0x0000000000000000 };
22
- static Packet2ul p2ul_CONJ_XOR2 = (Packet2ul) vec_sld((Packet4ui) p2l_ZERO, (Packet4ui) p2d_MZERO, 8);//{ 0x8000000000000000, 0x0000000000000000 };
26
+ inline Packet2ul p2ul_CONJ_XOR1() {
27
+ return (Packet2ul)vec_sld((Packet4ui)p2d_MZERO, (Packet4ui)p2l_ZERO,
28
+ 8); //{ 0x8000000000000000, 0x0000000000000000 };
29
+ }
30
+ inline Packet2ul p2ul_CONJ_XOR2() {
31
+ return (Packet2ul)vec_sld((Packet4ui)p2l_ZERO, (Packet4ui)p2d_MZERO,
32
+ 8); //{ 0x8000000000000000, 0x0000000000000000 };
33
+ }
23
34
  #else
24
- static Packet2ul p2ul_CONJ_XOR1 = (Packet2ul) vec_sld((Packet4ui) p2l_ZERO, (Packet4ui) p2d_MZERO, 8);//{ 0x8000000000000000, 0x0000000000000000 };
25
- static Packet2ul p2ul_CONJ_XOR2 = (Packet2ul) vec_sld((Packet4ui) p2d_MZERO, (Packet4ui) p2l_ZERO, 8);//{ 0x8000000000000000, 0x0000000000000000 };
35
+ inline Packet2ul p2ul_CONJ_XOR1() {
36
+ return (Packet2ul)vec_sld((Packet4ui)p2l_ZERO, (Packet4ui)p2d_MZERO,
37
+ 8); //{ 0x8000000000000000, 0x0000000000000000 };
38
+ }
39
+ inline Packet2ul p2ul_CONJ_XOR2() {
40
+ return (Packet2ul)vec_sld((Packet4ui)p2d_MZERO, (Packet4ui)p2l_ZERO,
41
+ 8); //{ 0x8000000000000000, 0x0000000000000000 };
42
+ }
26
43
  #endif
27
44
  #endif
28
45
 
29
46
  //---------- float ----------
30
- struct Packet2cf
31
- {
32
- EIGEN_STRONG_INLINE explicit Packet2cf() : v(p4f_ZERO) {}
47
+ struct Packet2cf {
48
+ EIGEN_STRONG_INLINE explicit Packet2cf() {}
33
49
  EIGEN_STRONG_INLINE explicit Packet2cf(const Packet4f& a) : v(a) {}
34
- Packet4f v;
50
+
51
+ EIGEN_STRONG_INLINE Packet2cf pmul(const Packet2cf& a, const Packet2cf& b) {
52
+ Packet4f v1, v2;
53
+
54
+ // Permute and multiply the real parts of a and b
55
+ v1 = vec_perm(a.v, a.v, p16uc_PSET32_WODD);
56
+ // Get the imaginary parts of a
57
+ v2 = vec_perm(a.v, a.v, p16uc_PSET32_WEVEN);
58
+ // multiply a_re * b
59
+ v1 = vec_madd(v1, b.v, p4f_ZERO);
60
+ // multiply a_im * b and get the conjugate result
61
+ v2 = vec_madd(v2, b.v, p4f_ZERO);
62
+ v2 = reinterpret_cast<Packet4f>(pxor(v2, reinterpret_cast<Packet4f>(p4ui_CONJ_XOR())));
63
+ // permute back to a proper order
64
+ v2 = vec_perm(v2, v2, p16uc_COMPLEX32_REV);
65
+
66
+ return Packet2cf(padd<Packet4f>(v1, v2));
67
+ }
68
+
69
+ EIGEN_STRONG_INLINE Packet2cf& operator*=(const Packet2cf& b) {
70
+ v = pmul(Packet2cf(*this), b).v;
71
+ return *this;
72
+ }
73
+ EIGEN_STRONG_INLINE Packet2cf operator*(const Packet2cf& b) const { return Packet2cf(*this) *= b; }
74
+
75
+ EIGEN_STRONG_INLINE Packet2cf& operator+=(const Packet2cf& b) {
76
+ v = padd(v, b.v);
77
+ return *this;
78
+ }
79
+ EIGEN_STRONG_INLINE Packet2cf operator+(const Packet2cf& b) const { return Packet2cf(*this) += b; }
80
+ EIGEN_STRONG_INLINE Packet2cf& operator-=(const Packet2cf& b) {
81
+ v = psub(v, b.v);
82
+ return *this;
83
+ }
84
+ EIGEN_STRONG_INLINE Packet2cf operator-(const Packet2cf& b) const { return Packet2cf(*this) -= b; }
85
+ EIGEN_STRONG_INLINE Packet2cf operator-(void) const { return Packet2cf(-v); }
86
+
87
+ Packet4f v;
35
88
  };
36
89
 
37
- template<> struct packet_traits<std::complex<float> > : default_packet_traits
38
- {
90
+ template <>
91
+ struct packet_traits<std::complex<float> > : default_packet_traits {
39
92
  typedef Packet2cf type;
40
93
  typedef Packet2cf half;
94
+ typedef Packet4f as_real;
41
95
  enum {
42
96
  Vectorizable = 1,
43
97
  AlignedOnScalar = 1,
44
98
  size = 2,
45
- HasHalfPacket = 0,
46
99
 
47
- HasAdd = 1,
48
- HasSub = 1,
49
- HasMul = 1,
50
- HasDiv = 1,
100
+ HasAdd = 1,
101
+ HasSub = 1,
102
+ HasMul = 1,
103
+ HasDiv = 1,
51
104
  HasNegate = 1,
52
- HasAbs = 0,
53
- HasAbs2 = 0,
54
- HasMin = 0,
55
- HasMax = 0,
56
- #ifdef __VSX__
57
- HasBlend = 1,
105
+ HasAbs = 0,
106
+ HasAbs2 = 0,
107
+ HasMin = 0,
108
+ HasMax = 0,
109
+ HasSqrt = 1,
110
+ HasLog = 1,
111
+ HasExp = 1,
112
+ #ifdef EIGEN_VECTORIZE_VSX
113
+ HasBlend = 1,
58
114
  #endif
59
115
  HasSetLinear = 0
60
116
  };
61
117
  };
62
118
 
63
- template<> struct unpacket_traits<Packet2cf> { typedef std::complex<float> type; enum {size=2, alignment=Aligned16}; typedef Packet2cf half; };
119
+ template <>
120
+ struct unpacket_traits<Packet2cf> {
121
+ typedef std::complex<float> type;
122
+ enum {
123
+ size = 2,
124
+ alignment = Aligned16,
125
+ vectorizable = true,
126
+ masked_load_available = false,
127
+ masked_store_available = false
128
+ };
129
+ typedef Packet2cf half;
130
+ typedef Packet4f as_real;
131
+ };
64
132
 
65
- template<> EIGEN_STRONG_INLINE Packet2cf pset1<Packet2cf>(const std::complex<float>& from)
66
- {
133
+ template <>
134
+ EIGEN_STRONG_INLINE Packet2cf pset1<Packet2cf>(const std::complex<float>& from) {
67
135
  Packet2cf res;
68
- if((std::ptrdiff_t(&from) % 16) == 0)
69
- res.v = pload<Packet4f>((const float *)&from);
136
+ #ifdef EIGEN_VECTORIZE_VSX
137
+ // Load a single std::complex<float> from memory and duplicate
138
+ //
139
+ // Using pload would read past the end of the reference in this case
140
+ // Using vec_xl_len + vec_splat, generates poor assembly
141
+ __asm__("lxvdsx %x0,%y1" : "=wa"(res.v) : "Z"(from));
142
+ #else
143
+ if ((std::ptrdiff_t(&from) % 16) == 0)
144
+ res.v = pload<Packet4f>((const float*)&from);
70
145
  else
71
- res.v = ploadu<Packet4f>((const float *)&from);
146
+ res.v = ploadu<Packet4f>((const float*)&from);
72
147
  res.v = vec_perm(res.v, res.v, p16uc_PSET64_HI);
148
+ #endif
73
149
  return res;
74
150
  }
75
151
 
76
- template<> EIGEN_STRONG_INLINE Packet2cf pload<Packet2cf>(const std::complex<float>* from) { return Packet2cf(pload<Packet4f>((const float *) from)); }
77
- template<> EIGEN_STRONG_INLINE Packet2cf ploadu<Packet2cf>(const std::complex<float>* from) { return Packet2cf(ploadu<Packet4f>((const float*) from)); }
78
- template<> EIGEN_STRONG_INLINE Packet2cf ploaddup<Packet2cf>(const std::complex<float>* from) { return pset1<Packet2cf>(*from); }
79
-
80
- template<> EIGEN_STRONG_INLINE void pstore <std::complex<float> >(std::complex<float> * to, const Packet2cf& from) { pstore((float*)to, from.v); }
81
- template<> EIGEN_STRONG_INLINE void pstoreu<std::complex<float> >(std::complex<float> * to, const Packet2cf& from) { pstoreu((float*)to, from.v); }
82
-
83
- template<> EIGEN_DEVICE_FUNC inline Packet2cf pgather<std::complex<float>, Packet2cf>(const std::complex<float>* from, Index stride)
84
- {
85
- std::complex<float> EIGEN_ALIGN16 af[2];
86
- af[0] = from[0*stride];
87
- af[1] = from[1*stride];
88
- return pload<Packet2cf>(af);
89
- }
90
- template<> EIGEN_DEVICE_FUNC inline void pscatter<std::complex<float>, Packet2cf>(std::complex<float>* to, const Packet2cf& from, Index stride)
91
- {
92
- std::complex<float> EIGEN_ALIGN16 af[2];
93
- pstore<std::complex<float> >((std::complex<float> *) af, from);
94
- to[0*stride] = af[0];
95
- to[1*stride] = af[1];
96
- }
97
-
98
- template<> EIGEN_STRONG_INLINE Packet2cf padd<Packet2cf>(const Packet2cf& a, const Packet2cf& b) { return Packet2cf(a.v + b.v); }
99
- template<> EIGEN_STRONG_INLINE Packet2cf psub<Packet2cf>(const Packet2cf& a, const Packet2cf& b) { return Packet2cf(a.v - b.v); }
100
- template<> EIGEN_STRONG_INLINE Packet2cf pnegate(const Packet2cf& a) { return Packet2cf(pnegate(a.v)); }
101
- template<> EIGEN_STRONG_INLINE Packet2cf pconj(const Packet2cf& a) { return Packet2cf(pxor<Packet4f>(a.v, reinterpret_cast<Packet4f>(p4ui_CONJ_XOR))); }
102
-
103
- template<> EIGEN_STRONG_INLINE Packet2cf pmul<Packet2cf>(const Packet2cf& a, const Packet2cf& b)
104
- {
105
- Packet4f v1, v2;
106
-
107
- // Permute and multiply the real parts of a and b
108
- v1 = vec_perm(a.v, a.v, p16uc_PSET32_WODD);
109
- // Get the imaginary parts of a
110
- v2 = vec_perm(a.v, a.v, p16uc_PSET32_WEVEN);
111
- // multiply a_re * b
112
- v1 = vec_madd(v1, b.v, p4f_ZERO);
113
- // multiply a_im * b and get the conjugate result
114
- v2 = vec_madd(v2, b.v, p4f_ZERO);
115
- v2 = reinterpret_cast<Packet4f>(pxor(v2, reinterpret_cast<Packet4f>(p4ui_CONJ_XOR)));
116
- // permute back to a proper order
117
- v2 = vec_perm(v2, v2, p16uc_COMPLEX32_REV);
118
-
119
- return Packet2cf(padd<Packet4f>(v1, v2));
120
- }
121
-
122
- template<> EIGEN_STRONG_INLINE Packet2cf pand <Packet2cf>(const Packet2cf& a, const Packet2cf& b) { return Packet2cf(pand<Packet4f>(a.v, b.v)); }
123
- template<> EIGEN_STRONG_INLINE Packet2cf por <Packet2cf>(const Packet2cf& a, const Packet2cf& b) { return Packet2cf(por<Packet4f>(a.v, b.v)); }
124
- template<> EIGEN_STRONG_INLINE Packet2cf pxor <Packet2cf>(const Packet2cf& a, const Packet2cf& b) { return Packet2cf(pxor<Packet4f>(a.v, b.v)); }
125
- template<> EIGEN_STRONG_INLINE Packet2cf pandnot<Packet2cf>(const Packet2cf& a, const Packet2cf& b) { return Packet2cf(pandnot<Packet4f>(a.v, b.v)); }
126
-
127
- template<> EIGEN_STRONG_INLINE void prefetch<std::complex<float> >(const std::complex<float> * addr) { EIGEN_PPC_PREFETCH(addr); }
128
-
129
- template<> EIGEN_STRONG_INLINE std::complex<float> pfirst<Packet2cf>(const Packet2cf& a)
130
- {
131
- std::complex<float> EIGEN_ALIGN16 res[2];
132
- pstore((float *)&res, a.v);
152
+ template <>
153
+ EIGEN_STRONG_INLINE Packet2cf pload<Packet2cf>(const std::complex<float>* from) {
154
+ return Packet2cf(pload<Packet4f>((const float*)from));
155
+ }
156
+ template <>
157
+ EIGEN_STRONG_INLINE Packet2cf ploadu<Packet2cf>(const std::complex<float>* from) {
158
+ return Packet2cf(ploadu<Packet4f>((const float*)from));
159
+ }
160
+ template <>
161
+ EIGEN_ALWAYS_INLINE Packet2cf pload_partial<Packet2cf>(const std::complex<float>* from, const Index n,
162
+ const Index offset) {
163
+ return Packet2cf(pload_partial<Packet4f>((const float*)from, n * 2, offset * 2));
164
+ }
165
+ template <>
166
+ EIGEN_ALWAYS_INLINE Packet2cf ploadu_partial<Packet2cf>(const std::complex<float>* from, const Index n,
167
+ const Index offset) {
168
+ return Packet2cf(ploadu_partial<Packet4f>((const float*)from, n * 2, offset * 2));
169
+ }
170
+ template <>
171
+ EIGEN_STRONG_INLINE Packet2cf ploaddup<Packet2cf>(const std::complex<float>* from) {
172
+ return pset1<Packet2cf>(*from);
173
+ }
174
+
175
+ template <>
176
+ EIGEN_STRONG_INLINE void pstore<std::complex<float> >(std::complex<float>* to, const Packet2cf& from) {
177
+ pstore((float*)to, from.v);
178
+ }
179
+ template <>
180
+ EIGEN_STRONG_INLINE void pstoreu<std::complex<float> >(std::complex<float>* to, const Packet2cf& from) {
181
+ pstoreu((float*)to, from.v);
182
+ }
183
+ template <>
184
+ EIGEN_ALWAYS_INLINE void pstore_partial<std::complex<float> >(std::complex<float>* to, const Packet2cf& from,
185
+ const Index n, const Index offset) {
186
+ pstore_partial((float*)to, from.v, n * 2, offset * 2);
187
+ }
188
+ template <>
189
+ EIGEN_ALWAYS_INLINE void pstoreu_partial<std::complex<float> >(std::complex<float>* to, const Packet2cf& from,
190
+ const Index n, const Index offset) {
191
+ pstoreu_partial((float*)to, from.v, n * 2, offset * 2);
192
+ }
193
+
194
+ EIGEN_STRONG_INLINE Packet2cf pload2(const std::complex<float>& from0, const std::complex<float>& from1) {
195
+ Packet4f res0, res1;
196
+ #ifdef EIGEN_VECTORIZE_VSX
197
+ // Load two std::complex<float> from memory and combine
198
+ __asm__("lxsdx %x0,%y1" : "=wa"(res0) : "Z"(from0));
199
+ __asm__("lxsdx %x0,%y1" : "=wa"(res1) : "Z"(from1));
200
+ #ifdef _BIG_ENDIAN
201
+ __asm__("xxpermdi %x0, %x1, %x2, 0" : "=wa"(res0) : "wa"(res0), "wa"(res1));
202
+ #else
203
+ __asm__("xxpermdi %x0, %x2, %x1, 0" : "=wa"(res0) : "wa"(res0), "wa"(res1));
204
+ #endif
205
+ #else
206
+ *reinterpret_cast<std::complex<float>*>(&res0) = from0;
207
+ *reinterpret_cast<std::complex<float>*>(&res1) = from1;
208
+ res0 = vec_perm(res0, res1, p16uc_TRANSPOSE64_HI);
209
+ #endif
210
+ return Packet2cf(res0);
211
+ }
212
+
213
+ template <>
214
+ EIGEN_ALWAYS_INLINE Packet2cf pload_ignore<Packet2cf>(const std::complex<float>* from) {
215
+ Packet2cf res;
216
+ res.v = pload_ignore<Packet4f>(reinterpret_cast<const float*>(from));
217
+ return res;
218
+ }
219
+
220
+ template <typename Scalar, typename Packet>
221
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet pgather_complex_size2(const Scalar* from, Index stride,
222
+ const Index n = 2) {
223
+ eigen_internal_assert(n <= unpacket_traits<Packet>::size && "number of elements will gather past end of packet");
224
+ EIGEN_ALIGN16 Scalar af[2];
225
+ for (Index i = 0; i < n; i++) {
226
+ af[i] = from[i * stride];
227
+ }
228
+ return pload_ignore<Packet>(af);
229
+ }
230
+ template <>
231
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet2cf pgather<std::complex<float>, Packet2cf>(const std::complex<float>* from,
232
+ Index stride) {
233
+ return pgather_complex_size2<std::complex<float>, Packet2cf>(from, stride);
234
+ }
235
+ template <>
236
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet2cf
237
+ pgather_partial<std::complex<float>, Packet2cf>(const std::complex<float>* from, Index stride, const Index n) {
238
+ return pgather_complex_size2<std::complex<float>, Packet2cf>(from, stride, n);
239
+ }
240
+ template <typename Scalar, typename Packet>
241
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter_complex_size2(Scalar* to, const Packet& from, Index stride,
242
+ const Index n = 2) {
243
+ eigen_internal_assert(n <= unpacket_traits<Packet>::size && "number of elements will scatter past end of packet");
244
+ EIGEN_ALIGN16 Scalar af[2];
245
+ pstore<Scalar>((Scalar*)af, from);
246
+ for (Index i = 0; i < n; i++) {
247
+ to[i * stride] = af[i];
248
+ }
249
+ }
250
+ template <>
251
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter<std::complex<float>, Packet2cf>(std::complex<float>* to,
252
+ const Packet2cf& from,
253
+ Index stride) {
254
+ pscatter_complex_size2<std::complex<float>, Packet2cf>(to, from, stride);
255
+ }
256
+ template <>
257
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter_partial<std::complex<float>, Packet2cf>(std::complex<float>* to,
258
+ const Packet2cf& from,
259
+ Index stride,
260
+ const Index n) {
261
+ pscatter_complex_size2<std::complex<float>, Packet2cf>(to, from, stride, n);
262
+ }
263
+
264
+ template <>
265
+ EIGEN_STRONG_INLINE Packet2cf padd<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
266
+ return Packet2cf(a.v + b.v);
267
+ }
268
+ template <>
269
+ EIGEN_STRONG_INLINE Packet2cf psub<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
270
+ return Packet2cf(a.v - b.v);
271
+ }
272
+ template <>
273
+ EIGEN_STRONG_INLINE Packet2cf pnegate(const Packet2cf& a) {
274
+ return Packet2cf(pnegate(a.v));
275
+ }
276
+ template <>
277
+ EIGEN_STRONG_INLINE Packet2cf pconj(const Packet2cf& a) {
278
+ return Packet2cf(pxor<Packet4f>(a.v, reinterpret_cast<Packet4f>(p4ui_CONJ_XOR())));
279
+ }
280
+
281
+ template <>
282
+ EIGEN_STRONG_INLINE Packet2cf pand<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
283
+ return Packet2cf(pand<Packet4f>(a.v, b.v));
284
+ }
285
+ template <>
286
+ EIGEN_STRONG_INLINE Packet2cf por<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
287
+ return Packet2cf(por<Packet4f>(a.v, b.v));
288
+ }
289
+ template <>
290
+ EIGEN_STRONG_INLINE Packet2cf pxor<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
291
+ return Packet2cf(pxor<Packet4f>(a.v, b.v));
292
+ }
293
+ template <>
294
+ EIGEN_STRONG_INLINE Packet2cf pandnot<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
295
+ return Packet2cf(pandnot<Packet4f>(a.v, b.v));
296
+ }
297
+
298
+ template <>
299
+ EIGEN_STRONG_INLINE void prefetch<std::complex<float> >(const std::complex<float>* addr) {
300
+ EIGEN_PPC_PREFETCH(addr);
301
+ }
302
+
303
+ template <>
304
+ EIGEN_STRONG_INLINE std::complex<float> pfirst<Packet2cf>(const Packet2cf& a) {
305
+ EIGEN_ALIGN16 std::complex<float> res[2];
306
+ pstore((float*)&res, a.v);
133
307
 
134
308
  return res[0];
135
309
  }
136
310
 
137
- template<> EIGEN_STRONG_INLINE Packet2cf preverse(const Packet2cf& a)
138
- {
311
+ template <>
312
+ EIGEN_STRONG_INLINE Packet2cf preverse(const Packet2cf& a) {
139
313
  Packet4f rev_a;
140
- rev_a = vec_perm(a.v, a.v, p16uc_COMPLEX32_REV2);
314
+ rev_a = vec_sld(a.v, a.v, 8);
141
315
  return Packet2cf(rev_a);
142
316
  }
143
317
 
144
- template<> EIGEN_STRONG_INLINE std::complex<float> predux<Packet2cf>(const Packet2cf& a)
145
- {
318
+ template <>
319
+ EIGEN_STRONG_INLINE std::complex<float> predux<Packet2cf>(const Packet2cf& a) {
146
320
  Packet4f b;
147
321
  b = vec_sld(a.v, a.v, 8);
148
322
  b = padd<Packet4f>(a.v, b);
149
323
  return pfirst<Packet2cf>(Packet2cf(b));
150
324
  }
151
325
 
152
- template<> EIGEN_STRONG_INLINE Packet2cf preduxp<Packet2cf>(const Packet2cf* vecs)
153
- {
154
- Packet4f b1, b2;
155
- #ifdef _BIG_ENDIAN
156
- b1 = vec_sld(vecs[0].v, vecs[1].v, 8);
157
- b2 = vec_sld(vecs[1].v, vecs[0].v, 8);
158
- #else
159
- b1 = vec_sld(vecs[1].v, vecs[0].v, 8);
160
- b2 = vec_sld(vecs[0].v, vecs[1].v, 8);
161
- #endif
162
- b2 = vec_sld(b2, b2, 8);
163
- b2 = padd<Packet4f>(b1, b2);
164
-
165
- return Packet2cf(b2);
166
- }
167
-
168
- template<> EIGEN_STRONG_INLINE std::complex<float> predux_mul<Packet2cf>(const Packet2cf& a)
169
- {
326
+ template <>
327
+ EIGEN_STRONG_INLINE std::complex<float> predux_mul<Packet2cf>(const Packet2cf& a) {
170
328
  Packet4f b;
171
329
  Packet2cf prod;
172
330
  b = vec_sld(a.v, a.v, 8);
@@ -175,256 +333,321 @@ template<> EIGEN_STRONG_INLINE std::complex<float> predux_mul<Packet2cf>(const P
175
333
  return pfirst<Packet2cf>(prod);
176
334
  }
177
335
 
178
- template<int Offset>
179
- struct palign_impl<Offset,Packet2cf>
180
- {
181
- static EIGEN_STRONG_INLINE void run(Packet2cf& first, const Packet2cf& second)
182
- {
183
- if (Offset==1)
184
- {
185
- #ifdef _BIG_ENDIAN
186
- first.v = vec_sld(first.v, second.v, 8);
187
- #else
188
- first.v = vec_sld(second.v, first.v, 8);
189
- #endif
190
- }
191
- }
192
- };
193
-
194
- template<> struct conj_helper<Packet2cf, Packet2cf, false,true>
195
- {
196
- EIGEN_STRONG_INLINE Packet2cf pmadd(const Packet2cf& x, const Packet2cf& y, const Packet2cf& c) const
197
- { return padd(pmul(x,y),c); }
198
-
199
- EIGEN_STRONG_INLINE Packet2cf pmul(const Packet2cf& a, const Packet2cf& b) const
200
- {
201
- return internal::pmul(a, pconj(b));
202
- }
203
- };
204
-
205
- template<> struct conj_helper<Packet2cf, Packet2cf, true,false>
206
- {
207
- EIGEN_STRONG_INLINE Packet2cf pmadd(const Packet2cf& x, const Packet2cf& y, const Packet2cf& c) const
208
- { return padd(pmul(x,y),c); }
336
+ EIGEN_MAKE_CONJ_HELPER_CPLX_REAL(Packet2cf, Packet4f)
209
337
 
210
- EIGEN_STRONG_INLINE Packet2cf pmul(const Packet2cf& a, const Packet2cf& b) const
211
- {
212
- return internal::pmul(pconj(a), b);
213
- }
214
- };
215
-
216
- template<> struct conj_helper<Packet2cf, Packet2cf, true,true>
217
- {
218
- EIGEN_STRONG_INLINE Packet2cf pmadd(const Packet2cf& x, const Packet2cf& y, const Packet2cf& c) const
219
- { return padd(pmul(x,y),c); }
220
-
221
- EIGEN_STRONG_INLINE Packet2cf pmul(const Packet2cf& a, const Packet2cf& b) const
222
- {
223
- return pconj(internal::pmul(a, b));
224
- }
225
- };
226
-
227
- EIGEN_MAKE_CONJ_HELPER_CPLX_REAL(Packet2cf,Packet4f)
228
-
229
- template<> EIGEN_STRONG_INLINE Packet2cf pdiv<Packet2cf>(const Packet2cf& a, const Packet2cf& b)
230
- {
231
- // TODO optimize it for AltiVec
232
- Packet2cf res = conj_helper<Packet2cf,Packet2cf,false,true>().pmul(a, b);
233
- Packet4f s = pmul<Packet4f>(b.v, b.v);
234
- return Packet2cf(pdiv(res.v, padd<Packet4f>(s, vec_perm(s, s, p16uc_COMPLEX32_REV))));
338
+ template <>
339
+ EIGEN_STRONG_INLINE Packet2cf pdiv<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
340
+ return pdiv_complex(a, b);
235
341
  }
236
342
 
237
- template<> EIGEN_STRONG_INLINE Packet2cf pcplxflip<Packet2cf>(const Packet2cf& x)
238
- {
343
+ template <>
344
+ EIGEN_STRONG_INLINE Packet2cf pcplxflip<Packet2cf>(const Packet2cf& x) {
239
345
  return Packet2cf(vec_perm(x.v, x.v, p16uc_COMPLEX32_REV));
240
346
  }
241
347
 
242
- EIGEN_STRONG_INLINE void ptranspose(PacketBlock<Packet2cf,2>& kernel)
243
- {
348
+ EIGEN_STRONG_INLINE void ptranspose(PacketBlock<Packet2cf, 2>& kernel) {
349
+ #ifdef EIGEN_VECTORIZE_VSX
350
+ Packet4f tmp = reinterpret_cast<Packet4f>(
351
+ vec_mergeh(reinterpret_cast<Packet2d>(kernel.packet[0].v), reinterpret_cast<Packet2d>(kernel.packet[1].v)));
352
+ kernel.packet[1].v = reinterpret_cast<Packet4f>(
353
+ vec_mergel(reinterpret_cast<Packet2d>(kernel.packet[0].v), reinterpret_cast<Packet2d>(kernel.packet[1].v)));
354
+ #else
244
355
  Packet4f tmp = vec_perm(kernel.packet[0].v, kernel.packet[1].v, p16uc_TRANSPOSE64_HI);
245
356
  kernel.packet[1].v = vec_perm(kernel.packet[0].v, kernel.packet[1].v, p16uc_TRANSPOSE64_LO);
357
+ #endif
246
358
  kernel.packet[0].v = tmp;
247
359
  }
248
360
 
249
- #ifdef __VSX__
250
- template<> EIGEN_STRONG_INLINE Packet2cf pblend(const Selector<2>& ifPacket, const Packet2cf& thenPacket, const Packet2cf& elsePacket) {
361
+ template <>
362
+ EIGEN_STRONG_INLINE Packet2cf pcmp_eq(const Packet2cf& a, const Packet2cf& b) {
363
+ Packet4f eq = reinterpret_cast<Packet4f>(vec_cmpeq(a.v, b.v));
364
+ return Packet2cf(vec_and(eq, vec_perm(eq, eq, p16uc_COMPLEX32_REV)));
365
+ }
366
+
367
+ #ifdef EIGEN_VECTORIZE_VSX
368
+ template <>
369
+ EIGEN_STRONG_INLINE Packet2cf pblend(const Selector<2>& ifPacket, const Packet2cf& thenPacket,
370
+ const Packet2cf& elsePacket) {
251
371
  Packet2cf result;
252
- result.v = reinterpret_cast<Packet4f>(pblend<Packet2d>(ifPacket, reinterpret_cast<Packet2d>(thenPacket.v), reinterpret_cast<Packet2d>(elsePacket.v)));
372
+ result.v = reinterpret_cast<Packet4f>(
373
+ pblend<Packet2d>(ifPacket, reinterpret_cast<Packet2d>(thenPacket.v), reinterpret_cast<Packet2d>(elsePacket.v)));
253
374
  return result;
254
375
  }
255
376
  #endif
256
377
 
378
+ template <>
379
+ EIGEN_STRONG_INLINE Packet2cf psqrt<Packet2cf>(const Packet2cf& a) {
380
+ return psqrt_complex<Packet2cf>(a);
381
+ }
382
+
383
+ template <>
384
+ EIGEN_STRONG_INLINE Packet2cf plog<Packet2cf>(const Packet2cf& a) {
385
+ return plog_complex<Packet2cf>(a);
386
+ }
387
+
388
+ template <>
389
+ EIGEN_STRONG_INLINE Packet2cf pexp<Packet2cf>(const Packet2cf& a) {
390
+ return pexp_complex<Packet2cf>(a);
391
+ }
392
+
257
393
  //---------- double ----------
258
- #ifdef __VSX__
259
- struct Packet1cd
260
- {
394
+ #ifdef EIGEN_VECTORIZE_VSX
395
+ struct Packet1cd {
261
396
  EIGEN_STRONG_INLINE Packet1cd() {}
262
397
  EIGEN_STRONG_INLINE explicit Packet1cd(const Packet2d& a) : v(a) {}
398
+
399
+ EIGEN_STRONG_INLINE Packet1cd pmul(const Packet1cd& a, const Packet1cd& b) {
400
+ Packet2d a_re, a_im, v1, v2;
401
+
402
+ // Permute and multiply the real parts of a and b
403
+ a_re = vec_perm(a.v, a.v, p16uc_PSET64_HI);
404
+ // Get the imaginary parts of a
405
+ a_im = vec_perm(a.v, a.v, p16uc_PSET64_LO);
406
+ // multiply a_re * b
407
+ v1 = vec_madd(a_re, b.v, p2d_ZERO);
408
+ // multiply a_im * b and get the conjugate result
409
+ v2 = vec_madd(a_im, b.v, p2d_ZERO);
410
+ v2 = reinterpret_cast<Packet2d>(vec_sld(reinterpret_cast<Packet4ui>(v2), reinterpret_cast<Packet4ui>(v2), 8));
411
+ v2 = pxor(v2, reinterpret_cast<Packet2d>(p2ul_CONJ_XOR1()));
412
+
413
+ return Packet1cd(padd<Packet2d>(v1, v2));
414
+ }
415
+
416
+ EIGEN_STRONG_INLINE Packet1cd& operator*=(const Packet1cd& b) {
417
+ v = pmul(Packet1cd(*this), b).v;
418
+ return *this;
419
+ }
420
+ EIGEN_STRONG_INLINE Packet1cd operator*(const Packet1cd& b) const { return Packet1cd(*this) *= b; }
421
+
422
+ EIGEN_STRONG_INLINE Packet1cd& operator+=(const Packet1cd& b) {
423
+ v = padd(v, b.v);
424
+ return *this;
425
+ }
426
+ EIGEN_STRONG_INLINE Packet1cd operator+(const Packet1cd& b) const { return Packet1cd(*this) += b; }
427
+ EIGEN_STRONG_INLINE Packet1cd& operator-=(const Packet1cd& b) {
428
+ v = psub(v, b.v);
429
+ return *this;
430
+ }
431
+ EIGEN_STRONG_INLINE Packet1cd operator-(const Packet1cd& b) const { return Packet1cd(*this) -= b; }
432
+ EIGEN_STRONG_INLINE Packet1cd operator-(void) const { return Packet1cd(-v); }
433
+
263
434
  Packet2d v;
264
435
  };
265
436
 
266
- template<> struct packet_traits<std::complex<double> > : default_packet_traits
267
- {
437
+ template <>
438
+ struct packet_traits<std::complex<double> > : default_packet_traits {
268
439
  typedef Packet1cd type;
269
440
  typedef Packet1cd half;
441
+ typedef Packet2d as_real;
270
442
  enum {
271
443
  Vectorizable = 1,
272
444
  AlignedOnScalar = 0,
273
445
  size = 1,
274
- HasHalfPacket = 0,
275
446
 
276
- HasAdd = 1,
277
- HasSub = 1,
278
- HasMul = 1,
279
- HasDiv = 1,
447
+ HasAdd = 1,
448
+ HasSub = 1,
449
+ HasMul = 1,
450
+ HasDiv = 1,
280
451
  HasNegate = 1,
281
- HasAbs = 0,
282
- HasAbs2 = 0,
283
- HasMin = 0,
284
- HasMax = 0,
452
+ HasAbs = 0,
453
+ HasAbs2 = 0,
454
+ HasMin = 0,
455
+ HasMax = 0,
456
+ HasSqrt = 1,
457
+ HasLog = 1,
285
458
  HasSetLinear = 0
286
459
  };
287
460
  };
288
461
 
289
- template<> struct unpacket_traits<Packet1cd> { typedef std::complex<double> type; enum {size=1, alignment=Aligned16}; typedef Packet1cd half; };
290
-
291
- template<> EIGEN_STRONG_INLINE Packet1cd pload <Packet1cd>(const std::complex<double>* from) { return Packet1cd(pload<Packet2d>((const double*)from)); }
292
- template<> EIGEN_STRONG_INLINE Packet1cd ploadu<Packet1cd>(const std::complex<double>* from) { return Packet1cd(ploadu<Packet2d>((const double*)from)); }
293
- template<> EIGEN_STRONG_INLINE void pstore <std::complex<double> >(std::complex<double> * to, const Packet1cd& from) { pstore((double*)to, from.v); }
294
- template<> EIGEN_STRONG_INLINE void pstoreu<std::complex<double> >(std::complex<double> * to, const Packet1cd& from) { pstoreu((double*)to, from.v); }
295
-
296
- template<> EIGEN_STRONG_INLINE Packet1cd pset1<Packet1cd>(const std::complex<double>& from)
297
- { /* here we really have to use unaligned loads :( */ return ploadu<Packet1cd>(&from); }
462
+ template <>
463
+ struct unpacket_traits<Packet1cd> {
464
+ typedef std::complex<double> type;
465
+ enum {
466
+ size = 1,
467
+ alignment = Aligned16,
468
+ vectorizable = true,
469
+ masked_load_available = false,
470
+ masked_store_available = false
471
+ };
472
+ typedef Packet1cd half;
473
+ typedef Packet2d as_real;
474
+ };
298
475
 
299
- template<> EIGEN_DEVICE_FUNC inline Packet1cd pgather<std::complex<double>, Packet1cd>(const std::complex<double>* from, Index stride)
300
- {
301
- std::complex<double> EIGEN_ALIGN16 af[2];
302
- af[0] = from[0*stride];
303
- af[1] = from[1*stride];
304
- return pload<Packet1cd>(af);
476
+ template <>
477
+ EIGEN_STRONG_INLINE Packet1cd pload<Packet1cd>(const std::complex<double>* from) {
478
+ return Packet1cd(pload<Packet2d>((const double*)from));
305
479
  }
306
- template<> EIGEN_DEVICE_FUNC inline void pscatter<std::complex<double>, Packet1cd>(std::complex<double>* to, const Packet1cd& from, Index stride)
307
- {
308
- std::complex<double> EIGEN_ALIGN16 af[2];
309
- pstore<std::complex<double> >(af, from);
310
- to[0*stride] = af[0];
311
- to[1*stride] = af[1];
480
+ template <>
481
+ EIGEN_STRONG_INLINE Packet1cd ploadu<Packet1cd>(const std::complex<double>* from) {
482
+ return Packet1cd(ploadu<Packet2d>((const double*)from));
483
+ }
484
+ template <>
485
+ EIGEN_ALWAYS_INLINE Packet1cd pload_partial<Packet1cd>(const std::complex<double>* from, const Index n,
486
+ const Index offset) {
487
+ return Packet1cd(pload_partial<Packet2d>((const double*)from, n * 2, offset * 2));
488
+ }
489
+ template <>
490
+ EIGEN_ALWAYS_INLINE Packet1cd ploadu_partial<Packet1cd>(const std::complex<double>* from, const Index n,
491
+ const Index offset) {
492
+ return Packet1cd(ploadu_partial<Packet2d>((const double*)from, n * 2, offset * 2));
493
+ }
494
+ template <>
495
+ EIGEN_STRONG_INLINE void pstore<std::complex<double> >(std::complex<double>* to, const Packet1cd& from) {
496
+ pstore((double*)to, from.v);
497
+ }
498
+ template <>
499
+ EIGEN_STRONG_INLINE void pstoreu<std::complex<double> >(std::complex<double>* to, const Packet1cd& from) {
500
+ pstoreu((double*)to, from.v);
501
+ }
502
+ template <>
503
+ EIGEN_ALWAYS_INLINE void pstore_partial<std::complex<double> >(std::complex<double>* to, const Packet1cd& from,
504
+ const Index n, const Index offset) {
505
+ pstore_partial((double*)to, from.v, n * 2, offset * 2);
506
+ }
507
+ template <>
508
+ EIGEN_ALWAYS_INLINE void pstoreu_partial<std::complex<double> >(std::complex<double>* to, const Packet1cd& from,
509
+ const Index n, const Index offset) {
510
+ pstoreu_partial((double*)to, from.v, n * 2, offset * 2);
312
511
  }
313
512
 
314
- template<> EIGEN_STRONG_INLINE Packet1cd padd<Packet1cd>(const Packet1cd& a, const Packet1cd& b) { return Packet1cd(a.v + b.v); }
315
- template<> EIGEN_STRONG_INLINE Packet1cd psub<Packet1cd>(const Packet1cd& a, const Packet1cd& b) { return Packet1cd(a.v - b.v); }
316
- template<> EIGEN_STRONG_INLINE Packet1cd pnegate(const Packet1cd& a) { return Packet1cd(pnegate(Packet2d(a.v))); }
317
- template<> EIGEN_STRONG_INLINE Packet1cd pconj(const Packet1cd& a) { return Packet1cd(pxor(a.v, reinterpret_cast<Packet2d>(p2ul_CONJ_XOR2))); }
318
-
319
- template<> EIGEN_STRONG_INLINE Packet1cd pmul<Packet1cd>(const Packet1cd& a, const Packet1cd& b)
320
- {
321
- Packet2d a_re, a_im, v1, v2;
513
+ template <>
514
+ EIGEN_STRONG_INLINE Packet1cd
515
+ pset1<Packet1cd>(const std::complex<double>& from) { /* here we really have to use unaligned loads :( */
516
+ return ploadu<Packet1cd>(&from);
517
+ }
322
518
 
323
- // Permute and multiply the real parts of a and b
324
- a_re = vec_perm(a.v, a.v, p16uc_PSET64_HI);
325
- // Get the imaginary parts of a
326
- a_im = vec_perm(a.v, a.v, p16uc_PSET64_LO);
327
- // multiply a_re * b
328
- v1 = vec_madd(a_re, b.v, p2d_ZERO);
329
- // multiply a_im * b and get the conjugate result
330
- v2 = vec_madd(a_im, b.v, p2d_ZERO);
331
- v2 = reinterpret_cast<Packet2d>(vec_sld(reinterpret_cast<Packet4ui>(v2), reinterpret_cast<Packet4ui>(v2), 8));
332
- v2 = pxor(v2, reinterpret_cast<Packet2d>(p2ul_CONJ_XOR1));
519
+ template <>
520
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet1cd
521
+ pgather<std::complex<double>, Packet1cd>(const std::complex<double>* from, Index) {
522
+ return pload<Packet1cd>(from);
523
+ }
524
+ template <>
525
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet1cd
526
+ pgather_partial<std::complex<double>, Packet1cd>(const std::complex<double>* from, Index, const Index) {
527
+ return pload<Packet1cd>(from);
528
+ }
529
+ template <>
530
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter<std::complex<double>, Packet1cd>(std::complex<double>* to,
531
+ const Packet1cd& from, Index) {
532
+ pstore<std::complex<double> >(to, from);
533
+ }
534
+ template <>
535
+ EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter_partial<std::complex<double>, Packet1cd>(std::complex<double>* to,
536
+ const Packet1cd& from,
537
+ Index, const Index) {
538
+ pstore<std::complex<double> >(to, from);
539
+ }
333
540
 
334
- return Packet1cd(padd<Packet2d>(v1, v2));
541
+ template <>
542
+ EIGEN_STRONG_INLINE Packet1cd padd<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
543
+ return Packet1cd(a.v + b.v);
544
+ }
545
+ template <>
546
+ EIGEN_STRONG_INLINE Packet1cd psub<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
547
+ return Packet1cd(a.v - b.v);
548
+ }
549
+ template <>
550
+ EIGEN_STRONG_INLINE Packet1cd pnegate(const Packet1cd& a) {
551
+ return Packet1cd(pnegate(Packet2d(a.v)));
552
+ }
553
+ template <>
554
+ EIGEN_STRONG_INLINE Packet1cd pconj(const Packet1cd& a) {
555
+ return Packet1cd(pxor(a.v, reinterpret_cast<Packet2d>(p2ul_CONJ_XOR2())));
335
556
  }
336
557
 
337
- template<> EIGEN_STRONG_INLINE Packet1cd pand <Packet1cd>(const Packet1cd& a, const Packet1cd& b) { return Packet1cd(pand(a.v,b.v)); }
338
- template<> EIGEN_STRONG_INLINE Packet1cd por <Packet1cd>(const Packet1cd& a, const Packet1cd& b) { return Packet1cd(por(a.v,b.v)); }
339
- template<> EIGEN_STRONG_INLINE Packet1cd pxor <Packet1cd>(const Packet1cd& a, const Packet1cd& b) { return Packet1cd(pxor(a.v,b.v)); }
340
- template<> EIGEN_STRONG_INLINE Packet1cd pandnot<Packet1cd>(const Packet1cd& a, const Packet1cd& b) { return Packet1cd(pandnot(a.v, b.v)); }
558
+ template <>
559
+ EIGEN_STRONG_INLINE Packet1cd pand<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
560
+ return Packet1cd(pand(a.v, b.v));
561
+ }
562
+ template <>
563
+ EIGEN_STRONG_INLINE Packet1cd por<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
564
+ return Packet1cd(por(a.v, b.v));
565
+ }
566
+ template <>
567
+ EIGEN_STRONG_INLINE Packet1cd pxor<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
568
+ return Packet1cd(pxor(a.v, b.v));
569
+ }
570
+ template <>
571
+ EIGEN_STRONG_INLINE Packet1cd pandnot<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
572
+ return Packet1cd(pandnot(a.v, b.v));
573
+ }
341
574
 
342
- template<> EIGEN_STRONG_INLINE Packet1cd ploaddup<Packet1cd>(const std::complex<double>* from) { return pset1<Packet1cd>(*from); }
575
+ template <>
576
+ EIGEN_STRONG_INLINE Packet1cd ploaddup<Packet1cd>(const std::complex<double>* from) {
577
+ return pset1<Packet1cd>(*from);
578
+ }
343
579
 
344
- template<> EIGEN_STRONG_INLINE void prefetch<std::complex<double> >(const std::complex<double> * addr) { EIGEN_PPC_PREFETCH(addr); }
580
+ template <>
581
+ EIGEN_STRONG_INLINE void prefetch<std::complex<double> >(const std::complex<double>* addr) {
582
+ EIGEN_PPC_PREFETCH(addr);
583
+ }
345
584
 
346
- template<> EIGEN_STRONG_INLINE std::complex<double> pfirst<Packet1cd>(const Packet1cd& a)
347
- {
348
- std::complex<double> EIGEN_ALIGN16 res[2];
585
+ template <>
586
+ EIGEN_STRONG_INLINE std::complex<double> pfirst<Packet1cd>(const Packet1cd& a) {
587
+ EIGEN_ALIGN16 std::complex<double> res[1];
349
588
  pstore<std::complex<double> >(res, a);
350
589
 
351
590
  return res[0];
352
591
  }
353
592
 
354
- template<> EIGEN_STRONG_INLINE Packet1cd preverse(const Packet1cd& a) { return a; }
355
-
356
- template<> EIGEN_STRONG_INLINE std::complex<double> predux<Packet1cd>(const Packet1cd& a) { return pfirst(a); }
357
- template<> EIGEN_STRONG_INLINE Packet1cd preduxp<Packet1cd>(const Packet1cd* vecs) { return vecs[0]; }
358
-
359
- template<> EIGEN_STRONG_INLINE std::complex<double> predux_mul<Packet1cd>(const Packet1cd& a) { return pfirst(a); }
360
-
361
- template<int Offset>
362
- struct palign_impl<Offset,Packet1cd>
363
- {
364
- static EIGEN_STRONG_INLINE void run(Packet1cd& /*first*/, const Packet1cd& /*second*/)
365
- {
366
- // FIXME is it sure we never have to align a Packet1cd?
367
- // Even though a std::complex<double> has 16 bytes, it is not necessarily aligned on a 16 bytes boundary...
368
- }
369
- };
370
-
371
- template<> struct conj_helper<Packet1cd, Packet1cd, false,true>
372
- {
373
- EIGEN_STRONG_INLINE Packet1cd pmadd(const Packet1cd& x, const Packet1cd& y, const Packet1cd& c) const
374
- { return padd(pmul(x,y),c); }
593
+ template <>
594
+ EIGEN_STRONG_INLINE Packet1cd preverse(const Packet1cd& a) {
595
+ return a;
596
+ }
375
597
 
376
- EIGEN_STRONG_INLINE Packet1cd pmul(const Packet1cd& a, const Packet1cd& b) const
377
- {
378
- return internal::pmul(a, pconj(b));
379
- }
380
- };
598
+ template <>
599
+ EIGEN_STRONG_INLINE std::complex<double> predux<Packet1cd>(const Packet1cd& a) {
600
+ return pfirst(a);
601
+ }
381
602
 
382
- template<> struct conj_helper<Packet1cd, Packet1cd, true,false>
383
- {
384
- EIGEN_STRONG_INLINE Packet1cd pmadd(const Packet1cd& x, const Packet1cd& y, const Packet1cd& c) const
385
- { return padd(pmul(x,y),c); }
603
+ template <>
604
+ EIGEN_STRONG_INLINE std::complex<double> predux_mul<Packet1cd>(const Packet1cd& a) {
605
+ return pfirst(a);
606
+ }
386
607
 
387
- EIGEN_STRONG_INLINE Packet1cd pmul(const Packet1cd& a, const Packet1cd& b) const
388
- {
389
- return internal::pmul(pconj(a), b);
390
- }
391
- };
608
+ EIGEN_MAKE_CONJ_HELPER_CPLX_REAL(Packet1cd, Packet2d)
392
609
 
393
- template<> struct conj_helper<Packet1cd, Packet1cd, true,true>
394
- {
395
- EIGEN_STRONG_INLINE Packet1cd pmadd(const Packet1cd& x, const Packet1cd& y, const Packet1cd& c) const
396
- { return padd(pmul(x,y),c); }
610
+ template <>
611
+ EIGEN_STRONG_INLINE Packet1cd pdiv<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
612
+ return pdiv_complex(a, b);
613
+ }
397
614
 
398
- EIGEN_STRONG_INLINE Packet1cd pmul(const Packet1cd& a, const Packet1cd& b) const
399
- {
400
- return pconj(internal::pmul(a, b));
401
- }
402
- };
615
+ EIGEN_STRONG_INLINE Packet1cd pcplxflip /*<Packet1cd>*/ (const Packet1cd& x) {
616
+ return Packet1cd(preverse(Packet2d(x.v)));
617
+ }
403
618
 
404
- EIGEN_MAKE_CONJ_HELPER_CPLX_REAL(Packet1cd,Packet2d)
619
+ EIGEN_STRONG_INLINE void ptranspose(PacketBlock<Packet1cd, 2>& kernel) {
620
+ Packet2d tmp = vec_mergeh(kernel.packet[0].v, kernel.packet[1].v);
621
+ kernel.packet[1].v = vec_mergel(kernel.packet[0].v, kernel.packet[1].v);
622
+ kernel.packet[0].v = tmp;
623
+ }
405
624
 
406
- template<> EIGEN_STRONG_INLINE Packet1cd pdiv<Packet1cd>(const Packet1cd& a, const Packet1cd& b)
407
- {
408
- // TODO optimize it for AltiVec
409
- Packet1cd res = conj_helper<Packet1cd,Packet1cd,false,true>().pmul(a,b);
410
- Packet2d s = pmul<Packet2d>(b.v, b.v);
411
- return Packet1cd(pdiv(res.v, padd<Packet2d>(s, vec_perm(s, s, p16uc_REVERSE64))));
625
+ template <>
626
+ EIGEN_STRONG_INLINE Packet1cd pcmp_eq(const Packet1cd& a, const Packet1cd& b) {
627
+ // Compare real and imaginary parts of a and b to get the mask vector:
628
+ // [re(a)==re(b), im(a)==im(b)]
629
+ Packet2d eq = reinterpret_cast<Packet2d>(vec_cmpeq(a.v, b.v));
630
+ // Swap real/imag elements in the mask in to get:
631
+ // [im(a)==im(b), re(a)==re(b)]
632
+ Packet2d eq_swapped =
633
+ reinterpret_cast<Packet2d>(vec_sld(reinterpret_cast<Packet4ui>(eq), reinterpret_cast<Packet4ui>(eq), 8));
634
+ // Return re(a)==re(b) & im(a)==im(b) by computing bitwise AND of eq and eq_swapped
635
+ return Packet1cd(vec_and(eq, eq_swapped));
412
636
  }
413
637
 
414
- EIGEN_STRONG_INLINE Packet1cd pcplxflip/*<Packet1cd>*/(const Packet1cd& x)
415
- {
416
- return Packet1cd(preverse(Packet2d(x.v)));
638
+ template <>
639
+ EIGEN_STRONG_INLINE Packet1cd psqrt<Packet1cd>(const Packet1cd& a) {
640
+ return psqrt_complex<Packet1cd>(a);
417
641
  }
418
642
 
419
- EIGEN_STRONG_INLINE void ptranspose(PacketBlock<Packet1cd,2>& kernel)
420
- {
421
- Packet2d tmp = vec_perm(kernel.packet[0].v, kernel.packet[1].v, p16uc_TRANSPOSE64_HI);
422
- kernel.packet[1].v = vec_perm(kernel.packet[0].v, kernel.packet[1].v, p16uc_TRANSPOSE64_LO);
423
- kernel.packet[0].v = tmp;
643
+ template <>
644
+ EIGEN_STRONG_INLINE Packet1cd plog<Packet1cd>(const Packet1cd& a) {
645
+ return plog_complex<Packet1cd>(a);
424
646
  }
425
- #endif // __VSX__
426
- } // end namespace internal
427
647
 
428
- } // end namespace Eigen
648
+ #endif // __VSX__
649
+ } // end namespace internal
650
+
651
+ } // end namespace Eigen
429
652
 
430
- #endif // EIGEN_COMPLEX32_ALTIVEC_H
653
+ #endif // EIGEN_COMPLEX32_ALTIVEC_H