@smake/eigen 1.0.2 → 1.1.1
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +1 -1
- package/eigen/Eigen/AccelerateSupport +52 -0
- package/eigen/Eigen/Cholesky +18 -21
- package/eigen/Eigen/CholmodSupport +28 -28
- package/eigen/Eigen/Core +235 -326
- package/eigen/Eigen/Eigenvalues +16 -14
- package/eigen/Eigen/Geometry +21 -24
- package/eigen/Eigen/Householder +9 -8
- package/eigen/Eigen/IterativeLinearSolvers +8 -4
- package/eigen/Eigen/Jacobi +14 -14
- package/eigen/Eigen/KLUSupport +43 -0
- package/eigen/Eigen/LU +16 -20
- package/eigen/Eigen/MetisSupport +12 -12
- package/eigen/Eigen/OrderingMethods +54 -54
- package/eigen/Eigen/PaStiXSupport +23 -20
- package/eigen/Eigen/PardisoSupport +17 -14
- package/eigen/Eigen/QR +18 -21
- package/eigen/Eigen/QtAlignedMalloc +5 -13
- package/eigen/Eigen/SPQRSupport +21 -14
- package/eigen/Eigen/SVD +23 -18
- package/eigen/Eigen/Sparse +1 -4
- package/eigen/Eigen/SparseCholesky +18 -23
- package/eigen/Eigen/SparseCore +18 -17
- package/eigen/Eigen/SparseLU +12 -8
- package/eigen/Eigen/SparseQR +16 -14
- package/eigen/Eigen/StdDeque +5 -2
- package/eigen/Eigen/StdList +5 -2
- package/eigen/Eigen/StdVector +5 -2
- package/eigen/Eigen/SuperLUSupport +30 -24
- package/eigen/Eigen/ThreadPool +80 -0
- package/eigen/Eigen/UmfPackSupport +19 -17
- package/eigen/Eigen/Version +14 -0
- package/eigen/Eigen/src/AccelerateSupport/AccelerateSupport.h +423 -0
- package/eigen/Eigen/src/AccelerateSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/Cholesky/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/Cholesky/LDLT.h +377 -401
- package/eigen/Eigen/src/Cholesky/LLT.h +332 -360
- package/eigen/Eigen/src/Cholesky/LLT_LAPACKE.h +81 -56
- package/eigen/Eigen/src/CholmodSupport/CholmodSupport.h +620 -521
- package/eigen/Eigen/src/CholmodSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/Core/ArithmeticSequence.h +239 -0
- package/eigen/Eigen/src/Core/Array.h +341 -294
- package/eigen/Eigen/src/Core/ArrayBase.h +190 -203
- package/eigen/Eigen/src/Core/ArrayWrapper.h +127 -171
- package/eigen/Eigen/src/Core/Assign.h +30 -40
- package/eigen/Eigen/src/Core/AssignEvaluator.h +711 -589
- package/eigen/Eigen/src/Core/Assign_MKL.h +130 -125
- package/eigen/Eigen/src/Core/BandMatrix.h +268 -283
- package/eigen/Eigen/src/Core/Block.h +375 -398
- package/eigen/Eigen/src/Core/CommaInitializer.h +86 -97
- package/eigen/Eigen/src/Core/ConditionEstimator.h +51 -53
- package/eigen/Eigen/src/Core/CoreEvaluators.h +1356 -1026
- package/eigen/Eigen/src/Core/CoreIterators.h +73 -59
- package/eigen/Eigen/src/Core/CwiseBinaryOp.h +114 -132
- package/eigen/Eigen/src/Core/CwiseNullaryOp.h +726 -617
- package/eigen/Eigen/src/Core/CwiseTernaryOp.h +77 -103
- package/eigen/Eigen/src/Core/CwiseUnaryOp.h +56 -68
- package/eigen/Eigen/src/Core/CwiseUnaryView.h +132 -95
- package/eigen/Eigen/src/Core/DenseBase.h +632 -571
- package/eigen/Eigen/src/Core/DenseCoeffsBase.h +511 -624
- package/eigen/Eigen/src/Core/DenseStorage.h +512 -509
- package/eigen/Eigen/src/Core/DeviceWrapper.h +153 -0
- package/eigen/Eigen/src/Core/Diagonal.h +169 -210
- package/eigen/Eigen/src/Core/DiagonalMatrix.h +351 -274
- package/eigen/Eigen/src/Core/DiagonalProduct.h +12 -10
- package/eigen/Eigen/src/Core/Dot.h +172 -222
- package/eigen/Eigen/src/Core/EigenBase.h +75 -85
- package/eigen/Eigen/src/Core/Fill.h +138 -0
- package/eigen/Eigen/src/Core/FindCoeff.h +464 -0
- package/eigen/Eigen/src/Core/ForceAlignedAccess.h +90 -109
- package/eigen/Eigen/src/Core/Fuzzy.h +82 -105
- package/eigen/Eigen/src/Core/GeneralProduct.h +327 -263
- package/eigen/Eigen/src/Core/GenericPacketMath.h +1472 -360
- package/eigen/Eigen/src/Core/GlobalFunctions.h +194 -151
- package/eigen/Eigen/src/Core/IO.h +147 -139
- package/eigen/Eigen/src/Core/IndexedView.h +321 -0
- package/eigen/Eigen/src/Core/InnerProduct.h +260 -0
- package/eigen/Eigen/src/Core/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/Core/Inverse.h +56 -66
- package/eigen/Eigen/src/Core/Map.h +124 -142
- package/eigen/Eigen/src/Core/MapBase.h +256 -281
- package/eigen/Eigen/src/Core/MathFunctions.h +1620 -938
- package/eigen/Eigen/src/Core/MathFunctionsImpl.h +233 -71
- package/eigen/Eigen/src/Core/Matrix.h +491 -416
- package/eigen/Eigen/src/Core/MatrixBase.h +468 -453
- package/eigen/Eigen/src/Core/NestByValue.h +66 -85
- package/eigen/Eigen/src/Core/NoAlias.h +79 -85
- package/eigen/Eigen/src/Core/NumTraits.h +235 -148
- package/eigen/Eigen/src/Core/PartialReduxEvaluator.h +253 -0
- package/eigen/Eigen/src/Core/PermutationMatrix.h +461 -511
- package/eigen/Eigen/src/Core/PlainObjectBase.h +871 -894
- package/eigen/Eigen/src/Core/Product.h +260 -139
- package/eigen/Eigen/src/Core/ProductEvaluators.h +863 -714
- package/eigen/Eigen/src/Core/Random.h +161 -136
- package/eigen/Eigen/src/Core/RandomImpl.h +262 -0
- package/eigen/Eigen/src/Core/RealView.h +250 -0
- package/eigen/Eigen/src/Core/Redux.h +366 -336
- package/eigen/Eigen/src/Core/Ref.h +308 -209
- package/eigen/Eigen/src/Core/Replicate.h +94 -106
- package/eigen/Eigen/src/Core/Reshaped.h +398 -0
- package/eigen/Eigen/src/Core/ReturnByValue.h +49 -55
- package/eigen/Eigen/src/Core/Reverse.h +136 -145
- package/eigen/Eigen/src/Core/Select.h +70 -140
- package/eigen/Eigen/src/Core/SelfAdjointView.h +262 -285
- package/eigen/Eigen/src/Core/SelfCwiseBinaryOp.h +23 -20
- package/eigen/Eigen/src/Core/SkewSymmetricMatrix3.h +382 -0
- package/eigen/Eigen/src/Core/Solve.h +97 -111
- package/eigen/Eigen/src/Core/SolveTriangular.h +131 -129
- package/eigen/Eigen/src/Core/SolverBase.h +138 -101
- package/eigen/Eigen/src/Core/StableNorm.h +156 -160
- package/eigen/Eigen/src/Core/StlIterators.h +619 -0
- package/eigen/Eigen/src/Core/Stride.h +91 -88
- package/eigen/Eigen/src/Core/Swap.h +70 -38
- package/eigen/Eigen/src/Core/Transpose.h +295 -273
- package/eigen/Eigen/src/Core/Transpositions.h +272 -317
- package/eigen/Eigen/src/Core/TriangularMatrix.h +670 -755
- package/eigen/Eigen/src/Core/VectorBlock.h +59 -72
- package/eigen/Eigen/src/Core/VectorwiseOp.h +668 -630
- package/eigen/Eigen/src/Core/Visitor.h +480 -216
- package/eigen/Eigen/src/Core/arch/AVX/Complex.h +407 -293
- package/eigen/Eigen/src/Core/arch/AVX/MathFunctions.h +79 -388
- package/eigen/Eigen/src/Core/arch/AVX/PacketMath.h +2935 -491
- package/eigen/Eigen/src/Core/arch/AVX/Reductions.h +353 -0
- package/eigen/Eigen/src/Core/arch/AVX/TypeCasting.h +279 -22
- package/eigen/Eigen/src/Core/arch/AVX512/Complex.h +472 -0
- package/eigen/Eigen/src/Core/arch/AVX512/GemmKernel.h +1245 -0
- package/eigen/Eigen/src/Core/arch/AVX512/MathFunctions.h +85 -333
- package/eigen/Eigen/src/Core/arch/AVX512/MathFunctionsFP16.h +75 -0
- package/eigen/Eigen/src/Core/arch/AVX512/PacketMath.h +2490 -649
- package/eigen/Eigen/src/Core/arch/AVX512/PacketMathFP16.h +1413 -0
- package/eigen/Eigen/src/Core/arch/AVX512/Reductions.h +297 -0
- package/eigen/Eigen/src/Core/arch/AVX512/TrsmKernel.h +1167 -0
- package/eigen/Eigen/src/Core/arch/AVX512/TrsmUnrolls.inc +1219 -0
- package/eigen/Eigen/src/Core/arch/AVX512/TypeCasting.h +277 -0
- package/eigen/Eigen/src/Core/arch/AVX512/TypeCastingFP16.h +130 -0
- package/eigen/Eigen/src/Core/arch/AltiVec/Complex.h +521 -298
- package/eigen/Eigen/src/Core/arch/AltiVec/MathFunctions.h +39 -280
- package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProduct.h +3686 -0
- package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProductCommon.h +205 -0
- package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProductMMA.h +901 -0
- package/eigen/Eigen/src/Core/arch/AltiVec/MatrixProductMMAbfloat16.h +742 -0
- package/eigen/Eigen/src/Core/arch/AltiVec/MatrixVectorProduct.inc +2818 -0
- package/eigen/Eigen/src/Core/arch/AltiVec/PacketMath.h +3391 -723
- package/eigen/Eigen/src/Core/arch/AltiVec/TypeCasting.h +153 -0
- package/eigen/Eigen/src/Core/arch/Default/BFloat16.h +866 -0
- package/eigen/Eigen/src/Core/arch/Default/ConjHelper.h +113 -14
- package/eigen/Eigen/src/Core/arch/Default/GenericPacketMathFunctions.h +2634 -0
- package/eigen/Eigen/src/Core/arch/Default/GenericPacketMathFunctionsFwd.h +227 -0
- package/eigen/Eigen/src/Core/arch/Default/Half.h +1091 -0
- package/eigen/Eigen/src/Core/arch/Default/Settings.h +11 -13
- package/eigen/Eigen/src/Core/arch/GPU/Complex.h +244 -0
- package/eigen/Eigen/src/Core/arch/GPU/MathFunctions.h +104 -0
- package/eigen/Eigen/src/Core/arch/GPU/PacketMath.h +1712 -0
- package/eigen/Eigen/src/Core/arch/GPU/Tuple.h +268 -0
- package/eigen/Eigen/src/Core/arch/GPU/TypeCasting.h +77 -0
- package/eigen/Eigen/src/Core/arch/HIP/hcc/math_constants.h +23 -0
- package/eigen/Eigen/src/Core/arch/HVX/PacketMath.h +1088 -0
- package/eigen/Eigen/src/Core/arch/LSX/Complex.h +520 -0
- package/eigen/Eigen/src/Core/arch/LSX/GeneralBlockPanelKernel.h +23 -0
- package/eigen/Eigen/src/Core/arch/LSX/MathFunctions.h +43 -0
- package/eigen/Eigen/src/Core/arch/LSX/PacketMath.h +2866 -0
- package/eigen/Eigen/src/Core/arch/LSX/TypeCasting.h +526 -0
- package/eigen/Eigen/src/Core/arch/MSA/Complex.h +620 -0
- package/eigen/Eigen/src/Core/arch/MSA/MathFunctions.h +379 -0
- package/eigen/Eigen/src/Core/arch/MSA/PacketMath.h +1237 -0
- package/eigen/Eigen/src/Core/arch/NEON/Complex.h +531 -289
- package/eigen/Eigen/src/Core/arch/NEON/GeneralBlockPanelKernel.h +243 -0
- package/eigen/Eigen/src/Core/arch/NEON/MathFunctions.h +50 -73
- package/eigen/Eigen/src/Core/arch/NEON/PacketMath.h +5915 -579
- package/eigen/Eigen/src/Core/arch/NEON/TypeCasting.h +1642 -0
- package/eigen/Eigen/src/Core/arch/NEON/UnaryFunctors.h +57 -0
- package/eigen/Eigen/src/Core/arch/SSE/Complex.h +366 -334
- package/eigen/Eigen/src/Core/arch/SSE/MathFunctions.h +40 -514
- package/eigen/Eigen/src/Core/arch/SSE/PacketMath.h +2164 -675
- package/eigen/Eigen/src/Core/arch/SSE/Reductions.h +324 -0
- package/eigen/Eigen/src/Core/arch/SSE/TypeCasting.h +188 -35
- package/eigen/Eigen/src/Core/arch/SVE/MathFunctions.h +48 -0
- package/eigen/Eigen/src/Core/arch/SVE/PacketMath.h +674 -0
- package/eigen/Eigen/src/Core/arch/SVE/TypeCasting.h +52 -0
- package/eigen/Eigen/src/Core/arch/SYCL/InteropHeaders.h +227 -0
- package/eigen/Eigen/src/Core/arch/SYCL/MathFunctions.h +303 -0
- package/eigen/Eigen/src/Core/arch/SYCL/PacketMath.h +576 -0
- package/eigen/Eigen/src/Core/arch/SYCL/TypeCasting.h +83 -0
- package/eigen/Eigen/src/Core/arch/ZVector/Complex.h +434 -261
- package/eigen/Eigen/src/Core/arch/ZVector/MathFunctions.h +160 -53
- package/eigen/Eigen/src/Core/arch/ZVector/PacketMath.h +1073 -605
- package/eigen/Eigen/src/Core/functors/AssignmentFunctors.h +123 -117
- package/eigen/Eigen/src/Core/functors/BinaryFunctors.h +594 -322
- package/eigen/Eigen/src/Core/functors/NullaryFunctors.h +204 -118
- package/eigen/Eigen/src/Core/functors/StlFunctors.h +110 -97
- package/eigen/Eigen/src/Core/functors/TernaryFunctors.h +34 -7
- package/eigen/Eigen/src/Core/functors/UnaryFunctors.h +1158 -530
- package/eigen/Eigen/src/Core/products/GeneralBlockPanelKernel.h +2329 -1333
- package/eigen/Eigen/src/Core/products/GeneralMatrixMatrix.h +328 -364
- package/eigen/Eigen/src/Core/products/GeneralMatrixMatrixTriangular.h +191 -178
- package/eigen/Eigen/src/Core/products/GeneralMatrixMatrixTriangular_BLAS.h +85 -82
- package/eigen/Eigen/src/Core/products/GeneralMatrixMatrix_BLAS.h +154 -73
- package/eigen/Eigen/src/Core/products/GeneralMatrixVector.h +396 -542
- package/eigen/Eigen/src/Core/products/GeneralMatrixVector_BLAS.h +80 -77
- package/eigen/Eigen/src/Core/products/Parallelizer.h +208 -92
- package/eigen/Eigen/src/Core/products/SelfadjointMatrixMatrix.h +331 -375
- package/eigen/Eigen/src/Core/products/SelfadjointMatrixMatrix_BLAS.h +206 -224
- package/eigen/Eigen/src/Core/products/SelfadjointMatrixVector.h +139 -146
- package/eigen/Eigen/src/Core/products/SelfadjointMatrixVector_BLAS.h +58 -61
- package/eigen/Eigen/src/Core/products/SelfadjointProduct.h +71 -71
- package/eigen/Eigen/src/Core/products/SelfadjointRank2Update.h +48 -46
- package/eigen/Eigen/src/Core/products/TriangularMatrixMatrix.h +294 -369
- package/eigen/Eigen/src/Core/products/TriangularMatrixMatrix_BLAS.h +246 -238
- package/eigen/Eigen/src/Core/products/TriangularMatrixVector.h +244 -247
- package/eigen/Eigen/src/Core/products/TriangularMatrixVector_BLAS.h +212 -192
- package/eigen/Eigen/src/Core/products/TriangularSolverMatrix.h +328 -275
- package/eigen/Eigen/src/Core/products/TriangularSolverMatrix_BLAS.h +108 -109
- package/eigen/Eigen/src/Core/products/TriangularSolverVector.h +70 -93
- package/eigen/Eigen/src/Core/util/Assert.h +158 -0
- package/eigen/Eigen/src/Core/util/BlasUtil.h +413 -290
- package/eigen/Eigen/src/Core/util/ConfigureVectorization.h +543 -0
- package/eigen/Eigen/src/Core/util/Constants.h +314 -263
- package/eigen/Eigen/src/Core/util/DisableStupidWarnings.h +130 -78
- package/eigen/Eigen/src/Core/util/EmulateArray.h +270 -0
- package/eigen/Eigen/src/Core/util/ForwardDeclarations.h +450 -224
- package/eigen/Eigen/src/Core/util/GpuHipCudaDefines.inc +101 -0
- package/eigen/Eigen/src/Core/util/GpuHipCudaUndefines.inc +45 -0
- package/eigen/Eigen/src/Core/util/IndexedViewHelper.h +487 -0
- package/eigen/Eigen/src/Core/util/IntegralConstant.h +279 -0
- package/eigen/Eigen/src/Core/util/MKL_support.h +39 -30
- package/eigen/Eigen/src/Core/util/Macros.h +939 -646
- package/eigen/Eigen/src/Core/util/MaxSizeVector.h +139 -0
- package/eigen/Eigen/src/Core/util/Memory.h +1042 -650
- package/eigen/Eigen/src/Core/util/Meta.h +618 -426
- package/eigen/Eigen/src/Core/util/MoreMeta.h +638 -0
- package/eigen/Eigen/src/Core/util/ReenableStupidWarnings.h +32 -19
- package/eigen/Eigen/src/Core/util/ReshapedHelper.h +51 -0
- package/eigen/Eigen/src/Core/util/Serializer.h +209 -0
- package/eigen/Eigen/src/Core/util/StaticAssert.h +51 -164
- package/eigen/Eigen/src/Core/util/SymbolicIndex.h +445 -0
- package/eigen/Eigen/src/Core/util/XprHelper.h +793 -538
- package/eigen/Eigen/src/Eigenvalues/ComplexEigenSolver.h +246 -277
- package/eigen/Eigen/src/Eigenvalues/ComplexSchur.h +299 -319
- package/eigen/Eigen/src/Eigenvalues/ComplexSchur_LAPACKE.h +52 -48
- package/eigen/Eigen/src/Eigenvalues/EigenSolver.h +413 -456
- package/eigen/Eigen/src/Eigenvalues/GeneralizedEigenSolver.h +309 -325
- package/eigen/Eigen/src/Eigenvalues/GeneralizedSelfAdjointEigenSolver.h +157 -171
- package/eigen/Eigen/src/Eigenvalues/HessenbergDecomposition.h +292 -310
- package/eigen/Eigen/src/Eigenvalues/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/Eigenvalues/MatrixBaseEigenvalues.h +91 -107
- package/eigen/Eigen/src/Eigenvalues/RealQZ.h +539 -606
- package/eigen/Eigen/src/Eigenvalues/RealSchur.h +348 -382
- package/eigen/Eigen/src/Eigenvalues/RealSchur_LAPACKE.h +41 -35
- package/eigen/Eigen/src/Eigenvalues/SelfAdjointEigenSolver.h +579 -600
- package/eigen/Eigen/src/Eigenvalues/SelfAdjointEigenSolver_LAPACKE.h +47 -44
- package/eigen/Eigen/src/Eigenvalues/Tridiagonalization.h +434 -461
- package/eigen/Eigen/src/Geometry/AlignedBox.h +307 -214
- package/eigen/Eigen/src/Geometry/AngleAxis.h +135 -137
- package/eigen/Eigen/src/Geometry/EulerAngles.h +163 -74
- package/eigen/Eigen/src/Geometry/Homogeneous.h +289 -333
- package/eigen/Eigen/src/Geometry/Hyperplane.h +152 -161
- package/eigen/Eigen/src/Geometry/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/Geometry/OrthoMethods.h +168 -145
- package/eigen/Eigen/src/Geometry/ParametrizedLine.h +141 -104
- package/eigen/Eigen/src/Geometry/Quaternion.h +595 -497
- package/eigen/Eigen/src/Geometry/Rotation2D.h +110 -108
- package/eigen/Eigen/src/Geometry/RotationBase.h +148 -145
- package/eigen/Eigen/src/Geometry/Scaling.h +115 -90
- package/eigen/Eigen/src/Geometry/Transform.h +896 -953
- package/eigen/Eigen/src/Geometry/Translation.h +100 -98
- package/eigen/Eigen/src/Geometry/Umeyama.h +79 -84
- package/eigen/Eigen/src/Geometry/arch/Geometry_SIMD.h +154 -0
- package/eigen/Eigen/src/Householder/BlockHouseholder.h +54 -42
- package/eigen/Eigen/src/Householder/Householder.h +104 -122
- package/eigen/Eigen/src/Householder/HouseholderSequence.h +416 -382
- package/eigen/Eigen/src/Householder/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/IterativeLinearSolvers/BasicPreconditioners.h +153 -166
- package/eigen/Eigen/src/IterativeLinearSolvers/BiCGSTAB.h +127 -138
- package/eigen/Eigen/src/IterativeLinearSolvers/ConjugateGradient.h +95 -124
- package/eigen/Eigen/src/IterativeLinearSolvers/IncompleteCholesky.h +269 -267
- package/eigen/Eigen/src/IterativeLinearSolvers/IncompleteLUT.h +246 -259
- package/eigen/Eigen/src/IterativeLinearSolvers/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/IterativeLinearSolvers/IterativeSolverBase.h +218 -217
- package/eigen/Eigen/src/IterativeLinearSolvers/LeastSquareConjugateGradient.h +80 -103
- package/eigen/Eigen/src/IterativeLinearSolvers/SolveWithGuess.h +59 -63
- package/eigen/Eigen/src/Jacobi/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/Jacobi/Jacobi.h +256 -291
- package/eigen/Eigen/src/KLUSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/KLUSupport/KLUSupport.h +339 -0
- package/eigen/Eigen/src/LU/Determinant.h +60 -63
- package/eigen/Eigen/src/LU/FullPivLU.h +561 -626
- package/eigen/Eigen/src/LU/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/LU/InverseImpl.h +213 -275
- package/eigen/Eigen/src/LU/PartialPivLU.h +407 -435
- package/eigen/Eigen/src/LU/PartialPivLU_LAPACKE.h +54 -40
- package/eigen/Eigen/src/LU/arch/InverseSize4.h +353 -0
- package/eigen/Eigen/src/MetisSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/MetisSupport/MetisSupport.h +81 -93
- package/eigen/Eigen/src/OrderingMethods/Amd.h +250 -282
- package/eigen/Eigen/src/OrderingMethods/Eigen_Colamd.h +950 -1103
- package/eigen/Eigen/src/OrderingMethods/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/OrderingMethods/Ordering.h +111 -122
- package/eigen/Eigen/src/PaStiXSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/PaStiXSupport/PaStiXSupport.h +524 -570
- package/eigen/Eigen/src/PardisoSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/PardisoSupport/PardisoSupport.h +385 -429
- package/eigen/Eigen/src/QR/ColPivHouseholderQR.h +494 -473
- package/eigen/Eigen/src/QR/ColPivHouseholderQR_LAPACKE.h +120 -56
- package/eigen/Eigen/src/QR/CompleteOrthogonalDecomposition.h +223 -137
- package/eigen/Eigen/src/QR/FullPivHouseholderQR.h +517 -460
- package/eigen/Eigen/src/QR/HouseholderQR.h +412 -278
- package/eigen/Eigen/src/QR/HouseholderQR_LAPACKE.h +32 -23
- package/eigen/Eigen/src/QR/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SPQRSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SPQRSupport/SuiteSparseQRSupport.h +263 -261
- package/eigen/Eigen/src/SVD/BDCSVD.h +872 -679
- package/eigen/Eigen/src/SVD/BDCSVD_LAPACKE.h +174 -0
- package/eigen/Eigen/src/SVD/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SVD/JacobiSVD.h +585 -543
- package/eigen/Eigen/src/SVD/JacobiSVD_LAPACKE.h +85 -49
- package/eigen/Eigen/src/SVD/SVDBase.h +281 -160
- package/eigen/Eigen/src/SVD/UpperBidiagonalization.h +202 -237
- package/eigen/Eigen/src/SparseCholesky/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SparseCholesky/SimplicialCholesky.h +769 -590
- package/eigen/Eigen/src/SparseCholesky/SimplicialCholesky_impl.h +318 -129
- package/eigen/Eigen/src/SparseCore/AmbiVector.h +202 -251
- package/eigen/Eigen/src/SparseCore/CompressedStorage.h +184 -236
- package/eigen/Eigen/src/SparseCore/ConservativeSparseSparseProduct.h +140 -184
- package/eigen/Eigen/src/SparseCore/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SparseCore/SparseAssign.h +174 -111
- package/eigen/Eigen/src/SparseCore/SparseBlock.h +408 -477
- package/eigen/Eigen/src/SparseCore/SparseColEtree.h +100 -112
- package/eigen/Eigen/src/SparseCore/SparseCompressedBase.h +531 -280
- package/eigen/Eigen/src/SparseCore/SparseCwiseBinaryOp.h +559 -347
- package/eigen/Eigen/src/SparseCore/SparseCwiseUnaryOp.h +100 -108
- package/eigen/Eigen/src/SparseCore/SparseDenseProduct.h +185 -191
- package/eigen/Eigen/src/SparseCore/SparseDiagonalProduct.h +71 -71
- package/eigen/Eigen/src/SparseCore/SparseDot.h +49 -47
- package/eigen/Eigen/src/SparseCore/SparseFuzzy.h +13 -11
- package/eigen/Eigen/src/SparseCore/SparseMap.h +243 -253
- package/eigen/Eigen/src/SparseCore/SparseMatrix.h +1614 -1142
- package/eigen/Eigen/src/SparseCore/SparseMatrixBase.h +403 -357
- package/eigen/Eigen/src/SparseCore/SparsePermutation.h +186 -115
- package/eigen/Eigen/src/SparseCore/SparseProduct.h +100 -91
- package/eigen/Eigen/src/SparseCore/SparseRedux.h +22 -24
- package/eigen/Eigen/src/SparseCore/SparseRef.h +268 -295
- package/eigen/Eigen/src/SparseCore/SparseSelfAdjointView.h +371 -414
- package/eigen/Eigen/src/SparseCore/SparseSolverBase.h +78 -87
- package/eigen/Eigen/src/SparseCore/SparseSparseProductWithPruning.h +81 -95
- package/eigen/Eigen/src/SparseCore/SparseTranspose.h +62 -71
- package/eigen/Eigen/src/SparseCore/SparseTriangularView.h +132 -144
- package/eigen/Eigen/src/SparseCore/SparseUtil.h +146 -115
- package/eigen/Eigen/src/SparseCore/SparseVector.h +426 -372
- package/eigen/Eigen/src/SparseCore/SparseView.h +164 -193
- package/eigen/Eigen/src/SparseCore/TriangularSolver.h +129 -170
- package/eigen/Eigen/src/SparseLU/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SparseLU/SparseLU.h +814 -618
- package/eigen/Eigen/src/SparseLU/SparseLUImpl.h +61 -48
- package/eigen/Eigen/src/SparseLU/SparseLU_Memory.h +102 -118
- package/eigen/Eigen/src/SparseLU/SparseLU_Structs.h +38 -35
- package/eigen/Eigen/src/SparseLU/SparseLU_SupernodalMatrix.h +273 -255
- package/eigen/Eigen/src/SparseLU/SparseLU_Utils.h +44 -49
- package/eigen/Eigen/src/SparseLU/SparseLU_column_bmod.h +104 -108
- package/eigen/Eigen/src/SparseLU/SparseLU_column_dfs.h +90 -101
- package/eigen/Eigen/src/SparseLU/SparseLU_copy_to_ucol.h +57 -58
- package/eigen/Eigen/src/SparseLU/SparseLU_heap_relax_snode.h +43 -55
- package/eigen/Eigen/src/SparseLU/SparseLU_kernel_bmod.h +74 -71
- package/eigen/Eigen/src/SparseLU/SparseLU_panel_bmod.h +125 -133
- package/eigen/Eigen/src/SparseLU/SparseLU_panel_dfs.h +136 -159
- package/eigen/Eigen/src/SparseLU/SparseLU_pivotL.h +51 -52
- package/eigen/Eigen/src/SparseLU/SparseLU_pruneL.h +67 -73
- package/eigen/Eigen/src/SparseLU/SparseLU_relax_snode.h +24 -26
- package/eigen/Eigen/src/SparseQR/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SparseQR/SparseQR.h +451 -490
- package/eigen/Eigen/src/StlSupport/StdDeque.h +28 -105
- package/eigen/Eigen/src/StlSupport/StdList.h +28 -84
- package/eigen/Eigen/src/StlSupport/StdVector.h +28 -108
- package/eigen/Eigen/src/StlSupport/details.h +48 -50
- package/eigen/Eigen/src/SuperLUSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/SuperLUSupport/SuperLUSupport.h +634 -732
- package/eigen/Eigen/src/ThreadPool/Barrier.h +70 -0
- package/eigen/Eigen/src/ThreadPool/CoreThreadPoolDevice.h +336 -0
- package/eigen/Eigen/src/ThreadPool/EventCount.h +241 -0
- package/eigen/Eigen/src/ThreadPool/ForkJoin.h +140 -0
- package/eigen/Eigen/src/ThreadPool/InternalHeaderCheck.h +4 -0
- package/eigen/Eigen/src/ThreadPool/NonBlockingThreadPool.h +587 -0
- package/eigen/Eigen/src/ThreadPool/RunQueue.h +230 -0
- package/eigen/Eigen/src/ThreadPool/ThreadCancel.h +21 -0
- package/eigen/Eigen/src/ThreadPool/ThreadEnvironment.h +43 -0
- package/eigen/Eigen/src/ThreadPool/ThreadLocal.h +289 -0
- package/eigen/Eigen/src/ThreadPool/ThreadPoolInterface.h +50 -0
- package/eigen/Eigen/src/ThreadPool/ThreadYield.h +16 -0
- package/eigen/Eigen/src/UmfPackSupport/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/UmfPackSupport/UmfPackSupport.h +480 -380
- package/eigen/Eigen/src/misc/Image.h +41 -43
- package/eigen/Eigen/src/misc/InternalHeaderCheck.h +3 -0
- package/eigen/Eigen/src/misc/Kernel.h +39 -41
- package/eigen/Eigen/src/misc/RealSvd2x2.h +19 -21
- package/eigen/Eigen/src/misc/blas.h +83 -426
- package/eigen/Eigen/src/misc/lapacke.h +9976 -16182
- package/eigen/Eigen/src/misc/lapacke_helpers.h +163 -0
- package/eigen/Eigen/src/misc/lapacke_mangling.h +4 -5
- package/eigen/Eigen/src/plugins/ArrayCwiseBinaryOps.inc +344 -0
- package/eigen/Eigen/src/plugins/ArrayCwiseUnaryOps.inc +544 -0
- package/eigen/Eigen/src/plugins/BlockMethods.inc +1370 -0
- package/eigen/Eigen/src/plugins/CommonCwiseBinaryOps.inc +116 -0
- package/eigen/Eigen/src/plugins/CommonCwiseUnaryOps.inc +167 -0
- package/eigen/Eigen/src/plugins/IndexedViewMethods.inc +192 -0
- package/eigen/Eigen/src/plugins/InternalHeaderCheck.inc +3 -0
- package/eigen/Eigen/src/plugins/MatrixCwiseBinaryOps.inc +331 -0
- package/eigen/Eigen/src/plugins/MatrixCwiseUnaryOps.inc +118 -0
- package/eigen/Eigen/src/plugins/ReshapedMethods.inc +133 -0
- package/lib/LibEigen.d.ts +4 -0
- package/lib/LibEigen.js +14 -0
- package/lib/index.d.ts +1 -1
- package/lib/index.js +7 -3
- package/package.json +2 -10
- package/eigen/Eigen/CMakeLists.txt +0 -19
- package/eigen/Eigen/src/Core/BooleanRedux.h +0 -164
- package/eigen/Eigen/src/Core/arch/CUDA/Complex.h +0 -103
- package/eigen/Eigen/src/Core/arch/CUDA/Half.h +0 -675
- package/eigen/Eigen/src/Core/arch/CUDA/MathFunctions.h +0 -91
- package/eigen/Eigen/src/Core/arch/CUDA/PacketMath.h +0 -333
- package/eigen/Eigen/src/Core/arch/CUDA/PacketMathHalf.h +0 -1124
- package/eigen/Eigen/src/Core/arch/CUDA/TypeCasting.h +0 -212
- package/eigen/Eigen/src/Core/util/NonMPL2.h +0 -3
- package/eigen/Eigen/src/Geometry/arch/Geometry_SSE.h +0 -161
- package/eigen/Eigen/src/LU/arch/Inverse_SSE.h +0 -338
- package/eigen/Eigen/src/SparseCore/MappedSparseMatrix.h +0 -67
- package/eigen/Eigen/src/SparseLU/SparseLU_gemm_kernel.h +0 -280
- package/eigen/Eigen/src/misc/lapack.h +0 -152
- package/eigen/Eigen/src/plugins/ArrayCwiseBinaryOps.h +0 -332
- package/eigen/Eigen/src/plugins/ArrayCwiseUnaryOps.h +0 -552
- package/eigen/Eigen/src/plugins/BlockMethods.h +0 -1058
- package/eigen/Eigen/src/plugins/CommonCwiseBinaryOps.h +0 -115
- package/eigen/Eigen/src/plugins/CommonCwiseUnaryOps.h +0 -163
- package/eigen/Eigen/src/plugins/MatrixCwiseBinaryOps.h +0 -152
- package/eigen/Eigen/src/plugins/MatrixCwiseUnaryOps.h +0 -85
- package/lib/eigen.d.ts +0 -2
- package/lib/eigen.js +0 -15
|
@@ -11,162 +11,320 @@
|
|
|
11
11
|
#ifndef EIGEN_COMPLEX32_ALTIVEC_H
|
|
12
12
|
#define EIGEN_COMPLEX32_ALTIVEC_H
|
|
13
13
|
|
|
14
|
+
// IWYU pragma: private
|
|
15
|
+
#include "../../InternalHeaderCheck.h"
|
|
16
|
+
|
|
14
17
|
namespace Eigen {
|
|
15
18
|
|
|
16
19
|
namespace internal {
|
|
17
20
|
|
|
18
|
-
|
|
19
|
-
|
|
21
|
+
inline Packet4ui p4ui_CONJ_XOR() {
|
|
22
|
+
return vec_mergeh((Packet4ui)p4i_ZERO, (Packet4ui)p4f_MZERO); //{ 0x00000000, 0x80000000, 0x00000000, 0x80000000 };
|
|
23
|
+
}
|
|
24
|
+
#ifdef EIGEN_VECTORIZE_VSX
|
|
20
25
|
#if defined(_BIG_ENDIAN)
|
|
21
|
-
|
|
22
|
-
|
|
26
|
+
inline Packet2ul p2ul_CONJ_XOR1() {
|
|
27
|
+
return (Packet2ul)vec_sld((Packet4ui)p2d_MZERO, (Packet4ui)p2l_ZERO,
|
|
28
|
+
8); //{ 0x8000000000000000, 0x0000000000000000 };
|
|
29
|
+
}
|
|
30
|
+
inline Packet2ul p2ul_CONJ_XOR2() {
|
|
31
|
+
return (Packet2ul)vec_sld((Packet4ui)p2l_ZERO, (Packet4ui)p2d_MZERO,
|
|
32
|
+
8); //{ 0x8000000000000000, 0x0000000000000000 };
|
|
33
|
+
}
|
|
23
34
|
#else
|
|
24
|
-
|
|
25
|
-
|
|
35
|
+
inline Packet2ul p2ul_CONJ_XOR1() {
|
|
36
|
+
return (Packet2ul)vec_sld((Packet4ui)p2l_ZERO, (Packet4ui)p2d_MZERO,
|
|
37
|
+
8); //{ 0x8000000000000000, 0x0000000000000000 };
|
|
38
|
+
}
|
|
39
|
+
inline Packet2ul p2ul_CONJ_XOR2() {
|
|
40
|
+
return (Packet2ul)vec_sld((Packet4ui)p2d_MZERO, (Packet4ui)p2l_ZERO,
|
|
41
|
+
8); //{ 0x8000000000000000, 0x0000000000000000 };
|
|
42
|
+
}
|
|
26
43
|
#endif
|
|
27
44
|
#endif
|
|
28
45
|
|
|
29
46
|
//---------- float ----------
|
|
30
|
-
struct Packet2cf
|
|
31
|
-
{
|
|
32
|
-
EIGEN_STRONG_INLINE explicit Packet2cf() : v(p4f_ZERO) {}
|
|
47
|
+
struct Packet2cf {
|
|
48
|
+
EIGEN_STRONG_INLINE explicit Packet2cf() {}
|
|
33
49
|
EIGEN_STRONG_INLINE explicit Packet2cf(const Packet4f& a) : v(a) {}
|
|
34
|
-
|
|
50
|
+
|
|
51
|
+
EIGEN_STRONG_INLINE Packet2cf pmul(const Packet2cf& a, const Packet2cf& b) {
|
|
52
|
+
Packet4f v1, v2;
|
|
53
|
+
|
|
54
|
+
// Permute and multiply the real parts of a and b
|
|
55
|
+
v1 = vec_perm(a.v, a.v, p16uc_PSET32_WODD);
|
|
56
|
+
// Get the imaginary parts of a
|
|
57
|
+
v2 = vec_perm(a.v, a.v, p16uc_PSET32_WEVEN);
|
|
58
|
+
// multiply a_re * b
|
|
59
|
+
v1 = vec_madd(v1, b.v, p4f_ZERO);
|
|
60
|
+
// multiply a_im * b and get the conjugate result
|
|
61
|
+
v2 = vec_madd(v2, b.v, p4f_ZERO);
|
|
62
|
+
v2 = reinterpret_cast<Packet4f>(pxor(v2, reinterpret_cast<Packet4f>(p4ui_CONJ_XOR())));
|
|
63
|
+
// permute back to a proper order
|
|
64
|
+
v2 = vec_perm(v2, v2, p16uc_COMPLEX32_REV);
|
|
65
|
+
|
|
66
|
+
return Packet2cf(padd<Packet4f>(v1, v2));
|
|
67
|
+
}
|
|
68
|
+
|
|
69
|
+
EIGEN_STRONG_INLINE Packet2cf& operator*=(const Packet2cf& b) {
|
|
70
|
+
v = pmul(Packet2cf(*this), b).v;
|
|
71
|
+
return *this;
|
|
72
|
+
}
|
|
73
|
+
EIGEN_STRONG_INLINE Packet2cf operator*(const Packet2cf& b) const { return Packet2cf(*this) *= b; }
|
|
74
|
+
|
|
75
|
+
EIGEN_STRONG_INLINE Packet2cf& operator+=(const Packet2cf& b) {
|
|
76
|
+
v = padd(v, b.v);
|
|
77
|
+
return *this;
|
|
78
|
+
}
|
|
79
|
+
EIGEN_STRONG_INLINE Packet2cf operator+(const Packet2cf& b) const { return Packet2cf(*this) += b; }
|
|
80
|
+
EIGEN_STRONG_INLINE Packet2cf& operator-=(const Packet2cf& b) {
|
|
81
|
+
v = psub(v, b.v);
|
|
82
|
+
return *this;
|
|
83
|
+
}
|
|
84
|
+
EIGEN_STRONG_INLINE Packet2cf operator-(const Packet2cf& b) const { return Packet2cf(*this) -= b; }
|
|
85
|
+
EIGEN_STRONG_INLINE Packet2cf operator-(void) const { return Packet2cf(-v); }
|
|
86
|
+
|
|
87
|
+
Packet4f v;
|
|
35
88
|
};
|
|
36
89
|
|
|
37
|
-
template<>
|
|
38
|
-
{
|
|
90
|
+
template <>
|
|
91
|
+
struct packet_traits<std::complex<float> > : default_packet_traits {
|
|
39
92
|
typedef Packet2cf type;
|
|
40
93
|
typedef Packet2cf half;
|
|
94
|
+
typedef Packet4f as_real;
|
|
41
95
|
enum {
|
|
42
96
|
Vectorizable = 1,
|
|
43
97
|
AlignedOnScalar = 1,
|
|
44
98
|
size = 2,
|
|
45
|
-
HasHalfPacket = 0,
|
|
46
99
|
|
|
47
|
-
HasAdd
|
|
48
|
-
HasSub
|
|
49
|
-
HasMul
|
|
50
|
-
HasDiv
|
|
100
|
+
HasAdd = 1,
|
|
101
|
+
HasSub = 1,
|
|
102
|
+
HasMul = 1,
|
|
103
|
+
HasDiv = 1,
|
|
51
104
|
HasNegate = 1,
|
|
52
|
-
HasAbs
|
|
53
|
-
HasAbs2
|
|
54
|
-
HasMin
|
|
55
|
-
HasMax
|
|
56
|
-
|
|
57
|
-
|
|
105
|
+
HasAbs = 0,
|
|
106
|
+
HasAbs2 = 0,
|
|
107
|
+
HasMin = 0,
|
|
108
|
+
HasMax = 0,
|
|
109
|
+
HasSqrt = 1,
|
|
110
|
+
HasLog = 1,
|
|
111
|
+
HasExp = 1,
|
|
112
|
+
#ifdef EIGEN_VECTORIZE_VSX
|
|
113
|
+
HasBlend = 1,
|
|
58
114
|
#endif
|
|
59
115
|
HasSetLinear = 0
|
|
60
116
|
};
|
|
61
117
|
};
|
|
62
118
|
|
|
63
|
-
template<>
|
|
119
|
+
template <>
|
|
120
|
+
struct unpacket_traits<Packet2cf> {
|
|
121
|
+
typedef std::complex<float> type;
|
|
122
|
+
enum {
|
|
123
|
+
size = 2,
|
|
124
|
+
alignment = Aligned16,
|
|
125
|
+
vectorizable = true,
|
|
126
|
+
masked_load_available = false,
|
|
127
|
+
masked_store_available = false
|
|
128
|
+
};
|
|
129
|
+
typedef Packet2cf half;
|
|
130
|
+
typedef Packet4f as_real;
|
|
131
|
+
};
|
|
64
132
|
|
|
65
|
-
template<>
|
|
66
|
-
{
|
|
133
|
+
template <>
|
|
134
|
+
EIGEN_STRONG_INLINE Packet2cf pset1<Packet2cf>(const std::complex<float>& from) {
|
|
67
135
|
Packet2cf res;
|
|
68
|
-
|
|
69
|
-
|
|
136
|
+
#ifdef EIGEN_VECTORIZE_VSX
|
|
137
|
+
// Load a single std::complex<float> from memory and duplicate
|
|
138
|
+
//
|
|
139
|
+
// Using pload would read past the end of the reference in this case
|
|
140
|
+
// Using vec_xl_len + vec_splat, generates poor assembly
|
|
141
|
+
__asm__("lxvdsx %x0,%y1" : "=wa"(res.v) : "Z"(from));
|
|
142
|
+
#else
|
|
143
|
+
if ((std::ptrdiff_t(&from) % 16) == 0)
|
|
144
|
+
res.v = pload<Packet4f>((const float*)&from);
|
|
70
145
|
else
|
|
71
|
-
res.v = ploadu<Packet4f>((const float
|
|
146
|
+
res.v = ploadu<Packet4f>((const float*)&from);
|
|
72
147
|
res.v = vec_perm(res.v, res.v, p16uc_PSET64_HI);
|
|
148
|
+
#endif
|
|
73
149
|
return res;
|
|
74
150
|
}
|
|
75
151
|
|
|
76
|
-
template<>
|
|
77
|
-
|
|
78
|
-
|
|
79
|
-
|
|
80
|
-
template<>
|
|
81
|
-
|
|
82
|
-
|
|
83
|
-
|
|
84
|
-
|
|
85
|
-
|
|
86
|
-
|
|
87
|
-
|
|
88
|
-
|
|
89
|
-
|
|
90
|
-
|
|
91
|
-
{
|
|
92
|
-
|
|
93
|
-
|
|
94
|
-
|
|
95
|
-
|
|
96
|
-
|
|
97
|
-
|
|
98
|
-
|
|
99
|
-
template<>
|
|
100
|
-
|
|
101
|
-
|
|
102
|
-
|
|
103
|
-
template<>
|
|
104
|
-
{
|
|
105
|
-
|
|
106
|
-
|
|
107
|
-
|
|
108
|
-
|
|
109
|
-
|
|
110
|
-
|
|
111
|
-
|
|
112
|
-
|
|
113
|
-
|
|
114
|
-
|
|
115
|
-
|
|
116
|
-
|
|
117
|
-
|
|
118
|
-
|
|
119
|
-
|
|
120
|
-
|
|
121
|
-
|
|
122
|
-
|
|
123
|
-
|
|
124
|
-
|
|
125
|
-
|
|
126
|
-
|
|
127
|
-
|
|
128
|
-
|
|
129
|
-
|
|
130
|
-
|
|
131
|
-
std::complex<float
|
|
132
|
-
|
|
152
|
+
template <>
|
|
153
|
+
EIGEN_STRONG_INLINE Packet2cf pload<Packet2cf>(const std::complex<float>* from) {
|
|
154
|
+
return Packet2cf(pload<Packet4f>((const float*)from));
|
|
155
|
+
}
|
|
156
|
+
template <>
|
|
157
|
+
EIGEN_STRONG_INLINE Packet2cf ploadu<Packet2cf>(const std::complex<float>* from) {
|
|
158
|
+
return Packet2cf(ploadu<Packet4f>((const float*)from));
|
|
159
|
+
}
|
|
160
|
+
template <>
|
|
161
|
+
EIGEN_ALWAYS_INLINE Packet2cf pload_partial<Packet2cf>(const std::complex<float>* from, const Index n,
|
|
162
|
+
const Index offset) {
|
|
163
|
+
return Packet2cf(pload_partial<Packet4f>((const float*)from, n * 2, offset * 2));
|
|
164
|
+
}
|
|
165
|
+
template <>
|
|
166
|
+
EIGEN_ALWAYS_INLINE Packet2cf ploadu_partial<Packet2cf>(const std::complex<float>* from, const Index n,
|
|
167
|
+
const Index offset) {
|
|
168
|
+
return Packet2cf(ploadu_partial<Packet4f>((const float*)from, n * 2, offset * 2));
|
|
169
|
+
}
|
|
170
|
+
template <>
|
|
171
|
+
EIGEN_STRONG_INLINE Packet2cf ploaddup<Packet2cf>(const std::complex<float>* from) {
|
|
172
|
+
return pset1<Packet2cf>(*from);
|
|
173
|
+
}
|
|
174
|
+
|
|
175
|
+
template <>
|
|
176
|
+
EIGEN_STRONG_INLINE void pstore<std::complex<float> >(std::complex<float>* to, const Packet2cf& from) {
|
|
177
|
+
pstore((float*)to, from.v);
|
|
178
|
+
}
|
|
179
|
+
template <>
|
|
180
|
+
EIGEN_STRONG_INLINE void pstoreu<std::complex<float> >(std::complex<float>* to, const Packet2cf& from) {
|
|
181
|
+
pstoreu((float*)to, from.v);
|
|
182
|
+
}
|
|
183
|
+
template <>
|
|
184
|
+
EIGEN_ALWAYS_INLINE void pstore_partial<std::complex<float> >(std::complex<float>* to, const Packet2cf& from,
|
|
185
|
+
const Index n, const Index offset) {
|
|
186
|
+
pstore_partial((float*)to, from.v, n * 2, offset * 2);
|
|
187
|
+
}
|
|
188
|
+
template <>
|
|
189
|
+
EIGEN_ALWAYS_INLINE void pstoreu_partial<std::complex<float> >(std::complex<float>* to, const Packet2cf& from,
|
|
190
|
+
const Index n, const Index offset) {
|
|
191
|
+
pstoreu_partial((float*)to, from.v, n * 2, offset * 2);
|
|
192
|
+
}
|
|
193
|
+
|
|
194
|
+
EIGEN_STRONG_INLINE Packet2cf pload2(const std::complex<float>& from0, const std::complex<float>& from1) {
|
|
195
|
+
Packet4f res0, res1;
|
|
196
|
+
#ifdef EIGEN_VECTORIZE_VSX
|
|
197
|
+
// Load two std::complex<float> from memory and combine
|
|
198
|
+
__asm__("lxsdx %x0,%y1" : "=wa"(res0) : "Z"(from0));
|
|
199
|
+
__asm__("lxsdx %x0,%y1" : "=wa"(res1) : "Z"(from1));
|
|
200
|
+
#ifdef _BIG_ENDIAN
|
|
201
|
+
__asm__("xxpermdi %x0, %x1, %x2, 0" : "=wa"(res0) : "wa"(res0), "wa"(res1));
|
|
202
|
+
#else
|
|
203
|
+
__asm__("xxpermdi %x0, %x2, %x1, 0" : "=wa"(res0) : "wa"(res0), "wa"(res1));
|
|
204
|
+
#endif
|
|
205
|
+
#else
|
|
206
|
+
*reinterpret_cast<std::complex<float>*>(&res0) = from0;
|
|
207
|
+
*reinterpret_cast<std::complex<float>*>(&res1) = from1;
|
|
208
|
+
res0 = vec_perm(res0, res1, p16uc_TRANSPOSE64_HI);
|
|
209
|
+
#endif
|
|
210
|
+
return Packet2cf(res0);
|
|
211
|
+
}
|
|
212
|
+
|
|
213
|
+
template <>
|
|
214
|
+
EIGEN_ALWAYS_INLINE Packet2cf pload_ignore<Packet2cf>(const std::complex<float>* from) {
|
|
215
|
+
Packet2cf res;
|
|
216
|
+
res.v = pload_ignore<Packet4f>(reinterpret_cast<const float*>(from));
|
|
217
|
+
return res;
|
|
218
|
+
}
|
|
219
|
+
|
|
220
|
+
template <typename Scalar, typename Packet>
|
|
221
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet pgather_complex_size2(const Scalar* from, Index stride,
|
|
222
|
+
const Index n = 2) {
|
|
223
|
+
eigen_internal_assert(n <= unpacket_traits<Packet>::size && "number of elements will gather past end of packet");
|
|
224
|
+
EIGEN_ALIGN16 Scalar af[2];
|
|
225
|
+
for (Index i = 0; i < n; i++) {
|
|
226
|
+
af[i] = from[i * stride];
|
|
227
|
+
}
|
|
228
|
+
return pload_ignore<Packet>(af);
|
|
229
|
+
}
|
|
230
|
+
template <>
|
|
231
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet2cf pgather<std::complex<float>, Packet2cf>(const std::complex<float>* from,
|
|
232
|
+
Index stride) {
|
|
233
|
+
return pgather_complex_size2<std::complex<float>, Packet2cf>(from, stride);
|
|
234
|
+
}
|
|
235
|
+
template <>
|
|
236
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet2cf
|
|
237
|
+
pgather_partial<std::complex<float>, Packet2cf>(const std::complex<float>* from, Index stride, const Index n) {
|
|
238
|
+
return pgather_complex_size2<std::complex<float>, Packet2cf>(from, stride, n);
|
|
239
|
+
}
|
|
240
|
+
template <typename Scalar, typename Packet>
|
|
241
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter_complex_size2(Scalar* to, const Packet& from, Index stride,
|
|
242
|
+
const Index n = 2) {
|
|
243
|
+
eigen_internal_assert(n <= unpacket_traits<Packet>::size && "number of elements will scatter past end of packet");
|
|
244
|
+
EIGEN_ALIGN16 Scalar af[2];
|
|
245
|
+
pstore<Scalar>((Scalar*)af, from);
|
|
246
|
+
for (Index i = 0; i < n; i++) {
|
|
247
|
+
to[i * stride] = af[i];
|
|
248
|
+
}
|
|
249
|
+
}
|
|
250
|
+
template <>
|
|
251
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter<std::complex<float>, Packet2cf>(std::complex<float>* to,
|
|
252
|
+
const Packet2cf& from,
|
|
253
|
+
Index stride) {
|
|
254
|
+
pscatter_complex_size2<std::complex<float>, Packet2cf>(to, from, stride);
|
|
255
|
+
}
|
|
256
|
+
template <>
|
|
257
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter_partial<std::complex<float>, Packet2cf>(std::complex<float>* to,
|
|
258
|
+
const Packet2cf& from,
|
|
259
|
+
Index stride,
|
|
260
|
+
const Index n) {
|
|
261
|
+
pscatter_complex_size2<std::complex<float>, Packet2cf>(to, from, stride, n);
|
|
262
|
+
}
|
|
263
|
+
|
|
264
|
+
template <>
|
|
265
|
+
EIGEN_STRONG_INLINE Packet2cf padd<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
|
|
266
|
+
return Packet2cf(a.v + b.v);
|
|
267
|
+
}
|
|
268
|
+
template <>
|
|
269
|
+
EIGEN_STRONG_INLINE Packet2cf psub<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
|
|
270
|
+
return Packet2cf(a.v - b.v);
|
|
271
|
+
}
|
|
272
|
+
template <>
|
|
273
|
+
EIGEN_STRONG_INLINE Packet2cf pnegate(const Packet2cf& a) {
|
|
274
|
+
return Packet2cf(pnegate(a.v));
|
|
275
|
+
}
|
|
276
|
+
template <>
|
|
277
|
+
EIGEN_STRONG_INLINE Packet2cf pconj(const Packet2cf& a) {
|
|
278
|
+
return Packet2cf(pxor<Packet4f>(a.v, reinterpret_cast<Packet4f>(p4ui_CONJ_XOR())));
|
|
279
|
+
}
|
|
280
|
+
|
|
281
|
+
template <>
|
|
282
|
+
EIGEN_STRONG_INLINE Packet2cf pand<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
|
|
283
|
+
return Packet2cf(pand<Packet4f>(a.v, b.v));
|
|
284
|
+
}
|
|
285
|
+
template <>
|
|
286
|
+
EIGEN_STRONG_INLINE Packet2cf por<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
|
|
287
|
+
return Packet2cf(por<Packet4f>(a.v, b.v));
|
|
288
|
+
}
|
|
289
|
+
template <>
|
|
290
|
+
EIGEN_STRONG_INLINE Packet2cf pxor<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
|
|
291
|
+
return Packet2cf(pxor<Packet4f>(a.v, b.v));
|
|
292
|
+
}
|
|
293
|
+
template <>
|
|
294
|
+
EIGEN_STRONG_INLINE Packet2cf pandnot<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
|
|
295
|
+
return Packet2cf(pandnot<Packet4f>(a.v, b.v));
|
|
296
|
+
}
|
|
297
|
+
|
|
298
|
+
template <>
|
|
299
|
+
EIGEN_STRONG_INLINE void prefetch<std::complex<float> >(const std::complex<float>* addr) {
|
|
300
|
+
EIGEN_PPC_PREFETCH(addr);
|
|
301
|
+
}
|
|
302
|
+
|
|
303
|
+
template <>
|
|
304
|
+
EIGEN_STRONG_INLINE std::complex<float> pfirst<Packet2cf>(const Packet2cf& a) {
|
|
305
|
+
EIGEN_ALIGN16 std::complex<float> res[2];
|
|
306
|
+
pstore((float*)&res, a.v);
|
|
133
307
|
|
|
134
308
|
return res[0];
|
|
135
309
|
}
|
|
136
310
|
|
|
137
|
-
template<>
|
|
138
|
-
{
|
|
311
|
+
template <>
|
|
312
|
+
EIGEN_STRONG_INLINE Packet2cf preverse(const Packet2cf& a) {
|
|
139
313
|
Packet4f rev_a;
|
|
140
|
-
rev_a =
|
|
314
|
+
rev_a = vec_sld(a.v, a.v, 8);
|
|
141
315
|
return Packet2cf(rev_a);
|
|
142
316
|
}
|
|
143
317
|
|
|
144
|
-
template<>
|
|
145
|
-
{
|
|
318
|
+
template <>
|
|
319
|
+
EIGEN_STRONG_INLINE std::complex<float> predux<Packet2cf>(const Packet2cf& a) {
|
|
146
320
|
Packet4f b;
|
|
147
321
|
b = vec_sld(a.v, a.v, 8);
|
|
148
322
|
b = padd<Packet4f>(a.v, b);
|
|
149
323
|
return pfirst<Packet2cf>(Packet2cf(b));
|
|
150
324
|
}
|
|
151
325
|
|
|
152
|
-
template<>
|
|
153
|
-
{
|
|
154
|
-
Packet4f b1, b2;
|
|
155
|
-
#ifdef _BIG_ENDIAN
|
|
156
|
-
b1 = vec_sld(vecs[0].v, vecs[1].v, 8);
|
|
157
|
-
b2 = vec_sld(vecs[1].v, vecs[0].v, 8);
|
|
158
|
-
#else
|
|
159
|
-
b1 = vec_sld(vecs[1].v, vecs[0].v, 8);
|
|
160
|
-
b2 = vec_sld(vecs[0].v, vecs[1].v, 8);
|
|
161
|
-
#endif
|
|
162
|
-
b2 = vec_sld(b2, b2, 8);
|
|
163
|
-
b2 = padd<Packet4f>(b1, b2);
|
|
164
|
-
|
|
165
|
-
return Packet2cf(b2);
|
|
166
|
-
}
|
|
167
|
-
|
|
168
|
-
template<> EIGEN_STRONG_INLINE std::complex<float> predux_mul<Packet2cf>(const Packet2cf& a)
|
|
169
|
-
{
|
|
326
|
+
template <>
|
|
327
|
+
EIGEN_STRONG_INLINE std::complex<float> predux_mul<Packet2cf>(const Packet2cf& a) {
|
|
170
328
|
Packet4f b;
|
|
171
329
|
Packet2cf prod;
|
|
172
330
|
b = vec_sld(a.v, a.v, 8);
|
|
@@ -175,256 +333,321 @@ template<> EIGEN_STRONG_INLINE std::complex<float> predux_mul<Packet2cf>(const P
|
|
|
175
333
|
return pfirst<Packet2cf>(prod);
|
|
176
334
|
}
|
|
177
335
|
|
|
178
|
-
|
|
179
|
-
struct palign_impl<Offset,Packet2cf>
|
|
180
|
-
{
|
|
181
|
-
static EIGEN_STRONG_INLINE void run(Packet2cf& first, const Packet2cf& second)
|
|
182
|
-
{
|
|
183
|
-
if (Offset==1)
|
|
184
|
-
{
|
|
185
|
-
#ifdef _BIG_ENDIAN
|
|
186
|
-
first.v = vec_sld(first.v, second.v, 8);
|
|
187
|
-
#else
|
|
188
|
-
first.v = vec_sld(second.v, first.v, 8);
|
|
189
|
-
#endif
|
|
190
|
-
}
|
|
191
|
-
}
|
|
192
|
-
};
|
|
193
|
-
|
|
194
|
-
template<> struct conj_helper<Packet2cf, Packet2cf, false,true>
|
|
195
|
-
{
|
|
196
|
-
EIGEN_STRONG_INLINE Packet2cf pmadd(const Packet2cf& x, const Packet2cf& y, const Packet2cf& c) const
|
|
197
|
-
{ return padd(pmul(x,y),c); }
|
|
198
|
-
|
|
199
|
-
EIGEN_STRONG_INLINE Packet2cf pmul(const Packet2cf& a, const Packet2cf& b) const
|
|
200
|
-
{
|
|
201
|
-
return internal::pmul(a, pconj(b));
|
|
202
|
-
}
|
|
203
|
-
};
|
|
204
|
-
|
|
205
|
-
template<> struct conj_helper<Packet2cf, Packet2cf, true,false>
|
|
206
|
-
{
|
|
207
|
-
EIGEN_STRONG_INLINE Packet2cf pmadd(const Packet2cf& x, const Packet2cf& y, const Packet2cf& c) const
|
|
208
|
-
{ return padd(pmul(x,y),c); }
|
|
336
|
+
EIGEN_MAKE_CONJ_HELPER_CPLX_REAL(Packet2cf, Packet4f)
|
|
209
337
|
|
|
210
|
-
|
|
211
|
-
|
|
212
|
-
|
|
213
|
-
}
|
|
214
|
-
};
|
|
215
|
-
|
|
216
|
-
template<> struct conj_helper<Packet2cf, Packet2cf, true,true>
|
|
217
|
-
{
|
|
218
|
-
EIGEN_STRONG_INLINE Packet2cf pmadd(const Packet2cf& x, const Packet2cf& y, const Packet2cf& c) const
|
|
219
|
-
{ return padd(pmul(x,y),c); }
|
|
220
|
-
|
|
221
|
-
EIGEN_STRONG_INLINE Packet2cf pmul(const Packet2cf& a, const Packet2cf& b) const
|
|
222
|
-
{
|
|
223
|
-
return pconj(internal::pmul(a, b));
|
|
224
|
-
}
|
|
225
|
-
};
|
|
226
|
-
|
|
227
|
-
EIGEN_MAKE_CONJ_HELPER_CPLX_REAL(Packet2cf,Packet4f)
|
|
228
|
-
|
|
229
|
-
template<> EIGEN_STRONG_INLINE Packet2cf pdiv<Packet2cf>(const Packet2cf& a, const Packet2cf& b)
|
|
230
|
-
{
|
|
231
|
-
// TODO optimize it for AltiVec
|
|
232
|
-
Packet2cf res = conj_helper<Packet2cf,Packet2cf,false,true>().pmul(a, b);
|
|
233
|
-
Packet4f s = pmul<Packet4f>(b.v, b.v);
|
|
234
|
-
return Packet2cf(pdiv(res.v, padd<Packet4f>(s, vec_perm(s, s, p16uc_COMPLEX32_REV))));
|
|
338
|
+
template <>
|
|
339
|
+
EIGEN_STRONG_INLINE Packet2cf pdiv<Packet2cf>(const Packet2cf& a, const Packet2cf& b) {
|
|
340
|
+
return pdiv_complex(a, b);
|
|
235
341
|
}
|
|
236
342
|
|
|
237
|
-
template<>
|
|
238
|
-
{
|
|
343
|
+
template <>
|
|
344
|
+
EIGEN_STRONG_INLINE Packet2cf pcplxflip<Packet2cf>(const Packet2cf& x) {
|
|
239
345
|
return Packet2cf(vec_perm(x.v, x.v, p16uc_COMPLEX32_REV));
|
|
240
346
|
}
|
|
241
347
|
|
|
242
|
-
EIGEN_STRONG_INLINE void ptranspose(PacketBlock<Packet2cf,2>& kernel)
|
|
243
|
-
|
|
348
|
+
EIGEN_STRONG_INLINE void ptranspose(PacketBlock<Packet2cf, 2>& kernel) {
|
|
349
|
+
#ifdef EIGEN_VECTORIZE_VSX
|
|
350
|
+
Packet4f tmp = reinterpret_cast<Packet4f>(
|
|
351
|
+
vec_mergeh(reinterpret_cast<Packet2d>(kernel.packet[0].v), reinterpret_cast<Packet2d>(kernel.packet[1].v)));
|
|
352
|
+
kernel.packet[1].v = reinterpret_cast<Packet4f>(
|
|
353
|
+
vec_mergel(reinterpret_cast<Packet2d>(kernel.packet[0].v), reinterpret_cast<Packet2d>(kernel.packet[1].v)));
|
|
354
|
+
#else
|
|
244
355
|
Packet4f tmp = vec_perm(kernel.packet[0].v, kernel.packet[1].v, p16uc_TRANSPOSE64_HI);
|
|
245
356
|
kernel.packet[1].v = vec_perm(kernel.packet[0].v, kernel.packet[1].v, p16uc_TRANSPOSE64_LO);
|
|
357
|
+
#endif
|
|
246
358
|
kernel.packet[0].v = tmp;
|
|
247
359
|
}
|
|
248
360
|
|
|
249
|
-
|
|
250
|
-
|
|
361
|
+
template <>
|
|
362
|
+
EIGEN_STRONG_INLINE Packet2cf pcmp_eq(const Packet2cf& a, const Packet2cf& b) {
|
|
363
|
+
Packet4f eq = reinterpret_cast<Packet4f>(vec_cmpeq(a.v, b.v));
|
|
364
|
+
return Packet2cf(vec_and(eq, vec_perm(eq, eq, p16uc_COMPLEX32_REV)));
|
|
365
|
+
}
|
|
366
|
+
|
|
367
|
+
#ifdef EIGEN_VECTORIZE_VSX
|
|
368
|
+
template <>
|
|
369
|
+
EIGEN_STRONG_INLINE Packet2cf pblend(const Selector<2>& ifPacket, const Packet2cf& thenPacket,
|
|
370
|
+
const Packet2cf& elsePacket) {
|
|
251
371
|
Packet2cf result;
|
|
252
|
-
result.v = reinterpret_cast<Packet4f>(
|
|
372
|
+
result.v = reinterpret_cast<Packet4f>(
|
|
373
|
+
pblend<Packet2d>(ifPacket, reinterpret_cast<Packet2d>(thenPacket.v), reinterpret_cast<Packet2d>(elsePacket.v)));
|
|
253
374
|
return result;
|
|
254
375
|
}
|
|
255
376
|
#endif
|
|
256
377
|
|
|
378
|
+
template <>
|
|
379
|
+
EIGEN_STRONG_INLINE Packet2cf psqrt<Packet2cf>(const Packet2cf& a) {
|
|
380
|
+
return psqrt_complex<Packet2cf>(a);
|
|
381
|
+
}
|
|
382
|
+
|
|
383
|
+
template <>
|
|
384
|
+
EIGEN_STRONG_INLINE Packet2cf plog<Packet2cf>(const Packet2cf& a) {
|
|
385
|
+
return plog_complex<Packet2cf>(a);
|
|
386
|
+
}
|
|
387
|
+
|
|
388
|
+
template <>
|
|
389
|
+
EIGEN_STRONG_INLINE Packet2cf pexp<Packet2cf>(const Packet2cf& a) {
|
|
390
|
+
return pexp_complex<Packet2cf>(a);
|
|
391
|
+
}
|
|
392
|
+
|
|
257
393
|
//---------- double ----------
|
|
258
|
-
#ifdef
|
|
259
|
-
struct Packet1cd
|
|
260
|
-
{
|
|
394
|
+
#ifdef EIGEN_VECTORIZE_VSX
|
|
395
|
+
struct Packet1cd {
|
|
261
396
|
EIGEN_STRONG_INLINE Packet1cd() {}
|
|
262
397
|
EIGEN_STRONG_INLINE explicit Packet1cd(const Packet2d& a) : v(a) {}
|
|
398
|
+
|
|
399
|
+
EIGEN_STRONG_INLINE Packet1cd pmul(const Packet1cd& a, const Packet1cd& b) {
|
|
400
|
+
Packet2d a_re, a_im, v1, v2;
|
|
401
|
+
|
|
402
|
+
// Permute and multiply the real parts of a and b
|
|
403
|
+
a_re = vec_perm(a.v, a.v, p16uc_PSET64_HI);
|
|
404
|
+
// Get the imaginary parts of a
|
|
405
|
+
a_im = vec_perm(a.v, a.v, p16uc_PSET64_LO);
|
|
406
|
+
// multiply a_re * b
|
|
407
|
+
v1 = vec_madd(a_re, b.v, p2d_ZERO);
|
|
408
|
+
// multiply a_im * b and get the conjugate result
|
|
409
|
+
v2 = vec_madd(a_im, b.v, p2d_ZERO);
|
|
410
|
+
v2 = reinterpret_cast<Packet2d>(vec_sld(reinterpret_cast<Packet4ui>(v2), reinterpret_cast<Packet4ui>(v2), 8));
|
|
411
|
+
v2 = pxor(v2, reinterpret_cast<Packet2d>(p2ul_CONJ_XOR1()));
|
|
412
|
+
|
|
413
|
+
return Packet1cd(padd<Packet2d>(v1, v2));
|
|
414
|
+
}
|
|
415
|
+
|
|
416
|
+
EIGEN_STRONG_INLINE Packet1cd& operator*=(const Packet1cd& b) {
|
|
417
|
+
v = pmul(Packet1cd(*this), b).v;
|
|
418
|
+
return *this;
|
|
419
|
+
}
|
|
420
|
+
EIGEN_STRONG_INLINE Packet1cd operator*(const Packet1cd& b) const { return Packet1cd(*this) *= b; }
|
|
421
|
+
|
|
422
|
+
EIGEN_STRONG_INLINE Packet1cd& operator+=(const Packet1cd& b) {
|
|
423
|
+
v = padd(v, b.v);
|
|
424
|
+
return *this;
|
|
425
|
+
}
|
|
426
|
+
EIGEN_STRONG_INLINE Packet1cd operator+(const Packet1cd& b) const { return Packet1cd(*this) += b; }
|
|
427
|
+
EIGEN_STRONG_INLINE Packet1cd& operator-=(const Packet1cd& b) {
|
|
428
|
+
v = psub(v, b.v);
|
|
429
|
+
return *this;
|
|
430
|
+
}
|
|
431
|
+
EIGEN_STRONG_INLINE Packet1cd operator-(const Packet1cd& b) const { return Packet1cd(*this) -= b; }
|
|
432
|
+
EIGEN_STRONG_INLINE Packet1cd operator-(void) const { return Packet1cd(-v); }
|
|
433
|
+
|
|
263
434
|
Packet2d v;
|
|
264
435
|
};
|
|
265
436
|
|
|
266
|
-
template<>
|
|
267
|
-
{
|
|
437
|
+
template <>
|
|
438
|
+
struct packet_traits<std::complex<double> > : default_packet_traits {
|
|
268
439
|
typedef Packet1cd type;
|
|
269
440
|
typedef Packet1cd half;
|
|
441
|
+
typedef Packet2d as_real;
|
|
270
442
|
enum {
|
|
271
443
|
Vectorizable = 1,
|
|
272
444
|
AlignedOnScalar = 0,
|
|
273
445
|
size = 1,
|
|
274
|
-
HasHalfPacket = 0,
|
|
275
446
|
|
|
276
|
-
HasAdd
|
|
277
|
-
HasSub
|
|
278
|
-
HasMul
|
|
279
|
-
HasDiv
|
|
447
|
+
HasAdd = 1,
|
|
448
|
+
HasSub = 1,
|
|
449
|
+
HasMul = 1,
|
|
450
|
+
HasDiv = 1,
|
|
280
451
|
HasNegate = 1,
|
|
281
|
-
HasAbs
|
|
282
|
-
HasAbs2
|
|
283
|
-
HasMin
|
|
284
|
-
HasMax
|
|
452
|
+
HasAbs = 0,
|
|
453
|
+
HasAbs2 = 0,
|
|
454
|
+
HasMin = 0,
|
|
455
|
+
HasMax = 0,
|
|
456
|
+
HasSqrt = 1,
|
|
457
|
+
HasLog = 1,
|
|
285
458
|
HasSetLinear = 0
|
|
286
459
|
};
|
|
287
460
|
};
|
|
288
461
|
|
|
289
|
-
template<>
|
|
290
|
-
|
|
291
|
-
|
|
292
|
-
|
|
293
|
-
|
|
294
|
-
|
|
295
|
-
|
|
296
|
-
|
|
297
|
-
|
|
462
|
+
template <>
|
|
463
|
+
struct unpacket_traits<Packet1cd> {
|
|
464
|
+
typedef std::complex<double> type;
|
|
465
|
+
enum {
|
|
466
|
+
size = 1,
|
|
467
|
+
alignment = Aligned16,
|
|
468
|
+
vectorizable = true,
|
|
469
|
+
masked_load_available = false,
|
|
470
|
+
masked_store_available = false
|
|
471
|
+
};
|
|
472
|
+
typedef Packet1cd half;
|
|
473
|
+
typedef Packet2d as_real;
|
|
474
|
+
};
|
|
298
475
|
|
|
299
|
-
template<>
|
|
300
|
-
{
|
|
301
|
-
|
|
302
|
-
af[0] = from[0*stride];
|
|
303
|
-
af[1] = from[1*stride];
|
|
304
|
-
return pload<Packet1cd>(af);
|
|
476
|
+
template <>
|
|
477
|
+
EIGEN_STRONG_INLINE Packet1cd pload<Packet1cd>(const std::complex<double>* from) {
|
|
478
|
+
return Packet1cd(pload<Packet2d>((const double*)from));
|
|
305
479
|
}
|
|
306
|
-
template<>
|
|
307
|
-
{
|
|
308
|
-
|
|
309
|
-
|
|
310
|
-
|
|
311
|
-
|
|
480
|
+
template <>
|
|
481
|
+
EIGEN_STRONG_INLINE Packet1cd ploadu<Packet1cd>(const std::complex<double>* from) {
|
|
482
|
+
return Packet1cd(ploadu<Packet2d>((const double*)from));
|
|
483
|
+
}
|
|
484
|
+
template <>
|
|
485
|
+
EIGEN_ALWAYS_INLINE Packet1cd pload_partial<Packet1cd>(const std::complex<double>* from, const Index n,
|
|
486
|
+
const Index offset) {
|
|
487
|
+
return Packet1cd(pload_partial<Packet2d>((const double*)from, n * 2, offset * 2));
|
|
488
|
+
}
|
|
489
|
+
template <>
|
|
490
|
+
EIGEN_ALWAYS_INLINE Packet1cd ploadu_partial<Packet1cd>(const std::complex<double>* from, const Index n,
|
|
491
|
+
const Index offset) {
|
|
492
|
+
return Packet1cd(ploadu_partial<Packet2d>((const double*)from, n * 2, offset * 2));
|
|
493
|
+
}
|
|
494
|
+
template <>
|
|
495
|
+
EIGEN_STRONG_INLINE void pstore<std::complex<double> >(std::complex<double>* to, const Packet1cd& from) {
|
|
496
|
+
pstore((double*)to, from.v);
|
|
497
|
+
}
|
|
498
|
+
template <>
|
|
499
|
+
EIGEN_STRONG_INLINE void pstoreu<std::complex<double> >(std::complex<double>* to, const Packet1cd& from) {
|
|
500
|
+
pstoreu((double*)to, from.v);
|
|
501
|
+
}
|
|
502
|
+
template <>
|
|
503
|
+
EIGEN_ALWAYS_INLINE void pstore_partial<std::complex<double> >(std::complex<double>* to, const Packet1cd& from,
|
|
504
|
+
const Index n, const Index offset) {
|
|
505
|
+
pstore_partial((double*)to, from.v, n * 2, offset * 2);
|
|
506
|
+
}
|
|
507
|
+
template <>
|
|
508
|
+
EIGEN_ALWAYS_INLINE void pstoreu_partial<std::complex<double> >(std::complex<double>* to, const Packet1cd& from,
|
|
509
|
+
const Index n, const Index offset) {
|
|
510
|
+
pstoreu_partial((double*)to, from.v, n * 2, offset * 2);
|
|
312
511
|
}
|
|
313
512
|
|
|
314
|
-
template<>
|
|
315
|
-
|
|
316
|
-
|
|
317
|
-
|
|
318
|
-
|
|
319
|
-
template<> EIGEN_STRONG_INLINE Packet1cd pmul<Packet1cd>(const Packet1cd& a, const Packet1cd& b)
|
|
320
|
-
{
|
|
321
|
-
Packet2d a_re, a_im, v1, v2;
|
|
513
|
+
template <>
|
|
514
|
+
EIGEN_STRONG_INLINE Packet1cd
|
|
515
|
+
pset1<Packet1cd>(const std::complex<double>& from) { /* here we really have to use unaligned loads :( */
|
|
516
|
+
return ploadu<Packet1cd>(&from);
|
|
517
|
+
}
|
|
322
518
|
|
|
323
|
-
|
|
324
|
-
|
|
325
|
-
|
|
326
|
-
|
|
327
|
-
|
|
328
|
-
|
|
329
|
-
|
|
330
|
-
|
|
331
|
-
|
|
332
|
-
|
|
519
|
+
template <>
|
|
520
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet1cd
|
|
521
|
+
pgather<std::complex<double>, Packet1cd>(const std::complex<double>* from, Index) {
|
|
522
|
+
return pload<Packet1cd>(from);
|
|
523
|
+
}
|
|
524
|
+
template <>
|
|
525
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE Packet1cd
|
|
526
|
+
pgather_partial<std::complex<double>, Packet1cd>(const std::complex<double>* from, Index, const Index) {
|
|
527
|
+
return pload<Packet1cd>(from);
|
|
528
|
+
}
|
|
529
|
+
template <>
|
|
530
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter<std::complex<double>, Packet1cd>(std::complex<double>* to,
|
|
531
|
+
const Packet1cd& from, Index) {
|
|
532
|
+
pstore<std::complex<double> >(to, from);
|
|
533
|
+
}
|
|
534
|
+
template <>
|
|
535
|
+
EIGEN_DEVICE_FUNC EIGEN_ALWAYS_INLINE void pscatter_partial<std::complex<double>, Packet1cd>(std::complex<double>* to,
|
|
536
|
+
const Packet1cd& from,
|
|
537
|
+
Index, const Index) {
|
|
538
|
+
pstore<std::complex<double> >(to, from);
|
|
539
|
+
}
|
|
333
540
|
|
|
334
|
-
|
|
541
|
+
template <>
|
|
542
|
+
EIGEN_STRONG_INLINE Packet1cd padd<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
|
|
543
|
+
return Packet1cd(a.v + b.v);
|
|
544
|
+
}
|
|
545
|
+
template <>
|
|
546
|
+
EIGEN_STRONG_INLINE Packet1cd psub<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
|
|
547
|
+
return Packet1cd(a.v - b.v);
|
|
548
|
+
}
|
|
549
|
+
template <>
|
|
550
|
+
EIGEN_STRONG_INLINE Packet1cd pnegate(const Packet1cd& a) {
|
|
551
|
+
return Packet1cd(pnegate(Packet2d(a.v)));
|
|
552
|
+
}
|
|
553
|
+
template <>
|
|
554
|
+
EIGEN_STRONG_INLINE Packet1cd pconj(const Packet1cd& a) {
|
|
555
|
+
return Packet1cd(pxor(a.v, reinterpret_cast<Packet2d>(p2ul_CONJ_XOR2())));
|
|
335
556
|
}
|
|
336
557
|
|
|
337
|
-
template<>
|
|
338
|
-
|
|
339
|
-
|
|
340
|
-
|
|
558
|
+
template <>
|
|
559
|
+
EIGEN_STRONG_INLINE Packet1cd pand<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
|
|
560
|
+
return Packet1cd(pand(a.v, b.v));
|
|
561
|
+
}
|
|
562
|
+
template <>
|
|
563
|
+
EIGEN_STRONG_INLINE Packet1cd por<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
|
|
564
|
+
return Packet1cd(por(a.v, b.v));
|
|
565
|
+
}
|
|
566
|
+
template <>
|
|
567
|
+
EIGEN_STRONG_INLINE Packet1cd pxor<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
|
|
568
|
+
return Packet1cd(pxor(a.v, b.v));
|
|
569
|
+
}
|
|
570
|
+
template <>
|
|
571
|
+
EIGEN_STRONG_INLINE Packet1cd pandnot<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
|
|
572
|
+
return Packet1cd(pandnot(a.v, b.v));
|
|
573
|
+
}
|
|
341
574
|
|
|
342
|
-
template<>
|
|
575
|
+
template <>
|
|
576
|
+
EIGEN_STRONG_INLINE Packet1cd ploaddup<Packet1cd>(const std::complex<double>* from) {
|
|
577
|
+
return pset1<Packet1cd>(*from);
|
|
578
|
+
}
|
|
343
579
|
|
|
344
|
-
template<>
|
|
580
|
+
template <>
|
|
581
|
+
EIGEN_STRONG_INLINE void prefetch<std::complex<double> >(const std::complex<double>* addr) {
|
|
582
|
+
EIGEN_PPC_PREFETCH(addr);
|
|
583
|
+
}
|
|
345
584
|
|
|
346
|
-
template<>
|
|
347
|
-
{
|
|
348
|
-
std::complex<double>
|
|
585
|
+
template <>
|
|
586
|
+
EIGEN_STRONG_INLINE std::complex<double> pfirst<Packet1cd>(const Packet1cd& a) {
|
|
587
|
+
EIGEN_ALIGN16 std::complex<double> res[1];
|
|
349
588
|
pstore<std::complex<double> >(res, a);
|
|
350
589
|
|
|
351
590
|
return res[0];
|
|
352
591
|
}
|
|
353
592
|
|
|
354
|
-
template<>
|
|
355
|
-
|
|
356
|
-
|
|
357
|
-
|
|
358
|
-
|
|
359
|
-
template<> EIGEN_STRONG_INLINE std::complex<double> predux_mul<Packet1cd>(const Packet1cd& a) { return pfirst(a); }
|
|
360
|
-
|
|
361
|
-
template<int Offset>
|
|
362
|
-
struct palign_impl<Offset,Packet1cd>
|
|
363
|
-
{
|
|
364
|
-
static EIGEN_STRONG_INLINE void run(Packet1cd& /*first*/, const Packet1cd& /*second*/)
|
|
365
|
-
{
|
|
366
|
-
// FIXME is it sure we never have to align a Packet1cd?
|
|
367
|
-
// Even though a std::complex<double> has 16 bytes, it is not necessarily aligned on a 16 bytes boundary...
|
|
368
|
-
}
|
|
369
|
-
};
|
|
370
|
-
|
|
371
|
-
template<> struct conj_helper<Packet1cd, Packet1cd, false,true>
|
|
372
|
-
{
|
|
373
|
-
EIGEN_STRONG_INLINE Packet1cd pmadd(const Packet1cd& x, const Packet1cd& y, const Packet1cd& c) const
|
|
374
|
-
{ return padd(pmul(x,y),c); }
|
|
593
|
+
template <>
|
|
594
|
+
EIGEN_STRONG_INLINE Packet1cd preverse(const Packet1cd& a) {
|
|
595
|
+
return a;
|
|
596
|
+
}
|
|
375
597
|
|
|
376
|
-
|
|
377
|
-
|
|
378
|
-
|
|
379
|
-
|
|
380
|
-
};
|
|
598
|
+
template <>
|
|
599
|
+
EIGEN_STRONG_INLINE std::complex<double> predux<Packet1cd>(const Packet1cd& a) {
|
|
600
|
+
return pfirst(a);
|
|
601
|
+
}
|
|
381
602
|
|
|
382
|
-
template<>
|
|
383
|
-
{
|
|
384
|
-
|
|
385
|
-
|
|
603
|
+
template <>
|
|
604
|
+
EIGEN_STRONG_INLINE std::complex<double> predux_mul<Packet1cd>(const Packet1cd& a) {
|
|
605
|
+
return pfirst(a);
|
|
606
|
+
}
|
|
386
607
|
|
|
387
|
-
|
|
388
|
-
{
|
|
389
|
-
return internal::pmul(pconj(a), b);
|
|
390
|
-
}
|
|
391
|
-
};
|
|
608
|
+
EIGEN_MAKE_CONJ_HELPER_CPLX_REAL(Packet1cd, Packet2d)
|
|
392
609
|
|
|
393
|
-
template<>
|
|
394
|
-
{
|
|
395
|
-
|
|
396
|
-
|
|
610
|
+
template <>
|
|
611
|
+
EIGEN_STRONG_INLINE Packet1cd pdiv<Packet1cd>(const Packet1cd& a, const Packet1cd& b) {
|
|
612
|
+
return pdiv_complex(a, b);
|
|
613
|
+
}
|
|
397
614
|
|
|
398
|
-
|
|
399
|
-
|
|
400
|
-
|
|
401
|
-
}
|
|
402
|
-
};
|
|
615
|
+
EIGEN_STRONG_INLINE Packet1cd pcplxflip /*<Packet1cd>*/ (const Packet1cd& x) {
|
|
616
|
+
return Packet1cd(preverse(Packet2d(x.v)));
|
|
617
|
+
}
|
|
403
618
|
|
|
404
|
-
|
|
619
|
+
EIGEN_STRONG_INLINE void ptranspose(PacketBlock<Packet1cd, 2>& kernel) {
|
|
620
|
+
Packet2d tmp = vec_mergeh(kernel.packet[0].v, kernel.packet[1].v);
|
|
621
|
+
kernel.packet[1].v = vec_mergel(kernel.packet[0].v, kernel.packet[1].v);
|
|
622
|
+
kernel.packet[0].v = tmp;
|
|
623
|
+
}
|
|
405
624
|
|
|
406
|
-
template<>
|
|
407
|
-
{
|
|
408
|
-
//
|
|
409
|
-
|
|
410
|
-
Packet2d
|
|
411
|
-
|
|
625
|
+
template <>
|
|
626
|
+
EIGEN_STRONG_INLINE Packet1cd pcmp_eq(const Packet1cd& a, const Packet1cd& b) {
|
|
627
|
+
// Compare real and imaginary parts of a and b to get the mask vector:
|
|
628
|
+
// [re(a)==re(b), im(a)==im(b)]
|
|
629
|
+
Packet2d eq = reinterpret_cast<Packet2d>(vec_cmpeq(a.v, b.v));
|
|
630
|
+
// Swap real/imag elements in the mask in to get:
|
|
631
|
+
// [im(a)==im(b), re(a)==re(b)]
|
|
632
|
+
Packet2d eq_swapped =
|
|
633
|
+
reinterpret_cast<Packet2d>(vec_sld(reinterpret_cast<Packet4ui>(eq), reinterpret_cast<Packet4ui>(eq), 8));
|
|
634
|
+
// Return re(a)==re(b) & im(a)==im(b) by computing bitwise AND of eq and eq_swapped
|
|
635
|
+
return Packet1cd(vec_and(eq, eq_swapped));
|
|
412
636
|
}
|
|
413
637
|
|
|
414
|
-
|
|
415
|
-
{
|
|
416
|
-
return Packet1cd(
|
|
638
|
+
template <>
|
|
639
|
+
EIGEN_STRONG_INLINE Packet1cd psqrt<Packet1cd>(const Packet1cd& a) {
|
|
640
|
+
return psqrt_complex<Packet1cd>(a);
|
|
417
641
|
}
|
|
418
642
|
|
|
419
|
-
|
|
420
|
-
{
|
|
421
|
-
|
|
422
|
-
kernel.packet[1].v = vec_perm(kernel.packet[0].v, kernel.packet[1].v, p16uc_TRANSPOSE64_LO);
|
|
423
|
-
kernel.packet[0].v = tmp;
|
|
643
|
+
template <>
|
|
644
|
+
EIGEN_STRONG_INLINE Packet1cd plog<Packet1cd>(const Packet1cd& a) {
|
|
645
|
+
return plog_complex<Packet1cd>(a);
|
|
424
646
|
}
|
|
425
|
-
#endif // __VSX__
|
|
426
|
-
} // end namespace internal
|
|
427
647
|
|
|
428
|
-
|
|
648
|
+
#endif // __VSX__
|
|
649
|
+
} // end namespace internal
|
|
650
|
+
|
|
651
|
+
} // end namespace Eigen
|
|
429
652
|
|
|
430
|
-
#endif
|
|
653
|
+
#endif // EIGEN_COMPLEX32_ALTIVEC_H
|