Correct cost model for vector shift on AVX2

- After moving logic recognizing vector shift with scalar amount from DAG combining into DAG lowering, we declare to customize all vector shifts even vector shift on AVX is legal. As a result, the cost model needs special tuning to identify these legal cases. git-svn-id: https://llvm.org/svn/llvm-project/llvm/trunk@177586 91177308-0d34-0410-b5e6-96231b3b80d8
author: Michael Liao <michael.liao@intel.com> 2013-03-20 22:01:10 +0000
committer: Michael Liao <michael.liao@intel.com> 2013-03-20 22:01:10 +0000
commit: f74e9bf650d7c40d595d3bb60e3c901e2bccec4b (patch)
tree: 26fc6ab996a2aa6e0516091cb0f5bee4483168c7 /lib/Target/X86/X86TargetTransformInfo.cpp
parent: 6178e5f50c0c8be26913cd93238a5035a39cdf37 (diff)
download: llvm-f74e9bf650d7c40d595d3bb60e3c901e2bccec4b.tar.gz
llvm-f74e9bf650d7c40d595d3bb60e3c901e2bccec4b.tar.bz2
llvm-f74e9bf650d7c40d595d3bb60e3c901e2bccec4b.tar.xz
1 files changed, 23 insertions, 0 deletions
diff --git a/lib/Target/X86/X86TargetTransformInfo.cpp b/lib/Target/X86/X86TargetTransformInfo.cpp
index 777ef508ec..3e3b86edbb 100644
--- a/lib/Target/X86/X86TargetTransformInfo.cpp
+++ b/lib/Target/X86/X86TargetTransformInfo.cpp
@@ -169,6 +169,29 @@ unsigned X86TTI::getArithmeticInstrCost(unsigned Opcode, Type *Ty) const {
   int ISD = TLI->InstructionOpcodeToISD(Opcode);
   assert(ISD && "Invalid opcode");
 
+  static const CostTblEntry<MVT> AVX2CostTable[] = {
+    // Shifts on v4i64/v8i32 on AVX2 is legal even though we declare to
+    // customize them to detect the cases where shift amount is a scalar one.
+    { ISD::SHL,     MVT::v4i32,    1 },
+    { ISD::SRL,     MVT::v4i32,    1 },
+    { ISD::SRA,     MVT::v4i32,    1 },
+    { ISD::SHL,     MVT::v8i32,    1 },
+    { ISD::SRL,     MVT::v8i32,    1 },
+    { ISD::SRA,     MVT::v8i32,    1 },
+    { ISD::SHL,     MVT::v2i64,    1 },
+    { ISD::SRL,     MVT::v2i64,    1 },
+    { ISD::SHL,     MVT::v4i64,    1 },
+    { ISD::SRL,     MVT::v4i64,    1 },
+  };
+
+  // Look for AVX2 lowering tricks.
+  if (ST->hasAVX2()) {
+    int Idx = CostTableLookup<MVT>(AVX2CostTable, array_lengthof(AVX2CostTable),
+                                   ISD, LT.second);
+    if (Idx != -1)
+      return LT.first * AVX2CostTable[Idx].Cost;
+  }
+
   static const CostTblEntry<MVT> AVX1CostTable[] = {
     // We don't have to scalarize unsupported ops. We can issue two half-sized
     // operations and we only need to extract the upper YMM half.
author	Michael Liao <michael.liao@intel.com>	2013-03-20 22:01:10 +0000
committer	Michael Liao <michael.liao@intel.com>	2013-03-20 22:01:10 +0000
commit	f74e9bf650d7c40d595d3bb60e3c901e2bccec4b (patch)
tree	26fc6ab996a2aa6e0516091cb0f5bee4483168c7 /lib/Target/X86/X86TargetTransformInfo.cpp
parent	6178e5f50c0c8be26913cd93238a5035a39cdf37 (diff)
download	llvm-f74e9bf650d7c40d595d3bb60e3c901e2bccec4b.tar.gz llvm-f74e9bf650d7c40d595d3bb60e3c901e2bccec4b.tar.bz2 llvm-f74e9bf650d7c40d595d3bb60e3c901e2bccec4b.tar.xz