ARM cost model: Address computation in vector mem ops not free

Adds a function to target transform info to query for the cost of address computation. The cost model analysis pass now also queries this interface. The code in LoopVectorize adds the cost of address computation as part of the memory instruction cost calculation. Only there, we know whether the instruction will be scalarized or not. Increase the penality for inserting in to D registers on swift. This becomes necessary because we now always assume that address computation has a cost and three is a closer value to the architecture. radar://13097204 git-svn-id: https://llvm.org/svn/llvm-project/llvm/trunk@174713 91177308-0d34-0410-b5e6-96231b3b80d8
author: Arnold Schwaighofer <aschwaighofer@apple.com> 2013-02-08 14:50:48 +0000
committer: Arnold Schwaighofer <aschwaighofer@apple.com> 2013-02-08 14:50:48 +0000
commit: fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d (patch)
tree: 6143816aee69c2cdf8ac4bd2414689216941b036 /lib/Target/ARM/ARMTargetTransformInfo.cpp
parent: f64edf8d802824202b50046638696a6b7897d4d6 (diff)
download: llvm-fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d.tar.gz
llvm-fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d.tar.bz2
llvm-fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d.tar.xz
1 files changed, 11 insertions, 2 deletions
diff --git a/lib/Target/ARM/ARMTargetTransformInfo.cpp b/lib/Target/ARM/ARMTargetTransformInfo.cpp
index 1f91e0ee36..f6fa319970 100644
--- a/lib/Target/ARM/ARMTargetTransformInfo.cpp
+++ b/lib/Target/ARM/ARMTargetTransformInfo.cpp
@@ -120,6 +120,8 @@ public:
   unsigned getCmpSelInstrCost(unsigned Opcode, Type *ValTy, Type *CondTy) const;
 
   unsigned getVectorInstrCost(unsigned Opcode, Type *Val, unsigned Index) const;
+
+  unsigned getAddressComputationCost(Type *Val) const;
   /// @}
 };
 
@@ -304,12 +306,13 @@ unsigned ARMTTI::getCastInstrCost(unsigned Opcode, Type *Dst,
 
 unsigned ARMTTI::getVectorInstrCost(unsigned Opcode, Type *ValTy,
                                     unsigned Index) const {
-  // Penalize inserting into an D-subregister.
+  // Penalize inserting into an D-subregister. We end up with a three times
+  // lower estimated throughput on swift.
   if (ST->isSwift() &&
       Opcode == Instruction::InsertElement &&
       ValTy->isVectorTy() &&
       ValTy->getScalarSizeInBits() <= 32)
-    return 2;
+    return 3;
 
   return TargetTransformInfo::getVectorInstrCost(Opcode, ValTy, Index);
 }
@@ -326,3 +329,9 @@ unsigned ARMTTI::getCmpSelInstrCost(unsigned Opcode, Type *ValTy,
 
   return TargetTransformInfo::getCmpSelInstrCost(Opcode, ValTy, CondTy);
 }
+
+unsigned ARMTTI::getAddressComputationCost(Type *Ty) const {
+  // In many cases the address computation is not merged into the instruction
+  // addressing mode.
+  return 1;
+}
author	Arnold Schwaighofer <aschwaighofer@apple.com>	2013-02-08 14:50:48 +0000
committer	Arnold Schwaighofer <aschwaighofer@apple.com>	2013-02-08 14:50:48 +0000
commit	fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d (patch)
tree	6143816aee69c2cdf8ac4bd2414689216941b036 /lib/Target/ARM/ARMTargetTransformInfo.cpp
parent	f64edf8d802824202b50046638696a6b7897d4d6 (diff)
download	llvm-fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d.tar.gz llvm-fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d.tar.bz2 llvm-fb55a8fd7c38aa09d9c243d48a8a72d890f36a3d.tar.xz