Remove the Function::getFnAttributes method in favor of using the AttributeSet

[oota-llvm.git] / lib / CodeGen / SelectionDAG / SelectionDAGBuilder.cpp
diff --git a/lib/CodeGen/SelectionDAG/SelectionDAGBuilder.cpp b/lib/CodeGen/SelectionDAG/SelectionDAGBuilder.cpp

index f3cf7582be3d5cde9e25d2c62292fc5d46145e12..23f277ae80eb3edddc13b915f7f73ab2e5c8d5e2 100644 (file)
--- a/lib/CodeGen/SelectionDAG/SelectionDAGBuilder.cpp
+++ b/lib/CodeGen/SelectionDAG/SelectionDAGBuilder.cpp
@@ -12,50 +12,51 @@
  //===----------------------------------------------------------------------===//
  
  #define DEBUG_TYPE "isel"
-#include "SDNodeDbgValue.h"
  #include "SelectionDAGBuilder.h"
+#include "SDNodeDbgValue.h"
  #include "llvm/ADT/BitVector.h"
  #include "llvm/ADT/PostOrderIterator.h"
  #include "llvm/ADT/SmallSet.h"
  #include "llvm/Analysis/AliasAnalysis.h"
  #include "llvm/Analysis/ConstantFolding.h"
-#include "llvm/Constants.h"
+#include "llvm/Analysis/ValueTracking.h"
  #include "llvm/CallingConv.h"
-#include "llvm/DebugInfo.h"
-#include "llvm/DerivedTypes.h"
-#include "llvm/Function.h"
-#include "llvm/GlobalVariable.h"
-#include "llvm/InlineAsm.h"
-#include "llvm/Instructions.h"
-#include "llvm/Intrinsics.h"
-#include "llvm/IntrinsicInst.h"
-#include "llvm/LLVMContext.h"
-#include "llvm/Module.h"
  #include "llvm/CodeGen/Analysis.h"
  #include "llvm/CodeGen/FastISel.h"
  #include "llvm/CodeGen/FunctionLoweringInfo.h"
-#include "llvm/CodeGen/GCStrategy.h"
  #include "llvm/CodeGen/GCMetadata.h"
-#include "llvm/CodeGen/MachineFunction.h"
+#include "llvm/CodeGen/GCStrategy.h"
  #include "llvm/CodeGen/MachineFrameInfo.h"
+#include "llvm/CodeGen/MachineFunction.h"
  #include "llvm/CodeGen/MachineInstrBuilder.h"
  #include "llvm/CodeGen/MachineJumpTableInfo.h"
  #include "llvm/CodeGen/MachineModuleInfo.h"
  #include "llvm/CodeGen/MachineRegisterInfo.h"
  #include "llvm/CodeGen/SelectionDAG.h"
-#include "llvm/Target/TargetData.h"
+#include "llvm/Constants.h"
+#include "llvm/DataLayout.h"
+#include "llvm/DebugInfo.h"
+#include "llvm/DerivedTypes.h"
+#include "llvm/Function.h"
+#include "llvm/GlobalVariable.h"
+#include "llvm/InlineAsm.h"
+#include "llvm/Instructions.h"
+#include "llvm/IntrinsicInst.h"
+#include "llvm/Intrinsics.h"
+#include "llvm/LLVMContext.h"
+#include "llvm/Module.h"
+#include "llvm/Support/CommandLine.h"
+#include "llvm/Support/Debug.h"
+#include "llvm/Support/ErrorHandling.h"
+#include "llvm/Support/IntegersSubsetMapping.h"
+#include "llvm/Support/MathExtras.h"
+#include "llvm/Support/raw_ostream.h"
  #include "llvm/Target/TargetFrameLowering.h"
  #include "llvm/Target/TargetInstrInfo.h"
  #include "llvm/Target/TargetIntrinsicInfo.h"
  #include "llvm/Target/TargetLibraryInfo.h"
  #include "llvm/Target/TargetLowering.h"
  #include "llvm/Target/TargetOptions.h"
-#include "llvm/Support/CommandLine.h"
-#include "llvm/Support/IntegersSubsetMapping.h"
-#include "llvm/Support/Debug.h"
-#include "llvm/Support/ErrorHandling.h"
-#include "llvm/Support/MathExtras.h"
-#include "llvm/Support/raw_ostream.h"
  #include <algorithm>
  using namespace llvm;
  
@@ -88,7 +89,7 @@ static const unsigned MaxParallelChains = 64;
  
  static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
                                        const SDValue *Parts, unsigned NumParts,
-                                      EVT PartVT, EVT ValueVT);
+                                      MVT PartVT, EVT ValueVT, const Value *V);
  
  /// getCopyFromParts - Create a value that contains the specified legal parts
  /// combined into the value they represent.  If the parts combine to a type
@@ -97,10 +98,12 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
  /// (ISD::AssertSext).
  static SDValue getCopyFromParts(SelectionDAG &DAG, DebugLoc DL,
                                  const SDValue *Parts,
-                                unsigned NumParts, EVT PartVT, EVT ValueVT,
+                                unsigned NumParts, MVT PartVT, EVT ValueVT,
+                                const Value *V,
                                  ISD::NodeType AssertOp = ISD::DELETED_NODE) {
    if (ValueVT.isVector())
-    return getCopyFromPartsVector(DAG, DL, Parts, NumParts, PartVT, ValueVT);
+    return getCopyFromPartsVector(DAG, DL, Parts, NumParts,
+                                  PartVT, ValueVT, V);
  
    assert(NumParts > 0 && "No parts to assemble!");
    const TargetLowering &TLI = DAG.getTargetLoweringInfo();
@@ -124,9 +127,9 @@ static SDValue getCopyFromParts(SelectionDAG &DAG, DebugLoc DL,
  
        if (RoundParts > 2) {
          Lo = getCopyFromParts(DAG, DL, Parts, RoundParts / 2,
-                              PartVT, HalfVT);
+                              PartVT, HalfVT, V);
          Hi = getCopyFromParts(DAG, DL, Parts + RoundParts / 2,
-                              RoundParts / 2, PartVT, HalfVT);
+                              RoundParts / 2, PartVT, HalfVT, V);
        } else {
          Lo = DAG.getNode(ISD::BITCAST, DL, HalfVT, Parts[0]);
          Hi = DAG.getNode(ISD::BITCAST, DL, HalfVT, Parts[1]);
@@ -142,7 +145,7 @@ static SDValue getCopyFromParts(SelectionDAG &DAG, DebugLoc DL,
          unsigned OddParts = NumParts - RoundParts;
          EVT OddVT = EVT::getIntegerVT(*DAG.getContext(), OddParts * PartBits);
          Hi = getCopyFromParts(DAG, DL,
-                              Parts + RoundParts, OddParts, PartVT, OddVT);
+                              Parts + RoundParts, OddParts, PartVT, OddVT, V);
  
          // Combine the round and odd parts.
          Lo = Val;
@@ -158,7 +161,7 @@ static SDValue getCopyFromParts(SelectionDAG &DAG, DebugLoc DL,
        }
      } else if (PartVT.isFloatingPoint()) {
        // FP split into multiple FP parts (for ppcf128)
-      assert(ValueVT == EVT(MVT::ppcf128) && PartVT == EVT(MVT::f64) &&
+      assert(ValueVT == EVT(MVT::ppcf128) && PartVT == MVT::f64 &&
               "Unexpected split");
        SDValue Lo, Hi;
        Lo = DAG.getNode(ISD::BITCAST, DL, EVT(MVT::f64), Parts[0]);
@@ -171,30 +174,30 @@ static SDValue getCopyFromParts(SelectionDAG &DAG, DebugLoc DL,
        assert(ValueVT.isFloatingPoint() && PartVT.isInteger() &&
               !PartVT.isVector() && "Unexpected split");
        EVT IntVT = EVT::getIntegerVT(*DAG.getContext(), ValueVT.getSizeInBits());
-      Val = getCopyFromParts(DAG, DL, Parts, NumParts, PartVT, IntVT);
+      Val = getCopyFromParts(DAG, DL, Parts, NumParts, PartVT, IntVT, V);
      }
    }
  
    // There is now one part, held in Val.  Correct it to match ValueVT.
-  PartVT = Val.getValueType();
+  EVT PartEVT = Val.getValueType();
  
-  if (PartVT == ValueVT)
+  if (PartEVT == ValueVT)
      return Val;
  
-  if (PartVT.isInteger() && ValueVT.isInteger()) {
-    if (ValueVT.bitsLT(PartVT)) {
+  if (PartEVT.isInteger() && ValueVT.isInteger()) {
+    if (ValueVT.bitsLT(PartEVT)) {
        // For a truncate, see if we have any information to
        // indicate whether the truncated bits will always be
        // zero or sign-extension.
        if (AssertOp != ISD::DELETED_NODE)
-        Val = DAG.getNode(AssertOp, DL, PartVT, Val,
+        Val = DAG.getNode(AssertOp, DL, PartEVT, Val,
                            DAG.getValueType(ValueVT));
        return DAG.getNode(ISD::TRUNCATE, DL, ValueVT, Val);
      }
      return DAG.getNode(ISD::ANY_EXTEND, DL, ValueVT, Val);
    }
  
-  if (PartVT.isFloatingPoint() && ValueVT.isFloatingPoint()) {
+  if (PartEVT.isFloatingPoint() && ValueVT.isFloatingPoint()) {
      // FP_ROUND's are always exact here.
      if (ValueVT.bitsLT(Val.getValueType()))
        return DAG.getNode(ISD::FP_ROUND, DL, ValueVT, Val,
@@ -203,20 +206,20 @@ static SDValue getCopyFromParts(SelectionDAG &DAG, DebugLoc DL,
      return DAG.getNode(ISD::FP_EXTEND, DL, ValueVT, Val);
    }
  
-  if (PartVT.getSizeInBits() == ValueVT.getSizeInBits())
+  if (PartEVT.getSizeInBits() == ValueVT.getSizeInBits())
      return DAG.getNode(ISD::BITCAST, DL, ValueVT, Val);
  
    llvm_unreachable("Unknown mismatch!");
  }
  
-/// getCopyFromParts - Create a value that contains the specified legal parts
-/// combined into the value they represent.  If the parts combine to a type
-/// larger then ValueVT then AssertOp can be used to specify whether the extra
-/// bits are known to be zero (ISD::AssertZext) or sign extended from ValueVT
-/// (ISD::AssertSext).
+/// getCopyFromPartsVector - Create a value that contains the specified legal
+/// parts combined into the value they represent.  If the parts combine to a
+/// type larger then ValueVT then AssertOp can be used to specify whether the
+/// extra bits are known to be zero (ISD::AssertZext) or sign extended from
+/// ValueVT (ISD::AssertSext).
  static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
                                        const SDValue *Parts, unsigned NumParts,
-                                      EVT PartVT, EVT ValueVT) {
+                                      MVT PartVT, EVT ValueVT, const Value *V) {
    assert(ValueVT.isVector() && "Not a vector value");
    assert(NumParts > 0 && "No parts to assemble!");
    const TargetLowering &TLI = DAG.getTargetLoweringInfo();
@@ -224,7 +227,8 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
  
    // Handle a multi-element vector.
    if (NumParts > 1) {
-    EVT IntermediateVT, RegisterVT;
+    EVT IntermediateVT;
+    MVT RegisterVT;
      unsigned NumIntermediates;
      unsigned NumRegs =
      TLI.getVectorTypeBreakdown(*DAG.getContext(), ValueVT, IntermediateVT,
@@ -232,7 +236,7 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
      assert(NumRegs == NumParts && "Part count doesn't match vector breakdown!");
      NumParts = NumRegs; // Silence a compiler warning.
      assert(RegisterVT == PartVT && "Part type doesn't match vector breakdown!");
-    assert(RegisterVT == Parts[0].getValueType() &&
+    assert(RegisterVT == Parts[0].getSimpleValueType() &&
             "Part type doesn't match part!");
  
      // Assemble the parts into intermediate operands.
@@ -242,7 +246,7 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
        // as appropriate.
        for (unsigned i = 0; i != NumParts; ++i)
          Ops[i] = getCopyFromParts(DAG, DL, &Parts[i], 1,
-                                  PartVT, IntermediateVT);
+                                  PartVT, IntermediateVT, V);
      } else if (NumParts > 0) {
        // If the intermediate type was expanded, build the intermediate
        // operands from the parts.
@@ -251,7 +255,7 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
        unsigned Factor = NumParts / NumIntermediates;
        for (unsigned i = 0; i != NumIntermediates; ++i)
          Ops[i] = getCopyFromParts(DAG, DL, &Parts[i * Factor], Factor,
-                                  PartVT, IntermediateVT);
+                                  PartVT, IntermediateVT, V);
      }
  
      // Build a vector with BUILD_VECTOR or CONCAT_VECTORS from the
@@ -262,31 +266,31 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
    }
  
    // There is now one part, held in Val.  Correct it to match ValueVT.
-  PartVT = Val.getValueType();
+  EVT PartEVT = Val.getValueType();
  
-  if (PartVT == ValueVT)
+  if (PartEVT == ValueVT)
      return Val;
  
-  if (PartVT.isVector()) {
+  if (PartEVT.isVector()) {
      // If the element type of the source/dest vectors are the same, but the
      // parts vector has more elements than the value vector, then we have a
      // vector widening case (e.g. <2 x float> -> <4 x float>).  Extract the
      // elements we want.
-    if (PartVT.getVectorElementType() == ValueVT.getVectorElementType()) {
-      assert(PartVT.getVectorNumElements() > ValueVT.getVectorNumElements() &&
+    if (PartEVT.getVectorElementType() == ValueVT.getVectorElementType()) {
+      assert(PartEVT.getVectorNumElements() > ValueVT.getVectorNumElements() &&
               "Cannot narrow, it would be a lossy transformation");
        return DAG.getNode(ISD::EXTRACT_SUBVECTOR, DL, ValueVT, Val,
                           DAG.getIntPtrConstant(0));
      }
  
      // Vector/Vector bitcast.
-    if (ValueVT.getSizeInBits() == PartVT.getSizeInBits())
+    if (ValueVT.getSizeInBits() == PartEVT.getSizeInBits())
        return DAG.getNode(ISD::BITCAST, DL, ValueVT, Val);
  
-    assert(PartVT.getVectorNumElements() == ValueVT.getVectorNumElements() &&
+    assert(PartEVT.getVectorNumElements() == ValueVT.getVectorNumElements() &&
        "Cannot handle this kind of promotion");
      // Promoted vector extract
-    bool Smaller = ValueVT.bitsLE(PartVT);
+    bool Smaller = ValueVT.bitsLE(PartEVT);
      return DAG.getNode((Smaller ? ISD::TRUNCATE : ISD::ANY_EXTEND),
                         DL, ValueVT, Val);
  
@@ -294,17 +298,28 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
  
    // Trivial bitcast if the types are the same size and the destination
    // vector type is legal.
-  if (PartVT.getSizeInBits() == ValueVT.getSizeInBits() &&
+  if (PartEVT.getSizeInBits() == ValueVT.getSizeInBits() &&
        TLI.isTypeLegal(ValueVT))
      return DAG.getNode(ISD::BITCAST, DL, ValueVT, Val);
  
    // Handle cases such as i8 -> <1 x i1>
-  assert(ValueVT.getVectorNumElements() == 1 &&
-         "Only trivial scalar-to-vector conversions should get here!");
+  if (ValueVT.getVectorNumElements() != 1) {
+    LLVMContext &Ctx = *DAG.getContext();
+    Twine ErrMsg("non-trivial scalar-to-vector conversion");
+    if (const Instruction *I = dyn_cast_or_null<Instruction>(V)) {
+      if (const CallInst *CI = dyn_cast<CallInst>(I))
+        if (isa<InlineAsm>(CI->getCalledValue()))
+          ErrMsg = ErrMsg + ", possible invalid constraint for vector type";
+      Ctx.emitError(I, ErrMsg);
+    } else {
+      Ctx.emitError(ErrMsg);
+    }
+    report_fatal_error("Cannot handle scalar-to-vector conversion!");
+  }
  
    if (ValueVT.getVectorNumElements() == 1 &&
-      ValueVT.getVectorElementType() != PartVT) {
-    bool Smaller = ValueVT.bitsLE(PartVT);
+      ValueVT.getVectorElementType() != PartEVT) {
+    bool Smaller = ValueVT.bitsLE(PartEVT);
      Val = DAG.getNode((Smaller ? ISD::TRUNCATE : ISD::ANY_EXTEND),
                         DL, ValueVT.getScalarType(), Val);
    }
@@ -312,25 +327,22 @@ static SDValue getCopyFromPartsVector(SelectionDAG &DAG, DebugLoc DL,
    return DAG.getNode(ISD::BUILD_VECTOR, DL, ValueVT, Val);
  }
  
-
-
-
  static void getCopyToPartsVector(SelectionDAG &DAG, DebugLoc dl,
                                   SDValue Val, SDValue *Parts, unsigned NumParts,
-                                 EVT PartVT);
+                                 MVT PartVT, const Value *V);
  
  /// getCopyToParts - Create a series of nodes that contain the specified value
  /// split into legal parts.  If the parts contain more bits than Val, then, for
  /// integers, ExtendKind can be used to specify how to generate the extra bits.
  static void getCopyToParts(SelectionDAG &DAG, DebugLoc DL,
                             SDValue Val, SDValue *Parts, unsigned NumParts,
-                           EVT PartVT,
+                           MVT PartVT, const Value *V,
                             ISD::NodeType ExtendKind = ISD::ANY_EXTEND) {
    EVT ValueVT = Val.getValueType();
  
    // Handle the vector case separately.
    if (ValueVT.isVector())
-    return getCopyToPartsVector(DAG, DL, Val, Parts, NumParts, PartVT);
+    return getCopyToPartsVector(DAG, DL, Val, Parts, NumParts, PartVT, V);
  
    const TargetLowering &TLI = DAG.getTargetLoweringInfo();
    unsigned PartBits = PartVT.getSizeInBits();
@@ -341,7 +353,8 @@ static void getCopyToParts(SelectionDAG &DAG, DebugLoc DL,
      return;
  
    assert(!ValueVT.isVector() && "Vector case handled elsewhere");
-  if (PartVT == ValueVT) {
+  EVT PartEVT = PartVT;
+  if (PartEVT == ValueVT) {
      assert(NumParts == 1 && "No-op copy with multiple parts!");
      Parts[0] = Val;
      return;
@@ -363,7 +376,7 @@ static void getCopyToParts(SelectionDAG &DAG, DebugLoc DL,
      }
    } else if (PartBits == ValueVT.getSizeInBits()) {
      // Different types of the same size.
-    assert(NumParts == 1 && PartVT != ValueVT);
+    assert(NumParts == 1 && PartEVT != ValueVT);
      Val = DAG.getNode(ISD::BITCAST, DL, PartVT, Val);
    } else if (NumParts * PartBits < ValueVT.getSizeInBits()) {
      // If the parts cover less bits than value has, truncate the value.
@@ -382,7 +395,19 @@ static void getCopyToParts(SelectionDAG &DAG, DebugLoc DL,
           "Failed to tile the value with PartVT!");
  
    if (NumParts == 1) {
-    assert(PartVT == ValueVT && "Type conversion failed!");
+    if (PartEVT != ValueVT) {
+      LLVMContext &Ctx = *DAG.getContext();
+      Twine ErrMsg("scalar-to-vector conversion failed");
+      if (const Instruction *I = dyn_cast_or_null<Instruction>(V)) {
+        if (const CallInst *CI = dyn_cast<CallInst>(I))
+          if (isa<InlineAsm>(CI->getCalledValue()))
+            ErrMsg = ErrMsg + ", possible invalid constraint for vector type";
+        Ctx.emitError(I, ErrMsg);
+      } else {
+        Ctx.emitError(ErrMsg);
+      }
+    }
+
      Parts[0] = Val;
      return;
    }
@@ -397,7 +422,7 @@ static void getCopyToParts(SelectionDAG &DAG, DebugLoc DL,
      unsigned OddParts = NumParts - RoundParts;
      SDValue OddVal = DAG.getNode(ISD::SRL, DL, ValueVT, Val,
                                   DAG.getIntPtrConstant(RoundBits));
-    getCopyToParts(DAG, DL, OddVal, Parts + RoundParts, OddParts, PartVT);
+    getCopyToParts(DAG, DL, OddVal, Parts + RoundParts, OddParts, PartVT, V);
  
      if (TLI.isBigEndian())
        // The odd parts were reversed by getCopyToParts - unreverse them.
@@ -443,20 +468,21 @@ static void getCopyToParts(SelectionDAG &DAG, DebugLoc DL,
  /// value split into legal parts.
  static void getCopyToPartsVector(SelectionDAG &DAG, DebugLoc DL,
                                   SDValue Val, SDValue *Parts, unsigned NumParts,
-                                 EVT PartVT) {
+                                 MVT PartVT, const Value *V) {
    EVT ValueVT = Val.getValueType();
    assert(ValueVT.isVector() && "Not a vector");
    const TargetLowering &TLI = DAG.getTargetLoweringInfo();
  
    if (NumParts == 1) {
-    if (PartVT == ValueVT) {
+    EVT PartEVT = PartVT;
+    if (PartEVT == ValueVT) {
        // Nothing to do.
      } else if (PartVT.getSizeInBits() == ValueVT.getSizeInBits()) {
        // Bitconvert vector->vector case.
        Val = DAG.getNode(ISD::BITCAST, DL, PartVT, Val);
      } else if (PartVT.isVector() &&
-               PartVT.getVectorElementType() == ValueVT.getVectorElementType() &&
-               PartVT.getVectorNumElements() > ValueVT.getVectorNumElements()) {
+               PartEVT.getVectorElementType() == ValueVT.getVectorElementType() &&
+               PartEVT.getVectorNumElements() > ValueVT.getVectorNumElements()) {
        EVT ElementVT = PartVT.getVectorElementType();
        // Vector widening case, e.g. <2 x float> -> <4 x float>.  Shuffle in
        // undef elements.
@@ -476,12 +502,12 @@ static void getCopyToPartsVector(SelectionDAG &DAG, DebugLoc DL,
        //SDValue UndefElts = DAG.getUNDEF(VectorTy);
        //Val = DAG.getNode(ISD::CONCAT_VECTORS, DL, PartVT, Val, UndefElts);
      } else if (PartVT.isVector() &&
-               PartVT.getVectorElementType().bitsGE(
+               PartEVT.getVectorElementType().bitsGE(
                   ValueVT.getVectorElementType()) &&
-               PartVT.getVectorNumElements() == ValueVT.getVectorNumElements()) {
+               PartEVT.getVectorNumElements() == ValueVT.getVectorNumElements()) {
  
        // Promoted vector extract
-      bool Smaller = PartVT.bitsLE(ValueVT);
+      bool Smaller = PartEVT.bitsLE(ValueVT);
        Val = DAG.getNode((Smaller ? ISD::TRUNCATE : ISD::ANY_EXTEND),
                          DL, PartVT, Val);
      } else{
@@ -501,7 +527,8 @@ static void getCopyToPartsVector(SelectionDAG &DAG, DebugLoc DL,
    }
  
    // Handle a multi-element vector.
-  EVT IntermediateVT, RegisterVT;
+  EVT IntermediateVT;
+  MVT RegisterVT;
    unsigned NumIntermediates;
    unsigned NumRegs = TLI.getVectorTypeBreakdown(*DAG.getContext(), ValueVT,
                                                  IntermediateVT,
@@ -529,7 +556,7 @@ static void getCopyToPartsVector(SelectionDAG &DAG, DebugLoc DL,
      // If the register was not expanded, promote or copy the value,
      // as appropriate.
      for (unsigned i = 0; i != NumParts; ++i)
-      getCopyToParts(DAG, DL, Ops[i], &Parts[i], 1, PartVT);
+      getCopyToParts(DAG, DL, Ops[i], &Parts[i], 1, PartVT, V);
    } else if (NumParts > 0) {
      // If the intermediate type was expanded, split each the value into
      // legal parts.
@@ -537,13 +564,10 @@ static void getCopyToPartsVector(SelectionDAG &DAG, DebugLoc DL,
             "Must expand into a divisible number of parts!");
      unsigned Factor = NumParts / NumIntermediates;
      for (unsigned i = 0; i != NumIntermediates; ++i)
-      getCopyToParts(DAG, DL, Ops[i], &Parts[i*Factor], Factor, PartVT);
+      getCopyToParts(DAG, DL, Ops[i], &Parts[i*Factor], Factor, PartVT, V);
    }
  }
  
-
-
-
  namespace {
    /// RegsForValue - This struct represents the registers (physical or virtual)
    /// that a particular set of values is assigned, and the type information
@@ -569,7 +593,7 @@ namespace {
      /// getRegisterType member function, however when with physical registers
      /// it is necessary to have a separate record of the types.
      ///
-    SmallVector<EVT, 4> RegVTs;
+    SmallVector<MVT, 4> RegVTs;
  
      /// Regs - This list holds the registers assigned to the values.
      /// Each legal or promoted value requires one register, and each
@@ -580,7 +604,7 @@ namespace {
      RegsForValue() {}
  
      RegsForValue(const SmallVector<unsigned, 4> &regs,
-                 EVT regvt, EVT valuevt)
+                 MVT regvt, EVT valuevt)
        : ValueVTs(1, valuevt), RegVTs(1, regvt), Regs(regs) {}
  
      RegsForValue(LLVMContext &Context, const TargetLowering &tli,
@@ -590,7 +614,7 @@ namespace {
        for (unsigned Value = 0, e = ValueVTs.size(); Value != e; ++Value) {
          EVT ValueVT = ValueVTs[Value];
          unsigned NumRegs = tli.getNumRegisters(Context, ValueVT);
-        EVT RegisterVT = tli.getRegisterType(Context, ValueVT);
+        MVT RegisterVT = tli.getRegisterType(Context, ValueVT);
          for (unsigned i = 0; i != NumRegs; ++i)
            Regs.push_back(Reg + i);
          RegVTs.push_back(RegisterVT);
@@ -601,7 +625,7 @@ namespace {
      /// areValueTypesLegal - Return true if types of all the values are legal.
      bool areValueTypesLegal(const TargetLowering &TLI) {
        for (unsigned Value = 0, e = ValueVTs.size(); Value != e; ++Value) {
-        EVT RegisterVT = RegVTs[Value];
+        MVT RegisterVT = RegVTs[Value];
          if (!TLI.isTypeLegal(RegisterVT))
            return false;
        }
@@ -621,14 +645,15 @@ namespace {
      /// If the Flag pointer is NULL, no flag is used.
      SDValue getCopyFromRegs(SelectionDAG &DAG, FunctionLoweringInfo &FuncInfo,
                              DebugLoc dl,
-                            SDValue &Chain, SDValue *Flag) const;
+                            SDValue &Chain, SDValue *Flag,
+                            const Value *V = 0) const;
  
      /// getCopyToRegs - Emit a series of CopyToReg nodes that copies the
      /// specified value into the registers specified by this object.  This uses
      /// Chain/Flag as the input and updates them for the output Chain/Flag.
      /// If the Flag pointer is NULL, no flag is used.
      void getCopyToRegs(SDValue Val, SelectionDAG &DAG, DebugLoc dl,
-                       SDValue &Chain, SDValue *Flag) const;
+                       SDValue &Chain, SDValue *Flag, const Value *V) const;
  
      /// AddInlineAsmOperands - Add this value to the specified inlineasm node
      /// operand list.  This adds the code marker, matching input operand index
@@ -647,7 +672,8 @@ namespace {
  SDValue RegsForValue::getCopyFromRegs(SelectionDAG &DAG,
                                        FunctionLoweringInfo &FuncInfo,
                                        DebugLoc dl,
-                                      SDValue &Chain, SDValue *Flag) const {
+                                      SDValue &Chain, SDValue *Flag,
+                                      const Value *V) const {
    // A Value with type {} or [0 x %t] needs no registers.
    if (ValueVTs.empty())
      return SDValue();
@@ -661,7 +687,7 @@ SDValue RegsForValue::getCopyFromRegs(SelectionDAG &DAG,
      // Copy the legal parts from the registers.
      EVT ValueVT = ValueVTs[Value];
      unsigned NumRegs = TLI.getNumRegisters(*DAG.getContext(), ValueVT);
-    EVT RegisterVT = RegVTs[Value];
+    MVT RegisterVT = RegVTs[Value];
  
      Parts.resize(NumRegs);
      for (unsigned i = 0; i != NumRegs; ++i) {
@@ -721,7 +747,7 @@ SDValue RegsForValue::getCopyFromRegs(SelectionDAG &DAG,
      }
  
      Values[Value] = getCopyFromParts(DAG, dl, Parts.begin(),
-                                     NumRegs, RegisterVT, ValueVT);
+                                     NumRegs, RegisterVT, ValueVT, V);
      Part += NumRegs;
      Parts.clear();
    }
@@ -736,7 +762,8 @@ SDValue RegsForValue::getCopyFromRegs(SelectionDAG &DAG,
  /// Chain/Flag as the input and updates them for the output Chain/Flag.
  /// If the Flag pointer is NULL, no flag is used.
  void RegsForValue::getCopyToRegs(SDValue Val, SelectionDAG &DAG, DebugLoc dl,
-                                 SDValue &Chain, SDValue *Flag) const {
+                                 SDValue &Chain, SDValue *Flag,
+                                 const Value *V) const {
    const TargetLowering &TLI = DAG.getTargetLoweringInfo();
  
    // Get the list of the values's legal parts.
@@ -745,10 +772,12 @@ void RegsForValue::getCopyToRegs(SDValue Val, SelectionDAG &DAG, DebugLoc dl,
    for (unsigned Value = 0, Part = 0, e = ValueVTs.size(); Value != e; ++Value) {
      EVT ValueVT = ValueVTs[Value];
      unsigned NumParts = TLI.getNumRegisters(*DAG.getContext(), ValueVT);
-    EVT RegisterVT = RegVTs[Value];
+    MVT RegisterVT = RegVTs[Value];
+    ISD::NodeType ExtendKind =
+      TLI.isZExtFree(Val, RegisterVT)? ISD::ZERO_EXTEND: ISD::ANY_EXTEND;
  
      getCopyToParts(DAG, dl, Val.getValue(Val.getResNo() + Value),
-                   &Parts[Part], NumParts, RegisterVT);
+                   &Parts[Part], NumParts, RegisterVT, V, ExtendKind);
      Part += NumParts;
    }
  
@@ -811,7 +840,7 @@ void RegsForValue::AddInlineAsmOperands(unsigned Code, bool HasMatching,
  
    for (unsigned Value = 0, Reg = 0, e = ValueVTs.size(); Value != e; ++Value) {
      unsigned NumRegs = TLI.getNumRegisters(*DAG.getContext(), ValueVTs[Value]);
-    EVT RegisterVT = RegVTs[Value];
+    MVT RegisterVT = RegVTs[Value];
      for (unsigned i = 0; i != NumRegs; ++i) {
        assert(Reg < Regs.size() && "Mismatch in # registers expected");
        Ops.push_back(DAG.getRegister(Regs[Reg++], RegisterVT));
@@ -824,7 +853,8 @@ void SelectionDAGBuilder::init(GCFunctionInfo *gfi, AliasAnalysis &aa,
    AA = &aa;
    GFI = gfi;
    LibInfo = li;
-  TD = DAG.getTarget().getTargetData();
+  TD = DAG.getTarget().getDataLayout();
+  Context = DAG.getContext();
    LPadToCallSiteMap.clear();
  }
  
@@ -992,7 +1022,7 @@ SDValue SelectionDAGBuilder::getValue(const Value *V) {
      unsigned InReg = It->second;
      RegsForValue RFV(*DAG.getContext(), TLI, InReg, V->getType());
      SDValue Chain = DAG.getEntryNode();
-    N = RFV.getCopyFromRegs(DAG, FuncInfo, getCurDebugLoc(), Chain, NULL);
+    N = RFV.getCopyFromRegs(DAG, FuncInfo, getCurDebugLoc(), Chain, NULL, V);
      resolveDanglingDebugInfo(V, N);
      return N;
    }
@@ -1147,7 +1177,7 @@ SDValue SelectionDAGBuilder::getValueImpl(const Value *V) {
      unsigned InReg = FuncInfo.InitializeRegForValue(Inst);
      RegsForValue RFV(*DAG.getContext(), TLI, InReg, Inst->getType());
      SDValue Chain = DAG.getEntryNode();
-    return RFV.getCopyFromRegs(DAG, FuncInfo, getCurDebugLoc(), Chain, NULL);
+    return RFV.getCopyFromRegs(DAG, FuncInfo, getCurDebugLoc(), Chain, NULL, V);
    }
  
    llvm_unreachable("Can't get register for value!");
@@ -1203,24 +1233,24 @@ void SelectionDAGBuilder::visitRet(const ReturnInst &I) {
          ISD::NodeType ExtendKind = ISD::ANY_EXTEND;
  
          const Function *F = I.getParent()->getParent();
-        if (F->paramHasAttr(0, Attribute::SExt))
+        if (F->getRetAttributes().hasAttribute(Attribute::SExt))
            ExtendKind = ISD::SIGN_EXTEND;
-        else if (F->paramHasAttr(0, Attribute::ZExt))
+        else if (F->getRetAttributes().hasAttribute(Attribute::ZExt))
            ExtendKind = ISD::ZERO_EXTEND;
  
          if (ExtendKind != ISD::ANY_EXTEND && VT.isInteger())
-          VT = TLI.getTypeForExtArgOrReturn(*DAG.getContext(), VT, ExtendKind);
+          VT = TLI.getTypeForExtArgOrReturn(VT.getSimpleVT(), ExtendKind);
  
          unsigned NumParts = TLI.getNumRegisters(*DAG.getContext(), VT);
-        EVT PartVT = TLI.getRegisterType(*DAG.getContext(), VT);
+        MVT PartVT = TLI.getRegisterType(*DAG.getContext(), VT);
          SmallVector<SDValue, 4> Parts(NumParts);
          getCopyToParts(DAG, getCurDebugLoc(),
                         SDValue(RetOp.getNode(), RetOp.getResNo() + j),
-                       &Parts[0], NumParts, PartVT, ExtendKind);
+                       &Parts[0], NumParts, PartVT, &I, ExtendKind);
  
          // 'inreg' on function refers to return value
          ISD::ArgFlagsTy Flags = ISD::ArgFlagsTy();
-        if (F->paramHasAttr(0, Attribute::InReg))
+        if (F->getRetAttributes().hasAttribute(Attribute::InReg))
            Flags.setInReg();
  
          // Propagate extension type if any
@@ -1231,7 +1261,7 @@ void SelectionDAGBuilder::visitRet(const ReturnInst &I) {
  
          for (unsigned i = 0; i < NumParts; ++i) {
            Outs.push_back(ISD::OutputArg(Flags, Parts[i].getValueType(),
-                                        /*isfixed=*/true));
+                                        /*isfixed=*/true, 0, 0));
            OutVals.push_back(Parts[i]);
          }
        }
@@ -1734,8 +1764,8 @@ void SelectionDAGBuilder::visitBitTestHeader(BitTestBlock &B,
      Sub = DAG.getZExtOrTrunc(Sub, getCurDebugLoc(), VT);
    }
  
-  B.RegVT = VT;
-  B.Reg = FuncInfo.CreateReg(VT);
+  B.RegVT = VT.getSimpleVT();
+  B.Reg = FuncInfo.CreateReg(B.RegVT);
    SDValue CopyTo = DAG.getCopyToReg(getControlRoot(), getCurDebugLoc(),
                                      B.Reg, Sub);
  
@@ -1765,10 +1795,11 @@ void SelectionDAGBuilder::visitBitTestHeader(BitTestBlock &B,
  /// visitBitTestCase - this function produces one "bit test"
  void SelectionDAGBuilder::visitBitTestCase(BitTestBlock &BB,
                                             MachineBasicBlock* NextMBB,
+                                           uint32_t BranchWeightToNext,
                                             unsigned Reg,
                                             BitTestCase &B,
                                             MachineBasicBlock *SwitchBB) {
-  EVT VT = BB.RegVT;
+  MVT VT = BB.RegVT;
    SDValue ShiftOp = DAG.getCopyFromReg(getControlRoot(), getCurDebugLoc(),
                                         Reg, VT);
    SDValue Cmp;
@@ -1802,8 +1833,10 @@ void SelectionDAGBuilder::visitBitTestCase(BitTestBlock &BB,
                         ISD::SETNE);
    }
  
-  addSuccessorWithWeight(SwitchBB, B.TargetBB);
-  addSuccessorWithWeight(SwitchBB, NextMBB);
+  // The branch weight from SwitchBB to B.TargetBB is B.ExtraWeight.
+  addSuccessorWithWeight(SwitchBB, B.TargetBB, B.ExtraWeight);
+  // The branch weight from SwitchBB to NextMBB is BranchWeightToNext.
+  addSuccessorWithWeight(SwitchBB, NextMBB, BranchWeightToNext);
  
    SDValue BrAnd = DAG.getNode(ISD::BRCOND, getCurDebugLoc(),
                                MVT::Other, getControlRoot(),
@@ -1926,6 +1959,7 @@ bool SelectionDAGBuilder::handleSmallSwitchRange(CaseRec& CR,
    if (++BBI != FuncInfo.MF->end())
      NextBlock = BBI;
  
+  BranchProbabilityInfo *BPI = FuncInfo.BPI;
    // If any two of the cases has the same destination, and if one value
    // is the same as the other, but has one bit unset that the other has set,
    // use bit manipulation to do two compares at once.  For example:
@@ -1959,8 +1993,12 @@ bool SelectionDAGBuilder::handleSmallSwitchRange(CaseRec& CR,
                                      ISD::SETEQ);
  
          // Update successor info.
-        addSuccessorWithWeight(SwitchBB, Small.BB);
-        addSuccessorWithWeight(SwitchBB, Default);
+        // Both Small and Big will jump to Small.BB, so we sum up the weights.
+        addSuccessorWithWeight(SwitchBB, Small.BB,
+                               Small.ExtraWeight + Big.ExtraWeight);
+        addSuccessorWithWeight(SwitchBB, Default,
+          // The default destination is the first successor in IR.
+          BPI ? BPI->getEdgeWeight(SwitchBB->getBasicBlock(), (unsigned)0) : 0);
  
          // Insert the true branch.
          SDValue BrCond = DAG.getNode(ISD::BRCOND, DL, MVT::Other,
@@ -1978,14 +2016,13 @@ bool SelectionDAGBuilder::handleSmallSwitchRange(CaseRec& CR,
    }
  
    // Order cases by weight so the most likely case will be checked first.
-  BranchProbabilityInfo *BPI = FuncInfo.BPI;
+  uint32_t UnhandledWeights = 0;
    if (BPI) {
      for (CaseItr I = CR.Range.first, IE = CR.Range.second; I != IE; ++I) {
-      uint32_t IWeight = BPI->getEdgeWeight(SwitchBB->getBasicBlock(),
-                                            I->BB->getBasicBlock());
+      uint32_t IWeight = I->ExtraWeight;
+      UnhandledWeights += IWeight;
        for (CaseItr J = CR.Range.first; J < I; ++J) {
-        uint32_t JWeight = BPI->getEdgeWeight(SwitchBB->getBasicBlock(),
-                                              J->BB->getBasicBlock());
+        uint32_t JWeight = J->ExtraWeight;
          if (IWeight > JWeight)
            std::swap(*I, *J);
        }
@@ -2034,10 +2071,12 @@ bool SelectionDAGBuilder::handleSmallSwitchRange(CaseRec& CR,
        LHS = I->Low; MHS = SV; RHS = I->High;
      }
  
-    uint32_t ExtraWeight = I->ExtraWeight;
+    // The false weight should be sum of all un-handled cases.
+    UnhandledWeights -= I->ExtraWeight;
      CaseBlock CB(CC, LHS, RHS, MHS, /* truebb */ I->BB, /* falsebb */ FallThrough,
                   /* me */ CurBlock,
-                 /* trueweight */ ExtraWeight / 2, /* falseweight */ ExtraWeight / 2);
+                 /* trueweight */ I->ExtraWeight,
+                 /* falseweight */ UnhandledWeights);
  
      // If emitting the first comparison, just call visitSwitchCase to emit the
      // code into the current block.  Otherwise, push the CaseBlock onto the
@@ -2082,7 +2121,7 @@ bool SelectionDAGBuilder::handleJTSwitchCase(CaseRec &CR,
    for (CaseItr I = CR.Range.first, E = CR.Range.second; I != E; ++I)
      TSize += I->size();
  
-  if (!areJTsAllowed(TLI) || TSize.ult(4))
+  if (!areJTsAllowed(TLI) || TSize.ult(TLI.getMinimumJumpTableEntries()))
      return false;
  
    APInt Range = ComputeRange(First, Last);
@@ -2137,13 +2176,28 @@ bool SelectionDAGBuilder::handleJTSwitchCase(CaseRec &CR,
      }
    }
  
+  // Calculate weight for each unique destination in CR.
+  DenseMap<MachineBasicBlock*, uint32_t> DestWeights;
+  if (FuncInfo.BPI)
+    for (CaseItr I = CR.Range.first, E = CR.Range.second; I != E; ++I) {
+      DenseMap<MachineBasicBlock*, uint32_t>::iterator Itr =
+          DestWeights.find(I->BB);
+      if (Itr != DestWeights.end()) 
+        Itr->second += I->ExtraWeight;
+      else
+        DestWeights[I->BB] = I->ExtraWeight;
+    }
+
    // Update successor info. Add one edge to each unique successor.
    BitVector SuccsHandled(CR.CaseBB->getParent()->getNumBlockIDs());
    for (std::vector<MachineBasicBlock*>::iterator I = DestBBs.begin(),
           E = DestBBs.end(); I != E; ++I) {
      if (!SuccsHandled[(*I)->getNumber()]) {
        SuccsHandled[(*I)->getNumber()] = true;
-      addSuccessorWithWeight(JumpTableBB, *I);
+      DenseMap<MachineBasicBlock*, uint32_t>::iterator Itr =
+          DestWeights.find(*I);
+      addSuccessorWithWeight(JumpTableBB, *I,
+                             Itr != DestWeights.end() ? Itr->second : 0);
      }
    }
  
@@ -2374,7 +2428,7 @@ bool SelectionDAGBuilder::handleBitTestsSwitchCase(CaseRec& CR,
  
      if (i == count) {
        assert((count < 3) && "Too much destinations to test!");
-      CasesBits.push_back(CaseBits(0, Dest, 0));
+      CasesBits.push_back(CaseBits(0, Dest, 0, 0/*Weight*/));
        count++;
      }
  
@@ -2383,6 +2437,7 @@ bool SelectionDAGBuilder::handleBitTestsSwitchCase(CaseRec& CR,
  
      uint64_t lo = (lowValue - lowBound).getZExtValue();
      uint64_t hi = (highValue - lowBound).getZExtValue();
+    CasesBits[i].ExtraWeight += I->ExtraWeight;
  
      for (uint64_t j = lo; j <= hi; j++) {
        CasesBits[i].Mask |=  1ULL << j;
@@ -2410,7 +2465,7 @@ bool SelectionDAGBuilder::handleBitTestsSwitchCase(CaseRec& CR,
      CurMF->insert(BBI, CaseBB);
      BTC.push_back(BitTestCase(CasesBits[i].Mask,
                                CaseBB,
-                              CasesBits[i].BB));
+                              CasesBits[i].BB, CasesBits[i].ExtraWeight));
  
      // Put SV in a virtual register to make it available from the new blocks.
      ExportFromCurrentBlock(SV);
@@ -2438,30 +2493,25 @@ size_t SelectionDAGBuilder::Clusterify(CaseVector& Cases,
    
    Clusterifier TheClusterifier;
  
+  BranchProbabilityInfo *BPI = FuncInfo.BPI;
    // Start with "simple" cases
    for (SwitchInst::ConstCaseIt i = SI.case_begin(), e = SI.case_end();
         i != e; ++i) {
      const BasicBlock *SuccBB = i.getCaseSuccessor();
      MachineBasicBlock *SMBB = FuncInfo.MBBMap[SuccBB];
  
-    TheClusterifier.add(i.getCaseValueEx(), SMBB);
+    TheClusterifier.add(i.getCaseValueEx(), SMBB, 
+        BPI ? BPI->getEdgeWeight(SI.getParent(), i.getSuccessorIndex()) : 0);
    }
    
    TheClusterifier.optimize();
    
-  BranchProbabilityInfo *BPI = FuncInfo.BPI;
    size_t numCmps = 0;
    for (Clusterifier::RangeIterator i = TheClusterifier.begin(),
         e = TheClusterifier.end(); i != e; ++i, ++numCmps) {
      Clusterifier::Cluster &C = *i;
-    unsigned W = 0;
-    if (BPI) {
-      W = BPI->getEdgeWeight(SI.getParent(), C.second->getBasicBlock());
-      if (!W)
-        W = 16;
-      W *= C.first.Weight;
-      BPI->setEdgeWeight(SI.getParent(), C.second->getBasicBlock(), W);  
-    }
+    // Update edge weight for the cluster.
+    unsigned W = C.first.Weight;
  
      // FIXME: Currently work with ConstantInt based numbers.
      // Changing it to APInt based is a pretty heavy for this commit.
@@ -2543,9 +2593,10 @@ void SelectionDAGBuilder::visitSwitch(const SwitchInst &SI) {
      if (handleSmallSwitchRange(CR, WorkList, SV, Default, SwitchMBB))
        continue;
  
-    // If the switch has more than 5 blocks, and at least 40% dense, and the
+    // If the switch has more than N blocks, and is at least 40% dense, and the
      // target supports indirect branches, then emit a jump table rather than
      // lowering the switch to a binary tree of conditional branches.
+    // N defaults to 4 and is controlled via TLS.getMinimumJumpTableEntries().
      if (handleJTSwitchCase(CR, WorkList, SV, Default, SwitchMBB))
        continue;
  
@@ -2559,14 +2610,14 @@ void SelectionDAGBuilder::visitIndirectBr(const IndirectBrInst &I) {
    MachineBasicBlock *IndirectBrMBB = FuncInfo.MBB;
  
    // Update machine-CFG edges with unique successors.
-  SmallVector<BasicBlock*, 32> succs;
-  succs.reserve(I.getNumSuccessors());
-  for (unsigned i = 0, e = I.getNumSuccessors(); i != e; ++i)
-    succs.push_back(I.getSuccessor(i));
-  array_pod_sort(succs.begin(), succs.end());
-  succs.erase(std::unique(succs.begin(), succs.end()), succs.end());
-  for (unsigned i = 0, e = succs.size(); i != e; ++i) {
-    MachineBasicBlock *Succ = FuncInfo.MBBMap[succs[i]];
+  SmallSet<BasicBlock*, 32> Done;
+  for (unsigned i = 0, e = I.getNumSuccessors(); i != e; ++i) {
+    BasicBlock *BB = I.getSuccessor(i);
+    bool Inserted = Done.insert(BB);
+    if (!Inserted)
+        continue;
+
+    MachineBasicBlock *Succ = FuncInfo.MBBMap[BB];
      addSuccessorWithWeight(IndirectBrMBB, Succ);
    }
  
@@ -3092,12 +3143,12 @@ void SelectionDAGBuilder::visitGetElementPtr(const User &I) {
         OI != E; ++OI) {
      const Value *Idx = *OI;
      if (StructType *StTy = dyn_cast<StructType>(Ty)) {
-      unsigned Field = cast<ConstantInt>(Idx)->getZExtValue();
+      unsigned Field = cast<Constant>(Idx)->getUniqueInteger().getZExtValue();
        if (Field) {
          // N = N + Offset
          uint64_t Offset = TD->getStructLayout(StTy)->getElementOffset(Field);
          N = DAG.getNode(ISD::ADD, getCurDebugLoc(), N.getValueType(), N,
-                        DAG.getIntPtrConstant(Offset));
+                        DAG.getConstant(Offset, N.getValueType()));
        }
  
        Ty = StTy->getElementType(Field);
@@ -3142,7 +3193,7 @@ void SelectionDAGBuilder::visitGetElementPtr(const User &I) {
                               N.getValueType(), IdxN,
                               DAG.getConstant(Amt, IdxN.getValueType()));
          } else {
-          SDValue Scale = DAG.getConstant(ElementSize, TLI.getPointerTy());
+          SDValue Scale = DAG.getConstant(ElementSize, IdxN.getValueType());
            IdxN = DAG.getNode(ISD::MUL, getCurDebugLoc(),
                               N.getValueType(), IdxN, Scale);
          }
@@ -3163,9 +3214,9 @@ void SelectionDAGBuilder::visitAlloca(const AllocaInst &I) {
      return;   // getValue will auto-populate this.
  
    Type *Ty = I.getAllocatedType();
-  uint64_t TySize = TLI.getTargetData()->getTypeAllocSize(Ty);
+  uint64_t TySize = TLI.getDataLayout()->getTypeAllocSize(Ty);
    unsigned Align =
-    std::max((unsigned)TLI.getTargetData()->getPrefTypeAlignment(Ty),
+    std::max((unsigned)TLI.getDataLayout()->getPrefTypeAlignment(Ty),
               I.getAlignment());
  
    SDValue AllocSize = getValue(I.getArraySize());
@@ -3642,16 +3693,12 @@ getF32Constant(SelectionDAG &DAG, unsigned Flt) {
    return DAG.getConstantFP(APFloat(APInt(32, Flt)), MVT::f32);
  }
  
-/// visitExp - Lower an exp intrinsic. Handles the special sequences for
+/// expandExp - Lower an exp intrinsic. Handles the special sequences for
  /// limited-precision mode.
-void
-SelectionDAGBuilder::visitExp(const CallInst &I) {
-  SDValue result;
-  DebugLoc dl = getCurDebugLoc();
-
-  if (getValue(I.getArgOperand(0)).getValueType() == MVT::f32 &&
+static SDValue expandExp(DebugLoc dl, SDValue Op, SelectionDAG &DAG,
+                         const TargetLowering &TLI) {
+  if (Op.getValueType() == MVT::f32 &&
        LimitFloatPrecision > 0 && LimitFloatPrecision <= 18) {
-    SDValue Op = getValue(I.getArgOperand(0));
  
      // Put the exponent in the right bit position for later addition to the
      // final result:
@@ -3670,6 +3717,7 @@ SelectionDAGBuilder::visitExp(const CallInst &I) {
      IntegerPartOfX = DAG.getNode(ISD::SHL, dl, MVT::i32, IntegerPartOfX,
                                   DAG.getConstant(23, TLI.getPointerTy()));
  
+    SDValue TwoToFracPartOfX;
      if (LimitFloatPrecision <= 6) {
        // For floating-point precision of 6:
        //
@@ -3683,16 +3731,9 @@ SelectionDAGBuilder::visitExp(const CallInst &I) {
        SDValue t3 = DAG.getNode(ISD::FADD, dl, MVT::f32, t2,
                                 getF32Constant(DAG, 0x3f3c50c8));
        SDValue t4 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t3, X);
-      SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
-                               getF32Constant(DAG, 0x3f7f5e7e));
-      SDValue TwoToFracPartOfX = DAG.getNode(ISD::BITCAST, dl,MVT::i32, t5);
-
-      // Add the exponent into the result in integer domain.
-      SDValue t6 = DAG.getNode(ISD::ADD, dl, MVT::i32,
-                               TwoToFracPartOfX, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl, MVT::f32, t6);
-    } else if (LimitFloatPrecision > 6 && LimitFloatPrecision <= 12) {
+      TwoToFracPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
+                                     getF32Constant(DAG, 0x3f7f5e7e));
+    } else if (LimitFloatPrecision <= 12) {
        // For floating-point precision of 12:
        //
        //   TwoToFractionalPartOfX =
@@ -3709,16 +3750,9 @@ SelectionDAGBuilder::visitExp(const CallInst &I) {
        SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
                                 getF32Constant(DAG, 0x3f324b07));
        SDValue t6 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t5, X);
-      SDValue t7 = DAG.getNode(ISD::FADD, dl, MVT::f32, t6,
-                               getF32Constant(DAG, 0x3f7ff8fd));
-      SDValue TwoToFracPartOfX = DAG.getNode(ISD::BITCAST, dl,MVT::i32, t7);
-
-      // Add the exponent into the result in integer domain.
-      SDValue t8 = DAG.getNode(ISD::ADD, dl, MVT::i32,
-                               TwoToFracPartOfX, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl, MVT::f32, t8);
-    } else { // LimitFloatPrecision > 12 && LimitFloatPrecision <= 18
+      TwoToFracPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t6,
+                                     getF32Constant(DAG, 0x3f7ff8fd));
+    } else { // LimitFloatPrecision <= 18
        // For floating-point precision of 18:
        //
        //   TwoToFractionalPartOfX =
@@ -3747,37 +3781,27 @@ SelectionDAGBuilder::visitExp(const CallInst &I) {
        SDValue t11 = DAG.getNode(ISD::FADD, dl, MVT::f32, t10,
                                  getF32Constant(DAG, 0x3f317234));
        SDValue t12 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t11, X);
-      SDValue t13 = DAG.getNode(ISD::FADD, dl, MVT::f32, t12,
-                                getF32Constant(DAG, 0x3f800000));
-      SDValue TwoToFracPartOfX = DAG.getNode(ISD::BITCAST, dl,
-                                             MVT::i32, t13);
-
-      // Add the exponent into the result in integer domain.
-      SDValue t14 = DAG.getNode(ISD::ADD, dl, MVT::i32,
-                                TwoToFracPartOfX, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl, MVT::f32, t14);
+      TwoToFracPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t12,
+                                     getF32Constant(DAG, 0x3f800000));
      }
-  } else {
-    // No special expansion.
-    result = DAG.getNode(ISD::FEXP, dl,
-                         getValue(I.getArgOperand(0)).getValueType(),
-                         getValue(I.getArgOperand(0)));
+
+    // Add the exponent into the result in integer domain.
+    SDValue t13 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, TwoToFracPartOfX);
+    return DAG.getNode(ISD::BITCAST, dl, MVT::f32,
+                       DAG.getNode(ISD::ADD, dl, MVT::i32,
+                                   t13, IntegerPartOfX));
    }
  
-  setValue(&I, result);
+  // No special expansion.
+  return DAG.getNode(ISD::FEXP, dl, Op.getValueType(), Op);
  }
  
-/// visitLog - Lower a log intrinsic. Handles the special sequences for
+/// expandLog - Lower a log intrinsic. Handles the special sequences for
  /// limited-precision mode.
-void
-SelectionDAGBuilder::visitLog(const CallInst &I) {
-  SDValue result;
-  DebugLoc dl = getCurDebugLoc();
-
-  if (getValue(I.getArgOperand(0)).getValueType() == MVT::f32 &&
+static SDValue expandLog(DebugLoc dl, SDValue Op, SelectionDAG &DAG,
+                         const TargetLowering &TLI) {
+  if (Op.getValueType() == MVT::f32 &&
        LimitFloatPrecision > 0 && LimitFloatPrecision <= 18) {
-    SDValue Op = getValue(I.getArgOperand(0));
      SDValue Op1 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, Op);
  
      // Scale the exponent by log(2) [0.69314718f].
@@ -3789,6 +3813,7 @@ SelectionDAGBuilder::visitLog(const CallInst &I) {
      // exponent of 1.
      SDValue X = GetSignificand(DAG, Op1, dl);
  
+    SDValue LogOfMantissa;
      if (LimitFloatPrecision <= 6) {
        // For floating-point precision of 6:
        //
@@ -3802,12 +3827,9 @@ SelectionDAGBuilder::visitLog(const CallInst &I) {
        SDValue t1 = DAG.getNode(ISD::FADD, dl, MVT::f32, t0,
                                 getF32Constant(DAG, 0x3fb3a2b1));
        SDValue t2 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t1, X);
-      SDValue LogOfMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t2,
-                                          getF32Constant(DAG, 0x3f949a29));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, LogOfMantissa);
-    } else if (LimitFloatPrecision > 6 && LimitFloatPrecision <= 12) {
+      LogOfMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t2,
+                                  getF32Constant(DAG, 0x3f949a29));
+    } else if (LimitFloatPrecision <= 12) {
        // For floating-point precision of 12:
        //
        //   LogOfMantissa =
@@ -3828,12 +3850,9 @@ SelectionDAGBuilder::visitLog(const CallInst &I) {
        SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
                                 getF32Constant(DAG, 0x40348e95));
        SDValue t6 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t5, X);
-      SDValue LogOfMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t6,
-                                          getF32Constant(DAG, 0x3fdef31a));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, LogOfMantissa);
-    } else { // LimitFloatPrecision > 12 && LimitFloatPrecision <= 18
+      LogOfMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t6,
+                                  getF32Constant(DAG, 0x3fdef31a));
+    } else { // LimitFloatPrecision <= 18
        // For floating-point precision of 18:
        //
        //   LogOfMantissa =
@@ -3862,32 +3881,23 @@ SelectionDAGBuilder::visitLog(const CallInst &I) {
        SDValue t9 = DAG.getNode(ISD::FADD, dl, MVT::f32, t8,
                                 getF32Constant(DAG, 0x408797cb));
        SDValue t10 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t9, X);
-      SDValue LogOfMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t10,
-                                          getF32Constant(DAG, 0x4006dcab));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, LogOfMantissa);
+      LogOfMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t10,
+                                  getF32Constant(DAG, 0x4006dcab));
      }
-  } else {
-    // No special expansion.
-    result = DAG.getNode(ISD::FLOG, dl,
-                         getValue(I.getArgOperand(0)).getValueType(),
-                         getValue(I.getArgOperand(0)));
+
+    return DAG.getNode(ISD::FADD, dl, MVT::f32, LogOfExponent, LogOfMantissa);
    }
  
-  setValue(&I, result);
+  // No special expansion.
+  return DAG.getNode(ISD::FLOG, dl, Op.getValueType(), Op);
  }
  
-/// visitLog2 - Lower a log2 intrinsic. Handles the special sequences for
+/// expandLog2 - Lower a log2 intrinsic. Handles the special sequences for
  /// limited-precision mode.
-void
-SelectionDAGBuilder::visitLog2(const CallInst &I) {
-  SDValue result;
-  DebugLoc dl = getCurDebugLoc();
-
-  if (getValue(I.getArgOperand(0)).getValueType() == MVT::f32 &&
+static SDValue expandLog2(DebugLoc dl, SDValue Op, SelectionDAG &DAG,
+                          const TargetLowering &TLI) {
+  if (Op.getValueType() == MVT::f32 &&
        LimitFloatPrecision > 0 && LimitFloatPrecision <= 18) {
-    SDValue Op = getValue(I.getArgOperand(0));
      SDValue Op1 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, Op);
  
      // Get the exponent.
@@ -3899,6 +3909,7 @@ SelectionDAGBuilder::visitLog2(const CallInst &I) {
  
      // Different possible minimax approximations of significand in
      // floating-point for various degrees of accuracy over [1,2].
+    SDValue Log2ofMantissa;
      if (LimitFloatPrecision <= 6) {
        // For floating-point precision of 6:
        //
@@ -3910,12 +3921,9 @@ SelectionDAGBuilder::visitLog2(const CallInst &I) {
        SDValue t1 = DAG.getNode(ISD::FADD, dl, MVT::f32, t0,
                                 getF32Constant(DAG, 0x40019463));
        SDValue t2 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t1, X);
-      SDValue Log2ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t2,
-                                           getF32Constant(DAG, 0x3fd6633d));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, Log2ofMantissa);
-    } else if (LimitFloatPrecision > 6 && LimitFloatPrecision <= 12) {
+      Log2ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t2,
+                                   getF32Constant(DAG, 0x3fd6633d));
+    } else if (LimitFloatPrecision <= 12) {
        // For floating-point precision of 12:
        //
        //   Log2ofMantissa =
@@ -3936,12 +3944,9 @@ SelectionDAGBuilder::visitLog2(const CallInst &I) {
        SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
                                 getF32Constant(DAG, 0x40823e2f));
        SDValue t6 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t5, X);
-      SDValue Log2ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t6,
-                                           getF32Constant(DAG, 0x4020d29c));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, Log2ofMantissa);
-    } else { // LimitFloatPrecision > 12 && LimitFloatPrecision <= 18
+      Log2ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t6,
+                                   getF32Constant(DAG, 0x4020d29c));
+    } else { // LimitFloatPrecision <= 18
        // For floating-point precision of 18:
        //
        //   Log2ofMantissa =
@@ -3971,32 +3976,23 @@ SelectionDAGBuilder::visitLog2(const CallInst &I) {
        SDValue t9 = DAG.getNode(ISD::FADD, dl, MVT::f32, t8,
                                 getF32Constant(DAG, 0x40c39dad));
        SDValue t10 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t9, X);
-      SDValue Log2ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t10,
-                                           getF32Constant(DAG, 0x4042902c));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, Log2ofMantissa);
+      Log2ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t10,
+                                   getF32Constant(DAG, 0x4042902c));
      }
-  } else {
-    // No special expansion.
-    result = DAG.getNode(ISD::FLOG2, dl,
-                         getValue(I.getArgOperand(0)).getValueType(),
-                         getValue(I.getArgOperand(0)));
+
+    return DAG.getNode(ISD::FADD, dl, MVT::f32, LogOfExponent, Log2ofMantissa);
    }
  
-  setValue(&I, result);
+  // No special expansion.
+  return DAG.getNode(ISD::FLOG2, dl, Op.getValueType(), Op);
  }
  
-/// visitLog10 - Lower a log10 intrinsic. Handles the special sequences for
+/// expandLog10 - Lower a log10 intrinsic. Handles the special sequences for
  /// limited-precision mode.
-void
-SelectionDAGBuilder::visitLog10(const CallInst &I) {
-  SDValue result;
-  DebugLoc dl = getCurDebugLoc();
-
-  if (getValue(I.getArgOperand(0)).getValueType() == MVT::f32 &&
+static SDValue expandLog10(DebugLoc dl, SDValue Op, SelectionDAG &DAG,
+                           const TargetLowering &TLI) {
+  if (Op.getValueType() == MVT::f32 &&
        LimitFloatPrecision > 0 && LimitFloatPrecision <= 18) {
-    SDValue Op = getValue(I.getArgOperand(0));
      SDValue Op1 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, Op);
  
      // Scale the exponent by log10(2) [0.30102999f].
@@ -4008,6 +4004,7 @@ SelectionDAGBuilder::visitLog10(const CallInst &I) {
      // exponent of 1.
      SDValue X = GetSignificand(DAG, Op1, dl);
  
+    SDValue Log10ofMantissa;
      if (LimitFloatPrecision <= 6) {
        // For floating-point precision of 6:
        //
@@ -4021,12 +4018,9 @@ SelectionDAGBuilder::visitLog10(const CallInst &I) {
        SDValue t1 = DAG.getNode(ISD::FADD, dl, MVT::f32, t0,
                                 getF32Constant(DAG, 0x3f1c0789));
        SDValue t2 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t1, X);
-      SDValue Log10ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t2,
-                                            getF32Constant(DAG, 0x3f011300));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, Log10ofMantissa);
-    } else if (LimitFloatPrecision > 6 && LimitFloatPrecision <= 12) {
+      Log10ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t2,
+                                    getF32Constant(DAG, 0x3f011300));
+    } else if (LimitFloatPrecision <= 12) {
        // For floating-point precision of 12:
        //
        //   Log10ofMantissa =
@@ -4043,12 +4037,9 @@ SelectionDAGBuilder::visitLog10(const CallInst &I) {
        SDValue t3 = DAG.getNode(ISD::FADD, dl, MVT::f32, t2,
                                 getF32Constant(DAG, 0x3f6ae232));
        SDValue t4 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t3, X);
-      SDValue Log10ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t4,
-                                            getF32Constant(DAG, 0x3f25f7c3));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, Log10ofMantissa);
-    } else { // LimitFloatPrecision > 12 && LimitFloatPrecision <= 18
+      Log10ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t4,
+                                    getF32Constant(DAG, 0x3f25f7c3));
+    } else { // LimitFloatPrecision <= 18
        // For floating-point precision of 18:
        //
        //   Log10ofMantissa =
@@ -4073,33 +4064,23 @@ SelectionDAGBuilder::visitLog10(const CallInst &I) {
        SDValue t7 = DAG.getNode(ISD::FADD, dl, MVT::f32, t6,
                                 getF32Constant(DAG, 0x3fc4316c));
        SDValue t8 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t7, X);
-      SDValue Log10ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t8,
-                                            getF32Constant(DAG, 0x3f57ce70));
-
-      result = DAG.getNode(ISD::FADD, dl,
-                           MVT::f32, LogOfExponent, Log10ofMantissa);
+      Log10ofMantissa = DAG.getNode(ISD::FSUB, dl, MVT::f32, t8,
+                                    getF32Constant(DAG, 0x3f57ce70));
      }
-  } else {
-    // No special expansion.
-    result = DAG.getNode(ISD::FLOG10, dl,
-                         getValue(I.getArgOperand(0)).getValueType(),
-                         getValue(I.getArgOperand(0)));
+
+    return DAG.getNode(ISD::FADD, dl, MVT::f32, LogOfExponent, Log10ofMantissa);
    }
  
-  setValue(&I, result);
+  // No special expansion.
+  return DAG.getNode(ISD::FLOG10, dl, Op.getValueType(), Op);
  }
  
-/// visitExp2 - Lower an exp2 intrinsic. Handles the special sequences for
+/// expandExp2 - Lower an exp2 intrinsic. Handles the special sequences for
  /// limited-precision mode.
-void
-SelectionDAGBuilder::visitExp2(const CallInst &I) {
-  SDValue result;
-  DebugLoc dl = getCurDebugLoc();
-
-  if (getValue(I.getArgOperand(0)).getValueType() == MVT::f32 &&
+static SDValue expandExp2(DebugLoc dl, SDValue Op, SelectionDAG &DAG,
+                          const TargetLowering &TLI) {
+  if (Op.getValueType() == MVT::f32 &&
        LimitFloatPrecision > 0 && LimitFloatPrecision <= 18) {
-    SDValue Op = getValue(I.getArgOperand(0));
-
      SDValue IntegerPartOfX = DAG.getNode(ISD::FP_TO_SINT, dl, MVT::i32, Op);
  
      //   FractionalPartOfX = x - (float)IntegerPartOfX;
@@ -4110,6 +4091,7 @@ SelectionDAGBuilder::visitExp2(const CallInst &I) {
      IntegerPartOfX = DAG.getNode(ISD::SHL, dl, MVT::i32, IntegerPartOfX,
                                   DAG.getConstant(23, TLI.getPointerTy()));
  
+    SDValue TwoToFractionalPartOfX;
      if (LimitFloatPrecision <= 6) {
        // For floating-point precision of 6:
        //
@@ -4123,15 +4105,9 @@ SelectionDAGBuilder::visitExp2(const CallInst &I) {
        SDValue t3 = DAG.getNode(ISD::FADD, dl, MVT::f32, t2,
                                 getF32Constant(DAG, 0x3f3c50c8));
        SDValue t4 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t3, X);
-      SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
-                               getF32Constant(DAG, 0x3f7f5e7e));
-      SDValue t6 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, t5);
-      SDValue TwoToFractionalPartOfX =
-        DAG.getNode(ISD::ADD, dl, MVT::i32, t6, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl,
-                           MVT::f32, TwoToFractionalPartOfX);
-    } else if (LimitFloatPrecision > 6 && LimitFloatPrecision <= 12) {
+      TwoToFractionalPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
+                                           getF32Constant(DAG, 0x3f7f5e7e));
+    } else if (LimitFloatPrecision <= 12) {
        // For floating-point precision of 12:
        //
        //   TwoToFractionalPartOfX =
@@ -4148,15 +4124,9 @@ SelectionDAGBuilder::visitExp2(const CallInst &I) {
        SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
                                 getF32Constant(DAG, 0x3f324b07));
        SDValue t6 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t5, X);
-      SDValue t7 = DAG.getNode(ISD::FADD, dl, MVT::f32, t6,
-                               getF32Constant(DAG, 0x3f7ff8fd));
-      SDValue t8 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, t7);
-      SDValue TwoToFractionalPartOfX =
-        DAG.getNode(ISD::ADD, dl, MVT::i32, t8, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl,
-                           MVT::f32, TwoToFractionalPartOfX);
-    } else { // LimitFloatPrecision > 12 && LimitFloatPrecision <= 18
+      TwoToFractionalPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t6,
+                                           getF32Constant(DAG, 0x3f7ff8fd));
+    } else { // LimitFloatPrecision <= 18
        // For floating-point precision of 18:
        //
        //   TwoToFractionalPartOfX =
@@ -4184,54 +4154,42 @@ SelectionDAGBuilder::visitExp2(const CallInst &I) {
        SDValue t11 = DAG.getNode(ISD::FADD, dl, MVT::f32, t10,
                                  getF32Constant(DAG, 0x3f317234));
        SDValue t12 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t11, X);
-      SDValue t13 = DAG.getNode(ISD::FADD, dl, MVT::f32, t12,
-                                getF32Constant(DAG, 0x3f800000));
-      SDValue t14 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, t13);
-      SDValue TwoToFractionalPartOfX =
-        DAG.getNode(ISD::ADD, dl, MVT::i32, t14, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl,
-                           MVT::f32, TwoToFractionalPartOfX);
+      TwoToFractionalPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t12,
+                                           getF32Constant(DAG, 0x3f800000));
      }
-  } else {
-    // No special expansion.
-    result = DAG.getNode(ISD::FEXP2, dl,
-                         getValue(I.getArgOperand(0)).getValueType(),
-                         getValue(I.getArgOperand(0)));
+
+    // Add the exponent into the result in integer domain.
+    SDValue t13 = DAG.getNode(ISD::BITCAST, dl, MVT::i32,
+                              TwoToFractionalPartOfX);
+    return DAG.getNode(ISD::BITCAST, dl, MVT::f32,
+                       DAG.getNode(ISD::ADD, dl, MVT::i32,
+                                   t13, IntegerPartOfX));
    }
  
-  setValue(&I, result);
+  // No special expansion.
+  return DAG.getNode(ISD::FEXP2, dl, Op.getValueType(), Op);
  }
  
  /// visitPow - Lower a pow intrinsic. Handles the special sequences for
  /// limited-precision mode with x == 10.0f.
-void
-SelectionDAGBuilder::visitPow(const CallInst &I) {
-  SDValue result;
-  const Value *Val = I.getArgOperand(0);
-  DebugLoc dl = getCurDebugLoc();
+static SDValue expandPow(DebugLoc dl, SDValue LHS, SDValue RHS,
+                         SelectionDAG &DAG, const TargetLowering &TLI) {
    bool IsExp10 = false;
-
-  if (getValue(Val).getValueType() == MVT::f32 &&
-      getValue(I.getArgOperand(1)).getValueType() == MVT::f32 &&
+  if (LHS.getValueType() == MVT::f32 && LHS.getValueType() == MVT::f32 &&
        LimitFloatPrecision > 0 && LimitFloatPrecision <= 18) {
-    if (Constant *C = const_cast<Constant*>(dyn_cast<Constant>(Val))) {
-      if (ConstantFP *CFP = dyn_cast<ConstantFP>(C)) {
-        APFloat Ten(10.0f);
-        IsExp10 = CFP->getValueAPF().bitwiseIsEqual(Ten);
-      }
+    if (ConstantFPSDNode *LHSC = dyn_cast<ConstantFPSDNode>(LHS)) {
+      APFloat Ten(10.0f);
+      IsExp10 = LHSC->isExactlyValue(Ten);
      }
    }
  
-  if (IsExp10 && LimitFloatPrecision > 0 && LimitFloatPrecision <= 18) {
-    SDValue Op = getValue(I.getArgOperand(1));
-
+  if (IsExp10) {
      // Put the exponent in the right bit position for later addition to the
      // final result:
      //
      //   #define LOG2OF10 3.3219281f
      //   IntegerPartOfX = (int32_t)(x * LOG2OF10);
-    SDValue t0 = DAG.getNode(ISD::FMUL, dl, MVT::f32, Op,
+    SDValue t0 = DAG.getNode(ISD::FMUL, dl, MVT::f32, RHS,
                               getF32Constant(DAG, 0x40549a78));
      SDValue IntegerPartOfX = DAG.getNode(ISD::FP_TO_SINT, dl, MVT::i32, t0);
  
@@ -4243,6 +4201,7 @@ SelectionDAGBuilder::visitPow(const CallInst &I) {
      IntegerPartOfX = DAG.getNode(ISD::SHL, dl, MVT::i32, IntegerPartOfX,
                                   DAG.getConstant(23, TLI.getPointerTy()));
  
+    SDValue TwoToFractionalPartOfX;
      if (LimitFloatPrecision <= 6) {
        // For floating-point precision of 6:
        //
@@ -4256,15 +4215,9 @@ SelectionDAGBuilder::visitPow(const CallInst &I) {
        SDValue t3 = DAG.getNode(ISD::FADD, dl, MVT::f32, t2,
                                 getF32Constant(DAG, 0x3f3c50c8));
        SDValue t4 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t3, X);
-      SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
-                               getF32Constant(DAG, 0x3f7f5e7e));
-      SDValue t6 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, t5);
-      SDValue TwoToFractionalPartOfX =
-        DAG.getNode(ISD::ADD, dl, MVT::i32, t6, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl,
-                           MVT::f32, TwoToFractionalPartOfX);
-    } else if (LimitFloatPrecision > 6 && LimitFloatPrecision <= 12) {
+      TwoToFractionalPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
+                                           getF32Constant(DAG, 0x3f7f5e7e));
+    } else if (LimitFloatPrecision <= 12) {
        // For floating-point precision of 12:
        //
        //   TwoToFractionalPartOfX =
@@ -4281,15 +4234,9 @@ SelectionDAGBuilder::visitPow(const CallInst &I) {
        SDValue t5 = DAG.getNode(ISD::FADD, dl, MVT::f32, t4,
                                 getF32Constant(DAG, 0x3f324b07));
        SDValue t6 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t5, X);
-      SDValue t7 = DAG.getNode(ISD::FADD, dl, MVT::f32, t6,
-                               getF32Constant(DAG, 0x3f7ff8fd));
-      SDValue t8 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, t7);
-      SDValue TwoToFractionalPartOfX =
-        DAG.getNode(ISD::ADD, dl, MVT::i32, t8, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl,
-                           MVT::f32, TwoToFractionalPartOfX);
-    } else { // LimitFloatPrecision > 12 && LimitFloatPrecision <= 18
+      TwoToFractionalPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t6,
+                                           getF32Constant(DAG, 0x3f7ff8fd));
+    } else { // LimitFloatPrecision <= 18
        // For floating-point precision of 18:
        //
        //   TwoToFractionalPartOfX =
@@ -4317,24 +4264,18 @@ SelectionDAGBuilder::visitPow(const CallInst &I) {
        SDValue t11 = DAG.getNode(ISD::FADD, dl, MVT::f32, t10,
                                  getF32Constant(DAG, 0x3f317234));
        SDValue t12 = DAG.getNode(ISD::FMUL, dl, MVT::f32, t11, X);
-      SDValue t13 = DAG.getNode(ISD::FADD, dl, MVT::f32, t12,
-                                getF32Constant(DAG, 0x3f800000));
-      SDValue t14 = DAG.getNode(ISD::BITCAST, dl, MVT::i32, t13);
-      SDValue TwoToFractionalPartOfX =
-        DAG.getNode(ISD::ADD, dl, MVT::i32, t14, IntegerPartOfX);
-
-      result = DAG.getNode(ISD::BITCAST, dl,
-                           MVT::f32, TwoToFractionalPartOfX);
+      TwoToFractionalPartOfX = DAG.getNode(ISD::FADD, dl, MVT::f32, t12,
+                                           getF32Constant(DAG, 0x3f800000));
      }
-  } else {
-    // No special expansion.
-    result = DAG.getNode(ISD::FPOW, dl,
-                         getValue(I.getArgOperand(0)).getValueType(),
-                         getValue(I.getArgOperand(0)),
-                         getValue(I.getArgOperand(1)));
+
+    SDValue t13 = DAG.getNode(ISD::BITCAST, dl,MVT::i32,TwoToFractionalPartOfX);
+    return DAG.getNode(ISD::BITCAST, dl, MVT::f32,
+                       DAG.getNode(ISD::ADD, dl, MVT::i32,
+                                   t13, IntegerPartOfX));
    }
  
-  setValue(&I, result);
+  // No special expansion.
+  return DAG.getNode(ISD::FPOW, dl, LHS.getValueType(), LHS, RHS);
  }
  
  
@@ -4355,7 +4296,8 @@ static SDValue ExpandPowI(DebugLoc DL, SDValue LHS, SDValue RHS,
        return DAG.getConstantFP(1.0, LHS.getValueType());
  
      const Function *F = DAG.getMachineFunction().getFunction();
-    if (!F->hasFnAttr(Attribute::OptimizeForSize) ||
+    if (!F->getAttributes().hasAttribute(AttributeSet::FunctionIndex,
+                                         Attribute::OptimizeForSize) ||
          // If optimizing for size, don't insert too many multiplies.  This
          // inserts up to 5 multiplies.
          CountPopulation_32(Val)+Log2_32(Val) < 7) {
@@ -4828,7 +4770,6 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
      // the sse2/mmx shift instructions reads 64 bits. Set the upper 32 bits
      // to be zero.
      // We must do this early because v2i32 is not a legal type.
-    DebugLoc dl = getCurDebugLoc();
      SDValue ShOps[2];
      ShOps[0] = ShAmt;
      ShOps[1] = DAG.getConstant(0, MVT::i32);
@@ -4845,7 +4786,6 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
    case Intrinsic::x86_avx_vinsertf128_ps_256:
    case Intrinsic::x86_avx_vinsertf128_si_256:
    case Intrinsic::x86_avx2_vinserti128: {
-    DebugLoc dl = getCurDebugLoc();
      EVT DestVT = TLI.getValueType(I.getType());
      EVT ElVT = TLI.getValueType(I.getArgOperand(1)->getType());
      uint64_t Idx = (cast<ConstantInt>(I.getArgOperand(2))->getZExtValue() & 1) *
@@ -4853,7 +4793,20 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
      Res = DAG.getNode(ISD::INSERT_SUBVECTOR, dl, DestVT,
                        getValue(I.getArgOperand(0)),
                        getValue(I.getArgOperand(1)),
-                      DAG.getConstant(Idx, MVT::i32));
+                      DAG.getIntPtrConstant(Idx));
+    setValue(&I, Res);
+    return 0;
+  }
+  case Intrinsic::x86_avx_vextractf128_pd_256:
+  case Intrinsic::x86_avx_vextractf128_ps_256:
+  case Intrinsic::x86_avx_vextractf128_si_256:
+  case Intrinsic::x86_avx2_vextracti128: {
+    EVT DestVT = TLI.getValueType(I.getType());
+    uint64_t Idx = (cast<ConstantInt>(I.getArgOperand(1))->getZExtValue() & 1) *
+                   DestVT.getVectorNumElements();
+    Res = DAG.getNode(ISD::EXTRACT_SUBVECTOR, dl, DestVT,
+                      getValue(I.getArgOperand(0)),
+                      DAG.getIntPtrConstant(Idx));
      setValue(&I, Res);
      return 0;
    }
@@ -4881,7 +4834,7 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
      }
      EVT DestVT = TLI.getValueType(I.getType());
      const Value *Op1 = I.getArgOperand(0);
-    Res = DAG.getConvertRndSat(DestVT, getCurDebugLoc(), getValue(Op1),
+    Res = DAG.getConvertRndSat(DestVT, dl, getValue(Op1),
                                 DAG.getValueType(DestVT),
                                 DAG.getValueType(getValue(Op1).getValueType()),
                                 getValue(I.getArgOperand(1)),
@@ -4890,53 +4843,57 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
      setValue(&I, Res);
      return 0;
    }
-  case Intrinsic::sqrt:
-    setValue(&I, DAG.getNode(ISD::FSQRT, dl,
-                             getValue(I.getArgOperand(0)).getValueType(),
-                             getValue(I.getArgOperand(0))));
-    return 0;
    case Intrinsic::powi:
      setValue(&I, ExpandPowI(dl, getValue(I.getArgOperand(0)),
                              getValue(I.getArgOperand(1)), DAG));
      return 0;
-  case Intrinsic::sin:
-    setValue(&I, DAG.getNode(ISD::FSIN, dl,
-                             getValue(I.getArgOperand(0)).getValueType(),
-                             getValue(I.getArgOperand(0))));
-    return 0;
-  case Intrinsic::cos:
-    setValue(&I, DAG.getNode(ISD::FCOS, dl,
-                             getValue(I.getArgOperand(0)).getValueType(),
-                             getValue(I.getArgOperand(0))));
-    return 0;
    case Intrinsic::log:
-    visitLog(I);
+    setValue(&I, expandLog(dl, getValue(I.getArgOperand(0)), DAG, TLI));
      return 0;
    case Intrinsic::log2:
-    visitLog2(I);
+    setValue(&I, expandLog2(dl, getValue(I.getArgOperand(0)), DAG, TLI));
      return 0;
    case Intrinsic::log10:
-    visitLog10(I);
+    setValue(&I, expandLog10(dl, getValue(I.getArgOperand(0)), DAG, TLI));
      return 0;
    case Intrinsic::exp:
-    visitExp(I);
+    setValue(&I, expandExp(dl, getValue(I.getArgOperand(0)), DAG, TLI));
      return 0;
    case Intrinsic::exp2:
-    visitExp2(I);
+    setValue(&I, expandExp2(dl, getValue(I.getArgOperand(0)), DAG, TLI));
      return 0;
    case Intrinsic::pow:
-    visitPow(I);
+    setValue(&I, expandPow(dl, getValue(I.getArgOperand(0)),
+                           getValue(I.getArgOperand(1)), DAG, TLI));
      return 0;
+  case Intrinsic::sqrt:
    case Intrinsic::fabs:
-    setValue(&I, DAG.getNode(ISD::FABS, dl,
-                             getValue(I.getArgOperand(0)).getValueType(),
-                             getValue(I.getArgOperand(0))));
-    return 0;
+  case Intrinsic::sin:
+  case Intrinsic::cos:
    case Intrinsic::floor:
-    setValue(&I, DAG.getNode(ISD::FFLOOR, dl,
+  case Intrinsic::ceil:
+  case Intrinsic::trunc:
+  case Intrinsic::rint:
+  case Intrinsic::nearbyint: {
+    unsigned Opcode;
+    switch (Intrinsic) {
+    default: llvm_unreachable("Impossible intrinsic");  // Can't reach here.
+    case Intrinsic::sqrt:      Opcode = ISD::FSQRT;      break;
+    case Intrinsic::fabs:      Opcode = ISD::FABS;       break;
+    case Intrinsic::sin:       Opcode = ISD::FSIN;       break;
+    case Intrinsic::cos:       Opcode = ISD::FCOS;       break;
+    case Intrinsic::floor:     Opcode = ISD::FFLOOR;     break;
+    case Intrinsic::ceil:      Opcode = ISD::FCEIL;      break;
+    case Intrinsic::trunc:     Opcode = ISD::FTRUNC;     break;
+    case Intrinsic::rint:      Opcode = ISD::FRINT;      break;
+    case Intrinsic::nearbyint: Opcode = ISD::FNEARBYINT; break;
+    }
+
+    setValue(&I, DAG.getNode(Opcode, dl,
                               getValue(I.getArgOperand(0)).getValueType(),
                               getValue(I.getArgOperand(0))));
      return 0;
+  }
    case Intrinsic::fma:
      setValue(&I, DAG.getNode(ISD::FMA, dl,
                               getValue(I.getArgOperand(0)).getValueType(),
@@ -4947,7 +4904,7 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
    case Intrinsic::fmuladd: {
      EVT VT = TLI.getValueType(I.getType());
      if (TM.Options.AllowFPOpFusion != FPOpFusion::Strict &&
-        TLI.isOperationLegal(ISD::FMA, VT) &&
+        TLI.isOperationLegalOrCustom(ISD::FMA, VT) &&
          TLI.isFMAFasterThanMulAndAdd(VT)){
        setValue(&I, DAG.getNode(ISD::FMA, dl,
                                 getValue(I.getArgOperand(0)).getValueType(),
@@ -5044,7 +5001,7 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
      SDValue FIN = DAG.getFrameIndex(FI, PtrTy);
  
      // Store the stack protector onto the stack.
-    Res = DAG.getStore(getRoot(), getCurDebugLoc(), Src, FIN,
+    Res = DAG.getStore(getRoot(), dl, Src, FIN,
                         MachinePointerInfo::getFixedStack(FI),
                         true, false, 0);
      setValue(&I, Res);
@@ -5116,10 +5073,13 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
      return 0;
    }
  
+  case Intrinsic::debugtrap:
    case Intrinsic::trap: {
      StringRef TrapFuncName = TM.Options.getTrapFunctionName();
      if (TrapFuncName.empty()) {
-      DAG.setRoot(DAG.getNode(ISD::TRAP, dl,MVT::Other, getRoot()));
+      ISD::NodeType Op = (Intrinsic == Intrinsic::trap) ? 
+        ISD::TRAP : ISD::DEBUGTRAP;
+      DAG.setRoot(DAG.getNode(Op, dl,MVT::Other, getRoot()));
        return 0;
      }
      TargetLowering::ArgListTy Args;
@@ -5129,15 +5089,12 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
                   /*isTailCall=*/false,
                   /*doesNotRet=*/false, /*isReturnValueUsed=*/true,
                   DAG.getExternalSymbol(TrapFuncName.data(), TLI.getPointerTy()),
-                 Args, DAG, getCurDebugLoc());
+                 Args, DAG, dl);
      std::pair<SDValue, SDValue> Result = TLI.LowerCallTo(CLI);
      DAG.setRoot(Result.second);
      return 0;
    }
-  case Intrinsic::debugtrap: {
-    DAG.setRoot(DAG.getNode(ISD::DEBUGTRAP, dl,MVT::Other, getRoot()));
-    return 0;
-  }
+
    case Intrinsic::uadd_with_overflow:
    case Intrinsic::sadd_with_overflow:
    case Intrinsic::usub_with_overflow:
@@ -5158,7 +5115,7 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
      SDValue Op2 = getValue(I.getArgOperand(1));
  
      SDVTList VTs = DAG.getVTList(Op1.getValueType(), MVT::i1);
-    setValue(&I, DAG.getNode(Op, getCurDebugLoc(), VTs, Op1, Op2));
+    setValue(&I, DAG.getNode(Op, dl, VTs, Op1, Op2));
      return 0;
    }
    case Intrinsic::prefetch: {
@@ -5180,14 +5137,40 @@ SelectionDAGBuilder::visitIntrinsicCall(const CallInst &I, unsigned Intrinsic) {
                                          rw==1)); /* write */
      return 0;
    }
+  case Intrinsic::lifetime_start:
+  case Intrinsic::lifetime_end: {
+    bool IsStart = (Intrinsic == Intrinsic::lifetime_start);
+    // Stack coloring is not enabled in O0, discard region information.
+    if (TM.getOptLevel() == CodeGenOpt::None)
+      return 0;
  
+    SmallVector<Value *, 4> Allocas;
+    GetUnderlyingObjects(I.getArgOperand(1), Allocas, TD);
+
+    for (SmallVector<Value*, 4>::iterator Object = Allocas.begin(),
+         E = Allocas.end(); Object != E; ++Object) {
+      AllocaInst *LifetimeObject = dyn_cast_or_null<AllocaInst>(*Object);
+
+      // Could not find an Alloca.
+      if (!LifetimeObject)
+        continue;
+
+      int FI = FuncInfo.StaticAllocaMap[LifetimeObject];
+
+      SDValue Ops[2];
+      Ops[0] = getRoot();
+      Ops[1] = DAG.getFrameIndex(FI, TLI.getPointerTy(), true);
+      unsigned Opcode = (IsStart ? ISD::LIFETIME_START : ISD::LIFETIME_END);
+
+      Res = DAG.getNode(Opcode, dl, MVT::Other, Ops, 2);
+      DAG.setRoot(Res);
+    }
+  }
    case Intrinsic::invariant_start:
-  case Intrinsic::lifetime_start:
      // Discard region information.
      setValue(&I, DAG.getUNDEF(TLI.getPointerTy()));
      return 0;
    case Intrinsic::invariant_end:
-  case Intrinsic::lifetime_end:
      // Discard region information.
      return 0;
    case Intrinsic::donothing:
@@ -5223,9 +5206,9 @@ void SelectionDAGBuilder::LowerCallTo(ImmutableCallSite CS, SDValue Callee,
    int DemoteStackIdx = -100;
  
    if (!CanLowerReturn) {
-    uint64_t TySize = TLI.getTargetData()->getTypeAllocSize(
+    uint64_t TySize = TLI.getDataLayout()->getTypeAllocSize(
                        FTy->getReturnType());
-    unsigned Align  = TLI.getTargetData()->getPrefTypeAlignment(
+    unsigned Align  = TLI.getDataLayout()->getPrefTypeAlignment(
                        FTy->getReturnType());
      MachineFunction &MF = DAG.getMachineFunction();
      DemoteStackIdx = MF.getFrameInfo()->CreateStackObject(TySize, Align, false);
@@ -5295,11 +5278,6 @@ void SelectionDAGBuilder::LowerCallTo(ImmutableCallSite CS, SDValue Callee,
        !isInTailCallPosition(CS, CS.getAttributes().getRetAttributes(), TLI))
      isTailCall = false;
  
-  // If there's a possibility that fast-isel has already selected some amount
-  // of the current basic block, don't emit a tail call.
-  if (isTailCall && TM.Options.EnableFastISel)
-    isTailCall = false;
-
    TargetLowering::
    CallLoweringInfo CLI(getRoot(), RetTy, FTy, isTailCall, Callee, Args, DAG,
                         getCurDebugLoc(), CS);
@@ -5690,7 +5668,7 @@ public:
    /// MVT::Other.
    EVT getCallOperandValEVT(LLVMContext &Context,
                             const TargetLowering &TLI,
-                           const TargetData *TD) const {
+                           const DataLayout *TD) const {
      if (CallOperandVal == 0) return MVT::Other;
  
      if (isa<BasicBlock>(CallOperandVal))
@@ -5771,7 +5749,7 @@ static void GetRegistersForValue(SelectionDAG &DAG,
        // Try to convert to the first EVT that the reg class contains.  If the
        // types are identical size, use a bitcast to convert (e.g. two differing
        // vector types).
-      EVT RegVT = *PhysReg.second->vt_begin();
+      MVT RegVT = *PhysReg.second->vt_begin();
        if (RegVT.getSizeInBits() == OpInfo.ConstraintVT.getSizeInBits()) {
          OpInfo.CallOperand = DAG.getNode(ISD::BITCAST, DL,
                                           RegVT, OpInfo.CallOperand);
@@ -5781,8 +5759,7 @@ static void GetRegistersForValue(SelectionDAG &DAG,
          // bitcast to the corresponding integer type.  This turns an f64 value
          // into i64, which can be passed with two i32 values on a 32-bit
          // machine.
-        RegVT = EVT::getIntegerVT(Context,
-                                  OpInfo.ConstraintVT.getSizeInBits());
+        RegVT = MVT::getIntegerVT(OpInfo.ConstraintVT.getSizeInBits());
          OpInfo.CallOperand = DAG.getNode(ISD::BITCAST, DL,
                                           RegVT, OpInfo.CallOperand);
          OpInfo.ConstraintVT = RegVT;
@@ -5792,7 +5769,7 @@ static void GetRegistersForValue(SelectionDAG &DAG,
      NumRegs = TLI.getNumRegisters(Context, OpInfo.ConstraintVT);
    }
  
-  EVT RegVT;
+  MVT RegVT;
    EVT ValueVT = OpInfo.ConstraintVT;
  
    // If this is a constraint for a specific physical register, like {r17},
@@ -5866,7 +5843,7 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
      ConstraintOperands.push_back(SDISelAsmOperandInfo(TargetConstraints[i]));
      SDISelAsmOperandInfo &OpInfo = ConstraintOperands.back();
  
-    EVT OpVT = MVT::Other;
+    MVT OpVT = MVT::Other;
  
      // Compute the value type for each operand.
      switch (OpInfo.Type) {
@@ -5881,10 +5858,10 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
        // corresponding argument.
        assert(!CS.getType()->isVoidTy() && "Bad inline asm!");
        if (StructType *STy = dyn_cast<StructType>(CS.getType())) {
-        OpVT = TLI.getValueType(STy->getElementType(ResNo));
+        OpVT = TLI.getSimpleValueType(STy->getElementType(ResNo));
        } else {
          assert(ResNo == 0 && "Asm only has one result!");
-        OpVT = TLI.getValueType(CS.getType());
+        OpVT = TLI.getSimpleValueType(CS.getType());
        }
        ++ResNo;
        break;
@@ -5905,7 +5882,8 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
          OpInfo.CallOperand = getValue(OpInfo.CallOperandVal);
        }
  
-      OpVT = OpInfo.getCallOperandValEVT(*DAG.getContext(), TLI, TD);
+      OpVT = OpInfo.getCallOperandValEVT(*DAG.getContext(), TLI, TD).
+        getSimpleVT();
      }
  
      OpInfo.ConstraintVT = OpVT;
@@ -5994,8 +5972,8 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
          // Otherwise, create a stack slot and emit a store to it before the
          // asm.
          Type *Ty = OpVal->getType();
-        uint64_t TySize = TLI.getTargetData()->getTypeAllocSize(Ty);
-        unsigned Align  = TLI.getTargetData()->getPrefTypeAlignment(Ty);
+        uint64_t TySize = TLI.getDataLayout()->getTypeAllocSize(Ty);
+        unsigned Align  = TLI.getDataLayout()->getPrefTypeAlignment(Ty);
          MachineFunction &MF = DAG.getMachineFunction();
          int SSFI = MF.getFrameInfo()->CreateStackObject(TySize, Align, false);
          SDValue StackSlot = DAG.getFrameIndex(SSFI, TLI.getPointerTy());
@@ -6043,12 +6021,36 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
    const MDNode *SrcLoc = CS.getInstruction()->getMetadata("srcloc");
    AsmNodeOperands.push_back(DAG.getMDNode(SrcLoc));
  
-  // Remember the HasSideEffect and AlignStack bits as operand 3.
+  // Remember the HasSideEffect, AlignStack, AsmDialect, MayLoad and MayStore
+  // bits as operand 3.
    unsigned ExtraInfo = 0;
    if (IA->hasSideEffects())
      ExtraInfo |= InlineAsm::Extra_HasSideEffects;
    if (IA->isAlignStack())
      ExtraInfo |= InlineAsm::Extra_IsAlignStack;
+  // Set the asm dialect.
+  ExtraInfo |= IA->getDialect() * InlineAsm::Extra_AsmDialect;
+
+  // Determine if this InlineAsm MayLoad or MayStore based on the constraints.
+  for (unsigned i = 0, e = TargetConstraints.size(); i != e; ++i) {
+    TargetLowering::AsmOperandInfo &OpInfo = TargetConstraints[i];
+
+    // Compute the constraint code and ConstraintType to use.
+    TLI.ComputeConstraintToUse(OpInfo, SDValue());
+
+    // Ideally, we would only check against memory constraints.  However, the
+    // meaning of an other constraint can be target-specific and we can't easily
+    // reason about it.  Therefore, be conservative and set MayLoad/MayStore
+    // for other constriants as well.
+    if (OpInfo.ConstraintType == TargetLowering::C_Memory ||
+        OpInfo.ConstraintType == TargetLowering::C_Other) {
+      if (OpInfo.Type == InlineAsm::isInput)
+        ExtraInfo |= InlineAsm::Extra_MayLoad;
+      else if (OpInfo.Type == InlineAsm::isOutput)
+        ExtraInfo |= InlineAsm::Extra_MayStore;
+    }
+  }
+
    AsmNodeOperands.push_back(DAG.getTargetConstant(ExtraInfo,
                                                    TLI.getPointerTy()));
  
@@ -6148,7 +6150,7 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
  
            RegsForValue MatchedRegs;
            MatchedRegs.ValueVTs.push_back(InOperandVal.getValueType());
-          EVT RegVT = AsmNodeOperands[CurOp+1].getValueType();
+          MVT RegVT = AsmNodeOperands[CurOp+1].getSimpleValueType();
            MatchedRegs.RegVTs.push_back(RegVT);
            MachineRegisterInfo &RegInfo = DAG.getMachineFunction().getRegInfo();
            for (unsigned i = 0, e = InlineAsm::getNumOperandRegisters(OpFlag);
@@ -6158,7 +6160,7 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
  
            // Use the produced MatchedRegs object to
            MatchedRegs.getCopyToRegs(InOperandVal, DAG, getCurDebugLoc(),
-                                    Chain, &Flag);
+                                    Chain, &Flag, CS.getInstruction());
            MatchedRegs.AddInlineAsmOperands(InlineAsm::Kind_RegUse,
                                             true, OpInfo.getMatchedOperand(),
                                             DAG, AsmNodeOperands);
@@ -6240,7 +6242,7 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
        }
  
        OpInfo.AssignedRegs.getCopyToRegs(InOperandVal, DAG, getCurDebugLoc(),
-                                        Chain, &Flag);
+                                        Chain, &Flag, CS.getInstruction());
  
        OpInfo.AssignedRegs.AddInlineAsmOperands(InlineAsm::Kind_RegUse, false, 0,
                                                 DAG, AsmNodeOperands);
@@ -6271,7 +6273,7 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
    // and set it as the value of the call.
    if (!RetValRegs.Regs.empty()) {
      SDValue Val = RetValRegs.getCopyFromRegs(DAG, FuncInfo, getCurDebugLoc(),
-                                             Chain, &Flag);
+                                             Chain, &Flag, CS.getInstruction());
  
      // FIXME: Why don't we do this for inline asms with MRVs?
      if (CS.getType()->isSingleValueType() && CS.getType()->isSized()) {
@@ -6311,7 +6313,7 @@ void SelectionDAGBuilder::visitInlineAsm(ImmutableCallSite CS) {
      RegsForValue &OutRegs = IndirectStoresToEmit[i].first;
      const Value *Ptr = IndirectStoresToEmit[i].second;
      SDValue OutVal = OutRegs.getCopyFromRegs(DAG, FuncInfo, getCurDebugLoc(),
-                                             Chain, &Flag);
+                                             Chain, &Flag, IA);
      StoresToEmit.push_back(std::make_pair(OutVal, Ptr));
    }
  
@@ -6341,7 +6343,7 @@ void SelectionDAGBuilder::visitVAStart(const CallInst &I) {
  }
  
  void SelectionDAGBuilder::visitVAArg(const VAArgInst &I) {
-  const TargetData &TD = *TLI.getTargetData();
+  const DataLayout &TD = *TLI.getDataLayout();
    SDValue V = DAG.getVAArg(TLI.getValueType(I.getType()), getCurDebugLoc(),
                             getRoot(), getValue(I.getOperand(0)),
                             DAG.getSrcValue(I.getOperand(0)),
@@ -6387,7 +6389,7 @@ TargetLowering::LowerCallTo(TargetLowering::CallLoweringInfo &CLI) const {
                             Args[i].Node.getResNo() + Value);
        ISD::ArgFlagsTy Flags;
        unsigned OriginalAlignment =
-        getTargetData()->getABITypeAlignment(ArgTy);
+        getDataLayout()->getABITypeAlignment(ArgTy);
  
        if (Args[i].isZExt)
          Flags.setZExt();
@@ -6401,7 +6403,7 @@ TargetLowering::LowerCallTo(TargetLowering::CallLoweringInfo &CLI) const {
          Flags.setByVal();
          PointerType *Ty = cast<PointerType>(Args[i].Ty);
          Type *ElementTy = Ty->getElementType();
-        Flags.setByValSize(getTargetData()->getTypeAllocSize(ElementTy));
+        Flags.setByValSize(getDataLayout()->getTypeAllocSize(ElementTy));
          // For ByVal, alignment should come from FE.  BE will guess if this
          // info is not there but there are cases it cannot get right.
          unsigned FrameAlign;
@@ -6415,7 +6417,7 @@ TargetLowering::LowerCallTo(TargetLowering::CallLoweringInfo &CLI) const {
          Flags.setNest();
        Flags.setOrigAlign(OriginalAlignment);
  
-      EVT PartVT = getRegisterType(CLI.RetTy->getContext(), VT);
+      MVT PartVT = getRegisterType(CLI.RetTy->getContext(), VT);
        unsigned NumParts = getNumRegisters(CLI.RetTy->getContext(), VT);
        SmallVector<SDValue, 4> Parts(NumParts);
        ISD::NodeType ExtendKind = ISD::ANY_EXTEND;
@@ -6426,12 +6428,13 @@ TargetLowering::LowerCallTo(TargetLowering::CallLoweringInfo &CLI) const {
          ExtendKind = ISD::ZERO_EXTEND;
  
        getCopyToParts(CLI.DAG, CLI.DL, Op, &Parts[0], NumParts,
-                     PartVT, ExtendKind);
+                     PartVT, CLI.CS ? CLI.CS->getInstruction() : 0, ExtendKind);
  
        for (unsigned j = 0; j != NumParts; ++j) {
          // if it isn't first piece, alignment must be 1
          ISD::OutputArg MyFlags(Flags, Parts[j].getValueType(),
-                               i < CLI.NumFixedArgs);
+                               i < CLI.NumFixedArgs,
+                               i, j*Parts[j].getValueType().getStoreSize());
          if (NumParts > 1 && j == 0)
            MyFlags.Flags.setSplit();
          else if (j != 0)
@@ -6449,11 +6452,11 @@ TargetLowering::LowerCallTo(TargetLowering::CallLoweringInfo &CLI) const {
    ComputeValueVTs(*this, CLI.RetTy, RetTys);
    for (unsigned I = 0, E = RetTys.size(); I != E; ++I) {
      EVT VT = RetTys[I];
-    EVT RegisterVT = getRegisterType(CLI.RetTy->getContext(), VT);
+    MVT RegisterVT = getRegisterType(CLI.RetTy->getContext(), VT);
      unsigned NumRegs = getNumRegisters(CLI.RetTy->getContext(), VT);
      for (unsigned i = 0; i != NumRegs; ++i) {
        ISD::InputArg MyFlags;
-      MyFlags.VT = RegisterVT.getSimpleVT();
+      MyFlags.VT = RegisterVT;
        MyFlags.Used = CLI.IsReturnValueUsed;
        if (CLI.RetSExt)
          MyFlags.Flags.setSExt();
@@ -6503,11 +6506,11 @@ TargetLowering::LowerCallTo(TargetLowering::CallLoweringInfo &CLI) const {
    unsigned CurReg = 0;
    for (unsigned I = 0, E = RetTys.size(); I != E; ++I) {
      EVT VT = RetTys[I];
-    EVT RegisterVT = getRegisterType(CLI.RetTy->getContext(), VT);
+    MVT RegisterVT = getRegisterType(CLI.RetTy->getContext(), VT);
      unsigned NumRegs = getNumRegisters(CLI.RetTy->getContext(), VT);
  
      ReturnValues.push_back(getCopyFromParts(CLI.DAG, CLI.DL, &InVals[CurReg],
-                                            NumRegs, RegisterVT, VT,
+                                            NumRegs, RegisterVT, VT, NULL,
                                              AssertOp));
      CurReg += NumRegs;
    }
@@ -6546,7 +6549,7 @@ SelectionDAGBuilder::CopyValueToVirtualRegister(const Value *V, unsigned Reg) {
  
    RegsForValue RFV(V->getContext(), TLI, Reg, V->getType());
    SDValue Chain = DAG.getEntryNode();
-  RFV.getCopyToRegs(Op, DAG, getCurDebugLoc(), Chain, 0);
+  RFV.getCopyToRegs(Op, DAG, getCurDebugLoc(), Chain, 0, V);
    PendingExports.push_back(Chain);
  }
  
@@ -6576,7 +6579,7 @@ void SelectionDAGISel::LowerArguments(const BasicBlock *LLVMBB) {
    const Function &F = *LLVMBB->getParent();
    SelectionDAG &DAG = SDB->DAG;
    DebugLoc dl = SDB->getCurDebugLoc();
-  const TargetData *TD = TLI.getTargetData();
+  const DataLayout *TD = TLI.getDataLayout();
    SmallVector<ISD::InputArg, 16> Ins;
  
    // Check whether the function can return without sret-demotion.
@@ -6593,8 +6596,8 @@ void SelectionDAGISel::LowerArguments(const BasicBlock *LLVMBB) {
      // or one register.
      ISD::ArgFlagsTy Flags;
      Flags.setSRet();
-    EVT RegisterVT = TLI.getRegisterType(*DAG.getContext(), ValueVTs[0]);
-    ISD::InputArg RetArg(Flags, RegisterVT, true);
+    MVT RegisterVT = TLI.getRegisterType(*DAG.getContext(), ValueVTs[0]);
+    ISD::InputArg RetArg(Flags, RegisterVT, true, 0, 0);
      Ins.push_back(RetArg);
    }
  
@@ -6613,15 +6616,15 @@ void SelectionDAGISel::LowerArguments(const BasicBlock *LLVMBB) {
        unsigned OriginalAlignment =
          TD->getABITypeAlignment(ArgTy);
  
-      if (F.paramHasAttr(Idx, Attribute::ZExt))
+      if (F.getParamAttributes(Idx).hasAttribute(Attribute::ZExt))
          Flags.setZExt();
-      if (F.paramHasAttr(Idx, Attribute::SExt))
+      if (F.getParamAttributes(Idx).hasAttribute(Attribute::SExt))
          Flags.setSExt();
-      if (F.paramHasAttr(Idx, Attribute::InReg))
+      if (F.getParamAttributes(Idx).hasAttribute(Attribute::InReg))
          Flags.setInReg();
-      if (F.paramHasAttr(Idx, Attribute::StructRet))
+      if (F.getParamAttributes(Idx).hasAttribute(Attribute::StructRet))
          Flags.setSRet();
-      if (F.paramHasAttr(Idx, Attribute::ByVal)) {
+      if (F.getParamAttributes(Idx).hasAttribute(Attribute::ByVal)) {
          Flags.setByVal();
          PointerType *Ty = cast<PointerType>(I->getType());
          Type *ElementTy = Ty->getElementType();
@@ -6635,14 +6638,15 @@ void SelectionDAGISel::LowerArguments(const BasicBlock *LLVMBB) {
            FrameAlign = TLI.getByValTypeAlignment(ElementTy);
          Flags.setByValAlign(FrameAlign);
        }
-      if (F.paramHasAttr(Idx, Attribute::Nest))
+      if (F.getParamAttributes(Idx).hasAttribute(Attribute::Nest))
          Flags.setNest();
        Flags.setOrigAlign(OriginalAlignment);
  
-      EVT RegisterVT = TLI.getRegisterType(*CurDAG->getContext(), VT);
+      MVT RegisterVT = TLI.getRegisterType(*CurDAG->getContext(), VT);
        unsigned NumRegs = TLI.getNumRegisters(*CurDAG->getContext(), VT);
        for (unsigned i = 0; i != NumRegs; ++i) {
-        ISD::InputArg MyFlags(Flags, RegisterVT, isArgValueUsed);
+        ISD::InputArg MyFlags(Flags, RegisterVT, isArgValueUsed,
+                              Idx-1, i*RegisterVT.getStoreSize());
          if (NumRegs > 1 && i == 0)
            MyFlags.Flags.setSplit();
          // if it isn't first piece, alignment must be 1
@@ -6684,11 +6688,11 @@ void SelectionDAGISel::LowerArguments(const BasicBlock *LLVMBB) {
      // from the sret argument into it.
      SmallVector<EVT, 1> ValueVTs;
      ComputeValueVTs(TLI, PointerType::getUnqual(F.getReturnType()), ValueVTs);
-    EVT VT = ValueVTs[0];
-    EVT RegVT = TLI.getRegisterType(*CurDAG->getContext(), VT);
+    MVT VT = ValueVTs[0].getSimpleVT();
+    MVT RegVT = TLI.getRegisterType(*CurDAG->getContext(), VT);
      ISD::NodeType AssertOp = ISD::DELETED_NODE;
      SDValue ArgValue = getCopyFromParts(DAG, dl, &InVals[0], 1,
-                                        RegVT, VT, AssertOp);
+                                        RegVT, VT, NULL, AssertOp);
  
      MachineFunction& MF = SDB->DAG.getMachineFunction();
      MachineRegisterInfo& RegInfo = MF.getRegInfo();
@@ -6717,19 +6721,19 @@ void SelectionDAGISel::LowerArguments(const BasicBlock *LLVMBB) {
  
      for (unsigned Val = 0; Val != NumValues; ++Val) {
        EVT VT = ValueVTs[Val];
-      EVT PartVT = TLI.getRegisterType(*CurDAG->getContext(), VT);
+      MVT PartVT = TLI.getRegisterType(*CurDAG->getContext(), VT);
        unsigned NumParts = TLI.getNumRegisters(*CurDAG->getContext(), VT);
  
        if (!I->use_empty()) {
          ISD::NodeType AssertOp = ISD::DELETED_NODE;
-        if (F.paramHasAttr(Idx, Attribute::SExt))
+        if (F.getParamAttributes(Idx).hasAttribute(Attribute::SExt))
            AssertOp = ISD::AssertSext;
-        else if (F.paramHasAttr(Idx, Attribute::ZExt))
+        else if (F.getParamAttributes(Idx).hasAttribute(Attribute::ZExt))
            AssertOp = ISD::AssertZext;
  
          ArgValues.push_back(getCopyFromParts(DAG, dl, &InVals[i],
                                               NumParts, PartVT, VT,
-                                             AssertOp));
+                                             NULL, AssertOp));
        }
  
        i += NumParts;