[PATCH] D32445: [LV] Handle external uses of floating-point induction variables
Matthew Simpson via Phabricator via llvm-commits
llvm-commits at lists.llvm.org
Mon Apr 24 10:56:40 PDT 2017
mssimpso created this revision.
Herald added a subscriber: mzolotukhin.
Reference: https://bugs.llvm.org/show_bug.cgi?id=32758
https://reviews.llvm.org/D32445
Files:
lib/Transforms/Vectorize/LoopVectorize.cpp
test/Transforms/LoopVectorize/float-induction.ll
Index: test/Transforms/LoopVectorize/float-induction.ll
===================================================================
--- test/Transforms/LoopVectorize/float-induction.ll
+++ test/Transforms/LoopVectorize/float-induction.ll
@@ -338,3 +338,44 @@
for.end:
ret void
}
+
+; VEC4_INTERL1-LABEL: @external_use(
+; VEC4_INTERL1-NEXT: entry:
+; VEC4_INTERL1-NEXT: [[TMP0:%.*]] = icmp sgt i64 %n, 1
+; VEC4_INTERL1-NEXT: [[SMAX:%.*]] = select i1 [[TMP0]], i64 %n, i64 1
+; VEC4_INTERL1: br i1 {{.*}}, label %scalar.ph, label %min.iters.checked
+; VEC4_INTERL1: min.iters.checked:
+; VEC4_INTERL1-NEXT: [[N_VEC:%.*]] = and i64 [[SMAX]], 9223372036854775804
+; VEC4_INTERL1: br i1 {{.*}}, label %scalar.ph, label %vector.ph
+; VEC4_INTERL1: vector.ph:
+; VEC4_INTERL1-NEXT: br label %vector.body
+; VEC4_INTERL1: vector.body:
+; VEC4_INTERL1-NEXT: [[INDEX:%.*]] = phi i64 [ 0, %vector.ph ], [ [[INDEX_NEXT:%.*]], %vector.body ]
+; VEC4_INTERL1-NEXT: [[INDEX_NEXT]] = add i64 [[INDEX]], 4
+; VEC4_INTERL1-NEXT: [[TMP1:%.*]] = icmp eq i64 [[INDEX_NEXT]], [[N_VEC]]
+; VEC4_INTERL1-NEXT: br i1 [[TMP1]], label %middle.block, label %vector.body
+; VEC4_INTERL1: middle.block:
+; VEC4_INTERL1: [[TMP2:%.*]] = add nsw i64 [[N_VEC]], -1
+; VEC4_INTERL1-NEXT: [[CAST_CMO:%.*]] = sitofp i64 [[TMP2]] to double
+; VEC4_INTERL1-NEXT: [[IND_ESCAPE:%.*]] = fadd fast double [[CAST_CMO]], %m
+; VEC4_INTERL1-NEXT: br i1 {{.*}}, label %for.end, label %scalar.ph
+; VEC4_INTERL1: for.end:
+; VEC4_INTERL1-NEXT: [[TMP0:%.*]] = phi double [ %j, %for.body ], [ [[IND_ESCAPE]], %middle.block ]
+; VEC4_INTERL1-NEXT: ret double [[TMP0]]
+
+define double @external_use(double %m, i64 %n) {
+entry:
+ br label %for.body
+
+for.body:
+ %i = phi i64 [ 0, %entry ], [%i.next, %for.body]
+ %j = phi double [ %m, %entry ], [ %j.next, %for.body ]
+ %i.next = add i64 %i, 1
+ %j.next = fadd fast double %j, 1.0
+ %cond = icmp slt i64 %i.next, %n
+ br i1 %cond, label %for.body, label %for.end
+
+for.end:
+ %tmp0 = phi double [ %j, %for.body ]
+ ret double %tmp0
+}
Index: lib/Transforms/Vectorize/LoopVectorize.cpp
===================================================================
--- lib/Transforms/Vectorize/LoopVectorize.cpp
+++ lib/Transforms/Vectorize/LoopVectorize.cpp
@@ -3586,8 +3586,12 @@
IRBuilder<> B(MiddleBlock->getTerminator());
Value *CountMinusOne = B.CreateSub(
CountRoundDown, ConstantInt::get(CountRoundDown->getType(), 1));
- Value *CMO = B.CreateSExtOrTrunc(CountMinusOne, II.getStep()->getType(),
- "cast.cmo");
+ Value *CMO =
+ !II.getStep()->getType()->isIntegerTy()
+ ? B.CreateCast(Instruction::SIToFP, CountMinusOne,
+ II.getStep()->getType())
+ : B.CreateSExtOrTrunc(CountMinusOne, II.getStep()->getType());
+ CMO->setName("cast.cmo");
Value *Escape = II.transform(B, CMO, PSE.getSE(), DL);
Escape->setName("ind.escape");
MissingVals[UI] = Escape;
-------------- next part --------------
A non-text attachment was scrubbed...
Name: D32445.96428.patch
Type: text/x-patch
Size: 3114 bytes
Desc: not available
URL: <http://lists.llvm.org/pipermail/llvm-commits/attachments/20170424/12a4d629/attachment.bin>
More information about the llvm-commits
mailing list