[llvm] e3dfa82 - [AArch64] Update abs cost with scalar costs and CSSC. NFC (#210900)

via llvm-commits llvm-commits at lists.llvm.org
Tue Jul 21 00:47:07 PDT 2026


Author: David Green
Date: 2026-07-21T08:47:01+01:00
New Revision: e3dfa82a46b6df7cdc92657b27e75470b74459cb

URL: https://github.com/llvm/llvm-project/commit/e3dfa82a46b6df7cdc92657b27e75470b74459cb
DIFF: https://github.com/llvm/llvm-project/commit/e3dfa82a46b6df7cdc92657b27e75470b74459cb.diff

LOG: [AArch64] Update abs cost with scalar costs and CSSC. NFC (#210900)

Added: 
    

Modified: 
    llvm/test/Analysis/CostModel/AArch64/abs.ll

Removed: 
    


################################################################################
diff  --git a/llvm/test/Analysis/CostModel/AArch64/abs.ll b/llvm/test/Analysis/CostModel/AArch64/abs.ll
index 1d57cbb346bf2..a42efa89e474f 100644
--- a/llvm/test/Analysis/CostModel/AArch64/abs.ll
+++ b/llvm/test/Analysis/CostModel/AArch64/abs.ll
@@ -1,69 +1,80 @@
 ; NOTE: Assertions have been autogenerated by utils/update_analyze_test_checks.py
-; RUN: opt -passes="print<cost-model>" 2>&1 -disable-output -cost-kind=all -mtriple=aarch64 < %s | FileCheck %s
+; RUN: opt -passes="print<cost-model>" 2>&1 -disable-output -cost-kind=all -mtriple=aarch64 < %s | FileCheck %s --check-prefixes=CHECK,CHECK-BASE
+; RUN: opt -passes="print<cost-model>" 2>&1 -disable-output -cost-kind=all -mtriple=aarch64 -mattr=+cssc < %s | FileCheck %s --check-prefixes=CHECK,CHECK-CSSC
 
 target datalayout = "e-m:e-i8:8:32-i16:16:32-i64:64-i128:128-n32:64-S128"
 
-declare <2 x i64>  @llvm.abs.v2i64(<2 x i64>, i1)
-declare <4 x i64>  @llvm.abs.v4i64(<4 x i64>, i1)
-declare <8 x i64>  @llvm.abs.v8i64(<8 x i64>, i1)
-
-declare <2 x i32>  @llvm.abs.v2i32(<2 x i32>, i1)
-declare <4 x i32>  @llvm.abs.v4i32(<4 x i32>, i1)
-declare <8 x i32>  @llvm.abs.v8i32(<8 x i32>, i1)
-declare <16 x i32> @llvm.abs.v16i32(<16 x i32>, i1)
-
-declare <2 x i16>  @llvm.abs.v2i16(<2 x i16>, i1)
-declare <4 x i16>  @llvm.abs.v4i16(<4 x i16>, i1)
-declare <8 x i16>  @llvm.abs.v8i16(<8 x i16>, i1)
-declare <16 x i16> @llvm.abs.v16i16(<16 x i16>, i1)
-declare <32 x i16> @llvm.abs.v32i16(<32 x i16>, i1)
-
-declare <2 x i8>   @llvm.abs.v2i8(<2 x i8>, i1)
-declare <4 x i8>   @llvm.abs.v4i8(<4 x i8>, i1)
-declare <8 x i8>   @llvm.abs.v8i8(<8 x i8>, i1)
-declare <16 x i8>  @llvm.abs.v16i8(<16 x i8>, i1)
-declare <32 x i8>  @llvm.abs.v32i8(<32 x i8>, i1)
-declare <64 x i8>  @llvm.abs.v64i8(<64 x i8>, i1)
+define void @abs() {
+; CHECK-BASE-LABEL: 'abs'
+; CHECK-BASE-NEXT:  Cost Model: Found costs of 2 for: %I8 = call i8 @llvm.abs.i8(i8 undef, i1 false)
+; CHECK-BASE-NEXT:  Cost Model: Found costs of 2 for: %I16 = call i16 @llvm.abs.i16(i16 undef, i1 false)
+; CHECK-BASE-NEXT:  Cost Model: Found costs of 2 for: %I32 = call i32 @llvm.abs.i32(i32 undef, i1 false)
+; CHECK-BASE-NEXT:  Cost Model: Found costs of 2 for: %I64 = call i64 @llvm.abs.i64(i64 undef, i1 false)
+; CHECK-BASE-NEXT:  Cost Model: Found costs of 4 for: %I128 = call i128 @llvm.abs.i128(i128 undef, i1 false)
+; CHECK-BASE-NEXT:  Cost Model: Found costs of RThru:0 CodeSize:1 Lat:1 SizeLat:1 for: ret void
+;
+; CHECK-CSSC-LABEL: 'abs'
+; CHECK-CSSC-NEXT:  Cost Model: Found costs of 1 for: %I8 = call i8 @llvm.abs.i8(i8 undef, i1 false)
+; CHECK-CSSC-NEXT:  Cost Model: Found costs of 1 for: %I16 = call i16 @llvm.abs.i16(i16 undef, i1 false)
+; CHECK-CSSC-NEXT:  Cost Model: Found costs of 1 for: %I32 = call i32 @llvm.abs.i32(i32 undef, i1 false)
+; CHECK-CSSC-NEXT:  Cost Model: Found costs of 1 for: %I64 = call i64 @llvm.abs.i64(i64 undef, i1 false)
+; CHECK-CSSC-NEXT:  Cost Model: Found costs of 4 for: %I128 = call i128 @llvm.abs.i128(i128 undef, i1 false)
+; CHECK-CSSC-NEXT:  Cost Model: Found costs of RThru:0 CodeSize:1 Lat:1 SizeLat:1 for: ret void
+;
+  %I8 = call i8 @llvm.abs.i8(i8 undef, i1 false)
+  %I16 = call i16 @llvm.abs.i16(i16 undef, i1 false)
+  %I32 = call i32 @llvm.abs.i32(i32 undef, i1 false)
+  %I64 = call i64 @llvm.abs.i64(i64 undef, i1 false)
+  %I128 = call i128 @llvm.abs.i128(i128 undef, i1 false)
+  ret void
+}
 
-define i32 @abs(i32 %arg) {
-; CHECK-LABEL: 'abs'
-; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V2I64 = call <2 x i64> @llvm.abs.v2i64(<2 x i64> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 2 for: %V4I64 = call <4 x i64> @llvm.abs.v4i64(<4 x i64> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 4 for: %V8I64 = call <8 x i64> @llvm.abs.v8i64(<8 x i64> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V2I32 = call <2 x i32> @llvm.abs.v2i32(<2 x i32> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V4I32 = call <4 x i32> @llvm.abs.v4i32(<4 x i32> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 2 for: %V8I32 = call <8 x i32> @llvm.abs.v8i32(<8 x i32> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 4 for: %V16I32 = call <16 x i32> @llvm.abs.v16i32(<16 x i32> undef, i1 false)
+define void @vec() {
+; CHECK-LABEL: 'vec'
+; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V2I8 = call <2 x i8> @llvm.abs.v2i8(<2 x i8> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V4I8 = call <4 x i8> @llvm.abs.v4i8(<4 x i8> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V8I8 = call <8 x i8> @llvm.abs.v8i8(<8 x i8> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V16I8 = call <16 x i8> @llvm.abs.v16i8(<16 x i8> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 2 for: %V32I8 = call <32 x i8> @llvm.abs.v32i8(<32 x i8> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 4 for: %V64I8 = call <64 x i8> @llvm.abs.v64i8(<64 x i8> undef, i1 false)
 ; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V2I16 = call <2 x i16> @llvm.abs.v2i16(<2 x i16> undef, i1 false)
 ; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V4I16 = call <4 x i16> @llvm.abs.v4i16(<4 x i16> undef, i1 false)
 ; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V8I16 = call <8 x i16> @llvm.abs.v8i16(<8 x i16> undef, i1 false)
 ; CHECK-NEXT:  Cost Model: Found costs of 2 for: %V16I16 = call <16 x i16> @llvm.abs.v16i16(<16 x i16> undef, i1 false)
 ; CHECK-NEXT:  Cost Model: Found costs of 4 for: %V32I16 = call <32 x i16> @llvm.abs.v32i16(<32 x i16> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V8I8 = call <8 x i8> @llvm.abs.v8i8(<8 x i8> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V16I8 = call <16 x i8> @llvm.abs.v16i8(<16 x i8> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 2 for: %V32I8 = call <32 x i8> @llvm.abs.v32i8(<32 x i8> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of 4 for: %V64I8 = call <64 x i8> @llvm.abs.v64i8(<64 x i8> undef, i1 false)
-; CHECK-NEXT:  Cost Model: Found costs of RThru:0 CodeSize:1 Lat:1 SizeLat:1 for: ret i32 undef
+; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V2I32 = call <2 x i32> @llvm.abs.v2i32(<2 x i32> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V4I32 = call <4 x i32> @llvm.abs.v4i32(<4 x i32> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 2 for: %V8I32 = call <8 x i32> @llvm.abs.v8i32(<8 x i32> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 4 for: %V16I32 = call <16 x i32> @llvm.abs.v16i32(<16 x i32> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 1 for: %V2I64 = call <2 x i64> @llvm.abs.v2i64(<2 x i64> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 2 for: %V4I64 = call <4 x i64> @llvm.abs.v4i64(<4 x i64> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 4 for: %V8I64 = call <8 x i64> @llvm.abs.v8i64(<8 x i64> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of 8 for: %V2I128 = call <2 x i128> @llvm.abs.v2i128(<2 x i128> undef, i1 false)
+; CHECK-NEXT:  Cost Model: Found costs of RThru:0 CodeSize:1 Lat:1 SizeLat:1 for: ret void
 ;
-  %V2I64 = call <2 x i64> @llvm.abs.v2i64(<2 x i64> undef, i1 false)
-  %V4I64 = call <4 x i64> @llvm.abs.v4i64(<4 x i64> undef, i1 false)
-  %V8I64 = call <8 x i64> @llvm.abs.v8i64(<8 x i64> undef, i1 false)
-
-  %V2I32  = call <2 x i32>  @llvm.abs.v2i32(<2 x i32> undef, i1 false)
-  %V4I32  = call <4 x i32>  @llvm.abs.v4i32(<4 x i32> undef, i1 false)
-  %V8I32  = call <8 x i32>  @llvm.abs.v8i32(<8 x i32> undef, i1 false)
-  %V16I32 = call <16 x i32> @llvm.abs.v16i32(<16 x i32> undef, i1 false)
+  %V2I8  = call <2 x i8> @llvm.abs.v2i8(<2 x i8> undef, i1 false)
+  %V4I8  = call <4 x i8> @llvm.abs.v4i8(<4 x i8> undef, i1 false)
+  %V8I8  = call <8 x i8> @llvm.abs.v8i8(<8 x i8> undef, i1 false)
+  %V16I8 = call <16 x i8> @llvm.abs.v16i8(<16 x i8> undef, i1 false)
+  %V32I8 = call <32 x i8> @llvm.abs.v32i8(<32 x i8> undef, i1 false)
+  %V64I8 = call <64 x i8> @llvm.abs.v64i8(<64 x i8> undef, i1 false)
 
-  %V2I16  = call <2 x i16>  @llvm.abs.v2i16(<2 x i16> undef, i1 false)
-  %V4I16  = call <4 x i16>  @llvm.abs.v4i16(<4 x i16> undef, i1 false)
-  %V8I16  = call <8 x i16>  @llvm.abs.v8i16(<8 x i16> undef, i1 false)
+  %V2I16  = call <2 x i16> @llvm.abs.v2i16(<2 x i16> undef, i1 false)
+  %V4I16  = call <4 x i16> @llvm.abs.v4i16(<4 x i16> undef, i1 false)
+  %V8I16  = call <8 x i16> @llvm.abs.v8i16(<8 x i16> undef, i1 false)
   %V16I16 = call <16 x i16> @llvm.abs.v16i16(<16 x i16> undef, i1 false)
   %V32I16 = call <32 x i16> @llvm.abs.v32i16(<32 x i16> undef, i1 false)
 
-  %V8I8  = call <8 x i8>  @llvm.abs.v8i8(<8 x i8> undef, i1 false)
-  %V16I8 = call <16 x i8> @llvm.abs.v16i8(<16 x i8> undef, i1 false)
-  %V32I8 = call <32 x i8> @llvm.abs.v32i8(<32 x i8> undef, i1 false)
-  %V64I8 = call <64 x i8> @llvm.abs.v64i8(<64 x i8> undef, i1 false)
+  %V2I32  = call <2 x i32> @llvm.abs.v2i32(<2 x i32> undef, i1 false)
+  %V4I32  = call <4 x i32> @llvm.abs.v4i32(<4 x i32> undef, i1 false)
+  %V8I32  = call <8 x i32> @llvm.abs.v8i32(<8 x i32> undef, i1 false)
+  %V16I32 = call <16 x i32> @llvm.abs.v16i32(<16 x i32> undef, i1 false)
+
+  %V2I64 = call <2 x i64> @llvm.abs.v2i64(<2 x i64> undef, i1 false)
+  %V4I64 = call <4 x i64> @llvm.abs.v4i64(<4 x i64> undef, i1 false)
+  %V8I64 = call <8 x i64> @llvm.abs.v8i64(<8 x i64> undef, i1 false)
+
+  %V2I128 = call <2 x i128> @llvm.abs.v2i128(<2 x i128> undef, i1 false)
 
-  ret i32 undef
+  ret void
 }


        


More information about the llvm-commits mailing list