xref: /aosp_15_r20/external/llvm/test/CodeGen/X86/pshufb-mask-comments.ll (revision 9880d6810fe72a1726cb53787c6711e909410d58)
1*9880d681SAndroid Build Coastguard Worker; NOTE: Assertions have been autogenerated by utils/update_llc_test_checks.py
2*9880d681SAndroid Build Coastguard Worker; RUN: llc < %s -mtriple=x86_64-unknown -mattr=+ssse3 | FileCheck %s
3*9880d681SAndroid Build Coastguard Worker
4*9880d681SAndroid Build Coastguard Worker; Test that the pshufb mask comment is correct.
5*9880d681SAndroid Build Coastguard Worker
6*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test1(<16 x i8> %V) {
7*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test1:
8*9880d681SAndroid Build Coastguard Worker; CHECK:       # BB#0:
9*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    pshufb {{.*#+}} xmm0 = xmm0[1,0,0,0,0,2,0,0,0,0,3,0,0,0,0,4]
10*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    retq
11*9880d681SAndroid Build Coastguard Worker  %1 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> <i8 1, i8 0, i8 0, i8 0, i8 0, i8 2, i8 0, i8 0, i8 0, i8 0, i8 3, i8 0, i8 0, i8 0, i8 0, i8 4>)
12*9880d681SAndroid Build Coastguard Worker  ret <16 x i8> %1
13*9880d681SAndroid Build Coastguard Worker}
14*9880d681SAndroid Build Coastguard Worker
15*9880d681SAndroid Build Coastguard Worker; Test that indexes larger than the size of the vector are shown masked (bottom 4 bits).
16*9880d681SAndroid Build Coastguard Worker
17*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test2(<16 x i8> %V) {
18*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test2:
19*9880d681SAndroid Build Coastguard Worker; CHECK:       # BB#0:
20*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    pshufb {{.*#+}} xmm0 = xmm0[15,0,0,0,0,0,0,0,0,0,1,0,0,0,0,2]
21*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    retq
22*9880d681SAndroid Build Coastguard Worker  %1 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> <i8 15, i8 0, i8 0, i8 0, i8 0, i8 16, i8 0, i8 0, i8 0, i8 0, i8 17, i8 0, i8 0, i8 0, i8 0, i8 50>)
23*9880d681SAndroid Build Coastguard Worker  ret <16 x i8> %1
24*9880d681SAndroid Build Coastguard Worker}
25*9880d681SAndroid Build Coastguard Worker
26*9880d681SAndroid Build Coastguard Worker; Test that indexes with bit seven set are shown as zero.
27*9880d681SAndroid Build Coastguard Worker
28*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test3(<16 x i8> %V) {
29*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test3:
30*9880d681SAndroid Build Coastguard Worker; CHECK:       # BB#0:
31*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    pshufb {{.*#+}} xmm0 = xmm0[1,0,0,15,0,2,0,0],zero,xmm0[0,3,0,0],zero,xmm0[0,4]
32*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    retq
33*9880d681SAndroid Build Coastguard Worker  %1 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> <i8 1, i8 0, i8 0, i8 127, i8 0, i8 2, i8 0, i8 0, i8 128, i8 0, i8 3, i8 0, i8 0, i8 255, i8 0, i8 4>)
34*9880d681SAndroid Build Coastguard Worker  ret <16 x i8> %1
35*9880d681SAndroid Build Coastguard Worker}
36*9880d681SAndroid Build Coastguard Worker
37*9880d681SAndroid Build Coastguard Worker; Test that we won't crash when the constant was reused for another instruction.
38*9880d681SAndroid Build Coastguard Worker
39*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test4(<16 x i8> %V, <2 x i64>* %P) {
40*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test4:
41*9880d681SAndroid Build Coastguard Worker; CHECK:       # BB#0:
42*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movaps {{.*#+}} xmm1 = [1084818905618843912,506097522914230528]
43*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movaps %xmm1, (%rdi)
44*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    pshufd {{.*#+}} xmm0 = xmm0[2,3,0,1]
45*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    retq
46*9880d681SAndroid Build Coastguard Worker  %1 = insertelement <2 x i64> undef, i64 1084818905618843912, i32 0
47*9880d681SAndroid Build Coastguard Worker  %2 = insertelement <2 x i64>    %1, i64  506097522914230528, i32 1
48*9880d681SAndroid Build Coastguard Worker  store <2 x i64> %2, <2 x i64>* %P, align 16
49*9880d681SAndroid Build Coastguard Worker  %3 = bitcast <2 x i64> %2 to <16 x i8>
50*9880d681SAndroid Build Coastguard Worker  %4 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> %3)
51*9880d681SAndroid Build Coastguard Worker  ret <16 x i8> %4
52*9880d681SAndroid Build Coastguard Worker}
53*9880d681SAndroid Build Coastguard Worker
54*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test5(<16 x i8> %V) {
55*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test5:
56*9880d681SAndroid Build Coastguard Worker; CHECK:       # BB#0:
57*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movl $1, %eax
58*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movd %rax, %xmm1
59*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movaps %xmm1, (%rax)
60*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movdqa {{.*#+}} xmm1 = [1,1]
61*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movdqa %xmm1, (%rax)
62*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    pshufb %xmm1, %xmm0
63*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    retq
64*9880d681SAndroid Build Coastguard Worker  store <2 x i64> <i64 1, i64 0>, <2 x i64>* undef, align 16
65*9880d681SAndroid Build Coastguard Worker  %l = load <2 x i64>, <2 x i64>* undef, align 16
66*9880d681SAndroid Build Coastguard Worker  %shuffle = shufflevector <2 x i64> %l, <2 x i64> undef, <2 x i32> zeroinitializer
67*9880d681SAndroid Build Coastguard Worker  store <2 x i64> %shuffle, <2 x i64>* undef, align 16
68*9880d681SAndroid Build Coastguard Worker  %1 = load <16 x i8>, <16 x i8>* undef, align 16
69*9880d681SAndroid Build Coastguard Worker  %2 = call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> %1)
70*9880d681SAndroid Build Coastguard Worker  ret <16 x i8> %2
71*9880d681SAndroid Build Coastguard Worker}
72*9880d681SAndroid Build Coastguard Worker
73*9880d681SAndroid Build Coastguard Worker; Test for a reused constant that would allow the pshufb to combine to a simpler instruction.
74*9880d681SAndroid Build Coastguard Worker
75*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test6(<16 x i8> %V, <2 x i64>* %P) {
76*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test6:
77*9880d681SAndroid Build Coastguard Worker; CHECK:       # BB#0:
78*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movaps {{.*#+}} xmm1 = [217019414673948672,506380106026255364]
79*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    movaps %xmm1, (%rdi)
80*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    punpcklbw {{.*#+}} xmm0 = xmm0[0,0,1,1,2,2,3,3,4,4,5,5,6,6,7,7]
81*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT:    retq
82*9880d681SAndroid Build Coastguard Worker  %1 = insertelement <2 x i64> undef, i64 217019414673948672, i32 0
83*9880d681SAndroid Build Coastguard Worker  %2 = insertelement <2 x i64>    %1, i64 506380106026255364, i32 1
84*9880d681SAndroid Build Coastguard Worker  store <2 x i64> %2, <2 x i64>* %P, align 16
85*9880d681SAndroid Build Coastguard Worker  %3 = bitcast <2 x i64> %2 to <16 x i8>
86*9880d681SAndroid Build Coastguard Worker  %4 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> %3)
87*9880d681SAndroid Build Coastguard Worker  ret <16 x i8> %4
88*9880d681SAndroid Build Coastguard Worker}
89*9880d681SAndroid Build Coastguard Worker
90*9880d681SAndroid Build Coastguard Workerdeclare <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8>, <16 x i8>) nounwind readnone
91