1*9880d681SAndroid Build Coastguard Worker; NOTE: Assertions have been autogenerated by utils/update_llc_test_checks.py 2*9880d681SAndroid Build Coastguard Worker; RUN: llc < %s -mtriple=x86_64-unknown -mattr=+ssse3 | FileCheck %s 3*9880d681SAndroid Build Coastguard Worker 4*9880d681SAndroid Build Coastguard Worker; Test that the pshufb mask comment is correct. 5*9880d681SAndroid Build Coastguard Worker 6*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test1(<16 x i8> %V) { 7*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test1: 8*9880d681SAndroid Build Coastguard Worker; CHECK: # BB#0: 9*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: pshufb {{.*#+}} xmm0 = xmm0[1,0,0,0,0,2,0,0,0,0,3,0,0,0,0,4] 10*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: retq 11*9880d681SAndroid Build Coastguard Worker %1 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> <i8 1, i8 0, i8 0, i8 0, i8 0, i8 2, i8 0, i8 0, i8 0, i8 0, i8 3, i8 0, i8 0, i8 0, i8 0, i8 4>) 12*9880d681SAndroid Build Coastguard Worker ret <16 x i8> %1 13*9880d681SAndroid Build Coastguard Worker} 14*9880d681SAndroid Build Coastguard Worker 15*9880d681SAndroid Build Coastguard Worker; Test that indexes larger than the size of the vector are shown masked (bottom 4 bits). 16*9880d681SAndroid Build Coastguard Worker 17*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test2(<16 x i8> %V) { 18*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test2: 19*9880d681SAndroid Build Coastguard Worker; CHECK: # BB#0: 20*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: pshufb {{.*#+}} xmm0 = xmm0[15,0,0,0,0,0,0,0,0,0,1,0,0,0,0,2] 21*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: retq 22*9880d681SAndroid Build Coastguard Worker %1 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> <i8 15, i8 0, i8 0, i8 0, i8 0, i8 16, i8 0, i8 0, i8 0, i8 0, i8 17, i8 0, i8 0, i8 0, i8 0, i8 50>) 23*9880d681SAndroid Build Coastguard Worker ret <16 x i8> %1 24*9880d681SAndroid Build Coastguard Worker} 25*9880d681SAndroid Build Coastguard Worker 26*9880d681SAndroid Build Coastguard Worker; Test that indexes with bit seven set are shown as zero. 27*9880d681SAndroid Build Coastguard Worker 28*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test3(<16 x i8> %V) { 29*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test3: 30*9880d681SAndroid Build Coastguard Worker; CHECK: # BB#0: 31*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: pshufb {{.*#+}} xmm0 = xmm0[1,0,0,15,0,2,0,0],zero,xmm0[0,3,0,0],zero,xmm0[0,4] 32*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: retq 33*9880d681SAndroid Build Coastguard Worker %1 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> <i8 1, i8 0, i8 0, i8 127, i8 0, i8 2, i8 0, i8 0, i8 128, i8 0, i8 3, i8 0, i8 0, i8 255, i8 0, i8 4>) 34*9880d681SAndroid Build Coastguard Worker ret <16 x i8> %1 35*9880d681SAndroid Build Coastguard Worker} 36*9880d681SAndroid Build Coastguard Worker 37*9880d681SAndroid Build Coastguard Worker; Test that we won't crash when the constant was reused for another instruction. 38*9880d681SAndroid Build Coastguard Worker 39*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test4(<16 x i8> %V, <2 x i64>* %P) { 40*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test4: 41*9880d681SAndroid Build Coastguard Worker; CHECK: # BB#0: 42*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movaps {{.*#+}} xmm1 = [1084818905618843912,506097522914230528] 43*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movaps %xmm1, (%rdi) 44*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: pshufd {{.*#+}} xmm0 = xmm0[2,3,0,1] 45*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: retq 46*9880d681SAndroid Build Coastguard Worker %1 = insertelement <2 x i64> undef, i64 1084818905618843912, i32 0 47*9880d681SAndroid Build Coastguard Worker %2 = insertelement <2 x i64> %1, i64 506097522914230528, i32 1 48*9880d681SAndroid Build Coastguard Worker store <2 x i64> %2, <2 x i64>* %P, align 16 49*9880d681SAndroid Build Coastguard Worker %3 = bitcast <2 x i64> %2 to <16 x i8> 50*9880d681SAndroid Build Coastguard Worker %4 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> %3) 51*9880d681SAndroid Build Coastguard Worker ret <16 x i8> %4 52*9880d681SAndroid Build Coastguard Worker} 53*9880d681SAndroid Build Coastguard Worker 54*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test5(<16 x i8> %V) { 55*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test5: 56*9880d681SAndroid Build Coastguard Worker; CHECK: # BB#0: 57*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movl $1, %eax 58*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movd %rax, %xmm1 59*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movaps %xmm1, (%rax) 60*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movdqa {{.*#+}} xmm1 = [1,1] 61*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movdqa %xmm1, (%rax) 62*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: pshufb %xmm1, %xmm0 63*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: retq 64*9880d681SAndroid Build Coastguard Worker store <2 x i64> <i64 1, i64 0>, <2 x i64>* undef, align 16 65*9880d681SAndroid Build Coastguard Worker %l = load <2 x i64>, <2 x i64>* undef, align 16 66*9880d681SAndroid Build Coastguard Worker %shuffle = shufflevector <2 x i64> %l, <2 x i64> undef, <2 x i32> zeroinitializer 67*9880d681SAndroid Build Coastguard Worker store <2 x i64> %shuffle, <2 x i64>* undef, align 16 68*9880d681SAndroid Build Coastguard Worker %1 = load <16 x i8>, <16 x i8>* undef, align 16 69*9880d681SAndroid Build Coastguard Worker %2 = call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> %1) 70*9880d681SAndroid Build Coastguard Worker ret <16 x i8> %2 71*9880d681SAndroid Build Coastguard Worker} 72*9880d681SAndroid Build Coastguard Worker 73*9880d681SAndroid Build Coastguard Worker; Test for a reused constant that would allow the pshufb to combine to a simpler instruction. 74*9880d681SAndroid Build Coastguard Worker 75*9880d681SAndroid Build Coastguard Workerdefine <16 x i8> @test6(<16 x i8> %V, <2 x i64>* %P) { 76*9880d681SAndroid Build Coastguard Worker; CHECK-LABEL: test6: 77*9880d681SAndroid Build Coastguard Worker; CHECK: # BB#0: 78*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movaps {{.*#+}} xmm1 = [217019414673948672,506380106026255364] 79*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: movaps %xmm1, (%rdi) 80*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: punpcklbw {{.*#+}} xmm0 = xmm0[0,0,1,1,2,2,3,3,4,4,5,5,6,6,7,7] 81*9880d681SAndroid Build Coastguard Worker; CHECK-NEXT: retq 82*9880d681SAndroid Build Coastguard Worker %1 = insertelement <2 x i64> undef, i64 217019414673948672, i32 0 83*9880d681SAndroid Build Coastguard Worker %2 = insertelement <2 x i64> %1, i64 506380106026255364, i32 1 84*9880d681SAndroid Build Coastguard Worker store <2 x i64> %2, <2 x i64>* %P, align 16 85*9880d681SAndroid Build Coastguard Worker %3 = bitcast <2 x i64> %2 to <16 x i8> 86*9880d681SAndroid Build Coastguard Worker %4 = tail call <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8> %V, <16 x i8> %3) 87*9880d681SAndroid Build Coastguard Worker ret <16 x i8> %4 88*9880d681SAndroid Build Coastguard Worker} 89*9880d681SAndroid Build Coastguard Worker 90*9880d681SAndroid Build Coastguard Workerdeclare <16 x i8> @llvm.x86.ssse3.pshuf.b.128(<16 x i8>, <16 x i8>) nounwind readnone 91