xref: /llvm-project/llvm/test/Transforms/VectorCombine/X86/shuffle-of-shuffles.ll (revision 611401c11594871aa5c7692cd17a7f12b6fbe660)
1; NOTE: Assertions have been autogenerated by utils/update_test_checks.py UTC_ARGS: --version 4
2; RUN: opt < %s -passes=vector-combine -S -mtriple=x86_64-- -mattr=sse2 | FileCheck %s --check-prefixes=CHECK,SSE
3; RUN: opt < %s -passes=vector-combine -S -mtriple=x86_64-- -mattr=avx2 | FileCheck %s --check-prefixes=CHECK,AVX
4
5; fold to identity
6
7define <8 x i32> @concat_extract_subvectors(<8 x i32> %x) {
8; CHECK-LABEL: define <8 x i32> @concat_extract_subvectors(
9; CHECK-SAME: <8 x i32> [[X:%.*]]) #[[ATTR0:[0-9]+]] {
10; CHECK-NEXT:    ret <8 x i32> [[X]]
11;
12  %lo = shufflevector <8 x i32> %x, <8 x i32> poison, <4 x i32> <i32 0, i32 1, i32 2, i32 3>
13  %hi = shufflevector <8 x i32> %x, <8 x i32> poison, <4 x i32> <i32 4, i32 5, i32 6, i32 7>
14  %concat = shufflevector <4 x i32> %lo, <4 x i32> %hi, <8 x i32> <i32 0, i32 1, i32 2, i32 3, i32 4, i32 5, i32 6, i32 7>
15  ret <8 x i32> %concat
16}
17
18; negative test - shuffle contains undef
19
20define <8 x i32> @concat_extract_subvectors_undef(<8 x i32> %x) {
21; CHECK-LABEL: define <8 x i32> @concat_extract_subvectors_undef(
22; CHECK-SAME: <8 x i32> [[X:%.*]]) #[[ATTR0]] {
23; CHECK-NEXT:    [[LO:%.*]] = shufflevector <8 x i32> [[X]], <8 x i32> undef, <4 x i32> <i32 0, i32 1, i32 2, i32 8>
24; CHECK-NEXT:    [[HI:%.*]] = shufflevector <8 x i32> [[X]], <8 x i32> undef, <4 x i32> <i32 4, i32 5, i32 6, i32 8>
25; CHECK-NEXT:    [[CONCAT:%.*]] = shufflevector <4 x i32> [[LO]], <4 x i32> [[HI]], <8 x i32> <i32 0, i32 1, i32 2, i32 3, i32 4, i32 5, i32 6, i32 7>
26; CHECK-NEXT:    ret <8 x i32> [[CONCAT]]
27;
28  %lo = shufflevector <8 x i32> %x, <8 x i32> undef, <4 x i32> <i32 0, i32 1, i32 2, i32 8>
29  %hi = shufflevector <8 x i32> %x, <8 x i32> undef, <4 x i32> <i32 4, i32 5, i32 6, i32 8>
30  %concat = shufflevector <4 x i32> %lo, <4 x i32> %hi, <8 x i32> <i32 0, i32 1, i32 2, i32 3, i32 4, i32 5, i32 6, i32 7>
31  ret <8 x i32> %concat
32}
33
34; shuffle contains poison
35
36define <8 x i32> @concat_extract_subvectors_poison(<8 x i32> %x) {
37; CHECK-LABEL: define <8 x i32> @concat_extract_subvectors_poison(
38; CHECK-SAME: <8 x i32> [[X:%.*]]) #[[ATTR0]] {
39; CHECK-NEXT:    ret <8 x i32> [[X]]
40;
41  %lo = shufflevector <8 x i32> %x, <8 x i32> poison, <4 x i32> <i32 0, i32 1, i32 2, i32 8>
42  %hi = shufflevector <8 x i32> %x, <8 x i32> poison, <4 x i32> <i32 4, i32 5, i32 6, i32 8>
43  %concat = shufflevector <4 x i32> %lo, <4 x i32> %hi, <8 x i32> <i32 0, i32 1, i32 2, i32 3, i32 4, i32 5, i32 6, i32 7>
44  ret <8 x i32> %concat
45}
46
47; broadcast loads are free on AVX (and blends are much cheap than general 2-operand shuffles)
48
49define  <4 x double> @blend_broadcasts_v4f64(ptr %p0, ptr %p1)  {
50; SSE-LABEL: define <4 x double> @blend_broadcasts_v4f64(
51; SSE-SAME: ptr [[P0:%.*]], ptr [[P1:%.*]]) #[[ATTR0]] {
52; SSE-NEXT:    [[LD0:%.*]] = load <4 x double>, ptr [[P0]], align 32
53; SSE-NEXT:    [[LD1:%.*]] = load <4 x double>, ptr [[P1]], align 32
54; SSE-NEXT:    [[BLEND:%.*]] = shufflevector <4 x double> [[LD0]], <4 x double> [[LD1]], <4 x i32> <i32 0, i32 4, i32 4, i32 0>
55; SSE-NEXT:    ret <4 x double> [[BLEND]]
56;
57; AVX-LABEL: define <4 x double> @blend_broadcasts_v4f64(
58; AVX-SAME: ptr [[P0:%.*]], ptr [[P1:%.*]]) #[[ATTR0]] {
59; AVX-NEXT:    [[LD0:%.*]] = load <4 x double>, ptr [[P0]], align 32
60; AVX-NEXT:    [[LD1:%.*]] = load <4 x double>, ptr [[P1]], align 32
61; AVX-NEXT:    [[BCST0:%.*]] = shufflevector <4 x double> [[LD0]], <4 x double> undef, <4 x i32> zeroinitializer
62; AVX-NEXT:    [[BCST1:%.*]] = shufflevector <4 x double> [[LD1]], <4 x double> undef, <4 x i32> zeroinitializer
63; AVX-NEXT:    [[BLEND:%.*]] = shufflevector <4 x double> [[BCST0]], <4 x double> [[BCST1]], <4 x i32> <i32 0, i32 5, i32 6, i32 3>
64; AVX-NEXT:    ret <4 x double> [[BLEND]]
65;
66  %ld0 = load <4 x double>, ptr %p0, align 32
67  %ld1 = load <4 x double>, ptr %p1, align 32
68  %bcst0 = shufflevector <4 x double> %ld0, <4 x double> undef, <4 x i32> zeroinitializer
69  %bcst1 = shufflevector <4 x double> %ld1, <4 x double> undef, <4 x i32> zeroinitializer
70  %blend = shufflevector <4 x double> %bcst0, <4 x double> %bcst1, <4 x i32> <i32 0, i32 5, i32 6, i32 3>
71  ret <4 x double> %blend
72}
73
74define <2 x float> @PR86068(<2 x float> %a0, <2 x float> %a1) {
75; CHECK-LABEL: define <2 x float> @PR86068(
76; CHECK-SAME: <2 x float> [[A0:%.*]], <2 x float> [[A1:%.*]]) #[[ATTR0]] {
77; CHECK-NEXT:    [[S2:%.*]] = shufflevector <2 x float> [[A1]], <2 x float> [[A0]], <2 x i32> <i32 1, i32 3>
78; CHECK-NEXT:    ret <2 x float> [[S2]]
79;
80  %s1 = shufflevector <2 x float> %a1, <2 x float> poison, <2 x i32> <i32 1, i32 poison>
81  %s2 = shufflevector <2 x float> %s1, <2 x float> %a0, <2 x i32> <i32 0, i32 3>
82  ret <2 x float> %s2
83}
84