1 | // RUN: %clang_cc1 %s -triple x86_64-unknown-linux-gnu -emit-llvm -target-cpu x86-64 -o - |opt -instnamer -S |FileCheck %s --check-prefix SSE |
2 | // RUN: %clang_cc1 %s -triple x86_64-unknown-linux-gnu -emit-llvm -target-cpu skylake -D AVX -o -|opt -instnamer -S | FileCheck %s --check-prefixes AVX,SSE |
3 | // RUN: %clang_cc1 %s -triple x86_64-unknown-linux-gnu -emit-llvm -target-cpu skylake-avx512 -D AVX512 -D AVX -o -|opt -instnamer -S | FileCheck %s --check-prefixes AVX512,AVX,SSE |
4 | // RUN: %clang_cc1 %s -triple x86_64-unknown-linux-gnu -emit-llvm -target-cpu knl -D AVX -D AVX512 -o - |opt -instnamer -S | FileCheck %s --check-prefixes AVX512,AVX,SSE |
5 | |
6 | typedef float __m128 __attribute__ ((vector_size (16))); |
7 | typedef float __m256 __attribute__ ((vector_size (32))); |
8 | typedef float __m512 __attribute__ ((vector_size (64))); |
9 | |
10 | // SSE: call <4 x float> asm "vmovhlps $1, $2, $0", "=v,v,v,~{dirflag},~{fpsr},~{flags}"(i64 %tmp, <4 x float> %tmp1) |
11 | __m128 testXMM(__m128 _xmm0, long _l) { |
12 | __asm__("vmovhlps %1, %2, %0" :"=v"(_xmm0) : "v"(_l), "v"(_xmm0)); |
13 | return _xmm0; |
14 | } |
15 | |
16 | // AVX: call <8 x float> asm "vmovsldup $1, $0", "=v,v,~{dirflag},~{fpsr},~{flags}"(<8 x float> %tmp) |
17 | __m256 testYMM(__m256 _ymm0) { |
18 | #ifdef AVX |
19 | __asm__("vmovsldup %1, %0" :"=v"(_ymm0) : "v"(_ymm0)); |
20 | #endif |
21 | return _ymm0; |
22 | } |
23 | |
24 | // AVX512: call <16 x float> asm "vpternlogd $$0, $1, $2, $0", "=v,v,v,~{dirflag},~{fpsr},~{flags}"(<16 x float> %tmp, <16 x float> %tmp1) |
25 | __m512 testZMM(__m512 _zmm0, __m512 _zmm1) { |
26 | #ifdef AVX512 |
27 | __asm__("vpternlogd $0, %1, %2, %0" :"=v"(_zmm0) : "v"(_zmm1), "v"(_zmm0)); |
28 | #endif |
29 | return _zmm0; |
30 | } |
31 | |