[v3,16/31] tcg/i386: Support vector variable shift opcodes

Message ID	20190504055300.18426-17-richard.henderson@linaro.org
State	Superseded
Headers	show Delivered-To: patch@linaro.org Received-SPF: pass (google.com: domain of qemu-devel-bounces+patch=linaro.org@nongnu.org designates 209.51.188.17 as permitted sender) client-ip=209.51.188.17; From: Richard Henderson <richard.henderson@linaro.org> To: qemu-devel@nongnu.org Date: Fri, 3 May 2019 22:52:45 -0700 Message-Id: <20190504055300.18426-17-richard.henderson@linaro.org> In-Reply-To: <20190504055300.18426-1-richard.henderson@linaro.org> References: <20190504055300.18426-1-richard.henderson@linaro.org> Subject: [Qemu-devel] [PATCH v3 16/31] tcg/i386: Support vector variable shift opcodes Precedence: list Cc: alex.bennee@linaro.org, david@redhat.com Errors-To: qemu-devel-bounces+patch=linaro.org@nongnu.org Sender: "Qemu-devel" <qemu-devel-bounces+patch=linaro.org@nongnu.org>
Series	tcg vector improvements \| expand [v3,00/31] tcg vector improvements [v3,01/31] tcg: Implement tcg_gen_gvec_3i() [v3,02/31] tcg: Do not recreate INDEX_op_neg_vec unless supported [v3,03/31] tcg: Allow add_vec, sub_vec, neg_vec, not_vec to be expanded [v3,04/31] tcg: Specify optional vector requirements with a list [v3,05/31] tcg: Assert fixed_reg is read-only [v3,06/31] tcg/arm: Use tcg_out_mov_reg in tcg_out_mov [v3,07/31] tcg: Return bool success from tcg_out_mov [v3,08/31] tcg: Support cross-class moves without instruction support [v3,09/31] tcg: Promote tcg_out_{dup, dupi}_vec to backend interface [v3,10/31] tcg: Manually expand INDEX_op_dup_vec [v3,11/31] tcg: Add tcg_out_dupm_vec to the backend interface [v3,12/31] tcg/i386: Implement tcg_out_dupm_vec [v3,13/31] tcg/aarch64: Implement tcg_out_dupm_vec [v3,14/31] tcg: Add INDEX_op_dupm_vec [v3,15/31] tcg: Add gvec expanders for variable shift [v3,16/31] tcg/i386: Support vector variable shift opcodes [v3,17/31] tcg/aarch64: Support vector variable shift opcodes [v3,18/31] tcg: Add gvec expanders for vector shift by scalar [v3,19/31] tcg/i386: Support vector scalar shift opcodes [v3,20/31] tcg: Add support for integer absolute value [v3,21/31] tcg: Add support for vector absolute value [v3,22/31] tcg/i386: Support vector absolute value [v3,23/31] tcg/aarch64: Support vector absolute value [v3,24/31] target/arm: Use tcg_gen_abs_i64 and tcg_gen_gvec_abs [v3,25/31] target/cris: Use tcg_gen_abs_tl [v3,26/31] target/ppc: Use tcg_gen_abs_i32 [v3,27/31] target/ppc: Use tcg_gen_abs_tl [v3,28/31] target/s390x: Use tcg_gen_abs_i64 [v3,29/31] target/tricore: Use tcg_gen_abs_tl [v3,30/31] target/xtensa: Use tcg_gen_abs_i32 [v3,31/31] tcg/aarch64: Do not advertise minmax for MO_64

Message ID

20190504055300.18426-17-richard.henderson@linaro.org

State

Superseded

Headers

Received-SPF: pass (google.com: domain of
	qemu-devel-bounces+patch=linaro.org@nongnu.org designates
	209.51.188.17 as permitted sender) client-ip=209.51.188.17; 
From: Richard Henderson <richard.henderson@linaro.org>
To: qemu-devel@nongnu.org
Date: Fri,  3 May 2019 22:52:45 -0700
Message-Id: <20190504055300.18426-17-richard.henderson@linaro.org>
In-Reply-To: <20190504055300.18426-1-richard.henderson@linaro.org>
References: <20190504055300.18426-1-richard.henderson@linaro.org>
Subject: [Qemu-devel] [PATCH v3 16/31] tcg/i386: Support vector variable
	shift opcodes
Precedence: list
Cc: alex.bennee@linaro.org, david@redhat.com
Errors-To: qemu-devel-bounces+patch=linaro.org@nongnu.org
Sender: "Qemu-devel" <qemu-devel-bounces+patch=linaro.org@nongnu.org>

Series

tcg vector improvements | expand

Commit Message

Richard Henderson May 4, 2019, 5:52 a.m. UTC

Signed-off-by: Richard Henderson <richard.henderson@linaro.org>

---
 tcg/i386/tcg-target.h     |  2 +-
 tcg/i386/tcg-target.inc.c | 35 +++++++++++++++++++++++++++++++++++
 2 files changed, 36 insertions(+), 1 deletion(-)

-- 
2.17.1

diff --git a/tcg/i386/tcg-target.h b/tcg/i386/tcg-target.h
index 241bf19413..b240633455 100644
--- a/tcg/i386/tcg-target.h
+++ b/tcg/i386/tcg-target.h
@@ -184,7 +184,7 @@  extern bool have_avx2;
 #define TCG_TARGET_HAS_neg_vec          0
 #define TCG_TARGET_HAS_shi_vec          1
 #define TCG_TARGET_HAS_shs_vec          0
-#define TCG_TARGET_HAS_shv_vec          0
+#define TCG_TARGET_HAS_shv_vec          have_avx2
 #define TCG_TARGET_HAS_cmp_vec          1
 #define TCG_TARGET_HAS_mul_vec          1
 #define TCG_TARGET_HAS_sat_vec          1
diff --git a/tcg/i386/tcg-target.inc.c b/tcg/i386/tcg-target.inc.c
index 5b33bbd99b..c9448b6d84 100644
--- a/tcg/i386/tcg-target.inc.c
+++ b/tcg/i386/tcg-target.inc.c
@@ -467,6 +467,11 @@  static inline int tcg_target_const_match(tcg_target_long val, TCGType type,
 #define OPC_VPBROADCASTQ (0x59 | P_EXT38 | P_DATA16)
 #define OPC_VPERMQ      (0x00 | P_EXT3A | P_DATA16 | P_REXW)
 #define OPC_VPERM2I128  (0x46 | P_EXT3A | P_DATA16 | P_VEXL)
+#define OPC_VPSLLVD     (0x47 | P_EXT38 | P_DATA16)
+#define OPC_VPSLLVQ     (0x47 | P_EXT38 | P_DATA16 | P_REXW)
+#define OPC_VPSRAVD     (0x46 | P_EXT38 | P_DATA16)
+#define OPC_VPSRLVD     (0x45 | P_EXT38 | P_DATA16)
+#define OPC_VPSRLVQ     (0x45 | P_EXT38 | P_DATA16 | P_REXW)
 #define OPC_VZEROUPPER  (0x77 | P_EXT)
 #define OPC_XCHG_ax_r32	(0x90)
 
@@ -2707,6 +2712,18 @@  static void tcg_out_vec_op(TCGContext *s, TCGOpcode opc,
     static int const umax_insn[4] = {
         OPC_PMAXUB, OPC_PMAXUW, OPC_PMAXUD, OPC_UD2
     };
+    static int const shlv_insn[4] = {
+        /* TODO: AVX512 adds support for MO_16.  */
+        OPC_UD2, OPC_UD2, OPC_VPSLLVD, OPC_VPSLLVQ
+    };
+    static int const shrv_insn[4] = {
+        /* TODO: AVX512 adds support for MO_16.  */
+        OPC_UD2, OPC_UD2, OPC_VPSRLVD, OPC_VPSRLVQ
+    };
+    static int const sarv_insn[4] = {
+        /* TODO: AVX512 adds support for MO_16, MO_64.  */
+        OPC_UD2, OPC_UD2, OPC_VPSRAVD, OPC_UD2
+    };
 
     TCGType type = vecl + TCG_TYPE_V64;
     int insn, sub;
@@ -2759,6 +2776,15 @@  static void tcg_out_vec_op(TCGContext *s, TCGOpcode opc,
     case INDEX_op_umax_vec:
         insn = umax_insn[vece];
         goto gen_simd;
+    case INDEX_op_shlv_vec:
+        insn = shlv_insn[vece];
+        goto gen_simd;
+    case INDEX_op_shrv_vec:
+        insn = shrv_insn[vece];
+        goto gen_simd;
+    case INDEX_op_sarv_vec:
+        insn = sarv_insn[vece];
+        goto gen_simd;
     case INDEX_op_x86_punpckl_vec:
         insn = punpckl_insn[vece];
         goto gen_simd;
@@ -3136,6 +3162,9 @@  static const TCGTargetOpDef *tcg_target_op_def(TCGOpcode op)
     case INDEX_op_umin_vec:
     case INDEX_op_smax_vec:
     case INDEX_op_umax_vec:
+    case INDEX_op_shlv_vec:
+    case INDEX_op_shrv_vec:
+    case INDEX_op_sarv_vec:
     case INDEX_op_cmp_vec:
     case INDEX_op_x86_shufps_vec:
     case INDEX_op_x86_blend_vec:
@@ -3193,6 +3222,12 @@  int tcg_can_emit_vec_op(TCGOpcode opc, TCGType type, unsigned vece)
         }
         return 1;
 
+    case INDEX_op_shlv_vec:
+    case INDEX_op_shrv_vec:
+        return have_avx2 && vece >= MO_32;
+    case INDEX_op_sarv_vec:
+        return have_avx2 && vece == MO_32;
+
     case INDEX_op_mul_vec:
         if (vece == MO_8) {
             /* We can expand the operation for MO_8.  */

[v3,16/31] tcg/i386: Support vector variable shift opcodes

Commit Message

Patch