Blame SOURCES/kvm-x86-add-support-for-KVM_CAP_XSAVE2-and-AMX-state-mig.patch

4841a6
From 28cf1b55f346a9f56e84fa57921f5a28a99cd59b Mon Sep 17 00:00:00 2001
4841a6
From: Jing Liu <jing2.liu@intel.com>
4841a6
Date: Wed, 16 Feb 2022 22:04:32 -0800
4841a6
Subject: [PATCH 10/24] x86: add support for KVM_CAP_XSAVE2 and AMX state
4841a6
 migration
4841a6
4841a6
RH-Author: Paul Lai <plai@redhat.com>
4841a6
RH-MergeRequest: 176: Enable KVM AMX support
4841a6
RH-Commit: [10/13] d584f455ba1ecd8a4a87f3470e6aac24ba9a1f5a
4841a6
RH-Bugzilla: 1916415
4841a6
RH-Acked-by: Cornelia Huck <cohuck@redhat.com>
4841a6
RH-Acked-by: Igor Mammedov <imammedo@redhat.com>
4841a6
RH-Acked-by: Paolo Bonzini <pbonzini@redhat.com>
4841a6
4841a6
When dynamic xfeatures (e.g. AMX) are used by the guest, the xsave
4841a6
area would be larger than 4KB. KVM_GET_XSAVE2 and KVM_SET_XSAVE
4841a6
under KVM_CAP_XSAVE2 works with a xsave buffer larger than 4KB.
4841a6
Always use the new ioctls under KVM_CAP_XSAVE2 when KVM supports it.
4841a6
4841a6
Signed-off-by: Jing Liu <jing2.liu@intel.com>
4841a6
Signed-off-by: Zeng Guang <guang.zeng@intel.com>
4841a6
Signed-off-by: Wei Wang <wei.w.wang@intel.com>
4841a6
Signed-off-by: Yang Zhong <yang.zhong@intel.com>
4841a6
Message-Id: <20220217060434.52460-7-yang.zhong@intel.com>
4841a6
Signed-off-by: Paolo Bonzini <pbonzini@redhat.com>
4841a6
(cherry picked from commit e56dd3c70abb31893c61ac834109fa7a38841330)
4841a6
Signed-off-by: Paul Lai <plai@redhat.com>
4841a6
---
4841a6
 target/i386/cpu.h          |  4 ++++
4841a6
 target/i386/kvm/kvm.c      | 42 ++++++++++++++++++++++++--------------
4841a6
 target/i386/xsave_helper.c | 28 +++++++++++++++++++++++++
4841a6
 3 files changed, 59 insertions(+), 15 deletions(-)
4841a6
4841a6
diff --git a/target/i386/cpu.h b/target/i386/cpu.h
4841a6
index f2bdef9c26..14a3501b87 100644
4841a6
--- a/target/i386/cpu.h
4841a6
+++ b/target/i386/cpu.h
4841a6
@@ -1522,6 +1522,10 @@ typedef struct CPUX86State {
4841a6
     uint64_t opmask_regs[NB_OPMASK_REGS];
4841a6
     YMMReg zmmh_regs[CPU_NB_REGS];
4841a6
     ZMMReg hi16_zmm_regs[CPU_NB_REGS];
4841a6
+#ifdef TARGET_X86_64
4841a6
+    uint8_t xtilecfg[64];
4841a6
+    uint8_t xtiledata[8192];
4841a6
+#endif
4841a6
 
4841a6
     /* sysenter registers */
4841a6
     uint32_t sysenter_cs;
4841a6
diff --git a/target/i386/kvm/kvm.c b/target/i386/kvm/kvm.c
4841a6
index a64a79d870..d3d476df27 100644
4841a6
--- a/target/i386/kvm/kvm.c
4841a6
+++ b/target/i386/kvm/kvm.c
4841a6
@@ -123,6 +123,7 @@ static uint32_t num_architectural_pmu_gp_counters;
4841a6
 static uint32_t num_architectural_pmu_fixed_counters;
4841a6
 
4841a6
 static int has_xsave;
4841a6
+static int has_xsave2;
4841a6
 static int has_xcrs;
4841a6
 static int has_pit_state2;
4841a6
 static int has_exception_payload;
4841a6
@@ -1585,6 +1586,26 @@ static Error *invtsc_mig_blocker;
4841a6
 
4841a6
 #define KVM_MAX_CPUID_ENTRIES  100
4841a6
 
4841a6
+static void kvm_init_xsave(CPUX86State *env)
4841a6
+{
4841a6
+    if (has_xsave2) {
4841a6
+        env->xsave_buf_len = QEMU_ALIGN_UP(has_xsave2, 4096);
4841a6
+    } else if (has_xsave) {
4841a6
+        env->xsave_buf_len = sizeof(struct kvm_xsave);
4841a6
+    } else {
4841a6
+        return;
4841a6
+    }
4841a6
+
4841a6
+    env->xsave_buf = qemu_memalign(4096, env->xsave_buf_len);
4841a6
+    memset(env->xsave_buf, 0, env->xsave_buf_len);
4841a6
+    /*
4841a6
+     * The allocated storage must be large enough for all of the
4841a6
+     * possible XSAVE state components.
4841a6
+     */
4841a6
+    assert(kvm_arch_get_supported_cpuid(kvm_state, 0xd, 0, R_ECX) <=
4841a6
+           env->xsave_buf_len);
4841a6
+}
4841a6
+
4841a6
 int kvm_arch_init_vcpu(CPUState *cs)
4841a6
 {
4841a6
     struct {
4841a6
@@ -1614,6 +1635,8 @@ int kvm_arch_init_vcpu(CPUState *cs)
4841a6
 
4841a6
     cpuid_i = 0;
4841a6
 
4841a6
+    has_xsave2 = kvm_check_extension(cs->kvm_state, KVM_CAP_XSAVE2);
4841a6
+
4841a6
     r = kvm_arch_set_tsc_khz(cs);
4841a6
     if (r < 0) {
4841a6
         return r;
4841a6
@@ -2003,19 +2026,7 @@ int kvm_arch_init_vcpu(CPUState *cs)
4841a6
     if (r) {
4841a6
         goto fail;
4841a6
     }
4841a6
-
4841a6
-    if (has_xsave) {
4841a6
-        env->xsave_buf_len = sizeof(struct kvm_xsave);
4841a6
-        env->xsave_buf = qemu_memalign(4096, env->xsave_buf_len);
4841a6
-        memset(env->xsave_buf, 0, env->xsave_buf_len);
4841a6
-
4841a6
-        /*
4841a6
-         * The allocated storage must be large enough for all of the
4841a6
-         * possible XSAVE state components.
4841a6
-         */
4841a6
-        assert(kvm_arch_get_supported_cpuid(kvm_state, 0xd, 0, R_ECX)
4841a6
-               <= env->xsave_buf_len);
4841a6
-    }
4841a6
+    kvm_init_xsave(env);
4841a6
 
4841a6
     max_nested_state_len = kvm_max_nested_state_length();
4841a6
     if (max_nested_state_len > 0) {
4841a6
@@ -3263,13 +3274,14 @@ static int kvm_get_xsave(X86CPU *cpu)
4841a6
 {
4841a6
     CPUX86State *env = &cpu->env;
4841a6
     void *xsave = env->xsave_buf;
4841a6
-    int ret;
4841a6
+    int type, ret;
4841a6
 
4841a6
     if (!has_xsave) {
4841a6
         return kvm_get_fpu(cpu);
4841a6
     }
4841a6
 
4841a6
-    ret = kvm_vcpu_ioctl(CPU(cpu), KVM_GET_XSAVE, xsave);
4841a6
+    type = has_xsave2 ? KVM_GET_XSAVE2 : KVM_GET_XSAVE;
4841a6
+    ret = kvm_vcpu_ioctl(CPU(cpu), type, xsave);
4841a6
     if (ret < 0) {
4841a6
         return ret;
4841a6
     }
4841a6
diff --git a/target/i386/xsave_helper.c b/target/i386/xsave_helper.c
4841a6
index ac61a96344..996e9f3bfe 100644
4841a6
--- a/target/i386/xsave_helper.c
4841a6
+++ b/target/i386/xsave_helper.c
4841a6
@@ -126,6 +126,20 @@ void x86_cpu_xsave_all_areas(X86CPU *cpu, void *buf, uint32_t buflen)
4841a6
 
4841a6
         memcpy(pkru, &env->pkru, sizeof(env->pkru));
4841a6
     }
4841a6
+
4841a6
+    e = &x86_ext_save_areas[XSTATE_XTILE_CFG_BIT];
4841a6
+    if (e->size && e->offset) {
4841a6
+        XSaveXTILECFG *tilecfg = buf + e->offset;
4841a6
+
4841a6
+        memcpy(tilecfg, &env->xtilecfg, sizeof(env->xtilecfg));
4841a6
+    }
4841a6
+
4841a6
+    e = &x86_ext_save_areas[XSTATE_XTILE_DATA_BIT];
4841a6
+    if (e->size && e->offset && buflen >= e->size + e->offset) {
4841a6
+        XSaveXTILEDATA *tiledata = buf + e->offset;
4841a6
+
4841a6
+        memcpy(tiledata, &env->xtiledata, sizeof(env->xtiledata));
4841a6
+    }
4841a6
 #endif
4841a6
 }
4841a6
 
4841a6
@@ -247,5 +261,19 @@ void x86_cpu_xrstor_all_areas(X86CPU *cpu, const void *buf, uint32_t buflen)
4841a6
         pkru = buf + e->offset;
4841a6
         memcpy(&env->pkru, pkru, sizeof(env->pkru));
4841a6
     }
4841a6
+
4841a6
+    e = &x86_ext_save_areas[XSTATE_XTILE_CFG_BIT];
4841a6
+    if (e->size && e->offset) {
4841a6
+        const XSaveXTILECFG *tilecfg = buf + e->offset;
4841a6
+
4841a6
+        memcpy(&env->xtilecfg, tilecfg, sizeof(env->xtilecfg));
4841a6
+    }
4841a6
+
4841a6
+    e = &x86_ext_save_areas[XSTATE_XTILE_DATA_BIT];
4841a6
+    if (e->size && e->offset && buflen >= e->size + e->offset) {
4841a6
+        const XSaveXTILEDATA *tiledata = buf + e->offset;
4841a6
+
4841a6
+        memcpy(&env->xtiledata, tiledata, sizeof(env->xtiledata));
4841a6
+    }
4841a6
 #endif
4841a6
 }
4841a6
-- 
4841a6
2.35.3
4841a6