summaryrefslogtreecommitdiff
diff options
context:
space:
mode:
authorMatthew Auld <matthew.auld@intel.com>2026-06-26 12:15:26 +0100
committerMatthew Auld <matthew.auld@intel.com>2026-07-10 12:50:12 +0100
commit16c8849297a8ca755d1ef24a407dc7815c361928 (patch)
treeb35711d86f2ac80f2fa980ecdc4773e748278033
parent08af1a0efd04301d094f9acc09ce0c1ba8ea26c5 (diff)
downloadlinux-next-16c8849297a8ca755d1ef24a407dc7815c361928.tar.gz
linux-next-16c8849297a8ca755d1ef24a407dc7815c361928.zip
drm/xe: refactor the paging engine setup
On newer platforms, the paging configuration is now configured by the PF via the ADS object, where VF side should ensure that everything configured as GUC_PAGING_CLASS is correctly mirrored on VF side. For example PF could in theory reserve two BCS instances, and we expect VF to mirror that. With that move towards having a logical mask of all the paging engines, and also generalise selecting those engines, based on the number of paging engines. Also cache the first designated paging engine, which will makes things a little cleaner here, and in later patches. No functional changes for existing platforms. v2 (Sashiko): - Rework the loop slightly so that we don't needlessly check for the paging engine, before we have correctly set the logical instance. - Add a proper error return, if we encounter a bogus paging config. Thinking ahead to VF where the config is defined by the PF, we should just gracefully exit the probe sequence. v3: - Move paging_engines > copy_engines engines check to VF patch. Signed-off-by: Matthew Auld <matthew.auld@intel.com> Cc: Daniele Ceraolo Spurio <daniele.ceraolospurio@intel.com> Cc: Thomas Hellström <thomas.hellstrom@linux.intel.com> Cc: Matthew Brost <matthew.brost@intel.com> Reviewed-by: Francois Dugast <francois.dugast@intel.com> Link: https://patch.msgid.link/20260626111520.487997-17-matthew.auld@intel.com
-rw-r--r--drivers/gpu/drm/xe/xe_exec_queue.c5
-rw-r--r--drivers/gpu/drm/xe/xe_gt.h6
-rw-r--r--drivers/gpu/drm/xe/xe_gt_types.h12
-rw-r--r--drivers/gpu/drm/xe/xe_hw_engine.c40
-rw-r--r--drivers/gpu/drm/xe/xe_migrate.c32
5 files changed, 46 insertions, 49 deletions
diff --git a/drivers/gpu/drm/xe/xe_exec_queue.c b/drivers/gpu/drm/xe/xe_exec_queue.c
index 1b5ca3ce578a..cfd2a4e6d4c7 100644
--- a/drivers/gpu/drm/xe/xe_exec_queue.c
+++ b/drivers/gpu/drm/xe/xe_exec_queue.c
@@ -530,10 +530,7 @@ struct xe_exec_queue *xe_exec_queue_create_bind(struct xe_device *xe,
migrate_vm = xe_migrate_get_vm(tile->migrate);
if (xe->info.has_usm) {
- struct xe_hw_engine *hwe = xe_gt_hw_engine(gt,
- XE_ENGINE_CLASS_COPY,
- gt->usm.reserved_bcs_instance,
- false);
+ struct xe_hw_engine *hwe = gt->usm.paging_hwe0;
if (!hwe) {
xe_vm_put(migrate_vm);
diff --git a/drivers/gpu/drm/xe/xe_gt.h b/drivers/gpu/drm/xe/xe_gt.h
index 4150aa594f05..a6cfaa1af23f 100644
--- a/drivers/gpu/drm/xe/xe_gt.h
+++ b/drivers/gpu/drm/xe/xe_gt.h
@@ -137,10 +137,8 @@ static inline bool xe_gt_is_media_type(struct xe_gt *gt)
static inline bool xe_gt_is_usm_hwe(struct xe_gt *gt, struct xe_hw_engine *hwe)
{
- struct xe_device *xe = gt_to_xe(gt);
-
- return xe->info.has_usm && hwe->class == XE_ENGINE_CLASS_COPY &&
- hwe->instance == gt->usm.reserved_bcs_instance;
+ return hwe->class == XE_ENGINE_CLASS_COPY &&
+ (gt->usm.paging_logical_mask & BIT(hwe->logical_instance));
}
/**
diff --git a/drivers/gpu/drm/xe/xe_gt_types.h b/drivers/gpu/drm/xe/xe_gt_types.h
index 0d234160ee3a..a8bbfbdf3849 100644
--- a/drivers/gpu/drm/xe/xe_gt_types.h
+++ b/drivers/gpu/drm/xe/xe_gt_types.h
@@ -235,10 +235,16 @@ struct xe_gt {
*/
struct xe_sa_manager *bb_pool;
/**
- * @usm.reserved_bcs_instance: reserved BCS instance used for USM
- * operations (e.g. migrations, fixing page tables)
+ * @usm.paging_hwe0: The first designated paging engine.
+ * This is some reserved BCS instance used for USM operations
+ * (e.g. migrations, fixing page tables)
*/
- u16 reserved_bcs_instance;
+ struct xe_hw_engine *paging_hwe0;
+ /**
+ * @usm.paging_logical_mask: logical mask of paging engines.
+ * Should be densely populated.
+ */
+ u32 paging_logical_mask;
} usm;
/** @ordered_wq: used to serialize GT resets and TDRs */
diff --git a/drivers/gpu/drm/xe/xe_hw_engine.c b/drivers/gpu/drm/xe/xe_hw_engine.c
index dd2d37e3d80c..d741d93601b2 100644
--- a/drivers/gpu/drm/xe/xe_hw_engine.c
+++ b/drivers/gpu/drm/xe/xe_hw_engine.c
@@ -647,10 +647,6 @@ static int hw_engine_init(struct xe_gt *gt, struct xe_hw_engine *hwe,
xe_hw_engine_enable_ring(hwe);
}
- /* We reserve the highest BCS instance for USM */
- if (xe->info.has_usm && hwe->class == XE_ENGINE_CLASS_COPY)
- gt->usm.reserved_bcs_instance = hwe->instance;
-
/* Ensure IDLEDLY is lower than MAXCNT */
adjust_idledly(hwe);
@@ -662,19 +658,43 @@ err_name:
return err;
}
-static void hw_engine_setup_logical_mapping(struct xe_gt *gt)
+static void hw_engine_setup_logical_and_paging_mapping(struct xe_gt *gt)
{
+ struct xe_device *xe = gt_to_xe(gt);
+ unsigned int num_copy_engines = 0, num_paging_engines = 0;
+ unsigned int reserved_logical_bcs_start;
+ struct xe_hw_engine *hwe;
+ enum xe_hw_engine_id id;
int class;
+ for_each_hw_engine(hwe, gt, id)
+ if (hwe->class == XE_ENGINE_CLASS_COPY)
+ num_copy_engines++;
+
+ /* We just reserve the highest BCS instance for USM */
+ if (num_copy_engines && xe->info.has_usm)
+ num_paging_engines = 1;
+
+ reserved_logical_bcs_start = num_copy_engines - num_paging_engines;
+
/* FIXME: Doing a simple logical mapping that works for most hardware */
for (class = 0; class < XE_ENGINE_CLASS_MAX; ++class) {
- struct xe_hw_engine *hwe;
- enum xe_hw_engine_id id;
int logical_instance = 0;
- for_each_hw_engine(hwe, gt, id)
- if (hwe->class == class)
+ for_each_hw_engine(hwe, gt, id) {
+ if (hwe->class == class) {
hwe->logical_instance = logical_instance++;
+
+ if (class == XE_ENGINE_CLASS_COPY &&
+ hwe->logical_instance >=
+ reserved_logical_bcs_start) {
+ if (!gt->usm.paging_hwe0)
+ gt->usm.paging_hwe0 = hwe;
+ gt->usm.paging_logical_mask |=
+ BIT(hwe->logical_instance);
+ }
+ }
+ }
}
}
@@ -894,7 +914,7 @@ int xe_hw_engines_init(struct xe_gt *gt)
return err;
}
- hw_engine_setup_logical_mapping(gt);
+ hw_engine_setup_logical_and_paging_mapping(gt);
err = xe_hw_engine_setup_groups(gt);
if (err)
return err;
diff --git a/drivers/gpu/drm/xe/xe_migrate.c b/drivers/gpu/drm/xe/xe_migrate.c
index 9428dd5e7760..92d5e81ceac2 100644
--- a/drivers/gpu/drm/xe/xe_migrate.c
+++ b/drivers/gpu/drm/xe/xe_migrate.c
@@ -383,27 +383,6 @@ static void xe_migrate_suballoc_manager_init(struct xe_migrate *m, u32 map_ofs)
NUM_VMUSA_UNIT_PER_PAGE, 0);
}
-/*
- * Including the reserved copy engine is required to avoid deadlocks due to
- * migrate jobs servicing the faults gets stuck behind the job that faulted.
- */
-static u32 xe_migrate_usm_logical_mask(struct xe_gt *gt)
-{
- u32 logical_mask = 0;
- struct xe_hw_engine *hwe;
- enum xe_hw_engine_id id;
-
- for_each_hw_engine(hwe, gt, id) {
- if (hwe->class != XE_ENGINE_CLASS_COPY)
- continue;
-
- if (xe_gt_is_usm_hwe(gt, hwe))
- logical_mask |= BIT(hwe->logical_instance);
- }
-
- return logical_mask;
-}
-
static bool xe_migrate_needs_ccs_emit(struct xe_device *xe)
{
return xe_device_has_flat_ccs(xe) && !(GRAPHICS_VER(xe) >= 20 && IS_DGFX(xe));
@@ -479,13 +458,10 @@ int xe_migrate_init(struct xe_migrate *m)
goto err_out;
if (xe->info.has_usm) {
- struct xe_hw_engine *hwe = xe_gt_hw_engine(primary_gt,
- XE_ENGINE_CLASS_COPY,
- primary_gt->usm.reserved_bcs_instance,
- false);
- u32 logical_mask = xe_migrate_usm_logical_mask(primary_gt);
+ struct xe_hw_engine *hwe0 = primary_gt->usm.paging_hwe0;
+ u32 logical_mask = primary_gt->usm.paging_logical_mask;
- if (!hwe || !logical_mask) {
+ if (!hwe0 || !logical_mask) {
err = -EINVAL;
goto err_out;
}
@@ -494,7 +470,7 @@ int xe_migrate_init(struct xe_migrate *m)
* XXX: Currently only reserving 1 (likely slow) BCS instance on
* PVC, may want to revisit if performance is needed.
*/
- m->q = xe_exec_queue_create(xe, vm, logical_mask, 1, hwe,
+ m->q = xe_exec_queue_create(xe, vm, logical_mask, 1, hwe0,
EXEC_QUEUE_FLAG_KERNEL |
EXEC_QUEUE_FLAG_PERMANENT |
EXEC_QUEUE_FLAG_HIGH_PRIORITY |