drm/xe: refactor the paging engine setup

On newer platforms, the paging configuration is now configured by the PF
via the ADS object, where VF side should ensure that everything
configured as GUC_PAGING_CLASS is correctly mirrored on VF side. For
example PF could in theory reserve two BCS instances, and we expect VF
to mirror that.

With that move towards having a logical mask of all the paging engines,
and also generalise selecting those engines, based on the number of
paging engines. Also cache the first designated paging engine, which
will makes things a little cleaner here, and in later patches.

No functional changes for existing platforms.

v2 (Sashiko):
  - Rework the loop slightly so that we don't needlessly check for the
    paging engine, before we have correctly set the logical instance.
  - Add a proper error return, if we encounter a bogus paging config.
    Thinking ahead to VF where the config is defined by the PF, we
    should just gracefully exit the probe sequence.
v3:
  - Move paging_engines > copy_engines engines check to VF patch.

Signed-off-by: Matthew Auld <matthew.auld@intel.com>
Cc: Daniele Ceraolo Spurio <daniele.ceraolospurio@intel.com>
Cc: Thomas Hellström <thomas.hellstrom@linux.intel.com>
Cc: Matthew Brost <matthew.brost@intel.com>
Reviewed-by: Francois Dugast <francois.dugast@intel.com>
Link: https://patch.msgid.link/20260626111520.487997-17-matthew.auld@intel.com
This commit is contained in:
Matthew Auld
2026-06-26 12:15:26 +01:00
parent 08af1a0efd
commit 16c8849297
5 changed files with 46 additions and 49 deletions

View File

@@ -530,10 +530,7 @@ struct xe_exec_queue *xe_exec_queue_create_bind(struct xe_device *xe,
migrate_vm = xe_migrate_get_vm(tile->migrate);
if (xe->info.has_usm) {
struct xe_hw_engine *hwe = xe_gt_hw_engine(gt,
XE_ENGINE_CLASS_COPY,
gt->usm.reserved_bcs_instance,
false);
struct xe_hw_engine *hwe = gt->usm.paging_hwe0;
if (!hwe) {
xe_vm_put(migrate_vm);

View File

@@ -137,10 +137,8 @@ static inline bool xe_gt_is_media_type(struct xe_gt *gt)
static inline bool xe_gt_is_usm_hwe(struct xe_gt *gt, struct xe_hw_engine *hwe)
{
struct xe_device *xe = gt_to_xe(gt);
return xe->info.has_usm && hwe->class == XE_ENGINE_CLASS_COPY &&
hwe->instance == gt->usm.reserved_bcs_instance;
return hwe->class == XE_ENGINE_CLASS_COPY &&
(gt->usm.paging_logical_mask & BIT(hwe->logical_instance));
}
/**

View File

@@ -235,10 +235,16 @@ struct xe_gt {
*/
struct xe_sa_manager *bb_pool;
/**
* @usm.reserved_bcs_instance: reserved BCS instance used for USM
* operations (e.g. migrations, fixing page tables)
* @usm.paging_hwe0: The first designated paging engine.
* This is some reserved BCS instance used for USM operations
* (e.g. migrations, fixing page tables)
*/
u16 reserved_bcs_instance;
struct xe_hw_engine *paging_hwe0;
/**
* @usm.paging_logical_mask: logical mask of paging engines.
* Should be densely populated.
*/
u32 paging_logical_mask;
} usm;
/** @ordered_wq: used to serialize GT resets and TDRs */

View File

@@ -647,10 +647,6 @@ static int hw_engine_init(struct xe_gt *gt, struct xe_hw_engine *hwe,
xe_hw_engine_enable_ring(hwe);
}
/* We reserve the highest BCS instance for USM */
if (xe->info.has_usm && hwe->class == XE_ENGINE_CLASS_COPY)
gt->usm.reserved_bcs_instance = hwe->instance;
/* Ensure IDLEDLY is lower than MAXCNT */
adjust_idledly(hwe);
@@ -662,19 +658,43 @@ static int hw_engine_init(struct xe_gt *gt, struct xe_hw_engine *hwe,
return err;
}
static void hw_engine_setup_logical_mapping(struct xe_gt *gt)
static void hw_engine_setup_logical_and_paging_mapping(struct xe_gt *gt)
{
struct xe_device *xe = gt_to_xe(gt);
unsigned int num_copy_engines = 0, num_paging_engines = 0;
unsigned int reserved_logical_bcs_start;
struct xe_hw_engine *hwe;
enum xe_hw_engine_id id;
int class;
for_each_hw_engine(hwe, gt, id)
if (hwe->class == XE_ENGINE_CLASS_COPY)
num_copy_engines++;
/* We just reserve the highest BCS instance for USM */
if (num_copy_engines && xe->info.has_usm)
num_paging_engines = 1;
reserved_logical_bcs_start = num_copy_engines - num_paging_engines;
/* FIXME: Doing a simple logical mapping that works for most hardware */
for (class = 0; class < XE_ENGINE_CLASS_MAX; ++class) {
struct xe_hw_engine *hwe;
enum xe_hw_engine_id id;
int logical_instance = 0;
for_each_hw_engine(hwe, gt, id)
if (hwe->class == class)
for_each_hw_engine(hwe, gt, id) {
if (hwe->class == class) {
hwe->logical_instance = logical_instance++;
if (class == XE_ENGINE_CLASS_COPY &&
hwe->logical_instance >=
reserved_logical_bcs_start) {
if (!gt->usm.paging_hwe0)
gt->usm.paging_hwe0 = hwe;
gt->usm.paging_logical_mask |=
BIT(hwe->logical_instance);
}
}
}
}
}
@@ -894,7 +914,7 @@ int xe_hw_engines_init(struct xe_gt *gt)
return err;
}
hw_engine_setup_logical_mapping(gt);
hw_engine_setup_logical_and_paging_mapping(gt);
err = xe_hw_engine_setup_groups(gt);
if (err)
return err;

View File

@@ -383,27 +383,6 @@ static void xe_migrate_suballoc_manager_init(struct xe_migrate *m, u32 map_ofs)
NUM_VMUSA_UNIT_PER_PAGE, 0);
}
/*
* Including the reserved copy engine is required to avoid deadlocks due to
* migrate jobs servicing the faults gets stuck behind the job that faulted.
*/
static u32 xe_migrate_usm_logical_mask(struct xe_gt *gt)
{
u32 logical_mask = 0;
struct xe_hw_engine *hwe;
enum xe_hw_engine_id id;
for_each_hw_engine(hwe, gt, id) {
if (hwe->class != XE_ENGINE_CLASS_COPY)
continue;
if (xe_gt_is_usm_hwe(gt, hwe))
logical_mask |= BIT(hwe->logical_instance);
}
return logical_mask;
}
static bool xe_migrate_needs_ccs_emit(struct xe_device *xe)
{
return xe_device_has_flat_ccs(xe) && !(GRAPHICS_VER(xe) >= 20 && IS_DGFX(xe));
@@ -479,13 +458,10 @@ int xe_migrate_init(struct xe_migrate *m)
goto err_out;
if (xe->info.has_usm) {
struct xe_hw_engine *hwe = xe_gt_hw_engine(primary_gt,
XE_ENGINE_CLASS_COPY,
primary_gt->usm.reserved_bcs_instance,
false);
u32 logical_mask = xe_migrate_usm_logical_mask(primary_gt);
struct xe_hw_engine *hwe0 = primary_gt->usm.paging_hwe0;
u32 logical_mask = primary_gt->usm.paging_logical_mask;
if (!hwe || !logical_mask) {
if (!hwe0 || !logical_mask) {
err = -EINVAL;
goto err_out;
}
@@ -494,7 +470,7 @@ int xe_migrate_init(struct xe_migrate *m)
* XXX: Currently only reserving 1 (likely slow) BCS instance on
* PVC, may want to revisit if performance is needed.
*/
m->q = xe_exec_queue_create(xe, vm, logical_mask, 1, hwe,
m->q = xe_exec_queue_create(xe, vm, logical_mask, 1, hwe0,
EXEC_QUEUE_FLAG_KERNEL |
EXEC_QUEUE_FLAG_PERMANENT |
EXEC_QUEUE_FLAG_HIGH_PRIORITY |