mirror of
https://github.com/torvalds/linux.git
synced 2026-09-22 12:44:03 +02:00
mm/page_alloc: use existing highatomic reserves on the buddy fastpath
ALLOC_HIGHATOMIC currently provides both access to MIGRATE_HIGHATOMIC free pages and permission to create new highatomic pageblock reserves. This makes it unsuitable for the fastpath. However, the fastpath can reach rmqueue_buddy() while MIGRATE_HIGHATOMIC reserves have free pages available. In this situation, the allocation can fall back to other migratetypes without trying those reserves first. Allow high-priority non-blocking allocations to use existing MIGRATE_HIGHATOMIC reserves on the buddy fastpath without growing them. First tighten the criteria for reserving pageblocks so that growth may only occur in the slowpath. Then allow fastpath usage by enabling ALLOC_HIGHATOMIC when the GFP mask describes a non-blocking high-priority allocation. This logic has been factored out from gfp_to_alloc_flags() to a new function gfp_to_alloc_flags_nonblocking(). A UDP receive workload was run with free MIGRATE_HIGHATOMIC pageblocks available in the target zone. Before this patch, the workload did not consume these blocks. With this patch, eligible order-1 allocations reaching the buddy path consumed existing MIGRATE_HIGHATOMIC pageblocks, with no highatomic misses observed. The workload did not grow highatomic reserves and NAPI page-frag allocations remained healthy with no failures or order-0 fallbacks. Link: https://lore.kernel.org/20260623004600.113347-1-jp.kobryn@linux.dev Signed-off-by: JP Kobryn <jp.kobryn@linux.dev> Reviewed-by: Vlastimil Babka (SUSE) <vbabka@kernel.org> Acked-by: Johannes Weiner <hannes@cmpxchg.org> Reviewed-by: Shakeel Butt <shakeel.butt@linux.dev> Cc: Brendan Jackman <jackmanb@google.com> Cc: David Hildenbrand <david@kernel.org> Cc: Frank van der Linden <fvdl@google.com> Cc: Liam R. Howlett <liam@infradead.org> Cc: Lorenzo Stoakes <ljs@kernel.org> Cc: Michal Hocko <mhocko@suse.com> Cc: Mike Rapoport <rppt@kernel.org> Cc: Suren Baghdasaryan <surenb@google.com> Cc: Zi Yan <ziy@nvidia.com> Signed-off-by: Andrew Morton <akpm@linux-foundation.org>
This commit is contained in:
parent
6028a90f33
commit
9f2fd03c9b
|
|
@ -3249,10 +3249,11 @@ struct page *rmqueue_buddy(struct zone *preferred_zone, struct zone *zone,
|
|||
} while (check_new_pages(page, order));
|
||||
|
||||
/*
|
||||
* If this is a high-order atomic allocation then check
|
||||
* if the pageblock should be reserved for the future
|
||||
* Slowpath (precarious) high-atomic allocations may reserve
|
||||
* a pageblock for future use.
|
||||
*/
|
||||
if (unlikely(alloc_flags & ALLOC_HIGHATOMIC))
|
||||
if (unlikely((alloc_flags & ALLOC_HIGHATOMIC) &&
|
||||
((alloc_flags & ALLOC_WMARK_MASK) == ALLOC_WMARK_MIN)))
|
||||
reserve_highatomic_pageblock(page, order, zone);
|
||||
|
||||
__count_zid_vm_events(PGALLOC, page_zonenum(page), 1 << order);
|
||||
|
|
@ -4472,6 +4473,29 @@ static void wake_all_kswapds(unsigned int order, gfp_t gfp_mask,
|
|||
}
|
||||
}
|
||||
|
||||
static inline unsigned int
|
||||
gfp_to_alloc_flags_nonblocking(gfp_t gfp_mask, unsigned int order)
|
||||
{
|
||||
unsigned int alloc_flags = 0;
|
||||
|
||||
if (gfp_mask & __GFP_DIRECT_RECLAIM)
|
||||
return 0;
|
||||
|
||||
/*
|
||||
* Not worth trying to allocate harder for __GFP_NOMEMALLOC even
|
||||
* if it can't schedule.
|
||||
*/
|
||||
if (gfp_mask & __GFP_NOMEMALLOC)
|
||||
return 0;
|
||||
|
||||
alloc_flags |= ALLOC_NON_BLOCK;
|
||||
|
||||
if (order > 0 && (gfp_mask & __GFP_HIGH))
|
||||
alloc_flags |= ALLOC_HIGHATOMIC;
|
||||
|
||||
return alloc_flags;
|
||||
}
|
||||
|
||||
static inline unsigned int
|
||||
gfp_to_alloc_flags(gfp_t gfp_mask, unsigned int order)
|
||||
{
|
||||
|
|
@ -4488,18 +4512,9 @@ gfp_to_alloc_flags(gfp_t gfp_mask, unsigned int order)
|
|||
if (gfp_mask & __GFP_KSWAPD_RECLAIM)
|
||||
alloc_flags |= ALLOC_KSWAPD;
|
||||
|
||||
alloc_flags |= gfp_to_alloc_flags_nonblocking(gfp_mask, order);
|
||||
|
||||
if (!(gfp_mask & __GFP_DIRECT_RECLAIM)) {
|
||||
/*
|
||||
* Not worth trying to allocate harder for __GFP_NOMEMALLOC even
|
||||
* if it can't schedule.
|
||||
*/
|
||||
if (!(gfp_mask & __GFP_NOMEMALLOC)) {
|
||||
alloc_flags |= ALLOC_NON_BLOCK;
|
||||
|
||||
if (order > 0 && (alloc_flags & ALLOC_MIN_RESERVE))
|
||||
alloc_flags |= ALLOC_HIGHATOMIC;
|
||||
}
|
||||
|
||||
/*
|
||||
* Ignore cpuset mems for non-blocking __GFP_HIGH (probably
|
||||
* GFP_ATOMIC) rather than fail, see the comment for
|
||||
|
|
@ -5292,6 +5307,7 @@ struct page *__alloc_frozen_pages_noprof(gfp_t gfp, unsigned int order,
|
|||
* memory until all local zones are considered.
|
||||
*/
|
||||
alloc_flags |= alloc_flags_nofragment(zonelist_zone(ac.preferred_zoneref), gfp);
|
||||
alloc_flags |= gfp_to_alloc_flags_nonblocking(gfp, order) & ALLOC_HIGHATOMIC;
|
||||
|
||||
/* First allocation attempt */
|
||||
page = get_page_from_freelist(alloc_gfp, order, alloc_flags, &ac);
|
||||
|
|
|
|||
Loading…
Reference in New Issue
Block a user