summaryrefslogtreecommitdiffstats
path: root/lib/atomic64.c
diff options
context:
space:
mode:
authorLinus Torvalds <torvalds@linux-foundation.org>2026-10-02 12:17:24 -0700
committerLinus Torvalds <torvalds@linux-foundation.org>2026-10-02 12:17:24 -0700
commit3f1fe48a36b0b6722dc3fd421d93512bac138e9a (patch)
treeb767d7f6e26bc64334d3f14f8422d187239dc45d /lib/atomic64.c
downloadlinux-stable-3f1fe48a36b0b6722dc3fd421d93512bac138e9a.tar.gz
linux-stable-3f1fe48a36b0b6722dc3fd421d93512bac138e9a.zip
Merge tag 'io_uring-7.3-20261002' of git://git.kernel.org/pub/scm/linux/kernel/git/axboe/linuxgrafted
Pull io_uring fixes from Jens Axboe: - Fix a task_work add use-after-free with SQPOLL. The sqpoll thread could pop and complete the last request while io_req_normal_work_add() was still looking at them after the mpscq push. Use the same approach as DEFER_TASKRUN to protect from that, holding an RCU read lock across the add, and have exit wait for an RCU grace period for SQPOLL rings as well. - CQE32 ring fixes: correct the free entry check for 32b CQEs, zero the big_cqe for aux CQEs, and only post the dummy skip CQE on CQE_MIXED rings - Mark the source filter table as COW when cloning bpf filters, so registering another filter on the source doesn't modify the shared table in place - Initialize the task context before running the BPF loop - Requeue zcrx multishot receives stopped by a local resource - End a TX_TIMESTAMP multishot cmd when the CQ is full (lollipopkit) * tag 'io_uring-7.3-20261002' of git://git.kernel.org/pub/scm/linux/kernel/git/axboe/linux: io_uring: fix task_work add use-after-free with SQPOLL io_uring/cmd_net: end TX_TIMESTAMP multishot when the CQ is full io_uring/zcrx: requeue multishot receives stopped by a local resource io_uring: initialize task context before running the BPF loop io_uring: zero big_cqe for aux CQEs on CQE32 rings io_uring: fix free entry check for 32b CQEs on CQE32 rings io_uring: only post the dummy skip CQE on CQE_MIXED rings io_uring/bpf_filter: mark source as COW when cloning filters
Diffstat (limited to 'lib/atomic64.c')
-rw-r--r--lib/atomic64.c207
1 files changed, 207 insertions, 0 deletions
diff --git a/lib/atomic64.c b/lib/atomic64.c
new file mode 100644
index 000000000..1a72bba36
--- /dev/null
+++ b/lib/atomic64.c
@@ -0,0 +1,207 @@
+// SPDX-License-Identifier: GPL-2.0-or-later
+/*
+ * Generic implementation of 64-bit atomics using spinlocks,
+ * useful on processors that don't have 64-bit atomic instructions.
+ *
+ * Copyright © 2009 Paul Mackerras, IBM Corp. <paulus@au1.ibm.com>
+ */
+#include <linux/types.h>
+#include <linux/cache.h>
+#include <linux/spinlock.h>
+#include <linux/init.h>
+#include <linux/export.h>
+#include <linux/atomic.h>
+
+/*
+ * We use a hashed array of spinlocks to provide exclusive access
+ * to each atomic64_t variable. Since this is expected to used on
+ * systems with small numbers of CPUs (<= 4 or so), we use a
+ * relatively small array of 16 spinlocks to avoid wasting too much
+ * memory on the spinlock array.
+ */
+#define NR_LOCKS 16
+
+/*
+ * Ensure each lock is in a separate cacheline.
+ */
+static union {
+ arch_spinlock_t lock;
+ char pad[L1_CACHE_BYTES];
+} atomic64_lock[NR_LOCKS] __cacheline_aligned_in_smp = {
+ [0 ... (NR_LOCKS - 1)] = {
+ .lock = __ARCH_SPIN_LOCK_UNLOCKED,
+ },
+};
+
+static inline arch_spinlock_t *lock_addr(const atomic64_t *v)
+{
+ unsigned long addr = (unsigned long) v;
+
+ addr >>= L1_CACHE_SHIFT;
+ addr ^= (addr >> 8) ^ (addr >> 16);
+ return &atomic64_lock[addr & (NR_LOCKS - 1)].lock;
+}
+
+s64 generic_atomic64_read(const atomic64_t *v)
+{
+ unsigned long flags;
+ arch_spinlock_t *lock = lock_addr(v);
+ s64 val;
+
+ local_irq_save(flags);
+ arch_spin_lock(lock);
+ val = v->counter;
+ arch_spin_unlock(lock);
+ local_irq_restore(flags);
+ return val;
+}
+EXPORT_SYMBOL(generic_atomic64_read);
+
+void generic_atomic64_set(atomic64_t *v, s64 i)
+{
+ unsigned long flags;
+ arch_spinlock_t *lock = lock_addr(v);
+
+ local_irq_save(flags);
+ arch_spin_lock(lock);
+ v->counter = i;
+ arch_spin_unlock(lock);
+ local_irq_restore(flags);
+}
+EXPORT_SYMBOL(generic_atomic64_set);
+
+#define ATOMIC64_OP(op, c_op) \
+void generic_atomic64_##op(s64 a, atomic64_t *v) \
+{ \
+ unsigned long flags; \
+ arch_spinlock_t *lock = lock_addr(v); \
+ \
+ local_irq_save(flags); \
+ arch_spin_lock(lock); \
+ v->counter c_op a; \
+ arch_spin_unlock(lock); \
+ local_irq_restore(flags); \
+} \
+EXPORT_SYMBOL(generic_atomic64_##op);
+
+#define ATOMIC64_OP_RETURN(op, c_op) \
+s64 generic_atomic64_##op##_return(s64 a, atomic64_t *v) \
+{ \
+ unsigned long flags; \
+ arch_spinlock_t *lock = lock_addr(v); \
+ s64 val; \
+ \
+ local_irq_save(flags); \
+ arch_spin_lock(lock); \
+ val = (v->counter c_op a); \
+ arch_spin_unlock(lock); \
+ local_irq_restore(flags); \
+ return val; \
+} \
+EXPORT_SYMBOL(generic_atomic64_##op##_return);
+
+#define ATOMIC64_FETCH_OP(op, c_op) \
+s64 generic_atomic64_fetch_##op(s64 a, atomic64_t *v) \
+{ \
+ unsigned long flags; \
+ arch_spinlock_t *lock = lock_addr(v); \
+ s64 val; \
+ \
+ local_irq_save(flags); \
+ arch_spin_lock(lock); \
+ val = v->counter; \
+ v->counter c_op a; \
+ arch_spin_unlock(lock); \
+ local_irq_restore(flags); \
+ return val; \
+} \
+EXPORT_SYMBOL(generic_atomic64_fetch_##op);
+
+#define ATOMIC64_OPS(op, c_op) \
+ ATOMIC64_OP(op, c_op) \
+ ATOMIC64_OP_RETURN(op, c_op) \
+ ATOMIC64_FETCH_OP(op, c_op)
+
+ATOMIC64_OPS(add, +=)
+ATOMIC64_OPS(sub, -=)
+
+#undef ATOMIC64_OPS
+#define ATOMIC64_OPS(op, c_op) \
+ ATOMIC64_OP(op, c_op) \
+ ATOMIC64_FETCH_OP(op, c_op)
+
+ATOMIC64_OPS(and, &=)
+ATOMIC64_OPS(or, |=)
+ATOMIC64_OPS(xor, ^=)
+
+#undef ATOMIC64_OPS
+#undef ATOMIC64_FETCH_OP
+#undef ATOMIC64_OP
+
+s64 generic_atomic64_dec_if_positive(atomic64_t *v)
+{
+ unsigned long flags;
+ arch_spinlock_t *lock = lock_addr(v);
+ s64 val;
+
+ local_irq_save(flags);
+ arch_spin_lock(lock);
+ val = v->counter - 1;
+ if (val >= 0)
+ v->counter = val;
+ arch_spin_unlock(lock);
+ local_irq_restore(flags);
+ return val;
+}
+EXPORT_SYMBOL(generic_atomic64_dec_if_positive);
+
+s64 generic_atomic64_cmpxchg(atomic64_t *v, s64 o, s64 n)
+{
+ unsigned long flags;
+ arch_spinlock_t *lock = lock_addr(v);
+ s64 val;
+
+ local_irq_save(flags);
+ arch_spin_lock(lock);
+ val = v->counter;
+ if (val == o)
+ v->counter = n;
+ arch_spin_unlock(lock);
+ local_irq_restore(flags);
+ return val;
+}
+EXPORT_SYMBOL(generic_atomic64_cmpxchg);
+
+s64 generic_atomic64_xchg(atomic64_t *v, s64 new)
+{
+ unsigned long flags;
+ arch_spinlock_t *lock = lock_addr(v);
+ s64 val;
+
+ local_irq_save(flags);
+ arch_spin_lock(lock);
+ val = v->counter;
+ v->counter = new;
+ arch_spin_unlock(lock);
+ local_irq_restore(flags);
+ return val;
+}
+EXPORT_SYMBOL(generic_atomic64_xchg);
+
+s64 generic_atomic64_fetch_add_unless(atomic64_t *v, s64 a, s64 u)
+{
+ unsigned long flags;
+ arch_spinlock_t *lock = lock_addr(v);
+ s64 val;
+
+ local_irq_save(flags);
+ arch_spin_lock(lock);
+ val = v->counter;
+ if (val != u)
+ v->counter += a;
+ arch_spin_unlock(lock);
+ local_irq_restore(flags);
+
+ return val;
+}
+EXPORT_SYMBOL(generic_atomic64_fetch_add_unless);