All of lore.kernel.org
 help / color / mirror / Atom feed
From: Puranjay Mohan <puranjay@kernel.org>
To: Alexei Starovoitov <alexei.starovoitov@gmail.com>
Cc: bpf <bpf@vger.kernel.org>,
	kkd@meta.com, Alexei Starovoitov <ast@kernel.org>,
	Andrii Nakryiko <andrii@kernel.org>,
	Daniel Borkmann <daniel@iogearbox.net>,
	Martin KaFai Lau <martin.lau@kernel.org>,
	Eduard Zingerman <eddyz87@gmail.com>,
	Kernel Team <kernel-team@fb.com>
Subject: Re: [PATCH bpf-next] bpf: support nested rcu critical sections
Date: Wed, 17 Sep 2025 13:43:26 +0000	[thread overview]
Message-ID: <mb61pfrclqjnl.fsf@kernel.org> (raw)
In-Reply-To: <CAADnVQKqjQKrSYTm8DO-GLYTFyaGaN8_RiuuJ8kj4zaAShQF0w@mail.gmail.com>

Alexei Starovoitov <alexei.starovoitov@gmail.com> writes:

> On Tue, Sep 16, 2025 at 4:36 AM Puranjay Mohan <puranjay@kernel.org> wrote:
>>
>> Currently, nested rcu critical sections are rejected by the verifier and
>> rcu_lock state is managed by a boolean variable. Add support for nested
>> rcu critical sections by make active_rcu_locks a counter similar to
>> active_preempt_locks. bpf_rcu_read_lock() increments this counter and
>> bpf_rcu_read_unlock() decrements it, MEM_RCU -> PTR_UNTRUSTED transition
>> happens when active_rcu_locks drops to 0.
>>
>> Signed-off-by: Puranjay Mohan <puranjay@kernel.org>
>> ---
>>  include/linux/bpf_verifier.h                  |  2 +-
>>  kernel/bpf/verifier.c                         | 34 ++++++++--------
>>  .../selftests/bpf/prog_tests/rcu_read_lock.c  |  4 +-
>>  .../selftests/bpf/progs/rcu_read_lock.c       | 40 +++++++++++++++++++
>>  4 files changed, 61 insertions(+), 19 deletions(-)
>>
>> diff --git a/include/linux/bpf_verifier.h b/include/linux/bpf_verifier.h
>> index 020de62bd09c..3fb4632d5eed 100644
>> --- a/include/linux/bpf_verifier.h
>> +++ b/include/linux/bpf_verifier.h
>> @@ -441,7 +441,7 @@ struct bpf_verifier_state {
>>         u32 active_irq_id;
>>         u32 active_lock_id;
>>         void *active_lock_ptr;
>> -       bool active_rcu_lock;
>> +       u32 active_rcu_locks;
>>
>>         bool speculative;
>>         bool in_sleepable;
>> diff --git a/kernel/bpf/verifier.c b/kernel/bpf/verifier.c
>> index 1029380f84db..645af66e29ab 100644
>> --- a/kernel/bpf/verifier.c
>> +++ b/kernel/bpf/verifier.c
>> @@ -1438,7 +1438,7 @@ static int copy_reference_state(struct bpf_verifier_state *dst, const struct bpf
>>         dst->acquired_refs = src->acquired_refs;
>>         dst->active_locks = src->active_locks;
>>         dst->active_preempt_locks = src->active_preempt_locks;
>> -       dst->active_rcu_lock = src->active_rcu_lock;
>> +       dst->active_rcu_locks = src->active_rcu_locks;
>>         dst->active_irq_id = src->active_irq_id;
>>         dst->active_lock_id = src->active_lock_id;
>>         dst->active_lock_ptr = src->active_lock_ptr;
>> @@ -5924,7 +5924,7 @@ static bool in_sleepable(struct bpf_verifier_env *env)
>>   */
>>  static bool in_rcu_cs(struct bpf_verifier_env *env)
>>  {
>> -       return env->cur_state->active_rcu_lock ||
>> +       return env->cur_state->active_rcu_locks ||
>>                env->cur_state->active_locks ||
>>                !in_sleepable(env);
>>  }
>> @@ -10684,7 +10684,7 @@ static int check_func_call(struct bpf_verifier_env *env, struct bpf_insn *insn,
>>                 }
>>
>>                 if (env->subprog_info[subprog].might_sleep &&
>> -                   (env->cur_state->active_rcu_lock || env->cur_state->active_preempt_locks ||
>> +                   (env->cur_state->active_rcu_locks || env->cur_state->active_preempt_locks ||
>>                      env->cur_state->active_irq_id || !in_sleepable(env))) {
>>                         verbose(env, "global functions that may sleep are not allowed in non-sleepable context,\n"
>>                                      "i.e., in a RCU/IRQ/preempt-disabled section, or in\n"
>> @@ -11231,7 +11231,7 @@ static int check_resource_leak(struct bpf_verifier_env *env, bool exception_exit
>>                 return -EINVAL;
>>         }
>>
>> -       if (check_lock && env->cur_state->active_rcu_lock) {
>> +       if (check_lock && env->cur_state->active_rcu_locks) {
>>                 verbose(env, "%s cannot be used inside bpf_rcu_read_lock-ed region\n", prefix);
>>                 return -EINVAL;
>>         }
>> @@ -11426,7 +11426,7 @@ static int check_helper_call(struct bpf_verifier_env *env, struct bpf_insn *insn
>>                 return err;
>>         }
>>
>> -       if (env->cur_state->active_rcu_lock) {
>> +       if (env->cur_state->active_rcu_locks) {
>>                 if (fn->might_sleep) {
>>                         verbose(env, "sleepable helper %s#%d in rcu_read_lock region\n",
>>                                 func_id_name(func_id), func_id);
>> @@ -13863,7 +13863,7 @@ static int check_kfunc_call(struct bpf_verifier_env *env, struct bpf_insn *insn,
>>         preempt_disable = is_kfunc_bpf_preempt_disable(&meta);
>>         preempt_enable = is_kfunc_bpf_preempt_enable(&meta);
>>
>> -       if (env->cur_state->active_rcu_lock) {
>> +       if (env->cur_state->active_rcu_locks) {
>>                 struct bpf_func_state *state;
>>                 struct bpf_reg_state *reg;
>>                 u32 clear_mask = (1 << STACK_SPILL) | (1 << STACK_ITER);
>> @@ -13874,22 +13874,22 @@ static int check_kfunc_call(struct bpf_verifier_env *env, struct bpf_insn *insn,
>>                 }
>>
>>                 if (rcu_lock) {
>> -                       verbose(env, "nested rcu read lock (kernel function %s)\n", func_name);
>> -                       return -EINVAL;
>> +                       env->cur_state->active_rcu_locks++;
>>                 } else if (rcu_unlock) {
>> -                       bpf_for_each_reg_in_vstate_mask(env->cur_state, state, reg, clear_mask, ({
>> -                               if (reg->type & MEM_RCU) {
>> -                                       reg->type &= ~(MEM_RCU | PTR_MAYBE_NULL);
>> -                                       reg->type |= PTR_UNTRUSTED;
>> -                               }
>> -                       }));
>> -                       env->cur_state->active_rcu_lock = false;
>> +                       if (--env->cur_state->active_rcu_locks == 0) {
>
> hmm. can it go negative ?
>
> nested_rcu_region_unbalanced_1 test suppose to check it,
> but what kind of error is returned?

It can't go regative as active_rcu_locks is checked above before the
line (-- == 0) is executed, nested_rcu_region_unbalanced_1 will return
the following error:

unmatched rcu read unlock (kernel function bpf_rcu_read_unlock)

>
>
>> +                               bpf_for_each_reg_in_vstate_mask(env-
>
> rewrite it to avoid adding extra ident?

Ack!

>>cur_state, state, reg, clear_mask, ({
>> +                                       if (reg->type & MEM_RCU) {
>> +                                               reg->type &= ~(MEM_RCU | PTR_MAYBE_NULL);
>> +                                               reg->type |= PTR_UNTRUSTED;
>> +                                       }
>> +                               }));
>> +                       }
>>                 } else if (sleepable) {
>>                         verbose(env, "kernel func %s is sleepable within rcu_read_lock region\n", func_name);
>>                         return -EACCES;
>>                 }
>>         } else if (rcu_lock) {
>> -               env->cur_state->active_rcu_lock = true;
>> +               env->cur_state->active_rcu_locks++;
>>         } else if (rcu_unlock) {
>>                 verbose(env, "unmatched rcu read unlock (kernel function %s)\n", func_name);
>>                 return -EINVAL;
>> @@ -18887,7 +18887,7 @@ static bool refsafe(struct bpf_verifier_state *old, struct bpf_verifier_state *c
>>         if (old->active_preempt_locks != cur->active_preempt_locks)
>>                 return false;
>>
>> -       if (old->active_rcu_lock != cur->active_rcu_lock)
>> +       if (old->active_rcu_locks != cur->active_rcu_locks)
>>                 return false;
>>
>>         if (!check_ids(old->active_irq_id, cur->active_irq_id, idmap))
>> diff --git a/tools/testing/selftests/bpf/prog_tests/rcu_read_lock.c b/tools/testing/selftests/bpf/prog_tests/rcu_read_lock.c
>> index c9f855e5da24..246eb259c08a 100644
>> --- a/tools/testing/selftests/bpf/prog_tests/rcu_read_lock.c
>> +++ b/tools/testing/selftests/bpf/prog_tests/rcu_read_lock.c
>> @@ -28,6 +28,7 @@ static void test_success(void)
>>         bpf_program__set_autoload(skel->progs.two_regions, true);
>>         bpf_program__set_autoload(skel->progs.non_sleepable_1, true);
>>         bpf_program__set_autoload(skel->progs.non_sleepable_2, true);
>> +       bpf_program__set_autoload(skel->progs.nested_rcu_region, true);
>>         bpf_program__set_autoload(skel->progs.task_trusted_non_rcuptr, true);
>>         bpf_program__set_autoload(skel->progs.rcu_read_lock_subprog, true);
>>         bpf_program__set_autoload(skel->progs.rcu_read_lock_global_subprog, true);
>> @@ -78,7 +79,8 @@ static const char * const inproper_region_tests[] = {
>>         "non_sleepable_rcu_mismatch",
>>         "inproper_sleepable_helper",
>>         "inproper_sleepable_kfunc",
>> -       "nested_rcu_region",
>
> should be deleted from progs/rcu_read_lock.c too ?

This is not deleted, just moved above into test_success() because now it
will succeed.

> This selftests hunk should in the main patch,
> but below new tests need to go into another patch.

Sure, I will split the changes into two patches.


Thanks,
Puranjay

  reply	other threads:[~2025-09-17 13:43 UTC|newest]

Thread overview: 9+ messages / expand[flat|nested]  mbox.gz  Atom feed  top
2025-09-16 11:36 [PATCH bpf-next] bpf: support nested rcu critical sections Puranjay Mohan
2025-09-16 20:34 ` Alexei Starovoitov
2025-09-17 13:43   ` Puranjay Mohan [this message]
2025-09-16 23:26 ` Eduard Zingerman
2025-09-17 13:44   ` Puranjay Mohan
2025-09-17  0:59 ` Kumar Kartikeya Dwivedi
2025-09-17  4:51 ` Leon Hwang
2025-09-17 14:03   ` Puranjay Mohan
2025-09-18  2:25     ` Leon Hwang

Reply instructions:

You may reply publicly to this message via plain-text email
using any one of the following methods:

* Save the following mbox file, import it into your mail client,
  and reply-to-all from there: mbox

  Avoid top-posting and favor interleaved quoting:
  https://en.wikipedia.org/wiki/Posting_style#Interleaved_style

* Reply using the --to, --cc, and --in-reply-to
  switches of git-send-email(1):

  git send-email \
    --in-reply-to=mb61pfrclqjnl.fsf@kernel.org \
    --to=puranjay@kernel.org \
    --cc=alexei.starovoitov@gmail.com \
    --cc=andrii@kernel.org \
    --cc=ast@kernel.org \
    --cc=bpf@vger.kernel.org \
    --cc=daniel@iogearbox.net \
    --cc=eddyz87@gmail.com \
    --cc=kernel-team@fb.com \
    --cc=kkd@meta.com \
    --cc=martin.lau@kernel.org \
    /path/to/YOUR_REPLY

  https://kernel.org/pub/software/scm/git/docs/git-send-email.html

* If your mail client supports setting the In-Reply-To header
  via mailto: links, try the mailto: link
Be sure your reply has a Subject: header at the top and a blank line before the message body.
This is an external index of several public inboxes,
see mirroring instructions on how to clone and mirror
all data and code used by this external index.