From: Breno Leitao <leitao@debian.org>
To: David Ahern <dsahern@kernel.org>,
Ido Schimmel <idosch@nvidia.com>,
"David S. Miller" <davem@davemloft.net>,
Eric Dumazet <edumazet@google.com>,
Jakub Kicinski <kuba@kernel.org>,
Paolo Abeni <pabeni@redhat.com>, Simon Horman <horms@kernel.org>,
Alexei Starovoitov <ast@kernel.org>,
Daniel Borkmann <daniel@iogearbox.net>,
Andrii Nakryiko <andrii@kernel.org>,
Eduard Zingerman <eddyz87@gmail.com>,
Kumar Kartikeya Dwivedi <memxor@gmail.com>,
Martin KaFai Lau <martin.lau@linux.dev>,
Song Liu <song@kernel.org>,
Yonghong Song <yonghong.song@linux.dev>,
Jiri Olsa <jolsa@kernel.org>,
Emil Tsalapatis <emil@etsalapatis.com>,
Ihor Solodrai <ihor.solodrai@linux.dev>,
John Fastabend <john.fastabend@gmail.com>,
Stanislav Fomichev <sdf@fomichev.me>,
Shuah Khan <shuah@kernel.org>
Cc: netdev@vger.kernel.org, linux-kernel@vger.kernel.org,
bpf@vger.kernel.org, linux-kselftest@vger.kernel.org,
david.laight.linux@gmail.com, Breno Leitao <leitao@debian.org>,
kernel-team@meta.com
Subject: [PATCH net-next 2/6] ipv6: mcast: convert ip6_mc_msfget() to sockopt_t
Date: Fri, 25 Sep 2026 08:55:17 -0700 [thread overview]
Message-ID: <20260925-sockopt_expand_out_v2-v1-2-c3ef2e3bb5c0@debian.org> (raw)
In-Reply-To: <20260925-sockopt_expand_out_v2-v1-0-c3ef2e3bb5c0@debian.org>
MCAST_MSFILTER reads its reply through ip6_mc_msfget(), reached from
do_ipv6_getsockopt() and from nowhere else. Convert it, and build the
sockopt_t at the call site for as long as the caller still carries a
sockptr_t pair.
optlen only has to cover the fixed part, and the real reply size comes
from the gf_numsrc field inside it. Userspace relies on that, so
sockopt_expand_out() grows optval past optlen, for a user address only
and only far enough for the sources the socket has.
ip6_mc_msfget() now advances over the fixed part and writes the source
list through iter_out. Its callers rewind by the reply length they
already compute, which is exactly what the callee consumed, and land
back where they used to write: offset 0 for the native reply, gf_fmode
for the compat one.
The *optlen store moves out to the call site, guarded by !err so the
-EINVAL, -EADDRNOTAVAIL and -EFAULT returns still leave the caller's
optlen word untouched.
Signed-off-by: Breno Leitao <leitao@debian.org>
---
include/net/ipv6.h | 2 +-
net/ipv6/ipv6_sockglue.c | 63 ++++++++++++++++++++++++++++++++----------------
net/ipv6/mcast.c | 19 ++++++++++++---
3 files changed, 58 insertions(+), 26 deletions(-)
diff --git a/include/net/ipv6.h b/include/net/ipv6.h
index 3de07e738538f7..9bb68d75890364 100644
--- a/include/net/ipv6.h
+++ b/include/net/ipv6.h
@@ -1195,7 +1195,7 @@ int ip6_mc_source(int add, int omode, struct sock *sk,
int ip6_mc_msfilter(struct sock *sk, struct group_filter *gsf,
struct sockaddr_storage *list);
int ip6_mc_msfget(struct sock *sk, struct group_filter *gsf,
- sockptr_t optval, size_t ss_offset);
+ sockopt_t *opt, size_t ss_offset);
#ifdef CONFIG_PROC_FS
int ac6_proc_init(struct net *net);
diff --git a/net/ipv6/ipv6_sockglue.c b/net/ipv6/ipv6_sockglue.c
index 5c6a0819a2aaff..1bdb3e001e4fe6 100644
--- a/net/ipv6/ipv6_sockglue.c
+++ b/net/ipv6/ipv6_sockglue.c
@@ -922,48 +922,51 @@ static int ipv6_getsockopt_sticky(struct sock *sk, struct ipv6_txoptions *opt,
return len;
}
-static int ipv6_get_msfilter(struct sock *sk, sockptr_t optval,
- sockptr_t optlen, int len)
+static int ipv6_get_msfilter(struct sock *sk, sockopt_t *opt)
{
const int size0 = offsetof(struct group_filter, gf_slist_flex);
struct group_filter gsf;
- int num;
+ int num, len;
int err;
- if (len < size0)
+ if (opt->optlen < size0)
return -EINVAL;
- if (copy_from_sockptr(&gsf, optval, size0))
+ if (copy_from_iter(&gsf, size0, &opt->iter_in) != size0)
return -EFAULT;
if (gsf.gf_group.ss_family != AF_INET6)
return -EADDRNOTAVAIL;
num = gsf.gf_numsrc;
sockopt_lock_sock(sk);
- err = ip6_mc_msfget(sk, &gsf, optval, size0);
+ err = ip6_mc_msfget(sk, &gsf, opt, size0);
if (!err) {
if (num > gsf.gf_numsrc)
num = gsf.gf_numsrc;
len = GROUP_FILTER_SIZE(num);
- if (copy_to_sockptr(optlen, &len, sizeof(int)) ||
- copy_to_sockptr(optval, &gsf, size0))
+ opt->optlen = len;
+
+ /* ip6_mc_msfget() consumed the whole reply; rewind to the
+ * fixed part.
+ */
+ iov_iter_revert(&opt->iter_out, len);
+ if (copy_to_iter(&gsf, size0, &opt->iter_out) != size0)
err = -EFAULT;
}
sockopt_release_sock(sk);
return err;
}
-static int compat_ipv6_get_msfilter(struct sock *sk, sockptr_t optval,
- sockptr_t optlen, int len)
+static int compat_ipv6_get_msfilter(struct sock *sk, sockopt_t *opt)
{
const int size0 = offsetof(struct compat_group_filter, gf_slist_flex);
struct compat_group_filter gf32;
struct group_filter gf;
int err;
- int num;
+ int num, len;
- if (len < size0)
+ if (opt->optlen < size0)
return -EINVAL;
- if (copy_from_sockptr(&gf32, optval, size0))
+ if (copy_from_iter(&gf32, size0, &opt->iter_in) != size0)
return -EFAULT;
gf.gf_interface = gf32.gf_interface;
gf.gf_fmode = gf32.gf_fmode;
@@ -974,18 +977,22 @@ static int compat_ipv6_get_msfilter(struct sock *sk, sockptr_t optval,
return -EADDRNOTAVAIL;
sockopt_lock_sock(sk);
- err = ip6_mc_msfget(sk, &gf, optval, size0);
+ err = ip6_mc_msfget(sk, &gf, opt, size0);
sockopt_release_sock(sk);
if (err)
return err;
if (num > gf.gf_numsrc)
num = gf.gf_numsrc;
len = GROUP_FILTER_SIZE(num) - (sizeof(gf)-sizeof(gf32));
- if (copy_to_sockptr(optlen, &len, sizeof(int)) ||
- copy_to_sockptr_offset(optval, offsetof(struct compat_group_filter, gf_fmode),
- &gf.gf_fmode, sizeof(gf32.gf_fmode)) ||
- copy_to_sockptr_offset(optval, offsetof(struct compat_group_filter, gf_numsrc),
- &gf.gf_numsrc, sizeof(gf32.gf_numsrc)))
+ opt->optlen = len;
+
+ /* Rewind to gf_fmode, which gf_numsrc follows. */
+ iov_iter_revert(&opt->iter_out,
+ len - offsetof(struct compat_group_filter, gf_fmode));
+ if (copy_to_iter(&gf.gf_fmode, sizeof(gf32.gf_fmode),
+ &opt->iter_out) != sizeof(gf32.gf_fmode) ||
+ copy_to_iter(&gf.gf_numsrc, sizeof(gf32.gf_numsrc),
+ &opt->iter_out) != sizeof(gf32.gf_numsrc))
return -EFAULT;
return 0;
}
@@ -1006,9 +1013,23 @@ int do_ipv6_getsockopt(struct sock *sk, int level, int optname,
return -EINVAL;
switch (optname) {
case MCAST_MSFILTER:
+ {
+ struct kvec kvec;
+ sockopt_t opt;
+ int err;
+
+ err = sockptr_to_sockopt(&opt, optval, optlen, &kvec);
+ if (err)
+ return err;
+
if (in_compat_syscall())
- return compat_ipv6_get_msfilter(sk, optval, optlen, len);
- return ipv6_get_msfilter(sk, optval, optlen, len);
+ err = compat_ipv6_get_msfilter(sk, &opt);
+ else
+ err = ipv6_get_msfilter(sk, &opt);
+ if (!err && copy_to_sockptr(optlen, &opt.optlen, sizeof(int)))
+ err = -EFAULT;
+ return err;
+ }
case IPV6_2292PKTOPTIONS:
{
struct msghdr msg;
diff --git a/net/ipv6/mcast.c b/net/ipv6/mcast.c
index ecef55f261890c..4ca2d77811f4eb 100644
--- a/net/ipv6/mcast.c
+++ b/net/ipv6/mcast.c
@@ -600,14 +600,14 @@ int ip6_mc_msfilter(struct sock *sk, struct group_filter *gsf,
}
int ip6_mc_msfget(struct sock *sk, struct group_filter *gsf,
- sockptr_t optval, size_t ss_offset)
+ sockopt_t *opt, size_t ss_offset)
{
struct ipv6_pinfo *inet6 = inet6_sk(sk);
const struct in6_addr *group;
struct ipv6_mc_socklist *pmc;
struct ip6_sf_socklist *psl;
+ int i, copycount, err;
unsigned int count;
- int i, copycount;
group = &((struct sockaddr_in6 *)&gsf->gf_group)->sin6_addr;
@@ -629,6 +629,18 @@ int ip6_mc_msfget(struct sock *sk, struct group_filter *gsf,
copycount = min(count, gsf->gf_numsrc);
gsf->gf_numsrc = count;
+
+ /* The source list is sized by the gf_numsrc the caller left in optval,
+ * not by optlen, which only has to cover the fixed part.
+ */
+ err = sockopt_expand_out(opt, ss_offset +
+ copycount * sizeof(struct sockaddr_storage));
+ if (err)
+ return err;
+
+ /* The caller fills the fixed part in once it knows gf_numsrc. */
+ iov_iter_advance(&opt->iter_out, ss_offset);
+
for (i = 0; i < copycount; i++) {
struct sockaddr_in6 *psin6;
struct sockaddr_storage ss;
@@ -637,9 +649,8 @@ int ip6_mc_msfget(struct sock *sk, struct group_filter *gsf,
memset(&ss, 0, sizeof(ss));
psin6->sin6_family = AF_INET6;
psin6->sin6_addr = psl->sl_addr[i];
- if (copy_to_sockptr_offset(optval, ss_offset, &ss, sizeof(ss)))
+ if (copy_to_iter(&ss, sizeof(ss), &opt->iter_out) != sizeof(ss))
return -EFAULT;
- ss_offset += sizeof(ss);
}
return 0;
}
--
2.53.0-Meta
next prev parent reply other threads:[~2026-09-25 15:56 UTC|newest]
Thread overview: 16+ messages / expand[flat|nested] mbox.gz Atom feed top
2026-09-25 15:55 [PATCH net-next 0/6] ipv4,ipv6: convert the getsockopt switches to sockopt_t Breno Leitao
2026-09-25 15:55 ` [PATCH net-next 1/6] ipv6: reject a negative optlen in do_ipv6_getsockopt() Breno Leitao
2026-09-25 19:02 ` Stanislav Fomichev
2026-09-27 6:53 ` David Laight
2026-09-29 12:14 ` Breno Leitao
2026-09-28 18:55 ` netdev-bot+sashiko
2026-09-25 15:55 ` Breno Leitao [this message]
2026-09-28 18:55 ` [PATCH net-next 2/6] ipv6: mcast: convert ip6_mc_msfget() to sockopt_t netdev-bot+sashiko
2026-09-25 15:55 ` [PATCH net-next 3/6] ipv4: igmp: convert ip_mc_gsfget() " Breno Leitao
2026-09-28 18:55 ` netdev-bot+sashiko
2026-09-25 15:55 ` [PATCH net-next 4/6] ipv4: convert do_ip_getsockopt() " Breno Leitao
2026-09-28 18:55 ` netdev-bot+sashiko
2026-09-25 15:55 ` [PATCH net-next 5/6] ipv6: convert do_ipv6_getsockopt() " Breno Leitao
2026-09-28 18:55 ` netdev-bot+sashiko
2026-09-25 15:55 ` [PATCH net-next 6/6] selftests: net: getsockopt_iter: cover ip and ipv6 Breno Leitao
2026-09-28 18:55 ` netdev-bot+sashiko
Reply instructions:
You may reply publicly to this message via plain-text email
using any one of the following methods:
* Save the following mbox file, import it into your mail client,
and reply-to-all from there: mbox
Avoid top-posting and favor interleaved quoting:
https://en.wikipedia.org/wiki/Posting_style#Interleaved_style
* Reply using the --to, --cc, and --in-reply-to
switches of git-send-email(1):
git send-email \
--in-reply-to=20260925-sockopt_expand_out_v2-v1-2-c3ef2e3bb5c0@debian.org \
--to=leitao@debian.org \
--cc=andrii@kernel.org \
--cc=ast@kernel.org \
--cc=bpf@vger.kernel.org \
--cc=daniel@iogearbox.net \
--cc=davem@davemloft.net \
--cc=david.laight.linux@gmail.com \
--cc=dsahern@kernel.org \
--cc=eddyz87@gmail.com \
--cc=edumazet@google.com \
--cc=emil@etsalapatis.com \
--cc=horms@kernel.org \
--cc=idosch@nvidia.com \
--cc=ihor.solodrai@linux.dev \
--cc=john.fastabend@gmail.com \
--cc=jolsa@kernel.org \
--cc=kernel-team@meta.com \
--cc=kuba@kernel.org \
--cc=linux-kernel@vger.kernel.org \
--cc=linux-kselftest@vger.kernel.org \
--cc=martin.lau@linux.dev \
--cc=memxor@gmail.com \
--cc=netdev@vger.kernel.org \
--cc=pabeni@redhat.com \
--cc=sdf@fomichev.me \
--cc=shuah@kernel.org \
--cc=song@kernel.org \
--cc=yonghong.song@linux.dev \
/path/to/YOUR_REPLY
https://kernel.org/pub/software/scm/git/docs/git-send-email.html
* If your mail client supports setting the In-Reply-To header
via mailto: links, try the mailto: link
Be sure your reply has a Subject: header at the top and a blank line
before the message body.
This is a public inbox, see mirroring instructions
for how to clone and mirror all data and code used for this inbox