netdev.vger.kernel.org archive mirror
 help / color / mirror / Atom feed
* [PATCH net-next-2.6] gro: __napi_gro_receive() optimizations
@ 2010-08-25 20:33 Eric Dumazet
  2010-08-25 20:45 ` Stephen Hemminger
  2010-08-25 20:57 ` David Miller
  0 siblings, 2 replies; 15+ messages in thread
From: Eric Dumazet @ 2010-08-25 20:33 UTC (permalink / raw)
  To: David Miller; +Cc: netdev, Herbert Xu

compare_ether_header() can have a special implementation on 64 bit
arches if CONFIG_HAVE_EFFICIENT_UNALIGNED_ACCESS is defined

__napi_gro_receive() can avoid a conditional branch to perform device
match.

__napi_gro_receive() can be used from vlan_gro_common() instead of being
duplicated.

Signed-off-by: Eric Dumazet <eric.dumazet@gmail.com>
---
 include/linux/etherdevice.h |   10 +++++++++-
 include/linux/netdevice.h   |    2 ++
 net/8021q/vlan_core.c       |   14 ++------------
 net/core/dev.c              |   12 +++++++-----
 4 files changed, 20 insertions(+), 18 deletions(-)

diff --git a/include/linux/etherdevice.h b/include/linux/etherdevice.h
index 2308fbb..02144fd 100644
--- a/include/linux/etherdevice.h
+++ b/include/linux/etherdevice.h
@@ -237,13 +237,21 @@ static inline bool is_etherdev_addr(const struct net_device *dev,
  * entry points.
  */
 
-static inline int compare_ether_header(const void *a, const void *b)
+static inline unsigned long compare_ether_header(const void *a, const void *b)
 {
+#if defined(CONFIG_HAVE_EFFICIENT_UNALIGNED_ACCESS) && BITS_PER_LONG == 64
+	unsigned long fold;
+
+	fold = *(unsigned long *)a ^ *(unsigned long *)b;
+	fold |= *(unsigned long *)(a + 6) ^ *(unsigned long *)(b + 6);
+	return fold;
+#else
 	u32 *a32 = (u32 *)((u8 *)a + 2);
 	u32 *b32 = (u32 *)((u8 *)b + 2);
 
 	return (*(u16 *)a ^ *(u16 *)b) | (a32[0] ^ b32[0]) |
 	       (a32[1] ^ b32[1]) | (a32[2] ^ b32[2]);
+#endif
 }
 
 #endif	/* _LINUX_ETHERDEVICE_H */
diff --git a/include/linux/netdevice.h b/include/linux/netdevice.h
index 59962db..629773e 100644
--- a/include/linux/netdevice.h
+++ b/include/linux/netdevice.h
@@ -1695,6 +1695,8 @@ extern gro_result_t	dev_gro_receive(struct napi_struct *napi,
 extern gro_result_t	napi_skb_finish(gro_result_t ret, struct sk_buff *skb);
 extern gro_result_t	napi_gro_receive(struct napi_struct *napi,
 					 struct sk_buff *skb);
+extern gro_result_t   __napi_gro_receive(struct napi_struct *napi,
+					 struct sk_buff *skb);
 extern void		napi_reuse_skb(struct napi_struct *napi,
 				       struct sk_buff *skb);
 extern struct sk_buff *	napi_get_frags(struct napi_struct *napi);
diff --git a/net/8021q/vlan_core.c b/net/8021q/vlan_core.c
index 07eeb5b..947650a 100644
--- a/net/8021q/vlan_core.c
+++ b/net/8021q/vlan_core.c
@@ -102,19 +102,9 @@ vlan_gro_common(struct napi_struct *napi, struct vlan_group *grp,
 	if (vlan_dev)
 		skb->dev = vlan_dev;
 	else if (vlan_id)
-		goto drop;
-
-	for (p = napi->gro_list; p; p = p->next) {
-		NAPI_GRO_CB(p)->same_flow =
-			p->dev == skb->dev && !compare_ether_header(
-				skb_mac_header(p), skb_gro_mac_header(skb));
-		NAPI_GRO_CB(p)->flush = 0;
-	}
-
-	return dev_gro_receive(napi, skb);
+		return GRO_DROP;
 
-drop:
-	return GRO_DROP;
+	return __napi_gro_receive(napi, skb);
 }
 
 gro_result_t vlan_gro_receive(struct napi_struct *napi, struct vlan_group *grp,
diff --git a/net/core/dev.c b/net/core/dev.c
index 859e30f..61e9c96 100644
--- a/net/core/dev.c
+++ b/net/core/dev.c
@@ -3169,21 +3169,23 @@ normal:
 }
 EXPORT_SYMBOL(dev_gro_receive);
 
-static gro_result_t
-__napi_gro_receive(struct napi_struct *napi, struct sk_buff *skb)
+gro_result_t __napi_gro_receive(struct napi_struct *napi, struct sk_buff *skb)
 {
 	struct sk_buff *p;
 
 	for (p = napi->gro_list; p; p = p->next) {
-		NAPI_GRO_CB(p)->same_flow =
-			(p->dev == skb->dev) &&
-			!compare_ether_header(skb_mac_header(p),
+		unsigned long diffs;
+
+		diffs = (unsigned long)p->dev ^ (unsigned long)skb->dev;
+		diffs |= compare_ether_header(skb_mac_header(p),
 					      skb_gro_mac_header(skb));
+		NAPI_GRO_CB(p)->same_flow = !diffs;
 		NAPI_GRO_CB(p)->flush = 0;
 	}
 
 	return dev_gro_receive(napi, skb);
 }
+EXPORT_SYMBOL(__napi_gro_receive);
 
 gro_result_t napi_skb_finish(gro_result_t ret, struct sk_buff *skb)
 {



^ permalink raw reply related	[flat|nested] 15+ messages in thread

end of thread, other threads:[~2010-08-27  5:03 UTC | newest]

Thread overview: 15+ messages (download: mbox.gz follow: Atom feed
-- links below jump to the message on this page --
2010-08-25 20:33 [PATCH net-next-2.6] gro: __napi_gro_receive() optimizations Eric Dumazet
2010-08-25 20:45 ` Stephen Hemminger
2010-08-25 20:54   ` Eric Dumazet
2010-08-25 20:57 ` David Miller
2010-08-25 21:15   ` Eric Dumazet
2010-08-26  7:51     ` David Miller
2010-08-26  9:03       ` Eric Dumazet
2010-08-27  4:35     ` [PATCH net-next-2.6 v3] " Eric Dumazet
2010-08-27  4:38       ` David Miller
2010-08-27  4:42         ` Herbert Xu
2010-08-27  5:02           ` David Miller
2010-08-27  4:43         ` Eric Dumazet
2010-08-27  5:01           ` David Miller
2010-08-27  5:01         ` [PATCH net-next-2.6 v4] " Eric Dumazet
2010-08-27  5:03           ` David Miller

This is a public inbox, see mirroring instructions
for how to clone and mirror all data and code used for this inbox;
as well as URLs for NNTP newsgroup(s).