Commit 1216ce9d authored by Vlad Buslov's avatar Vlad Buslov Committed by Saeed Mahameed

net/mlx5e: Extend neigh hash entry with rcu

To remove dependency on rtnl lock and to allow unlocked iteration over list
of neigh hash entries, extend nhe with rcu. Change operations on neigh list
to their rcu counterparts and free neigh hash entry with rcu timeout.

Introduce mlx5e_get_next_nhe() helper that is used to iterate over rcu
neigh list with reference to nhe taken.
Signed-off-by: default avatarVlad Buslov <vladbu@mellanox.com>
Reviewed-by: default avatarJianbo Liu <jianbol@mellanox.com>
Reviewed-by: default avatarRoi Dayan <roid@mellanox.com>
Signed-off-by: default avatarSaeed Mahameed <saeedm@mellanox.com>
parent 61081f9c
...@@ -535,28 +535,56 @@ static void mlx5e_rep_neigh_entry_release(struct mlx5e_neigh_hash_entry *nhe) ...@@ -535,28 +535,56 @@ static void mlx5e_rep_neigh_entry_release(struct mlx5e_neigh_hash_entry *nhe)
{ {
if (refcount_dec_and_test(&nhe->refcnt)) { if (refcount_dec_and_test(&nhe->refcnt)) {
mlx5e_rep_neigh_entry_remove(nhe); mlx5e_rep_neigh_entry_remove(nhe);
kfree(nhe); kfree_rcu(nhe, rcu);
} }
} }
static struct mlx5e_neigh_hash_entry *
mlx5e_get_next_nhe(struct mlx5e_rep_priv *rpriv,
struct mlx5e_neigh_hash_entry *nhe)
{
struct mlx5e_neigh_hash_entry *next = NULL;
rcu_read_lock();
for (next = nhe ?
list_next_or_null_rcu(&rpriv->neigh_update.neigh_list,
&nhe->neigh_list,
struct mlx5e_neigh_hash_entry,
neigh_list) :
list_first_or_null_rcu(&rpriv->neigh_update.neigh_list,
struct mlx5e_neigh_hash_entry,
neigh_list);
next;
next = list_next_or_null_rcu(&rpriv->neigh_update.neigh_list,
&next->neigh_list,
struct mlx5e_neigh_hash_entry,
neigh_list))
if (mlx5e_rep_neigh_entry_hold(next))
break;
rcu_read_unlock();
if (nhe)
mlx5e_rep_neigh_entry_release(nhe);
return next;
}
static void mlx5e_rep_neigh_stats_work(struct work_struct *work) static void mlx5e_rep_neigh_stats_work(struct work_struct *work)
{ {
struct mlx5e_rep_priv *rpriv = container_of(work, struct mlx5e_rep_priv, struct mlx5e_rep_priv *rpriv = container_of(work, struct mlx5e_rep_priv,
neigh_update.neigh_stats_work.work); neigh_update.neigh_stats_work.work);
struct net_device *netdev = rpriv->netdev; struct net_device *netdev = rpriv->netdev;
struct mlx5e_priv *priv = netdev_priv(netdev); struct mlx5e_priv *priv = netdev_priv(netdev);
struct mlx5e_neigh_hash_entry *nhe; struct mlx5e_neigh_hash_entry *nhe = NULL;
rtnl_lock(); rtnl_lock();
if (!list_empty(&rpriv->neigh_update.neigh_list)) if (!list_empty(&rpriv->neigh_update.neigh_list))
mlx5e_rep_queue_neigh_stats_work(priv); mlx5e_rep_queue_neigh_stats_work(priv);
list_for_each_entry(nhe, &rpriv->neigh_update.neigh_list, neigh_list) { while ((nhe = mlx5e_get_next_nhe(rpriv, nhe)) != NULL)
if (mlx5e_rep_neigh_entry_hold(nhe)) { mlx5e_tc_update_neigh_used_value(nhe);
mlx5e_tc_update_neigh_used_value(nhe);
mlx5e_rep_neigh_entry_release(nhe);
}
}
rtnl_unlock(); rtnl_unlock();
} }
...@@ -883,13 +911,9 @@ static int mlx5e_rep_netevent_event(struct notifier_block *nb, ...@@ -883,13 +911,9 @@ static int mlx5e_rep_netevent_event(struct notifier_block *nb,
m_neigh.family = n->ops->family; m_neigh.family = n->ops->family;
memcpy(&m_neigh.dst_ip, n->primary_key, n->tbl->key_len); memcpy(&m_neigh.dst_ip, n->primary_key, n->tbl->key_len);
/* We are in atomic context and can't take RTNL mutex, so use rcu_read_lock();
* spin_lock_bh to lookup the neigh table. bh is used since
* netevent can be called from a softirq context.
*/
spin_lock_bh(&neigh_update->encap_lock);
nhe = mlx5e_rep_neigh_entry_lookup(priv, &m_neigh); nhe = mlx5e_rep_neigh_entry_lookup(priv, &m_neigh);
spin_unlock_bh(&neigh_update->encap_lock); rcu_read_unlock();
if (!nhe) if (!nhe)
return NOTIFY_DONE; return NOTIFY_DONE;
...@@ -910,19 +934,15 @@ static int mlx5e_rep_netevent_event(struct notifier_block *nb, ...@@ -910,19 +934,15 @@ static int mlx5e_rep_netevent_event(struct notifier_block *nb,
#endif #endif
return NOTIFY_DONE; return NOTIFY_DONE;
/* We are in atomic context and can't take RTNL mutex, rcu_read_lock();
* so use spin_lock_bh to walk the neigh list and look for list_for_each_entry_rcu(nhe, &neigh_update->neigh_list,
* the relevant device. bh is used since netevent can be neigh_list) {
* called from a softirq context.
*/
spin_lock_bh(&neigh_update->encap_lock);
list_for_each_entry(nhe, &neigh_update->neigh_list, neigh_list) {
if (p->dev == nhe->m_neigh.dev) { if (p->dev == nhe->m_neigh.dev) {
found = true; found = true;
break; break;
} }
} }
spin_unlock_bh(&neigh_update->encap_lock); rcu_read_unlock();
if (!found) if (!found)
return NOTIFY_DONE; return NOTIFY_DONE;
...@@ -995,7 +1015,7 @@ static int mlx5e_rep_neigh_entry_insert(struct mlx5e_priv *priv, ...@@ -995,7 +1015,7 @@ static int mlx5e_rep_neigh_entry_insert(struct mlx5e_priv *priv,
if (err) if (err)
return err; return err;
list_add(&nhe->neigh_list, &rpriv->neigh_update.neigh_list); list_add_rcu(&nhe->neigh_list, &rpriv->neigh_update.neigh_list);
return err; return err;
} }
...@@ -1006,7 +1026,7 @@ static void mlx5e_rep_neigh_entry_remove(struct mlx5e_neigh_hash_entry *nhe) ...@@ -1006,7 +1026,7 @@ static void mlx5e_rep_neigh_entry_remove(struct mlx5e_neigh_hash_entry *nhe)
spin_lock_bh(&rpriv->neigh_update.encap_lock); spin_lock_bh(&rpriv->neigh_update.encap_lock);
list_del(&nhe->neigh_list); list_del_rcu(&nhe->neigh_list);
rhashtable_remove_fast(&rpriv->neigh_update.neigh_ht, rhashtable_remove_fast(&rpriv->neigh_update.neigh_ht,
&nhe->rhash_node, &nhe->rhash_node,
......
...@@ -138,6 +138,8 @@ struct mlx5e_neigh_hash_entry { ...@@ -138,6 +138,8 @@ struct mlx5e_neigh_hash_entry {
* 'used' value and avoid neigh deleting by the kernel. * 'used' value and avoid neigh deleting by the kernel.
*/ */
unsigned long reported_lastuse; unsigned long reported_lastuse;
struct rcu_head rcu;
}; };
enum { enum {
......
Markdown is supported
0%
or
You are about to add 0 people to the discussion. Proceed with caution.
Finish editing this message first!
Please register or to comment