if_virt.c revision 1.26 1 1.26 tls /* $NetBSD: if_virt.c,v 1.26 2011/11/19 22:51:31 tls Exp $ */
2 1.1 pooka
3 1.1 pooka /*
4 1.1 pooka * Copyright (c) 2008 Antti Kantee. All Rights Reserved.
5 1.1 pooka *
6 1.1 pooka * Redistribution and use in source and binary forms, with or without
7 1.1 pooka * modification, are permitted provided that the following conditions
8 1.1 pooka * are met:
9 1.1 pooka * 1. Redistributions of source code must retain the above copyright
10 1.1 pooka * notice, this list of conditions and the following disclaimer.
11 1.1 pooka * 2. Redistributions in binary form must reproduce the above copyright
12 1.1 pooka * notice, this list of conditions and the following disclaimer in the
13 1.1 pooka * documentation and/or other materials provided with the distribution.
14 1.1 pooka *
15 1.1 pooka * THIS SOFTWARE IS PROVIDED BY THE AUTHOR ``AS IS'' AND ANY EXPRESS
16 1.1 pooka * OR IMPLIED WARRANTIES, INCLUDING, BUT NOT LIMITED TO, THE IMPLIED
17 1.1 pooka * WARRANTIES OF MERCHANTABILITY AND FITNESS FOR A PARTICULAR PURPOSE ARE
18 1.1 pooka * DISCLAIMED. IN NO EVENT SHALL THE AUTHOR OR CONTRIBUTORS BE LIABLE
19 1.1 pooka * FOR ANY DIRECT, INDIRECT, INCIDENTAL, SPECIAL, EXEMPLARY, OR CONSEQUENTIAL
20 1.1 pooka * DAMAGES (INCLUDING, BUT NOT LIMITED TO, PROCUREMENT OF SUBSTITUTE GOODS OR
21 1.1 pooka * SERVICES; LOSS OF USE, DATA, OR PROFITS; OR BUSINESS INTERRUPTION)
22 1.1 pooka * HOWEVER CAUSED AND ON ANY THEORY OF LIABILITY, WHETHER IN CONTRACT, STRICT
23 1.1 pooka * LIABILITY, OR TORT (INCLUDING NEGLIGENCE OR OTHERWISE) ARISING IN ANY WAY
24 1.1 pooka * OUT OF THE USE OF THIS SOFTWARE, EVEN IF ADVISED OF THE POSSIBILITY OF
25 1.1 pooka * SUCH DAMAGE.
26 1.1 pooka */
27 1.1 pooka
28 1.5 pooka #include <sys/cdefs.h>
29 1.26 tls __KERNEL_RCSID(0, "$NetBSD: if_virt.c,v 1.26 2011/11/19 22:51:31 tls Exp $");
30 1.5 pooka
31 1.1 pooka #include <sys/param.h>
32 1.1 pooka #include <sys/condvar.h>
33 1.1 pooka #include <sys/fcntl.h>
34 1.21 pooka #include <sys/kernel.h>
35 1.1 pooka #include <sys/kmem.h>
36 1.1 pooka #include <sys/kthread.h>
37 1.1 pooka #include <sys/mutex.h>
38 1.11 pooka #include <sys/poll.h>
39 1.1 pooka #include <sys/sockio.h>
40 1.1 pooka #include <sys/socketvar.h>
41 1.26 tls #include <sys/cprng.h>
42 1.1 pooka
43 1.15 pooka #include <net/bpf.h>
44 1.1 pooka #include <net/if.h>
45 1.1 pooka #include <net/if_ether.h>
46 1.1 pooka #include <net/if_tap.h>
47 1.1 pooka
48 1.1 pooka #include <netinet/in.h>
49 1.1 pooka #include <netinet/in_var.h>
50 1.1 pooka
51 1.1 pooka #include <rump/rump.h>
52 1.1 pooka #include <rump/rumpuser.h>
53 1.1 pooka
54 1.1 pooka #include "rump_private.h"
55 1.10 pooka #include "rump_net_private.h"
56 1.1 pooka
57 1.1 pooka /*
58 1.1 pooka * Virtual interface for userspace purposes. Uses tap(4) to
59 1.1 pooka * interface with the kernel and just simply shovels data
60 1.1 pooka * to/from /dev/tap.
61 1.1 pooka */
62 1.1 pooka
63 1.1 pooka #define VIRTIF_BASE "virt"
64 1.1 pooka
65 1.1 pooka static int virtif_init(struct ifnet *);
66 1.1 pooka static int virtif_ioctl(struct ifnet *, u_long, void *);
67 1.1 pooka static void virtif_start(struct ifnet *);
68 1.1 pooka static void virtif_stop(struct ifnet *, int);
69 1.1 pooka
70 1.1 pooka struct virtif_sc {
71 1.7 pooka struct ethercom sc_ec;
72 1.1 pooka int sc_tapfd;
73 1.21 pooka bool sc_dying;
74 1.21 pooka struct lwp *sc_l_snd, *sc_l_rcv;
75 1.21 pooka kmutex_t sc_mtx;
76 1.21 pooka kcondvar_t sc_cv;
77 1.1 pooka };
78 1.1 pooka
79 1.21 pooka static void virtif_receiver(void *);
80 1.8 pooka static void virtif_sender(void *);
81 1.20 pooka static int virtif_clone(struct if_clone *, int);
82 1.20 pooka static int virtif_unclone(struct ifnet *);
83 1.1 pooka
84 1.20 pooka struct if_clone virtif_cloner =
85 1.20 pooka IF_CLONE_INITIALIZER(VIRTIF_BASE, virtif_clone, virtif_unclone);
86 1.1 pooka
87 1.8 pooka int
88 1.14 pooka rump_virtif_create(int num)
89 1.1 pooka {
90 1.1 pooka struct virtif_sc *sc;
91 1.1 pooka struct ifnet *ifp;
92 1.3 pooka uint8_t enaddr[ETHER_ADDR_LEN] = { 0xb2, 0x0a, 0x00, 0x0b, 0x0e, 0x01 };
93 1.8 pooka char tapdev[16];
94 1.21 pooka int fd, error = 0;
95 1.21 pooka
96 1.21 pooka if (num >= 0x100)
97 1.21 pooka return E2BIG;
98 1.1 pooka
99 1.8 pooka snprintf(tapdev, sizeof(tapdev), "/dev/tap%d", num);
100 1.8 pooka fd = rumpuser_open(tapdev, O_RDWR, &error);
101 1.1 pooka if (fd == -1) {
102 1.19 pooka printf("virtif_create: can't open /dev/tap%d: %d\n",
103 1.19 pooka num, error);
104 1.1 pooka return error;
105 1.1 pooka }
106 1.26 tls enaddr[2] = cprng_fast32() & 0xff;
107 1.8 pooka enaddr[5] = num;
108 1.1 pooka
109 1.1 pooka sc = kmem_zalloc(sizeof(*sc), KM_SLEEP);
110 1.21 pooka sc->sc_dying = false;
111 1.1 pooka sc->sc_tapfd = fd;
112 1.1 pooka
113 1.21 pooka mutex_init(&sc->sc_mtx, MUTEX_DEFAULT, IPL_NONE);
114 1.21 pooka cv_init(&sc->sc_cv, "virtsnd");
115 1.7 pooka ifp = &sc->sc_ec.ec_if;
116 1.8 pooka sprintf(ifp->if_xname, "%s%d", VIRTIF_BASE, num);
117 1.1 pooka ifp->if_softc = sc;
118 1.21 pooka
119 1.21 pooka if (rump_threads) {
120 1.24 rmind if ((error = kthread_create(PRI_NONE, KTHREAD_MUSTJOIN, NULL,
121 1.21 pooka virtif_receiver, ifp, &sc->sc_l_rcv, "virtifr")) != 0)
122 1.21 pooka goto out;
123 1.21 pooka
124 1.21 pooka if ((error = kthread_create(PRI_NONE,
125 1.24 rmind KTHREAD_MUSTJOIN | KTHREAD_MPSAFE, NULL,
126 1.21 pooka virtif_sender, ifp, &sc->sc_l_snd, "virtifs")) != 0)
127 1.21 pooka goto out;
128 1.21 pooka } else {
129 1.21 pooka printf("WARNING: threads not enabled, receive NOT working\n");
130 1.21 pooka }
131 1.21 pooka
132 1.1 pooka ifp->if_flags = IFF_BROADCAST | IFF_SIMPLEX | IFF_MULTICAST;
133 1.1 pooka ifp->if_init = virtif_init;
134 1.1 pooka ifp->if_ioctl = virtif_ioctl;
135 1.1 pooka ifp->if_start = virtif_start;
136 1.1 pooka ifp->if_stop = virtif_stop;
137 1.21 pooka IFQ_SET_READY(&ifp->if_snd);
138 1.1 pooka
139 1.1 pooka if_attach(ifp);
140 1.1 pooka ether_ifattach(ifp, enaddr);
141 1.1 pooka
142 1.21 pooka out:
143 1.21 pooka if (error) {
144 1.21 pooka virtif_unclone(ifp);
145 1.21 pooka }
146 1.21 pooka
147 1.21 pooka return error;
148 1.1 pooka }
149 1.1 pooka
150 1.1 pooka static int
151 1.20 pooka virtif_clone(struct if_clone *ifc, int unit)
152 1.20 pooka {
153 1.20 pooka
154 1.20 pooka return rump_virtif_create(unit);
155 1.20 pooka }
156 1.20 pooka
157 1.20 pooka static int
158 1.20 pooka virtif_unclone(struct ifnet *ifp)
159 1.20 pooka {
160 1.21 pooka struct virtif_sc *sc = ifp->if_softc;
161 1.20 pooka
162 1.21 pooka mutex_enter(&sc->sc_mtx);
163 1.21 pooka if (sc->sc_dying) {
164 1.21 pooka mutex_exit(&sc->sc_mtx);
165 1.21 pooka return EINPROGRESS;
166 1.21 pooka }
167 1.21 pooka sc->sc_dying = true;
168 1.21 pooka cv_broadcast(&sc->sc_cv);
169 1.21 pooka mutex_exit(&sc->sc_mtx);
170 1.21 pooka
171 1.21 pooka virtif_stop(ifp, 1);
172 1.21 pooka if_down(ifp);
173 1.21 pooka
174 1.21 pooka if (sc->sc_l_snd) {
175 1.21 pooka kthread_join(sc->sc_l_snd);
176 1.21 pooka sc->sc_l_snd = NULL;
177 1.21 pooka }
178 1.21 pooka if (sc->sc_l_rcv) {
179 1.21 pooka kthread_join(sc->sc_l_rcv);
180 1.21 pooka sc->sc_l_rcv = NULL;
181 1.21 pooka }
182 1.21 pooka
183 1.21 pooka rumpuser_close(sc->sc_tapfd, NULL);
184 1.21 pooka
185 1.21 pooka mutex_destroy(&sc->sc_mtx);
186 1.21 pooka cv_destroy(&sc->sc_cv);
187 1.21 pooka kmem_free(sc, sizeof(*sc));
188 1.21 pooka
189 1.21 pooka ether_ifdetach(ifp);
190 1.21 pooka if_detach(ifp);
191 1.21 pooka
192 1.21 pooka return 0;
193 1.20 pooka }
194 1.20 pooka
195 1.20 pooka static int
196 1.1 pooka virtif_init(struct ifnet *ifp)
197 1.1 pooka {
198 1.21 pooka struct virtif_sc *sc = ifp->if_softc;
199 1.1 pooka
200 1.1 pooka ifp->if_flags |= IFF_RUNNING;
201 1.21 pooka
202 1.21 pooka mutex_enter(&sc->sc_mtx);
203 1.21 pooka cv_broadcast(&sc->sc_cv);
204 1.21 pooka mutex_exit(&sc->sc_mtx);
205 1.8 pooka
206 1.1 pooka return 0;
207 1.1 pooka }
208 1.1 pooka
209 1.1 pooka static int
210 1.1 pooka virtif_ioctl(struct ifnet *ifp, u_long cmd, void *data)
211 1.1 pooka {
212 1.1 pooka int s, rv;
213 1.1 pooka
214 1.1 pooka s = splnet();
215 1.1 pooka rv = ether_ioctl(ifp, cmd, data);
216 1.9 pooka if (rv == ENETRESET)
217 1.9 pooka rv = 0;
218 1.1 pooka splx(s);
219 1.1 pooka
220 1.1 pooka return rv;
221 1.1 pooka }
222 1.1 pooka
223 1.1 pooka /* just send everything in-context */
224 1.1 pooka static void
225 1.1 pooka virtif_start(struct ifnet *ifp)
226 1.1 pooka {
227 1.1 pooka struct virtif_sc *sc = ifp->if_softc;
228 1.1 pooka
229 1.21 pooka mutex_enter(&sc->sc_mtx);
230 1.21 pooka ifp->if_flags |= IFF_OACTIVE;
231 1.21 pooka cv_broadcast(&sc->sc_cv);
232 1.21 pooka mutex_exit(&sc->sc_mtx);
233 1.1 pooka }
234 1.1 pooka
235 1.1 pooka static void
236 1.1 pooka virtif_stop(struct ifnet *ifp, int disable)
237 1.1 pooka {
238 1.21 pooka struct virtif_sc *sc = ifp->if_softc;
239 1.1 pooka
240 1.21 pooka ifp->if_flags &= ~IFF_RUNNING;
241 1.21 pooka
242 1.21 pooka mutex_enter(&sc->sc_mtx);
243 1.21 pooka cv_broadcast(&sc->sc_cv);
244 1.21 pooka mutex_exit(&sc->sc_mtx);
245 1.1 pooka }
246 1.1 pooka
247 1.21 pooka #define POLLTIMO_MS 1
248 1.1 pooka static void
249 1.21 pooka virtif_receiver(void *arg)
250 1.1 pooka {
251 1.1 pooka struct ifnet *ifp = arg;
252 1.1 pooka struct virtif_sc *sc = ifp->if_softc;
253 1.1 pooka struct mbuf *m;
254 1.1 pooka size_t plen = ETHER_MAX_LEN_JUMBO+1;
255 1.21 pooka struct pollfd pfd;
256 1.1 pooka ssize_t n;
257 1.21 pooka int error, rv;
258 1.21 pooka
259 1.21 pooka pfd.fd = sc->sc_tapfd;
260 1.21 pooka pfd.events = POLLIN;
261 1.21 pooka
262 1.1 pooka for (;;) {
263 1.1 pooka m = m_gethdr(M_WAIT, MT_DATA);
264 1.1 pooka MEXTMALLOC(m, plen, M_WAIT);
265 1.1 pooka
266 1.11 pooka again:
267 1.21 pooka /* poll, but periodically check if we should die */
268 1.21 pooka rv = rumpuser_poll(&pfd, 1, POLLTIMO_MS, &error);
269 1.21 pooka if (sc->sc_dying) {
270 1.21 pooka m_freem(m);
271 1.21 pooka break;
272 1.21 pooka }
273 1.21 pooka if (rv == 0)
274 1.21 pooka goto again;
275 1.21 pooka
276 1.1 pooka n = rumpuser_read(sc->sc_tapfd, mtod(m, void *), plen, &error);
277 1.1 pooka KASSERT(n < ETHER_MAX_LEN_JUMBO);
278 1.19 pooka if (__predict_false(n < 0)) {
279 1.11 pooka if (n == -1 && error == EAGAIN) {
280 1.11 pooka goto again;
281 1.11 pooka }
282 1.19 pooka
283 1.25 yamt printf("%s: read from /dev/tap failed. host is down?\n",
284 1.21 pooka ifp->if_xname);
285 1.21 pooka mutex_enter(&sc->sc_mtx);
286 1.21 pooka /* could check if need go, done soon anyway */
287 1.21 pooka cv_timedwait(&sc->sc_cv, &sc->sc_mtx, hz);
288 1.21 pooka mutex_exit(&sc->sc_mtx);
289 1.21 pooka goto again;
290 1.1 pooka }
291 1.19 pooka
292 1.19 pooka /* tap sometimes returns EOF. don't sweat it and plow on */
293 1.19 pooka if (__predict_false(n == 0))
294 1.19 pooka goto again;
295 1.19 pooka
296 1.21 pooka /* discard if we're not up */
297 1.21 pooka if ((ifp->if_flags & IFF_RUNNING) == 0)
298 1.21 pooka goto again;
299 1.21 pooka
300 1.1 pooka m->m_len = m->m_pkthdr.len = n;
301 1.1 pooka m->m_pkthdr.rcvif = ifp;
302 1.18 joerg bpf_mtap(ifp, m);
303 1.1 pooka ether_input(ifp, m);
304 1.1 pooka }
305 1.1 pooka
306 1.21 pooka kthread_exit(0);
307 1.1 pooka }
308 1.8 pooka
309 1.12 pooka /* lazy bum stetson-harrison magic value */
310 1.12 pooka #define LB_SH 32
311 1.8 pooka static void
312 1.8 pooka virtif_sender(void *arg)
313 1.8 pooka {
314 1.8 pooka struct ifnet *ifp = arg;
315 1.8 pooka struct virtif_sc *sc = ifp->if_softc;
316 1.8 pooka struct mbuf *m, *m0;
317 1.12 pooka struct rumpuser_iovec io[LB_SH];
318 1.8 pooka int i, error;
319 1.8 pooka
320 1.21 pooka mutex_enter(&sc->sc_mtx);
321 1.21 pooka KERNEL_LOCK(1, NULL);
322 1.21 pooka while (!sc->sc_dying) {
323 1.23 mrg if (!(ifp->if_flags & IFF_RUNNING)) {
324 1.21 pooka cv_wait(&sc->sc_cv, &sc->sc_mtx);
325 1.21 pooka continue;
326 1.21 pooka }
327 1.8 pooka IF_DEQUEUE(&ifp->if_snd, m0);
328 1.8 pooka if (!m0) {
329 1.21 pooka ifp->if_flags &= ~IFF_OACTIVE;
330 1.21 pooka cv_wait(&sc->sc_cv, &sc->sc_mtx);
331 1.8 pooka continue;
332 1.8 pooka }
333 1.21 pooka mutex_exit(&sc->sc_mtx);
334 1.8 pooka
335 1.8 pooka m = m0;
336 1.12 pooka for (i = 0; i < LB_SH && m; i++) {
337 1.8 pooka io[i].iov_base = mtod(m, void *);
338 1.8 pooka io[i].iov_len = m->m_len;
339 1.8 pooka m = m->m_next;
340 1.8 pooka }
341 1.12 pooka if (i == LB_SH)
342 1.8 pooka panic("lazy bum");
343 1.17 joerg bpf_mtap(ifp, m0);
344 1.21 pooka KERNEL_UNLOCK_LAST(curlwp);
345 1.21 pooka
346 1.8 pooka rumpuser_writev(sc->sc_tapfd, io, i, &error);
347 1.21 pooka
348 1.21 pooka KERNEL_LOCK(1, NULL);
349 1.8 pooka m_freem(m0);
350 1.21 pooka mutex_enter(&sc->sc_mtx);
351 1.8 pooka }
352 1.21 pooka KERNEL_UNLOCK_LAST(curlwp);
353 1.21 pooka
354 1.21 pooka mutex_exit(&sc->sc_mtx);
355 1.8 pooka
356 1.21 pooka kthread_exit(0);
357 1.8 pooka }
358 1.10 pooka
359 1.10 pooka /*
360 1.10 pooka * dummyif is a nada-interface.
361 1.10 pooka * As it requires nothing external, it can be used for testing
362 1.10 pooka * interface configuration.
363 1.10 pooka */
364 1.10 pooka static int dummyif_init(struct ifnet *);
365 1.10 pooka static void dummyif_start(struct ifnet *);
366 1.10 pooka
367 1.10 pooka void
368 1.10 pooka rump_dummyif_create()
369 1.10 pooka {
370 1.10 pooka struct ifnet *ifp;
371 1.10 pooka struct ethercom *ec;
372 1.10 pooka uint8_t enaddr[ETHER_ADDR_LEN] = { 0xb2, 0x0a, 0x00, 0x0b, 0x0e, 0x01 };
373 1.10 pooka
374 1.26 tls enaddr[2] = cprng_fast32() & 0xff;
375 1.26 tls enaddr[5] = cprng_fast32() & 0xff;
376 1.10 pooka
377 1.10 pooka ec = kmem_zalloc(sizeof(*ec), KM_SLEEP);
378 1.10 pooka
379 1.10 pooka ifp = &ec->ec_if;
380 1.10 pooka strlcpy(ifp->if_xname, "dummy0", sizeof(ifp->if_xname));
381 1.10 pooka ifp->if_softc = ifp;
382 1.10 pooka ifp->if_flags = IFF_BROADCAST | IFF_SIMPLEX | IFF_MULTICAST;
383 1.10 pooka ifp->if_init = dummyif_init;
384 1.10 pooka ifp->if_ioctl = virtif_ioctl;
385 1.10 pooka ifp->if_start = dummyif_start;
386 1.10 pooka
387 1.10 pooka if_attach(ifp);
388 1.10 pooka ether_ifattach(ifp, enaddr);
389 1.10 pooka }
390 1.10 pooka
391 1.10 pooka static int
392 1.10 pooka dummyif_init(struct ifnet *ifp)
393 1.10 pooka {
394 1.10 pooka
395 1.10 pooka ifp->if_flags |= IFF_RUNNING;
396 1.10 pooka return 0;
397 1.10 pooka }
398 1.10 pooka
399 1.10 pooka static void
400 1.10 pooka dummyif_start(struct ifnet *ifp)
401 1.10 pooka {
402 1.10 pooka
403 1.10 pooka }
404