blob: 3eaedc40ad3d18cb135b6592a5fd9701ae81f070 [file] [log] [blame]
Mark McLoughlinc28b1c12009-10-22 17:49:12 +01001/*
2 * QEMU System Emulator
3 *
4 * Copyright (c) 2003-2008 Fabrice Bellard
5 * Copyright (c) 2009 Red Hat, Inc.
6 *
7 * Permission is hereby granted, free of charge, to any person obtaining a copy
8 * of this software and associated documentation files (the "Software"), to deal
9 * in the Software without restriction, including without limitation the rights
10 * to use, copy, modify, merge, publish, distribute, sublicense, and/or sell
11 * copies of the Software, and to permit persons to whom the Software is
12 * furnished to do so, subject to the following conditions:
13 *
14 * The above copyright notice and this permission notice shall be included in
15 * all copies or substantial portions of the Software.
16 *
17 * THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR
18 * IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY,
19 * FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL
20 * THE AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER
21 * LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM,
22 * OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN
23 * THE SOFTWARE.
24 */
25
26#include "net/tap.h"
27#include "net/tap-linux.h"
28
29#include <net/if.h>
30#include <sys/ioctl.h>
31
32#include "sysemu.h"
33#include "qemu-common.h"
Markus Armbruster2f792012010-02-18 16:24:31 +010034#include "qemu-error.h"
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010035
Michael Tokarev91ca60e2010-06-02 14:33:01 -030036#define PATH_NET_TUN "/dev/net/tun"
37
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010038int tap_open(char *ifname, int ifname_size, int *vnet_hdr, int vnet_hdr_required)
39{
40 struct ifreq ifr;
41 int fd, ret;
Michael S. Tsirkin89e6d682012-11-12 09:13:04 +020042 int len = sizeof(struct virtio_net_hdr);
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010043
Michael Tokarev91ca60e2010-06-02 14:33:01 -030044 TFR(fd = open(PATH_NET_TUN, O_RDWR));
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010045 if (fd < 0) {
Michael Tokarev91ca60e2010-06-02 14:33:01 -030046 error_report("could not open %s: %m", PATH_NET_TUN);
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010047 return -1;
48 }
49 memset(&ifr, 0, sizeof(ifr));
50 ifr.ifr_flags = IFF_TAP | IFF_NO_PI;
51
52 if (*vnet_hdr) {
53 unsigned int features;
54
55 if (ioctl(fd, TUNGETFEATURES, &features) == 0 &&
56 features & IFF_VNET_HDR) {
57 *vnet_hdr = 1;
58 ifr.ifr_flags |= IFF_VNET_HDR;
Pierre Riteau6720b352009-11-25 18:49:34 +000059 } else {
60 *vnet_hdr = 0;
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010061 }
62
63 if (vnet_hdr_required && !*vnet_hdr) {
Markus Armbruster1ecda022010-02-18 17:25:24 +010064 error_report("vnet_hdr=1 requested, but no kernel "
65 "support for IFF_VNET_HDR available");
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010066 close(fd);
67 return -1;
68 }
Michael S. Tsirkin89e6d682012-11-12 09:13:04 +020069 /*
70 * Make sure vnet header size has the default value: for a persistent
71 * tap it might have been modified e.g. by another instance of qemu.
72 * Ignore errors since old kernels do not support this ioctl: in this
73 * case the header size implicitly has the correct value.
74 */
75 ioctl(fd, TUNSETVNETHDRSZ, &len);
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010076 }
77
78 if (ifname[0] != '\0')
79 pstrcpy(ifr.ifr_name, IFNAMSIZ, ifname);
80 else
81 pstrcpy(ifr.ifr_name, IFNAMSIZ, "tap%d");
82 ret = ioctl(fd, TUNSETIFF, (void *) &ifr);
83 if (ret != 0) {
Luiz Capitulino93a73202011-10-14 15:05:10 -030084 if (ifname[0] != '\0') {
85 error_report("could not configure %s (%s): %m", PATH_NET_TUN, ifr.ifr_name);
86 } else {
87 error_report("could not configure %s: %m", PATH_NET_TUN);
88 }
Mark McLoughlinc28b1c12009-10-22 17:49:12 +010089 close(fd);
90 return -1;
91 }
92 pstrcpy(ifname, ifname_size, ifr.ifr_name);
93 fcntl(fd, F_SETFL, O_NONBLOCK);
94 return fd;
95}
Mark McLoughlin15ac9132009-10-22 17:49:13 +010096
Michael S. Tsirkinf157ed22011-02-01 14:25:40 +020097/* sndbuf implements a kind of flow control for tap.
98 * Unfortunately when it's enabled, and packets are sent
99 * to other guests on the same host, the receiver
100 * can lock up the transmitter indefinitely.
101 *
102 * To avoid packet loss, sndbuf should be set to a value lower than the tx
103 * queue capacity of any destination network interface.
Mark McLoughlin15ac9132009-10-22 17:49:13 +0100104 * Ethernet NICs generally have txqueuelen=1000, so 1Mb is
Michael S. Tsirkinf157ed22011-02-01 14:25:40 +0200105 * a good value, given a 1500 byte MTU.
Mark McLoughlin15ac9132009-10-22 17:49:13 +0100106 */
Michael S. Tsirkinf157ed22011-02-01 14:25:40 +0200107#define TAP_DEFAULT_SNDBUF 0
Mark McLoughlin15ac9132009-10-22 17:49:13 +0100108
Laszlo Ersek08c573a2012-07-17 16:17:19 +0200109int tap_set_sndbuf(int fd, const NetdevTapOptions *tap)
Mark McLoughlin15ac9132009-10-22 17:49:13 +0100110{
111 int sndbuf;
112
Laszlo Ersek08c573a2012-07-17 16:17:19 +0200113 sndbuf = !tap->has_sndbuf ? TAP_DEFAULT_SNDBUF :
114 tap->sndbuf > INT_MAX ? INT_MAX :
115 tap->sndbuf;
116
Mark McLoughlin15ac9132009-10-22 17:49:13 +0100117 if (!sndbuf) {
118 sndbuf = INT_MAX;
119 }
120
Laszlo Ersek08c573a2012-07-17 16:17:19 +0200121 if (ioctl(fd, TUNSETSNDBUF, &sndbuf) == -1 && tap->has_sndbuf) {
Markus Armbruster1ecda022010-02-18 17:25:24 +0100122 error_report("TUNSETSNDBUF ioctl failed: %s", strerror(errno));
Mark McLoughlin15ac9132009-10-22 17:49:13 +0100123 return -1;
124 }
125 return 0;
126}
Mark McLoughlindc690042009-10-22 17:49:14 +0100127
128int tap_probe_vnet_hdr(int fd)
129{
130 struct ifreq ifr;
131
132 if (ioctl(fd, TUNGETIFF, &ifr) != 0) {
Markus Armbruster1ecda022010-02-18 17:25:24 +0100133 error_report("TUNGETIFF ioctl() failed: %s", strerror(errno));
Mark McLoughlindc690042009-10-22 17:49:14 +0100134 return 0;
135 }
136
137 return ifr.ifr_flags & IFF_VNET_HDR;
138}
Mark McLoughlin1faac1f2009-10-22 17:49:15 +0100139
Mark McLoughlin9c282712009-10-22 17:49:16 +0100140int tap_probe_has_ufo(int fd)
141{
142 unsigned offload;
143
144 offload = TUN_F_CSUM | TUN_F_UFO;
145
146 if (ioctl(fd, TUNSETOFFLOAD, offload) < 0)
147 return 0;
148
149 return 1;
150}
151
Michael S. Tsirkin445d8922010-07-16 11:16:06 +0300152/* Verify that we can assign given length */
153int tap_probe_vnet_hdr_len(int fd, int len)
154{
155 int orig;
156 if (ioctl(fd, TUNGETVNETHDRSZ, &orig) == -1) {
157 return 0;
158 }
159 if (ioctl(fd, TUNSETVNETHDRSZ, &len) == -1) {
160 return 0;
161 }
162 /* Restore original length: we can't handle failure. */
163 if (ioctl(fd, TUNSETVNETHDRSZ, &orig) == -1) {
164 fprintf(stderr, "TUNGETVNETHDRSZ ioctl() failed: %s. Exiting.\n",
165 strerror(errno));
166 assert(0);
167 return -errno;
168 }
169 return 1;
170}
171
172void tap_fd_set_vnet_hdr_len(int fd, int len)
173{
174 if (ioctl(fd, TUNSETVNETHDRSZ, &len) == -1) {
175 fprintf(stderr, "TUNSETVNETHDRSZ ioctl() failed: %s. Exiting.\n",
176 strerror(errno));
177 assert(0);
178 }
179}
180
Mark McLoughlin1faac1f2009-10-22 17:49:15 +0100181void tap_fd_set_offload(int fd, int csum, int tso4,
182 int tso6, int ecn, int ufo)
183{
184 unsigned int offload = 0;
185
Pierre Riteau2e503262009-11-25 18:49:35 +0000186 /* Check if our kernel supports TUNSETOFFLOAD */
187 if (ioctl(fd, TUNSETOFFLOAD, 0) != 0 && errno == EINVAL) {
188 return;
189 }
190
Mark McLoughlin1faac1f2009-10-22 17:49:15 +0100191 if (csum) {
192 offload |= TUN_F_CSUM;
193 if (tso4)
194 offload |= TUN_F_TSO4;
195 if (tso6)
196 offload |= TUN_F_TSO6;
197 if ((tso4 || tso6) && ecn)
198 offload |= TUN_F_TSO_ECN;
199 if (ufo)
200 offload |= TUN_F_UFO;
201 }
202
203 if (ioctl(fd, TUNSETOFFLOAD, offload) != 0) {
204 offload &= ~TUN_F_UFO;
205 if (ioctl(fd, TUNSETOFFLOAD, offload) != 0) {
206 fprintf(stderr, "TUNSETOFFLOAD ioctl() failed: %s\n",
207 strerror(errno));
208 }
209 }
210}