iommu.c 7.8 KB
Newer Older
B
Ben-Ami Yassour 已提交
1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18
/*
 * Copyright (c) 2006, Intel Corporation.
 *
 * This program is free software; you can redistribute it and/or modify it
 * under the terms and conditions of the GNU General Public License,
 * version 2, as published by the Free Software Foundation.
 *
 * This program is distributed in the hope it will be useful, but WITHOUT
 * ANY WARRANTY; without even the implied warranty of MERCHANTABILITY or
 * FITNESS FOR A PARTICULAR PURPOSE.  See the GNU General Public License for
 * more details.
 *
 * You should have received a copy of the GNU General Public License along with
 * this program; if not, write to the Free Software Foundation, Inc., 59 Temple
 * Place - Suite 330, Boston, MA 02111-1307 USA.
 *
 * Copyright (C) 2006-2008 Intel Corporation
 * Copyright IBM Corporation, 2008
A
Avi Kivity 已提交
19 20
 * Copyright 2010 Red Hat, Inc. and/or its affiliates.
 *
B
Ben-Ami Yassour 已提交
21 22 23 24 25 26 27 28 29
 * Author: Allen M. Kay <allen.m.kay@intel.com>
 * Author: Weidong Han <weidong.han@intel.com>
 * Author: Ben-Ami Yassour <benami@il.ibm.com>
 */

#include <linux/list.h>
#include <linux/kvm_host.h>
#include <linux/pci.h>
#include <linux/dmar.h>
J
Joerg Roedel 已提交
30
#include <linux/iommu.h>
B
Ben-Ami Yassour 已提交
31 32
#include <linux/intel-iommu.h>

33 34 35 36 37 38
static int allow_unsafe_assigned_interrupts;
module_param_named(allow_unsafe_assigned_interrupts,
		   allow_unsafe_assigned_interrupts, bool, S_IRUGO | S_IWUSR);
MODULE_PARM_DESC(allow_unsafe_assigned_interrupts,
 "Enable device assignment on platforms without interrupt remapping support.");

B
Ben-Ami Yassour 已提交
39 40 41 42
static int kvm_iommu_unmap_memslots(struct kvm *kvm);
static void kvm_iommu_put_pages(struct kvm *kvm,
				gfn_t base_gfn, unsigned long npages);

43 44 45 46 47 48 49 50 51 52 53 54 55 56 57 58 59 60 61
static pfn_t kvm_pin_pages(struct kvm *kvm, struct kvm_memory_slot *slot,
			   gfn_t gfn, unsigned long size)
{
	gfn_t end_gfn;
	pfn_t pfn;

	pfn     = gfn_to_pfn_memslot(kvm, slot, gfn);
	end_gfn = gfn + (size >> PAGE_SHIFT);
	gfn    += 1;

	if (is_error_pfn(pfn))
		return pfn;

	while (gfn < end_gfn)
		gfn_to_pfn_memslot(kvm, slot, gfn++);

	return pfn;
}

62
int kvm_iommu_map_pages(struct kvm *kvm, struct kvm_memory_slot *slot)
B
Ben-Ami Yassour 已提交
63
{
64
	gfn_t gfn, end_gfn;
B
Ben-Ami Yassour 已提交
65
	pfn_t pfn;
66
	int r = 0;
J
Joerg Roedel 已提交
67
	struct iommu_domain *domain = kvm->arch.iommu_domain;
68
	int flags;
B
Ben-Ami Yassour 已提交
69 70 71 72 73

	/* check if iommu exists and in use */
	if (!domain)
		return 0;

74 75 76
	gfn     = slot->base_gfn;
	end_gfn = gfn + slot->npages;

77 78 79 80
	flags = IOMMU_READ | IOMMU_WRITE;
	if (kvm->arch.iommu_flags & KVM_IOMMU_CACHE_COHERENCY)
		flags |= IOMMU_CACHE;

81 82 83 84 85 86 87 88 89 90 91 92 93 94 95 96 97 98 99 100 101 102 103 104 105 106 107 108

	while (gfn < end_gfn) {
		unsigned long page_size;

		/* Check if already mapped */
		if (iommu_iova_to_phys(domain, gfn_to_gpa(gfn))) {
			gfn += 1;
			continue;
		}

		/* Get the page size we could use to map */
		page_size = kvm_host_page_size(kvm, gfn);

		/* Make sure the page_size does not exceed the memslot */
		while ((gfn + (page_size >> PAGE_SHIFT)) > end_gfn)
			page_size >>= 1;

		/* Make sure gfn is aligned to the page size we want to map */
		while ((gfn << PAGE_SHIFT) & (page_size - 1))
			page_size >>= 1;

		/*
		 * Pin all pages we are about to map in memory. This is
		 * important because we unmap and unpin in 4kb steps later.
		 */
		pfn = kvm_pin_pages(kvm, slot, gfn, page_size);
		if (is_error_pfn(pfn)) {
			gfn += 1;
B
Ben-Ami Yassour 已提交
109
			continue;
110
		}
B
Ben-Ami Yassour 已提交
111

112 113 114
		/* Map into IO address space */
		r = iommu_map(domain, gfn_to_gpa(gfn), pfn_to_hpa(pfn),
			      get_order(page_size), flags);
115
		if (r) {
W
Weidong Han 已提交
116
			printk(KERN_ERR "kvm_iommu_map_address:"
J
Joerg Roedel 已提交
117
			       "iommu failed to map pfn=%llx\n", pfn);
B
Ben-Ami Yassour 已提交
118 119
			goto unmap_pages;
		}
120 121 122 123

		gfn += page_size >> PAGE_SHIFT;


B
Ben-Ami Yassour 已提交
124
	}
125

B
Ben-Ami Yassour 已提交
126 127 128
	return 0;

unmap_pages:
129
	kvm_iommu_put_pages(kvm, slot->base_gfn, gfn);
B
Ben-Ami Yassour 已提交
130 131 132 133 134
	return r;
}

static int kvm_iommu_map_memslots(struct kvm *kvm)
{
135
	int i, idx, r = 0;
136
	struct kvm_memslots *slots;
B
Ben-Ami Yassour 已提交
137

138
	idx = srcu_read_lock(&kvm->srcu);
139
	slots = kvm_memslots(kvm);
140 141

	for (i = 0; i < slots->nmemslots; i++) {
142
		r = kvm_iommu_map_pages(kvm, &slots->memslots[i]);
B
Ben-Ami Yassour 已提交
143 144 145
		if (r)
			break;
	}
146
	srcu_read_unlock(&kvm->srcu, idx);
147

B
Ben-Ami Yassour 已提交
148 149 150
	return r;
}

W
Weidong Han 已提交
151 152
int kvm_assign_device(struct kvm *kvm,
		      struct kvm_assigned_dev_kernel *assigned_dev)
B
Ben-Ami Yassour 已提交
153 154
{
	struct pci_dev *pdev = NULL;
J
Joerg Roedel 已提交
155
	struct iommu_domain *domain = kvm->arch.iommu_domain;
156
	int r, last_flags;
B
Ben-Ami Yassour 已提交
157

W
Weidong Han 已提交
158 159 160 161 162 163
	/* check if iommu exists and in use */
	if (!domain)
		return 0;

	pdev = assigned_dev->dev;
	if (pdev == NULL)
B
Ben-Ami Yassour 已提交
164
		return -ENODEV;
W
Weidong Han 已提交
165

J
Joerg Roedel 已提交
166
	r = iommu_attach_device(domain, &pdev->dev);
W
Weidong Han 已提交
167
	if (r) {
168 169
		printk(KERN_ERR "assign device %x:%x:%x.%x failed",
			pci_domain_nr(pdev->bus),
W
Weidong Han 已提交
170 171 172 173
			pdev->bus->number,
			PCI_SLOT(pdev->devfn),
			PCI_FUNC(pdev->devfn));
		return r;
B
Ben-Ami Yassour 已提交
174 175
	}

176 177 178 179 180 181 182 183 184 185 186 187 188 189
	last_flags = kvm->arch.iommu_flags;
	if (iommu_domain_has_cap(kvm->arch.iommu_domain,
				 IOMMU_CAP_CACHE_COHERENCY))
		kvm->arch.iommu_flags |= KVM_IOMMU_CACHE_COHERENCY;

	/* Check if need to update IOMMU page table for guest memory */
	if ((last_flags ^ kvm->arch.iommu_flags) ==
			KVM_IOMMU_CACHE_COHERENCY) {
		kvm_iommu_unmap_memslots(kvm);
		r = kvm_iommu_map_memslots(kvm);
		if (r)
			goto out_unmap;
	}

190 191
	printk(KERN_DEBUG "assign device %x:%x:%x.%x\n",
		assigned_dev->host_segnr,
W
Weidong Han 已提交
192 193 194
		assigned_dev->host_busnr,
		PCI_SLOT(assigned_dev->host_devfn),
		PCI_FUNC(assigned_dev->host_devfn));
B
Ben-Ami Yassour 已提交
195

W
Weidong Han 已提交
196
	return 0;
197 198 199
out_unmap:
	kvm_iommu_unmap_memslots(kvm);
	return r;
W
Weidong Han 已提交
200
}
B
Ben-Ami Yassour 已提交
201

W
Weidong Han 已提交
202 203 204
int kvm_deassign_device(struct kvm *kvm,
			struct kvm_assigned_dev_kernel *assigned_dev)
{
J
Joerg Roedel 已提交
205
	struct iommu_domain *domain = kvm->arch.iommu_domain;
W
Weidong Han 已提交
206 207 208 209 210 211 212 213 214 215
	struct pci_dev *pdev = NULL;

	/* check if iommu exists and in use */
	if (!domain)
		return 0;

	pdev = assigned_dev->dev;
	if (pdev == NULL)
		return -ENODEV;

J
Joerg Roedel 已提交
216
	iommu_detach_device(domain, &pdev->dev);
W
Weidong Han 已提交
217

218 219
	printk(KERN_DEBUG "deassign device %x:%x:%x.%x\n",
		assigned_dev->host_segnr,
W
Weidong Han 已提交
220 221 222 223 224 225 226
		assigned_dev->host_busnr,
		PCI_SLOT(assigned_dev->host_devfn),
		PCI_FUNC(assigned_dev->host_devfn));

	return 0;
}

W
Weidong Han 已提交
227 228 229 230
int kvm_iommu_map_guest(struct kvm *kvm)
{
	int r;

J
Joerg Roedel 已提交
231 232
	if (!iommu_found()) {
		printk(KERN_ERR "%s: iommu not found\n", __func__);
B
Ben-Ami Yassour 已提交
233 234 235
		return -ENODEV;
	}

J
Joerg Roedel 已提交
236 237
	kvm->arch.iommu_domain = iommu_domain_alloc();
	if (!kvm->arch.iommu_domain)
W
Weidong Han 已提交
238
		return -ENOMEM;
B
Ben-Ami Yassour 已提交
239

240 241 242 243 244 245 246 247 248 249 250 251
	if (!allow_unsafe_assigned_interrupts &&
	    !iommu_domain_has_cap(kvm->arch.iommu_domain,
				  IOMMU_CAP_INTR_REMAP)) {
		printk(KERN_WARNING "%s: No interrupt remapping support,"
		       " disallowing device assignment."
		       " Re-enble with \"allow_unsafe_assigned_interrupts=1\""
		       " module option.\n", __func__);
		iommu_domain_free(kvm->arch.iommu_domain);
		kvm->arch.iommu_domain = NULL;
		return -EPERM;
	}

B
Ben-Ami Yassour 已提交
252 253 254 255 256 257 258 259 260 261 262
	r = kvm_iommu_map_memslots(kvm);
	if (r)
		goto out_unmap;

	return 0;

out_unmap:
	kvm_iommu_unmap_memslots(kvm);
	return r;
}

263 264 265 266 267 268 269 270
static void kvm_unpin_pages(struct kvm *kvm, pfn_t pfn, unsigned long npages)
{
	unsigned long i;

	for (i = 0; i < npages; ++i)
		kvm_release_pfn_clean(pfn + i);
}

B
Ben-Ami Yassour 已提交
271
static void kvm_iommu_put_pages(struct kvm *kvm,
W
Weidong Han 已提交
272
				gfn_t base_gfn, unsigned long npages)
B
Ben-Ami Yassour 已提交
273
{
274 275
	struct iommu_domain *domain;
	gfn_t end_gfn, gfn;
B
Ben-Ami Yassour 已提交
276
	pfn_t pfn;
W
Weidong Han 已提交
277 278
	u64 phys;

279 280 281 282
	domain  = kvm->arch.iommu_domain;
	end_gfn = base_gfn + npages;
	gfn     = base_gfn;

W
Weidong Han 已提交
283 284 285
	/* check if iommu exists and in use */
	if (!domain)
		return;
B
Ben-Ami Yassour 已提交
286

287 288 289 290 291
	while (gfn < end_gfn) {
		unsigned long unmap_pages;
		int order;

		/* Get physical address */
J
Joerg Roedel 已提交
292
		phys = iommu_iova_to_phys(domain, gfn_to_gpa(gfn));
293 294 295
		pfn  = phys >> PAGE_SHIFT;

		/* Unmap address from IO address space */
296
		order       = iommu_unmap(domain, gfn_to_gpa(gfn), 0);
297
		unmap_pages = 1ULL << order;
W
Weidong Han 已提交
298

299 300 301 302 303
		/* Unpin all pages we just unmapped to not leak any memory */
		kvm_unpin_pages(kvm, pfn, unmap_pages);

		gfn += unmap_pages;
	}
B
Ben-Ami Yassour 已提交
304 305 306 307
}

static int kvm_iommu_unmap_memslots(struct kvm *kvm)
{
308
	int i, idx;
309 310
	struct kvm_memslots *slots;

311
	idx = srcu_read_lock(&kvm->srcu);
312
	slots = kvm_memslots(kvm);
313

314 315 316
	for (i = 0; i < slots->nmemslots; i++) {
		kvm_iommu_put_pages(kvm, slots->memslots[i].base_gfn,
				    slots->memslots[i].npages);
B
Ben-Ami Yassour 已提交
317
	}
318
	srcu_read_unlock(&kvm->srcu, idx);
B
Ben-Ami Yassour 已提交
319 320 321 322 323 324

	return 0;
}

int kvm_iommu_unmap_guest(struct kvm *kvm)
{
J
Joerg Roedel 已提交
325
	struct iommu_domain *domain = kvm->arch.iommu_domain;
B
Ben-Ami Yassour 已提交
326 327 328 329 330 331

	/* check if iommu exists and in use */
	if (!domain)
		return 0;

	kvm_iommu_unmap_memslots(kvm);
J
Joerg Roedel 已提交
332
	iommu_domain_free(domain);
B
Ben-Ami Yassour 已提交
333 334
	return 0;
}