iommu.c 8.2 KB
Newer Older
B
Ben-Ami Yassour 已提交
1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18
/*
 * Copyright (c) 2006, Intel Corporation.
 *
 * This program is free software; you can redistribute it and/or modify it
 * under the terms and conditions of the GNU General Public License,
 * version 2, as published by the Free Software Foundation.
 *
 * This program is distributed in the hope it will be useful, but WITHOUT
 * ANY WARRANTY; without even the implied warranty of MERCHANTABILITY or
 * FITNESS FOR A PARTICULAR PURPOSE.  See the GNU General Public License for
 * more details.
 *
 * You should have received a copy of the GNU General Public License along with
 * this program; if not, write to the Free Software Foundation, Inc., 59 Temple
 * Place - Suite 330, Boston, MA 02111-1307 USA.
 *
 * Copyright (C) 2006-2008 Intel Corporation
 * Copyright IBM Corporation, 2008
A
Avi Kivity 已提交
19 20
 * Copyright 2010 Red Hat, Inc. and/or its affiliates.
 *
B
Ben-Ami Yassour 已提交
21 22 23 24 25 26 27
 * Author: Allen M. Kay <allen.m.kay@intel.com>
 * Author: Weidong Han <weidong.han@intel.com>
 * Author: Ben-Ami Yassour <benami@il.ibm.com>
 */

#include <linux/list.h>
#include <linux/kvm_host.h>
28
#include <linux/module.h>
B
Ben-Ami Yassour 已提交
29
#include <linux/pci.h>
30
#include <linux/stat.h>
B
Ben-Ami Yassour 已提交
31
#include <linux/dmar.h>
J
Joerg Roedel 已提交
32
#include <linux/iommu.h>
B
Ben-Ami Yassour 已提交
33 34
#include <linux/intel-iommu.h>

35
static bool allow_unsafe_assigned_interrupts;
36 37 38 39 40
module_param_named(allow_unsafe_assigned_interrupts,
		   allow_unsafe_assigned_interrupts, bool, S_IRUGO | S_IWUSR);
MODULE_PARM_DESC(allow_unsafe_assigned_interrupts,
 "Enable device assignment on platforms without interrupt remapping support.");

B
Ben-Ami Yassour 已提交
41 42 43 44
static int kvm_iommu_unmap_memslots(struct kvm *kvm);
static void kvm_iommu_put_pages(struct kvm *kvm,
				gfn_t base_gfn, unsigned long npages);

45
static pfn_t kvm_pin_pages(struct kvm_memory_slot *slot, gfn_t gfn,
46
			   unsigned long npages)
47 48 49 50
{
	gfn_t end_gfn;
	pfn_t pfn;

51
	pfn     = gfn_to_pfn_memslot(slot, gfn);
52
	end_gfn = gfn + npages;
53 54
	gfn    += 1;

55
	if (is_error_noslot_pfn(pfn))
56 57 58
		return pfn;

	while (gfn < end_gfn)
59
		gfn_to_pfn_memslot(slot, gfn++);
60 61 62 63

	return pfn;
}

64 65 66 67 68 69 70 71
static void kvm_unpin_pages(struct kvm *kvm, pfn_t pfn, unsigned long npages)
{
	unsigned long i;

	for (i = 0; i < npages; ++i)
		kvm_release_pfn_clean(pfn + i);
}

72
int kvm_iommu_map_pages(struct kvm *kvm, struct kvm_memory_slot *slot)
B
Ben-Ami Yassour 已提交
73
{
74
	gfn_t gfn, end_gfn;
B
Ben-Ami Yassour 已提交
75
	pfn_t pfn;
76
	int r = 0;
J
Joerg Roedel 已提交
77
	struct iommu_domain *domain = kvm->arch.iommu_domain;
78
	int flags;
B
Ben-Ami Yassour 已提交
79 80 81 82 83

	/* check if iommu exists and in use */
	if (!domain)
		return 0;

84 85 86
	gfn     = slot->base_gfn;
	end_gfn = gfn + slot->npages;

87 88 89
	flags = IOMMU_READ;
	if (!(slot->flags & KVM_MEM_READONLY))
		flags |= IOMMU_WRITE;
90
	if (!kvm->arch.iommu_noncoherent)
91 92
		flags |= IOMMU_CACHE;

93 94 95 96 97 98 99 100 101 102 103 104 105 106 107 108 109 110 111 112 113

	while (gfn < end_gfn) {
		unsigned long page_size;

		/* Check if already mapped */
		if (iommu_iova_to_phys(domain, gfn_to_gpa(gfn))) {
			gfn += 1;
			continue;
		}

		/* Get the page size we could use to map */
		page_size = kvm_host_page_size(kvm, gfn);

		/* Make sure the page_size does not exceed the memslot */
		while ((gfn + (page_size >> PAGE_SHIFT)) > end_gfn)
			page_size >>= 1;

		/* Make sure gfn is aligned to the page size we want to map */
		while ((gfn << PAGE_SHIFT) & (page_size - 1))
			page_size >>= 1;

114 115 116 117
		/* Make sure hva is aligned to the page size we want to map */
		while (__gfn_to_hva_memslot(slot, gfn) & (page_size - 1))
			page_size >>= 1;

118 119 120 121
		/*
		 * Pin all pages we are about to map in memory. This is
		 * important because we unmap and unpin in 4kb steps later.
		 */
122
		pfn = kvm_pin_pages(slot, gfn, page_size >> PAGE_SHIFT);
123
		if (is_error_noslot_pfn(pfn)) {
124
			gfn += 1;
B
Ben-Ami Yassour 已提交
125
			continue;
126
		}
B
Ben-Ami Yassour 已提交
127

128 129
		/* Map into IO address space */
		r = iommu_map(domain, gfn_to_gpa(gfn), pfn_to_hpa(pfn),
130
			      page_size, flags);
131
		if (r) {
W
Weidong Han 已提交
132
			printk(KERN_ERR "kvm_iommu_map_address:"
J
Joerg Roedel 已提交
133
			       "iommu failed to map pfn=%llx\n", pfn);
134
			kvm_unpin_pages(kvm, pfn, page_size >> PAGE_SHIFT);
B
Ben-Ami Yassour 已提交
135 136
			goto unmap_pages;
		}
137 138 139 140

		gfn += page_size >> PAGE_SHIFT;


B
Ben-Ami Yassour 已提交
141
	}
142

B
Ben-Ami Yassour 已提交
143 144 145
	return 0;

unmap_pages:
146
	kvm_iommu_put_pages(kvm, slot->base_gfn, gfn - slot->base_gfn);
B
Ben-Ami Yassour 已提交
147 148 149 150 151
	return r;
}

static int kvm_iommu_map_memslots(struct kvm *kvm)
{
152
	int idx, r = 0;
153
	struct kvm_memslots *slots;
154
	struct kvm_memory_slot *memslot;
B
Ben-Ami Yassour 已提交
155

156 157 158
	if (kvm->arch.iommu_noncoherent)
		kvm_arch_register_noncoherent_dma(kvm);

159
	idx = srcu_read_lock(&kvm->srcu);
160
	slots = kvm_memslots(kvm);
161

162 163
	kvm_for_each_memslot(memslot, slots) {
		r = kvm_iommu_map_pages(kvm, memslot);
B
Ben-Ami Yassour 已提交
164 165 166
		if (r)
			break;
	}
167
	srcu_read_unlock(&kvm->srcu, idx);
168

B
Ben-Ami Yassour 已提交
169 170 171
	return r;
}

W
Weidong Han 已提交
172 173
int kvm_assign_device(struct kvm *kvm,
		      struct kvm_assigned_dev_kernel *assigned_dev)
B
Ben-Ami Yassour 已提交
174 175
{
	struct pci_dev *pdev = NULL;
J
Joerg Roedel 已提交
176
	struct iommu_domain *domain = kvm->arch.iommu_domain;
177 178
	int r;
	bool noncoherent;
B
Ben-Ami Yassour 已提交
179

W
Weidong Han 已提交
180 181 182 183 184 185
	/* check if iommu exists and in use */
	if (!domain)
		return 0;

	pdev = assigned_dev->dev;
	if (pdev == NULL)
B
Ben-Ami Yassour 已提交
186
		return -ENODEV;
W
Weidong Han 已提交
187

J
Joerg Roedel 已提交
188
	r = iommu_attach_device(domain, &pdev->dev);
W
Weidong Han 已提交
189
	if (r) {
190
		dev_err(&pdev->dev, "kvm assign device failed ret %d", r);
W
Weidong Han 已提交
191
		return r;
B
Ben-Ami Yassour 已提交
192 193
	}

194
	noncoherent = !iommu_capable(&pci_bus_type, IOMMU_CAP_CACHE_COHERENCY);
195 196

	/* Check if need to update IOMMU page table for guest memory */
197
	if (noncoherent != kvm->arch.iommu_noncoherent) {
198
		kvm_iommu_unmap_memslots(kvm);
199
		kvm->arch.iommu_noncoherent = noncoherent;
200 201 202 203 204
		r = kvm_iommu_map_memslots(kvm);
		if (r)
			goto out_unmap;
	}

205
	pci_set_dev_assigned(pdev);
206

207
	dev_info(&pdev->dev, "kvm assign device\n");
B
Ben-Ami Yassour 已提交
208

W
Weidong Han 已提交
209
	return 0;
210 211 212
out_unmap:
	kvm_iommu_unmap_memslots(kvm);
	return r;
W
Weidong Han 已提交
213
}
B
Ben-Ami Yassour 已提交
214

W
Weidong Han 已提交
215 216 217
int kvm_deassign_device(struct kvm *kvm,
			struct kvm_assigned_dev_kernel *assigned_dev)
{
J
Joerg Roedel 已提交
218
	struct iommu_domain *domain = kvm->arch.iommu_domain;
W
Weidong Han 已提交
219 220 221 222 223 224 225 226 227 228
	struct pci_dev *pdev = NULL;

	/* check if iommu exists and in use */
	if (!domain)
		return 0;

	pdev = assigned_dev->dev;
	if (pdev == NULL)
		return -ENODEV;

J
Joerg Roedel 已提交
229
	iommu_detach_device(domain, &pdev->dev);
W
Weidong Han 已提交
230

231
	pci_clear_dev_assigned(pdev);
232

233
	dev_info(&pdev->dev, "kvm deassign device\n");
W
Weidong Han 已提交
234 235 236 237

	return 0;
}

W
Weidong Han 已提交
238 239 240 241
int kvm_iommu_map_guest(struct kvm *kvm)
{
	int r;

242
	if (!iommu_present(&pci_bus_type)) {
J
Joerg Roedel 已提交
243
		printk(KERN_ERR "%s: iommu not found\n", __func__);
B
Ben-Ami Yassour 已提交
244 245 246
		return -ENODEV;
	}

247 248
	mutex_lock(&kvm->slots_lock);

249
	kvm->arch.iommu_domain = iommu_domain_alloc(&pci_bus_type);
250 251 252 253
	if (!kvm->arch.iommu_domain) {
		r = -ENOMEM;
		goto out_unlock;
	}
B
Ben-Ami Yassour 已提交
254

255
	if (!allow_unsafe_assigned_interrupts &&
256
	    !iommu_capable(&pci_bus_type, IOMMU_CAP_INTR_REMAP)) {
257 258 259 260 261 262
		printk(KERN_WARNING "%s: No interrupt remapping support,"
		       " disallowing device assignment."
		       " Re-enble with \"allow_unsafe_assigned_interrupts=1\""
		       " module option.\n", __func__);
		iommu_domain_free(kvm->arch.iommu_domain);
		kvm->arch.iommu_domain = NULL;
263 264
		r = -EPERM;
		goto out_unlock;
265 266
	}

B
Ben-Ami Yassour 已提交
267 268
	r = kvm_iommu_map_memslots(kvm);
	if (r)
269
		kvm_iommu_unmap_memslots(kvm);
B
Ben-Ami Yassour 已提交
270

271 272
out_unlock:
	mutex_unlock(&kvm->slots_lock);
B
Ben-Ami Yassour 已提交
273 274 275 276
	return r;
}

static void kvm_iommu_put_pages(struct kvm *kvm,
W
Weidong Han 已提交
277
				gfn_t base_gfn, unsigned long npages)
B
Ben-Ami Yassour 已提交
278
{
279 280
	struct iommu_domain *domain;
	gfn_t end_gfn, gfn;
B
Ben-Ami Yassour 已提交
281
	pfn_t pfn;
W
Weidong Han 已提交
282 283
	u64 phys;

284 285 286 287
	domain  = kvm->arch.iommu_domain;
	end_gfn = base_gfn + npages;
	gfn     = base_gfn;

W
Weidong Han 已提交
288 289 290
	/* check if iommu exists and in use */
	if (!domain)
		return;
B
Ben-Ami Yassour 已提交
291

292 293
	while (gfn < end_gfn) {
		unsigned long unmap_pages;
294
		size_t size;
295 296

		/* Get physical address */
J
Joerg Roedel 已提交
297
		phys = iommu_iova_to_phys(domain, gfn_to_gpa(gfn));
298 299 300 301 302 303

		if (!phys) {
			gfn++;
			continue;
		}

304 305 306
		pfn  = phys >> PAGE_SHIFT;

		/* Unmap address from IO address space */
307 308
		size       = iommu_unmap(domain, gfn_to_gpa(gfn), PAGE_SIZE);
		unmap_pages = 1ULL << get_order(size);
W
Weidong Han 已提交
309

310 311 312 313 314
		/* Unpin all pages we just unmapped to not leak any memory */
		kvm_unpin_pages(kvm, pfn, unmap_pages);

		gfn += unmap_pages;
	}
B
Ben-Ami Yassour 已提交
315 316
}

317 318 319 320 321
void kvm_iommu_unmap_pages(struct kvm *kvm, struct kvm_memory_slot *slot)
{
	kvm_iommu_put_pages(kvm, slot->base_gfn, slot->npages);
}

B
Ben-Ami Yassour 已提交
322 323
static int kvm_iommu_unmap_memslots(struct kvm *kvm)
{
324
	int idx;
325
	struct kvm_memslots *slots;
326
	struct kvm_memory_slot *memslot;
327

328
	idx = srcu_read_lock(&kvm->srcu);
329
	slots = kvm_memslots(kvm);
330

331
	kvm_for_each_memslot(memslot, slots)
332
		kvm_iommu_unmap_pages(kvm, memslot);
333

334
	srcu_read_unlock(&kvm->srcu, idx);
B
Ben-Ami Yassour 已提交
335

336 337 338
	if (kvm->arch.iommu_noncoherent)
		kvm_arch_unregister_noncoherent_dma(kvm);

B
Ben-Ami Yassour 已提交
339 340 341 342 343
	return 0;
}

int kvm_iommu_unmap_guest(struct kvm *kvm)
{
J
Joerg Roedel 已提交
344
	struct iommu_domain *domain = kvm->arch.iommu_domain;
B
Ben-Ami Yassour 已提交
345 346 347 348 349

	/* check if iommu exists and in use */
	if (!domain)
		return 0;

350
	mutex_lock(&kvm->slots_lock);
B
Ben-Ami Yassour 已提交
351
	kvm_iommu_unmap_memslots(kvm);
352
	kvm->arch.iommu_domain = NULL;
353
	kvm->arch.iommu_noncoherent = false;
354 355
	mutex_unlock(&kvm->slots_lock);

J
Joerg Roedel 已提交
356
	iommu_domain_free(domain);
B
Ben-Ami Yassour 已提交
357 358
	return 0;
}