coco.py 5.9 KB
Newer Older
1 2
import os
import numpy as np
3
import logging
4
from ppdet.core.workspace import register, serializable
5
from .dataset import DetDataset
6 7 8 9 10 11

logger = logging.getLogger(__name__)


@register
@serializable
12
class COCODataset(DetDataset):
13
    def __init__(self,
14
                 dataset_dir=None,
15 16
                 image_dir=None,
                 anno_path=None,
17 18 19 20
                 with_background=True,
                 sample_num=-1):
        super(COCODataset, self).__init__(dataset_dir, image_dir, anno_path,
                                          with_background, sample_num)
W
wangguanzhong 已提交
21
        self.load_image_only = False
22
        self.load_semantic = False
23

24
    def parse_dataset(self):
25 26 27 28 29
        anno_path = os.path.join(self.dataset_dir, self.anno_path)
        image_dir = os.path.join(self.dataset_dir, self.image_dir)

        assert anno_path.endswith('.json'), \
            'invalid coco annotation file: ' + anno_path
W
wangguanzhong 已提交
30
        from pycocotools.coco import COCO
31 32 33 34 35 36 37 38 39 40 41 42 43 44 45 46 47
        coco = COCO(anno_path)
        img_ids = coco.getImgIds()
        cat_ids = coco.getCatIds()
        records = []
        ct = 0

        # when with_background = True, mapping category to classid, like:
        #   background:0, first_class:1, second_class:2, ...
        catid2clsid = dict({
            catid: i + int(self.with_background)
            for i, catid in enumerate(cat_ids)
        })
        cname2cid = dict({
            coco.loadCats(catid)[0]['name']: clsid
            for catid, clsid in catid2clsid.items()
        })

W
wangguanzhong 已提交
48 49 50 51 52
        if 'annotations' not in coco.dataset:
            self.load_image_only = True
            logger.warn('Annotation file: {} does not contains ground truth '
                        'and load image information only.'.format(anno_path))

53 54 55 56 57 58
        for img_id in img_ids:
            img_anno = coco.loadImgs(img_id)[0]
            im_fname = img_anno['file_name']
            im_w = float(img_anno['width'])
            im_h = float(img_anno['height'])

59 60 61
            im_path = os.path.join(image_dir,
                                   im_fname) if image_dir else im_fname
            if not os.path.exists(im_path):
W
wangguanzhong 已提交
62
                logger.warn('Illegal image file: {}, and it will be '
63
                            'ignored'.format(im_path))
W
wangguanzhong 已提交
64 65 66 67 68 69 70
                continue

            if im_w < 0 or im_h < 0:
                logger.warn('Illegal width: {} or height: {} in annotation, '
                            'and im_id: {} will be ignored'.format(im_w, im_h,
                                                                   img_id))
                continue
W
wangguanzhong 已提交
71

72
            coco_rec = {
73
                'im_file': im_path,
74 75 76 77 78
                'im_id': np.array([img_id]),
                'h': im_h,
                'w': im_w,
            }

W
wangguanzhong 已提交
79 80 81
            if not self.load_image_only:
                ins_anno_ids = coco.getAnnIds(imgIds=img_id, iscrowd=False)
                instances = coco.loadAnns(ins_anno_ids)
82

W
wangguanzhong 已提交
83 84
                bboxes = []
                for inst in instances:
S
sunxl1988 已提交
85 86 87 88 89 90
                    # check gt bbox
                    if 'bbox' not in inst.keys():
                        continue
                    else:
                        if not any(np.array(inst['bbox'])):
                            continue
W
wangguanzhong 已提交
91 92 93 94 95 96 97 98 99 100 101 102 103 104 105 106 107 108 109 110 111 112 113 114 115 116 117
                    x, y, box_w, box_h = inst['bbox']
                    x1 = max(0, x)
                    y1 = max(0, y)
                    x2 = min(im_w - 1, x1 + max(0, box_w - 1))
                    y2 = min(im_h - 1, y1 + max(0, box_h - 1))
                    if inst['area'] > 0 and x2 >= x1 and y2 >= y1:
                        inst['clean_bbox'] = [x1, y1, x2, y2]
                        bboxes.append(inst)
                    else:
                        logger.warn(
                            'Found an invalid bbox in annotations: im_id: {}, '
                            'area: {} x1: {}, y1: {}, x2: {}, y2: {}.'.format(
                                img_id, float(inst['area']), x1, y1, x2, y2))
                num_bbox = len(bboxes)

                gt_bbox = np.zeros((num_bbox, 4), dtype=np.float32)
                gt_class = np.zeros((num_bbox, 1), dtype=np.int32)
                gt_score = np.ones((num_bbox, 1), dtype=np.float32)
                is_crowd = np.zeros((num_bbox, 1), dtype=np.int32)
                difficult = np.zeros((num_bbox, 1), dtype=np.int32)
                gt_poly = [None] * num_bbox

                for i, box in enumerate(bboxes):
                    catid = box['category_id']
                    gt_class[i][0] = catid2clsid[catid]
                    gt_bbox[i, :] = box['clean_bbox']
                    is_crowd[i][0] = box['iscrowd']
S
sunxl1988 已提交
118 119 120 121
                    # check RLE format 
                    if box['iscrowd'] == 1:
                        gt_poly[i] = [[0.0, 0.0], ]
                        continue
W
wangguanzhong 已提交
122 123 124
                    if 'segmentation' in box:
                        gt_poly[i] = box['segmentation']

S
sunxl1988 已提交
125 126 127
                if not any(gt_poly):
                    continue

W
wangguanzhong 已提交
128 129 130 131 132 133 134
                coco_rec.update({
                    'is_crowd': is_crowd,
                    'gt_class': gt_class,
                    'gt_bbox': gt_bbox,
                    'gt_score': gt_score,
                    'gt_poly': gt_poly,
                })
135 136 137 138 139
                # TODO: remove load_semantic
                if self.load_semantic:
                    seg_path = os.path.join(self.dataset_dir, 'stuffthingmaps',
                                            'train2017', im_fname[:-3] + 'png')
                    coco_rec.update({'semantic': seg_path})
W
wangguanzhong 已提交
140

141
            logger.debug('Load file: {}, im_id: {}, h: {}, w: {}.'.format(
142
                im_path, img_id, im_h, im_w))
143 144 145 146 147
            records.append(coco_rec)
            ct += 1
            if self.sample_num > 0 and ct >= self.sample_num:
                break
        assert len(records) > 0, 'not found any coco record in %s' % (anno_path)
Y
Yang Zhang 已提交
148
        logger.debug('{} samples in file {}'.format(ct, anno_path))
149
        self.roidbs, self.cname2cid = records, cname2cid