TFRecord的使用事例

TFRecord的使用事例

存储为.tfrecord文件

def store(tfrecord_file, filenames, labels):
    with tf.io.TFRecordWriter(tfrecord_file) as writer:
        for filename, label in zip(filenames, labels):
            image = open(filename, 'rb').read()
            feature = {
                'image': tf.train.Feature(bytes_list=tf.train.BytesList(value=[image])),
                'label': tf.train.Feature(int64_list=tf.train.Int64List(value=[label]))
            }
            example = tf.train.Example(features=tf.train.Features(feature=feature))
            writer.write(example.SerializeToString())
data_dir = "./cats_vs_dogs/sample"

train_cats_dir = data_dir + "/train/cats/"
train_dogs_dir = data_dir + "/train/dogs/"
train_tfrecord_file = "./cats_vs_dogs/train.tfrecords"

train_cat_filenames = [train_cats_dir + filename for filename in os.listdir(train_cats_dir)]
train_dog_filenames = [train_dogs_dir + filename for filename in os.listdir(train_dogs_dir)]
train_filenames = train_cat_filenames + train_dog_filenames
train_labels = [0] * len(train_cat_filenames) + [1] * len(train_dog_filenames)

test_cats_dir = data_dir + "/valid/cats/"
test_dogs_dir = data_dir + "/valid/dogs/"
test_tfrecord_file = "./cats_vs_dogs/test.tfrecords"

test_cat_filenames = [test_cats_dir + filename for filename in os.listdir(test_cats_dir)]
test_dog_filenames = [test_dogs_dir + filename for filename in os.listdir(test_dogs_dir)]
test_filenames = test_cat_filenames + test_dog_filenames
test_labels = [0] * len(test_cat_filenames) + [1] * len(test_dog_filenames)
store(train_tfrecord_file)
store(test_tfrecord_file)

解析.tfrecord文件

train_tfrecord_file = "./cats_vs_dogs/train.tfrecords"
test_tfrecord_file = "./cats_vs_dogs/test.tfrecords"

train_raw_dataset = tf.data.TFRecordDataset(train_tfrecord_file)
test_raw_dataset = tf.data.TFRecordDataset(test_tfrecord_file)

feature_description = {
    'image': tf.io.FixedLenFeature([], tf.string),
    'label': tf.io.FixedLenFeature([], tf.int64)
}

test_dataset = test_raw_dataset.map(_parse_example)
train_dataset = train_raw_dataset.map(_parse_example)
def _parse_example(example_string):
    feature_dict = tf.io.parse_single_example(example_string, feature_description)
    feature_dict['image'] = tf.io.decode_jpeg(feature_dict['image'])  # 解码JPEG图片
    feature_dict['image']=tf.image.resize(feature_dict['image'], [256, 256]) / 255.0
    return feature_dict['image'], feature_dict['label']

训练数据集的预处理

train_dataset = train_dataset.shuffle(buffer_size=buffer_size).batch(batch_size=batch_size).prefetch(buffer_size=tf.data.experimental.AUTOTUNE)
#测试数据集如果太大,也需要这样处理,否则会报错
test_dataset = test_dataset.shuffle(buffer_size=buffer_size).batch(batch_size=batch_size).prefetch(buffer_size=tf.data.experimental.AUTOTUNE)
  • 0
    点赞
  • 0
    收藏
    觉得还不错? 一键收藏
  • 0
    评论

“相关推荐”对你有帮助么?

  • 非常没帮助
  • 没帮助
  • 一般
  • 有帮助
  • 非常有帮助
提交
评论
添加红包

请填写红包祝福语或标题

红包个数最小为10个

红包金额最低5元

当前余额3.43前往充值 >
需支付:10.00
成就一亿技术人!
领取后你会自动成为博主和红包主的粉丝 规则
hope_wisdom
发出的红包
实付
使用余额支付
点击重新获取
扫码支付
钱包余额 0

抵扣说明:

1.余额是钱包充值的虚拟货币,按照1:1的比例进行支付金额的抵扣。
2.余额无法直接购买下载,可以购买VIP、付费专栏及课程。

余额充值