python词频统计结果写入csv_Python3,字典从csv文件中统计词频

import csv

from collections import Counter

columns = defaultdict(list) # each value in each column is appended to a list

with open('csv_file.csv') as f:

reader = csv.DictReader(f) # read rows into a dictionary format

for row in reader: # read a row as {column1: value1, column2: value2,...}

for (k,v) in row.items(): # go over each column name and value

columns[k].append(v) # append the value into the appropriate list

# based on column name k

选项1output_dict_counter_version = dict(Counter(degree_list_clean))

print(output_dict_counter_version)

选项2degree_frequency_dict = {}

for deg in degree_list_clean:

if deg in degree_frequency_dict:

degree_frequency_dict[deg] += 1

else:

degree_frequency_dict[deg] = 1

print(degree_frequency_dict)

使用import pandas as pd

from collections import Counter

data = pd.read_csv("csv_file.csv")

degree_list = data['degree'].tolist()

degree_list_clean = []

for cad_degrees in degree_list:

cad_degrees_lst = cad_degrees.split()

for degree in cad_degrees_lst:

degree_clean = degree.strip().replace('.','').lower()

degree_list_clean.append(degree_clean)

print(dict(Counter(degree_list_clean)))

'''

Input

name,degree,email

ABC,PhD. ,abd@gmail.com

CDE,Ph.D. ,cde@gmail.com

FGH, MD PHD ,fgh@gmail.com

Output

{'phd': 3, 'md': 1}

'''

  • 0
    点赞
  • 0
    收藏
    觉得还不错? 一键收藏
  • 0
    评论
评论
添加红包

请填写红包祝福语或标题

红包个数最小为10个

红包金额最低5元

当前余额3.43前往充值 >
需支付:10.00
成就一亿技术人!
领取后你会自动成为博主和红包主的粉丝 规则
hope_wisdom
发出的红包
实付
使用余额支付
点击重新获取
扫码支付
钱包余额 0

抵扣说明:

1.余额是钱包充值的虚拟货币,按照1:1的比例进行支付金额的抵扣。
2.余额无法直接购买下载,可以购买VIP、付费专栏及课程。

余额充值