@inproceedings{772300ad0cf64d3b9185c30b49b6ac6a,
title = "Towards scalable emotion classification in microblog based on noisy training data",
abstract = "The availability of labeled corpus is of great importance for emotion classification tasks. Because manual labeling is too timeconsuming, hashtags have been used as naturally annotated labels to obtain large amount of labeled training data from microblog. However, the inconsistency and noise in annotation can adversely affect the data quality and thus the performance when used to train a classifier. In this paper, we propose a classification framework which allows naturally annotated data to be used as additional training data and employs a k-NN graph based data cleaning method to remove noise after noisy data has certain accumulations. Evaluation on NLP&CC2013 Chinese Weibo emotion classification dataset shows that our approach achieves 15.8% better performance than directly using the noisy data without noise filtering. After adding the filtered data with hashtags into an existing high-quality training data, the performance increases 3.7% compared to using the high-quality training data alone.",
keywords = "Data cleaning, Emotion classification, Hashtag, K-NN",
author = "Minglei Li and Qin Lu and Lin Gui and Yunfei Long",
year = "2016",
month = jan,
day = "1",
doi = "10.1007/978-3-319-47674-2_33",
language = "English",
isbn = "9783319476735",
series = "Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics)",
publisher = "Springer Verlag",
pages = "399--410",
booktitle = "Chinese Computational Linguistics and Natural Language Processing Based on Naturally Annotated Big Data - 15th China National Conference, CCL 2016 and 4th International Symposium, NLP-NABD 2016, Proceedings",
address = "Germany",
note = "15th China National Conference on Chinese Computational Linguistics, CCL 2016 and 4th International Symposium on Natural Language Processing Based on Naturally Annotated Big Data, NLP-NABD 2016 ; Conference date: 15-10-2016 Through 16-10-2016",
}