@inbook{f84b8ab5b10d490b9281aa4e3d83f955,
title = "Incorporating virtual relevant documents for learning in text categorization",
abstract = "This paper proposes a virtual relevant document technique in the learning phase for text categorization. The method uses a simple transformation of relevant documents, i.e. making virtual documents by combining document pairs in the training set. The virtual document produced by this method has the enriched term vector space, with greater weights for the terms that co-occur in two relevant documents. The experimental results showed a significant improvement over the baseline, which proves the usefulness of the proposed method: 71\% improvement on TREC-11 filtering test collection and 11\% improvement on Reuters-21578 test set for the topics with less than 100 relevant documents in the micro average F1. The result analysis indicates that the addition of virtual relevant documents contributes to the steady improvement of the performance.",
author = "Lee, \{Kyung Soon\} and Kyo Kageura",
year = "2003",
doi = "10.1007/978-3-540-24594-0\_6",
language = "English",
series = "Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics)",
publisher = "Springer Verlag",
pages = "62--72",
editor = "\{Tengku Sembok\}, \{Tengku Mohd\} and Zaman, \{Halimah Badioze\} and Hsinchun Chen and Urs, \{Shalini R.\} and Myaeng, \{Sung Hyon\}",
booktitle = "Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics)",
}