@inproceedings{a04ea2463ade43d7922bb644b029b7c6,
title = "Keyword spotting techniques for sanskrit documents",
abstract = "With advances in the field of digitization of printed documents and several mass digitization projects underway, information retrieval and document search have emerged as key research areas. However, most of the current work in these areas is limited to English and a few oriental languages. The lack of efficient solutions for Indic scripts and languages such as Sanskrit has hampered information extraction from a large body of documents of cultural and historical importance. This chapter presents two relevant topics in this area. First, we describe the use of a script specific Keyword Spotting for Sanskrit documents that makes use of domain knowledge of the script. Second, we address the needs of a digital library to provide access to a collection of documents from multiple scripts. This requires intelligent solutions which scale across different scripts. We present a script independent Keyword Spotting approach for this purpose. Experimental results illustrate the efficacy of our methods.",
keywords = "Document analysis, Document retrieval, Indic scripts, Keyword spotting, Optical character recognition",
author = "Anurag Bhardwaj and Srirangaraj Setlur and Venu Govindaraju",
year = "2009",
doi = "10.1007/978-3-642-00155-0\_22",
language = "English",
isbn = "9783642001543",
series = "Lecture Notes in Computer Science (including subseries Lecture Notes in Artificial Intelligence and Lecture Notes in Bioinformatics)",
pages = "403--416",
booktitle = "Sanskrit Computational Linguistics - First and Second International Symposia, Revised Selected and Invited Papers",
note = "1st and 2nd International Symposia on Sanskrit Computational Linguistics ; Conference date: 15-05-2008 Through 17-05-2008",
}