%Aigaion2 BibTeX export from Idiap Publications %Monday 20 January 2025 07:44:18 PM @TECHREPORT{Motlicek_Idiap-RR-21-2010, author = {Motlicek, Petr and Valente, Fabio and Garner, Philip N.}, projects = {TA2, IM2}, month = {7}, title = {English Spoken Term Detection in Multilingual Recordings}, type = {Idiap-RR}, number = {Idiap-RR-21-2010}, year = {2010}, institution = {Idiap}, address = {Rue Marconi 19, Martigny, 1920, Switzerland}, crossref = {Motlicek_INTERSPEECH2010_2010}, abstract = {This paper investigates the automatic detection of English spoken terms in a multi-language scenario over real lecture recordings. Spoken Term Detection (STD) is based on an LVCSR where the output is represented in the form of word lattices. The lattices are then used to search the required terms. Processed lectures are mainly composed of English, French and Italian recordings where the language can also change within one recording. Therefore, the English STD system uses an Out-Of-Language (OOL) detection module to filter out non-English input segments. OOL detection is evaluated w.r.t. various confidence measures estimated from word lattices. Experimental studies of OOL detection followed by English STD are performed on several hours of multilingual recordings. Significant improvement of OOL+STD over a stand-alone STD system is achieved (relatively more than 50\% in EER). Finally, an additional modality (text slides in the form of PowerPoint presentations) is exploited to improve STD.}, pdf = {https://publications.idiap.ch/attachments/reports/2010/Motlicek_Idiap-RR-21-2010.pdf} } crossreferenced publications: @INPROCEEDINGS{Motlicek_INTERSPEECH2010_2010, author = {Motlicek, Petr and Valente, Fabio and Garner, Philip N.}, keywords = {Confidence Measure (CM), LVCSR, Out-Of-Language (OOL) detection, Spoken Term Detection (STD)}, projects = {Idiap, IM2, TA2}, month = {9}, title = {English Spoken Term Detection in Multilingual Recordings}, booktitle = {Proceedings of Interspeech, Makuhari, Japan, 2010}, year = {2010}, location = {Makuhari, Japan}, organization = {ISCA}, crossref = {Motlicek_Idiap-RR-21-2010}, abstract = {This paper investigates the automatic detection of English spoken terms in a multi-language scenario over real lecture recordings. Spoken Term Detection (STD) is based on an LVCSR where the output is represented in the form of word lattices. The lattices are then used to search the required terms. Processed lectures are mainly composed of English, French and Italian recordings where the language can also change within one recording. Therefore, the English STD system uses an Out-Of-Language (OOL) detection module to filter out non-English input segments. OOL detection is evaluated w.r.t. various confidence measures estimated from word lattices. Experimental studies of OOL detection followed by English STD are performed on several hours of multilingual recordings. Significant improvement of OOL+STD over a stand-alone STD system is achieved (relatively more than 50\% in EER). Finally, an additional modality (text slides in the form of PowerPoint presentations) is exploited to improve STD.}, pdf = {https://publications.idiap.ch/attachments/papers/2010/Motlicek_INTERSPEECH2010_2010.pdf} }