@inproceedings{bb387600,
AUTHOR = "Ng, H.W. and Guan, C.T.",
TITLE = "Efficient Representation Learning for Inner Speech Domain
Generalization",
BOOKTITLE = CAIP23,
YEAR = "2023",
PAGES = "I:131-141",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381663"}
@inproceedings{bb387601,
AUTHOR = "Oneata, D. and Cucu, H.",
TITLE = "Improving Multimodal Speech Recognition by Data Augmentation and
Speech Representations",
BOOKTITLE = MULA22,
YEAR = "2022",
PAGES = "4578-4587",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381664"}
@inproceedings{bb387602,
AUTHOR = "Tapia, L.S. and Gomez, A. and Esparza, M. and Jatla, V. and Pattichis, M. and Celedon Pattichis, S. and Lopez Leiva, C.",
TITLE = "Bilingual Speech Recognition by Estimating Speaker Geometry from Video
Data",
BOOKTITLE = CAIP21,
YEAR = "2021",
PAGES = "I:79-89",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381665"}
@inproceedings{bb387603,
AUTHOR = "Qiao, F.C. and Peng, X.",
TITLE = "Uncertainty-guided Model Generalization to Unseen Domains",
BOOKTITLE = CVPR21,
YEAR = "2021",
PAGES = "6786-6796",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381666"}
@inproceedings{bb387604,
AUTHOR = "Ngantcha, P. and Amith, M. and Tao, C. and Roberts, K.",
TITLE = "Patient-Provider Communication Training Models for Interactive Speech
Devices",
BOOKTITLE = DHM21,
YEAR = "2021",
PAGES = "I:250-268",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381667"}
@inproceedings{bb387605,
AUTHOR = "Wu, Y.C. and Liao, W.H.",
TITLE = "Toward Text-independent Cross-lingual Speaker Recognition Using
English-Mandarin-Taiwanese Dataset",
BOOKTITLE = ICPR21,
YEAR = "2021",
PAGES = "8515-8522",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381668"}
@inproceedings{bb387606,
AUTHOR = "Chen, Y.B. and Ma, Y. and Ko, T. and Wang, J.P. and Li, Q.",
TITLE = "MetaMix: Improved Meta-Learning with Interpolation-based Consistency
Regularization",
BOOKTITLE = ICPR21,
YEAR = "2021",
PAGES = "407-414",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381669"}
@inproceedings{bb387607,
AUTHOR = "Zhou, L.X. and Zhang, J.",
TITLE = "From Bottom to Top: A Coordinated Feature Representation Method for
Speech Recognition",
BOOKTITLE = MMDLCA20,
YEAR = "2020",
PAGES = "396-403",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381670"}
@inproceedings{bb387608,
AUTHOR = "Zhao, J. and Parry, C.J. and dos Anjos, R. and Anslow, C. and Rhee, T.",
TITLE = "Voice Interaction for Augmented Reality Navigation Interfaces with
Natural Language Understanding",
BOOKTITLE = IVCNZ20,
YEAR = "2020",
PAGES = "1-6",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381671"}
@inproceedings{bb387609,
AUTHOR = "ABAKARIM, F. and ABENAOU, A.",
TITLE = "Amazigh isolated word speech recognition system using the Adaptive
Orthogonal Transform Method.",
BOOKTITLE = ISCV20,
YEAR = "2020",
PAGES = "1-6",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381672"}
@inproceedings{bb387610,
AUTHOR = "Perez, A.F. and Sanguineti, V. and Morerio, P. and Murino, V.",
TITLE = "Audio-Visual Model Distillation Using Acoustic Images",
BOOKTITLE = WACV20,
YEAR = "2020",
PAGES = "2843-2852",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381673"}
@inproceedings{bb387611,
AUTHOR = "Tapu, R. and Mocanu, B. and Zaharia, T.",
TITLE = "Dynamic Subtitles: A Multimodal Video Accessibility Enhancement
Dedicated to Deaf and Hearing Impaired Users",
BOOKTITLE = ACVR19,
YEAR = "2019",
PAGES = "2558-2566",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381674"}
@inproceedings{bb387612,
AUTHOR = "Roberto, A. and Saggese, A. and Vento, M.",
TITLE = "A Challenging Voice Dataset for Robotic Applications in Noisy
Environments",
BOOKTITLE = CAIP19,
YEAR = "2019",
PAGES = "II:354-364",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381675"}
@inproceedings{bb387613,
AUTHOR = "Gauvain, J. and Lamel, L. and Le, V.B. and Despres, J. and Gauvain, J.L. and Messaoudi, A. and Vieru, B. and Ben Kheder, W.",
TITLE = "Challenges in Audio Processing of Terrorist-Related Data",
BOOKTITLE = "MMMod19",
YEAR = "2019",
PAGES = "II:80-92",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381676"}
@inproceedings{bb387614,
AUTHOR = "Jorrin, J. and Buera, L.",
TITLE = "DANTE Speaker Recognition Module. An Efficient and Robust Automatic
Speaker Searching Solution for Terrorism-Related Scenarios",
BOOKTITLE = "MMMod19",
YEAR = "2019",
PAGES = "I:704-715",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381677"}
@inproceedings{bb387615,
AUTHOR = "Galanopoulos, D. and Mezaris, V.",
TITLE = "Temporal Lecture Video Fragmentation Using Word Embeddings",
BOOKTITLE = "MMMod19",
YEAR = "2019",
PAGES = "II:254-265",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381678"}
@inproceedings{bb387616,
AUTHOR = "Mukherjee, H. and Obaidullah, S.M. and Phadikar, S. and Roy, K.",
TITLE = "A Dravidian Language Identification System",
BOOKTITLE = ICPR18,
YEAR = "2018",
PAGES = "2654-2657",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381679"}
@inproceedings{bb387617,
AUTHOR = "Galiotou, E. and Karanikolas, N. and Ralli, A.",
TITLE = "Preservation and Management of Greek Dialectal Data",
BOOKTITLE = EuroMed18,
YEAR = "2018",
PAGES = "I:752-761",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381680"}
@inproceedings{bb387618,
AUTHOR = "Li, R. and Yu, J.",
TITLE = "Multimodal 3D visible articulation system for syllable based Mandarin
Chinese training",
BOOKTITLE = VCIP17,
YEAR = "2017",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381681"}
@inproceedings{bb387619,
AUTHOR = "Le, N. and Odobez, J.M.",
TITLE = "Improving Speaker Turn Embedding by Crossmodal Transfer Learning from
Face Embedding",
BOOKTITLE = CVAVM17,
YEAR = "2017",
PAGES = "428-437",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381682"}
@inproceedings{bb387620,
AUTHOR = "Arandjelovic, R. and Zisserman, A.",
TITLE = "Look, Listen and Learn",
BOOKTITLE = ICCV17,
YEAR = "2017",
PAGES = "609-617",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381683"}
@inproceedings{bb387621,
AUTHOR = "Muniandy, T. and Alvar, T.A. and Boon, C.J.",
TITLE = "Mandarin Language Learning System for Nasal Voice User",
BOOKTITLE = IVIC17,
YEAR = "2017",
PAGES = "376-388",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381684"}
@inproceedings{bb387622,
AUTHOR = "Madhavi, M.C. and Patil, H.A. and Bhendawade, N.",
TITLE = "Spoken Keyword Retrieval Using Source and System Features",
BOOKTITLE = PReMI17,
YEAR = "2017",
PAGES = "333-341",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381685"}
@inproceedings{bb387623,
AUTHOR = "Kacprzak, S.",
TITLE = "Spoken language clustering in the i-vectors space",
BOOKTITLE = WSSIP17,
YEAR = "2017",
PAGES = "1-5",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381686"}
@inproceedings{bb387624,
AUTHOR = "Pironkov, G. and Dupont, S. and Dutoit, T.",
TITLE = "Speaker-aware Multi-Task Learning for automatic speech recognition",
BOOKTITLE = ICPR16,
YEAR = "2016",
PAGES = "2900-2905",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381687"}
@inproceedings{bb387625,
AUTHOR = "Zhao, Y. and Zhao, R. and Wang, X.Y. and Ji, Q.",
TITLE = "Multilingual articulatory features augmentation learning",
BOOKTITLE = ICPR16,
YEAR = "2016",
PAGES = "2895-2899",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381688"}
@inproceedings{bb387626,
AUTHOR = "Ogawa, T. and Mallidi, S.H. and Dupoux, E. and Cohen, J. and Feldman, N.H. and Hermansky, H.",
TITLE = "A new efficient measure for accuracy prediction and its application
to multistream-based unsupervised adaptation",
BOOKTITLE = ICPR16,
YEAR = "2016",
PAGES = "2222-2227",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381689"}
@inproceedings{bb387627,
AUTHOR = "Mzah, Y. and Ahfir, M. and Jaidane, M.",
TITLE = "Late pre-dereverberation for speech intelligibility enhancement in
public address systems",
BOOKTITLE = ISIVC16,
YEAR = "2016",
PAGES = "291-296",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381690"}
@inproceedings{bb387628,
AUTHOR = "Montalvo, A. and Calvo, J.R.",
TITLE = "Discriminative Capacity and Phonetic Information of Bottleneck Features
in Speech",
BOOKTITLE = CIARP16,
YEAR = "2016",
PAGES = "134-141",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381691"}
@inproceedings{bb387629,
AUTHOR = "Ondas, S. and Juhar, J.",
TITLE = "Towards human-machine dialog in Slovak",
BOOKTITLE = WSSIP16,
YEAR = "2016",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381692"}
@inproceedings{bb387630,
AUTHOR = "Calvo, M. and Hurtado, L.F. and Garcia, F. and Sanchis, E.",
TITLE = "Combining Several ASR Outputs in a Graph-Based SLU System",
BOOKTITLE = CIARP15,
YEAR = "2015",
PAGES = "551-558",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381693"}
@inproceedings{bb387631,
AUTHOR = "Rohrbach, A. and Rohrbach, M. and Schiele, B.",
TITLE = "The Long-Short Story of Movie Description",
BOOKTITLE = GCPR15,
YEAR = "2015",
PAGES = "209-221",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381694"}
@inproceedings{bb387632,
AUTHOR = "Rohrbach, A. and Rohrbach, M. and Tandon, N. and Schiele, B.",
TITLE = "A dataset for Movie Description",
BOOKTITLE = CVPR15,
YEAR = "2015",
PAGES = "3202-3212",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381695"}
@inproceedings{bb387633,
AUTHOR = "Zhao, H.Q. and Qin, Z.C. and Wang, Y. and Wang, Y.X.",
TITLE = "A Bag-of-phonemes Model for Homeplace Classification of Mandarin
Speakers",
BOOKTITLE = IbPRIA15,
YEAR = "2015",
PAGES = "683-690",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381696"}
@inproceedings{bb387634,
AUTHOR = "Yakubu, M.A. and Maddage, N.C. and Atrey, P.K.",
TITLE = "Audio Secret Management Scheme Using Shamir's Secret Sharing",
BOOKTITLE = MMMod15,
YEAR = "2015",
PAGES = "I: 396-407",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381697"}
@inproceedings{bb387635,
AUTHOR = "Bello, C. and Ribas, D. and Calvo, J.R. and Ferrer, C.A.",
TITLE = "From Speech Quality Measures to Speaker Recognition Performance",
BOOKTITLE = CIARP14,
YEAR = "2014",
PAGES = "199-206",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381698"}
@inproceedings{bb387636,
AUTHOR = "Oropeza Rodriguez, J.L. and Suarez Guerra, S. and Jimenez Hernandez, M.",
TITLE = "The Place Theory as an Alternative Solution in Automatic Speech
Recognition Tasks",
BOOKTITLE = CIARP14,
YEAR = "2014",
PAGES = "167-174",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381699"}
@article{bb387637,
AUTHOR = "Diez, M. and Varona, A. and Penagarikano, M. and Rodriguez Fuentes, L.J. and Bordel, G.",
TITLE = "On the Projection of PLLRs for Unbounded Feature Distributions in
Spoken Language Recognition",
JOURNAL = SPLetters,
VOLUME = "21",
YEAR = "2014",
NUMBER = "9",
MONTH = "September",
PAGES = "1073-1077",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381700"}
@inproceedings{bb387638,
AUTHOR = "Diez, M. and Varona, A. and Penagarikano, M. and Rodriguez Fuentes, L.J. and Bordel, G.",
TITLE = "Optimizing PLLR Features for Spoken Language Recognition",
BOOKTITLE = ICPR14,
YEAR = "2014",
PAGES = "779-784",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381701"}
@inproceedings{bb387639,
AUTHOR = "Missaoui, I. and Lachiri, Z.",
TITLE = "Gabor Filterbank Features for Robust Speech Recognition",
BOOKTITLE = ICISP14,
YEAR = "2014",
PAGES = "665-671",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381702"}
@inproceedings{bb387640,
AUTHOR = "Carletti, V. and Foggia, P. and Percannella, G. and Saggese, A. and Strisciuglio, N. and Vento, M.",
TITLE = "Audio surveillance using a bag of aural words classifier",
BOOKTITLE = AVSS13,
YEAR = "2013",
PAGES = "81-86",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381703"}
@inproceedings{bb387641,
AUTHOR = "Hurtado, L.F. and Calvo, M. and Gomez, J.A. and Garcia, F. and Sanchis, E.",
TITLE = "A Phonetic-Based Approach to Query-by-Example Spoken Term Detection",
BOOKTITLE = CIARP13,
YEAR = "2013",
PAGES = "I:504-511",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381704"}
@inproceedings{bb387642,
AUTHOR = "Chaloupka, J. and Nouza, J. and Kucharova, M.",
TITLE = "Using Various Types of Multimedia Resources to Train System for
Automatic Transcription of Czech Historical Oral Archives",
BOOKTITLE = MM4CH13,
YEAR = "2013",
PAGES = "228-237",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381705"}
@inproceedings{bb387643,
AUTHOR = "Nouza, J. and Cerva, P. and Silovsky, J.",
TITLE = "Dealing with Bilingualism in Automatic Transcription of Historical
Archive of Czech Radio",
BOOKTITLE = MM4CH13,
YEAR = "2013",
PAGES = "238-246",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381706"}
@inproceedings{bb387644,
AUTHOR = "Chan, K.Y. and Nordholm, S.E. and Yiu, C.K.F.",
TITLE = "Multichannel filters for speech recognition using a particle swarm
optimization",
BOOKTITLE = ICARCV12,
YEAR = "2012",
PAGES = "937-942",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381707"}
@inproceedings{bb387645,
AUTHOR = "Zhao, Y. and Xu, X.N. and Yang, G.S.",
TITLE = "Unsupervised Tibetan speech features Learning based on Dynamic Bayesian
Networks",
BOOKTITLE = ICPR12,
YEAR = "2012",
PAGES = "2319-2322",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381708"}
@inproceedings{bb387646,
AUTHOR = "Nour Eddine, L. and Abdelkader, A.",
TITLE = "Reduced Universal Background Model for Speech Recognition and
Identification System",
BOOKTITLE = MCPR12,
YEAR = "2012",
PAGES = "303-312",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381709"}
@inproceedings{bb387647,
AUTHOR = "Amrous, A.I. and Debyeche, M.",
TITLE = "Robust Arabic Multi-stream Speech Recognition System in Noisy
Environment",
BOOKTITLE = ICISP12,
YEAR = "2012",
PAGES = "571-578",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381710"}
@inproceedings{bb387648,
AUTHOR = "Touazi, A. and Debyeche, M.",
TITLE = "New Encoding Algorithm for Distributed Speech Recognition Based on DTFS
Transform",
BOOKTITLE = ICISP12,
YEAR = "2012",
PAGES = "547-554",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381711"}
@inproceedings{bb387649,
AUTHOR = "Ghigi, F. and Tamarit, V. and Martinez Hinarejos, C.D. and Benedi, J.M.",
TITLE = "Active Learning for Dialogue Act Labelling",
BOOKTITLE = IbPRIA11,
YEAR = "2011",
PAGES = "652-659",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381712"}
@inproceedings{bb387650,
AUTHOR = "Meng, L. and Xiang, J. and Zhao, D. and Zhao, H.",
TITLE = "A New Application of MEG and DTI on Word Recognition",
BOOKTITLE = ICPR10,
YEAR = "2010",
PAGES = "2472-2475",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381713"}
@inproceedings{bb387651,
AUTHOR = "O'Gorman, L.",
TITLE = "Latency in Speech Feature Analysis for Telepresence Event Coding",
BOOKTITLE = ICPR10,
YEAR = "2010",
PAGES = "4464-4467",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381714"}
@inproceedings{bb387652,
AUTHOR = "Zhang, S.L. and Shi, Q. and Qin, Y.",
TITLE = "Modeling Syllable-Based Pronunciation Variation for Accented Mandarin
Speech Recognition",
BOOKTITLE = ICPR10,
YEAR = "2010",
PAGES = "1606-1609",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381715"}
@inproceedings{bb387653,
AUTHOR = "Zhang, S.L. and Zhang, S.W. and Xu, B.",
TITLE = "A Two-level Method for Unsupervised Speaker-based Audio Segmentation",
BOOKTITLE = ICPR06,
YEAR = "2006",
PAGES = "IV: 298-301",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381716"}
@inproceedings{bb387654,
AUTHOR = "Krajewski, J. and Batliner, A. and Kessel, S.",
TITLE = "Comparing Multiple Classifiers for Speech-Based Detection of
Self-Confidence: A Pilot Study",
BOOKTITLE = ICPR10,
YEAR = "2010",
PAGES = "3716-3719",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381717"}
@inproceedings{bb387655,
AUTHOR = "Nolazco Flores, J.A. and Aceves L., R.A. and Garcia Perera, L.P.",
TITLE = "Speech Magnitude-Spectrum Information-Entropy (MSIE) for Automatic
Speech Recognition in Noisy Environments",
BOOKTITLE = ICPR10,
YEAR = "2010",
PAGES = "4364-4367",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381718"}
@inproceedings{bb387656,
AUTHOR = "Kelly, F. and Harte, N.",
TITLE = "Auditory Features Revisited for Robust Speech Recognition",
BOOKTITLE = ICPR10,
YEAR = "2010",
PAGES = "4456-4459",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381719"}
@inproceedings{bb387657,
AUTHOR = "Xie, Z.Q. and Miao, Z.J.",
TITLE = "Tone Recognition of Isolated Mandarin Syllables",
BOOKTITLE = ICISP10,
YEAR = "2010",
PAGES = "412-418",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381720"}
@inproceedings{bb387658,
AUTHOR = "Alotaibi, Y.A. and Alghamdi, M. and Alotaiby, F.",
TITLE = "Speech Recognition System of Arabic Alphabet Based on a Telephony
Arabic Corpus",
BOOKTITLE = ICISP10,
YEAR = "2010",
PAGES = "122-129",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381721"}
@inproceedings{bb387659,
AUTHOR = "Lu, G. and Yu, H.Z. and Li, Y.H. and Zhang, R.S.",
TITLE = "Study on SAMPA_ST for Lhasa Tibetan and realization of automatic
labelling system",
BOOKTITLE = IASP10,
YEAR = "2010",
PAGES = "133-137",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381722"}
@inproceedings{bb387660,
AUTHOR = "Chen, X.Y. and Jin, H.M. and Yu, H.Z.",
TITLE = "Acoustic research on long and short vowels in Tibetan Lhasa dialect",
BOOKTITLE = IASP10,
YEAR = "2010",
PAGES = "561-564",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381723"}
@inproceedings{bb387661,
AUTHOR = "Sahu, V.P. and Mishra, H.K. and Sekhar, C.C.",
TITLE = "Variational Bayes Adapted GMM Based Models for Audio Clip
Classification",
BOOKTITLE = PReMI09,
YEAR = "2009",
PAGES = "513-518",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381724"}
@inproceedings{bb387662,
AUTHOR = "Verteletskaya, E. and Sakhnov, K. and Simak, B.",
TITLE = "Pitch Detection Algorithms and Voiced/Unvoiced Classification for Noisy
Speech",
BOOKTITLE = WSSIP09,
YEAR = "2009",
PAGES = "1-5",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381725"}
@inproceedings{bb387663,
AUTHOR = "Vlaj, D. and Kos, M. and Grasic, M. and Kacic, Z.",
TITLE = "Influence of Hangover and Hangbefore Criteria on Automatic Speech
Recognition",
BOOKTITLE = WSSIP09,
YEAR = "2009",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381726"}
@inproceedings{bb387664,
AUTHOR = "Hanzl, V. and Pollak, P.",
TITLE = "Accuracy Analysis of Generalized Pronunciation Variant Selection in ASR
Systems",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "399-408",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381727"}
@inproceedings{bb387665,
AUTHOR = "Camarena Ibarrola, A. and Chavez, E. and Tellez, E.S.",
TITLE = "Robust Radio Broadcast Monitoring Using a Multi-Band Spectral Entropy
Signature",
BOOKTITLE = CIARP09,
YEAR = "2009",
PAGES = "587-594",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381728"}
@inproceedings{bb387666,
AUTHOR = "Mantilla Caeiros, A. and Miyatake, M.N. and Perez Meana, H.",
TITLE = "Isolate Speech Recognition Based on Time-Frequency Analysis Methods",
BOOKTITLE = CIARP09,
YEAR = "2009",
PAGES = "297-304",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381729"}
@inproceedings{bb387667,
AUTHOR = "Veronkova, J. and Palkova, Z.",
TITLE = "Perception of Czech in Noise: Stability of Vowels",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "149-161",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381730"}
@inproceedings{bb387668,
AUTHOR = "Skarnitzl, R.",
TITLE = "Challenges in Segmenting the Czech Lateral Liquid",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "162-172",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381731"}
@inproceedings{bb387669,
AUTHOR = "Machac, P.",
TITLE = "Implications of Acoustic Variation for the Segmentation of the Czech
Trill r",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "173-181",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381732"}
@inproceedings{bb387670,
AUTHOR = "Jorschick, A.B.",
TITLE = "Voicing in Labial Plosives in Czech",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "182-189",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381733"}
@inproceedings{bb387671,
AUTHOR = "Volin, J.",
TITLE = "Normalization of the Vocalic Space",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "190-200",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381734"}
@inproceedings{bb387672,
AUTHOR = "Rajnoha, J. and Pollak, P.",
TITLE = "Czech Spontaneous Speech Collection and Annotation:
The Database of Technical Lectures",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "377-385",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381735"}
@inproceedings{bb387673,
AUTHOR = "Janda, J.",
TITLE = "Quantitative Analysis of the Relative Local Speech Rate",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "368-376",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381736"}
@inproceedings{bb387674,
AUTHOR = "Zhang, B. and Zhuang, X. and Huang, P. and Feng, C. and Zhao, J.",
TITLE = "Application of Uni-Directional Microphone Array for Identifying English
Pronunciation Errors",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-5",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381737"}
@inproceedings{bb387675,
AUTHOR = "Kuremoto, T. and Komoto, T. and Kobayashi, K. and Obayashi, M.",
TITLE = "A Voice Instruction Learning System Using PL-T-SOM",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-6",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381738"}
@inproceedings{bb387676,
AUTHOR = "Espi, M. and Takeuchi, Y.",
TITLE = "Substitution of Vocal Folds for Voice Generation by Means of Intra-Oral
Pulse Generator",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-5",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381739"}
@inproceedings{bb387677,
AUTHOR = "Orhan, Z. and Gormez, Z.",
TITLE = "Evaluation of the Concatenative Turkish Text-to-Speech System",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-5",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381740"}
@inproceedings{bb387678,
AUTHOR = "Cai, Y. and Yuan, J.P. and Hou, C.H. and Yang, J. and Wu, B.",
TITLE = "Harmonic Enhancement with Noise Reduction of Speech Signal by Comb
Filtering",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381741"}
@inproceedings{bb387679,
AUTHOR = "Li, W.F. and Billard, A. and Bourlard, H.",
TITLE = "Keyword Detection for Spontaneous Speech",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-5",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381742"}
@inproceedings{bb387680,
AUTHOR = "Zhang, X.Y. and Yao, J.X. and He, Q.A.",
TITLE = "Research of STRAIGHT Spectrogram and Difference Subspace Algorithm for
Speech Recognition",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381743"}
@inproceedings{bb387681,
AUTHOR = "Lu, X. and Matsuda, S. and Unoki, M. and Nakamura, S.",
TITLE = "Temporal Modulation Normalization for Robust Speech Feature Extraction
and Recognition",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381744"}
@inproceedings{bb387682,
AUTHOR = "Jun, Y.Z. and Lei, W. and Hao, W.",
TITLE = "A New Parameter of Speech Character Based on the Bloomfield's Model",
BOOKTITLE = CISP09,
YEAR = "2009",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381745"}
@inproceedings{bb387683,
AUTHOR = "Qasemi Zadeh, B. and Shen, J.L. and O'Neill, I. and Miller, P. and Hanna, P. and Stewart, D. and Wang, H.B.",
TITLE = "A Speech Based Approach to Surveillance Video Retrieval",
BOOKTITLE = AVSBS09,
YEAR = "2009",
PAGES = "336-339",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381746"}
@inproceedings{bb387684,
AUTHOR = "Cristani, M. and Pesarin, A. and Drioli, C. and Tavano, A. and Perina, A. and Murino, V.",
TITLE = "Auditory dialog analysis and understanding by generative modelling of
interactional dynamics",
BOOKTITLE = CVPR4HB09,
YEAR = "2009",
PAGES = "103-109",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381747"}
@inproceedings{bb387685,
AUTHOR = "Chen, J.B. and Zhang, S.Q.",
TITLE = "Manifold learning-based phoneme recognition",
BOOKTITLE = IASP09,
YEAR = "2009",
PAGES = "308-312",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381748"}
@inproceedings{bb387686,
AUTHOR = "Mahdhaoui, A. and Chetouani, M. and Zong, C.",
TITLE = "Motherese detection based on segmental and supra-segmental features",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381749"}
@inproceedings{bb387687,
AUTHOR = "Zeng, Z. and Li, X. and Ma, X.H. and Ji, Q.A.",
TITLE = "Adaptive context recognition based on audio signal",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381750"}
@inproceedings{bb387688,
AUTHOR = "Luo, L. and Lu, P.F. and Wang, Z.F.",
TITLE = "A real-time accompaniment system based on sung voice recognition",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381751"}
@inproceedings{bb387689,
AUTHOR = "Pesarin, A. and Cristani, M. and Murino, V. and Drioli, C. and Perina, A. and Tavano, A.",
TITLE = "A statistical signature for automatic dialogue classification",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381752"}
@inproceedings{bb387690,
AUTHOR = "Choi, H. and Gutierrez Osuna, R. and Choi, S.J. and Choe, Y.",
TITLE = "Kernel oriented discriminant analysis for speaker-independent phoneme
spaces",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381753"}
@inproceedings{bb387691,
AUTHOR = "Terry, L. and Katsaggelos, A.K.",
TITLE = "A phone-viseme dynamic Bayesian network for audio-visual automatic
speech recognition",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381754"}
@inproceedings{bb387692,
AUTHOR = "Krajewski, J. and Batliner, A. and Wieland, R.",
TITLE = "Multiple classifier applied on predicting microsleep from speech",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381755"}
@inproceedings{bb387693,
AUTHOR = "Banerjee, P. and Garg, G. and Mitra, P. and Basu, A.",
TITLE = "Application of triphone clustering in acoustic modeling for continuous
speech recognition in Bengali",
BOOKTITLE = ICPR08,
YEAR = "2008",
PAGES = "1-4",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381756"}
@inproceedings{bb387694,
AUTHOR = "Bouzid, A. and Ellouze, N.",
TITLE = "Voicing Detection in Noisy Speech Signal",
BOOKTITLE = ICISP08,
YEAR = "2008",
PAGES = "544-551",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381757"}
@inproceedings{bb387695,
AUTHOR = "Turkmen, H.I. and Karsligil, M.E.",
TITLE = "Reconstruction of Dysphonic Speech by MELP",
BOOKTITLE = CIARP08,
YEAR = "2008",
PAGES = "767-774",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381758"}
@inproceedings{bb387696,
AUTHOR = "Maskeliunas, R. and Rudzionis, A. and Rudzionis, V.",
TITLE = "Analysis of the Possibilities to Adapt the Foreign Language Speech
Recognition Engines for the Lithuanian Spoken Commands Recognition",
BOOKTITLE = COST08,
YEAR = "2008",
PAGES = "409-422",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381759"}
@inproceedings{bb387697,
AUTHOR = "Hain, T. and Burget, L. and Dines, J. and Garau, G. and Karafiat, M. and van Leeuwen, D. and Lincoln, M. and Wan, V.",
TITLE = "The 2007 AMI(DA) System for Meeting Transcription",
BOOKTITLE = MTPH07,
YEAR = "2007",
PAGES = "xx-yy",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381760"}
@inproceedings{bb387698,
AUTHOR = "Lamel, L. and Bilinski, E. and Gauvain, J.L. and Adda, G. and Barras, C. and Zhu, X.",
TITLE = "The LIMSI RT07 Lecture Transcription System",
BOOKTITLE = MTPH07,
YEAR = "2007",
PAGES = "xx-yy",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381761"}
@inproceedings{bb387699,
AUTHOR = "Fiscus, J.G. and Ajot, J. and Garofolo, J.S.",
TITLE = "The Rich Transcription 2007 Meeting Recognition Evaluation",
BOOKTITLE = MTPH07,
YEAR = "2007",
PAGES = "xx-yy",
BIBSOURCE = "http://www.visionbib.com/bibliography/other1023.html#TT381762"}
Last update:Sep 30, 2026 at 11:45:00