{"id":"https://openalex.org/W197367430","doi":"https://doi.org/10.21437/eurospeech.1993-64","title":"Within class optimization of cepstra for speaker recognition","display_name":"Within class optimization of cepstra for speaker recognition","publication_year":1993,"publication_date":"1993-09-22","ids":{"openalex":"https://openalex.org/W197367430","doi":"https://doi.org/10.21437/eurospeech.1993-64","mag":"197367430"},"language":"en","primary_location":{"id":"doi:10.21437/eurospeech.1993-64","is_oa":false,"landing_page_url":"https://doi.org/10.21437/eurospeech.1993-64","pdf_url":null,"source":null,"license":null,"license_id":null,"version":"publishedVersion","is_accepted":true,"is_published":true,"raw_source_name":"3rd European Conference on Speech Communication and Technology (Eurospeech 1993)","raw_type":"proceedings-article"},"type":"conference-paper","indexed_in":["crossref"],"open_access":{"is_oa":false,"oa_status":"closed","oa_url":null,"any_repository_has_fulltext":false},"authorships":[{"author_position":"first","author":{"id":"https://openalex.org/A5011872593","display_name":"Jeff Thompson","orcid":"https://orcid.org/0000-0002-2720-1625"},"institutions":[],"countries":[],"is_corresponding":false,"raw_author_name":"J. Thompson","raw_affiliation_strings":[],"raw_orcid":null,"affiliations":[]},{"author_position":"last","author":{"id":"https://openalex.org/A5021470344","display_name":"J. S. Mason","orcid":"https://orcid.org/0000-0003-0062-2339"},"institutions":[],"countries":[],"is_corresponding":false,"raw_author_name":"J. S. Mason","raw_affiliation_strings":[],"raw_orcid":null,"affiliations":[]}],"institutions":[],"countries_distinct_count":0,"institutions_distinct_count":0,"corresponding_author_ids":[],"corresponding_institution_ids":[],"apc_list":null,"apc_paid":null,"fwci":0.0,"has_fulltext":false,"cited_by_count":8,"citation_normalized_percentile":{"value":0.04668675,"is_in_top_1_percent":false,"is_in_top_10_percent":false},"cited_by_percentile_year":null,"biblio":{"volume":null,"issue":null,"first_page":"165","last_page":"168"},"is_retracted":false,"is_paratext":false,"is_xpac":false,"primary_topic":{"id":"https://openalex.org/T11309","display_name":"Music and Audio Processing","score":0.996999979019165,"subfield":{"id":"https://openalex.org/subfields/1711","display_name":"Signal Processing"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}},"topics":[{"id":"https://openalex.org/T11309","display_name":"Music and Audio Processing","score":0.996999979019165,"subfield":{"id":"https://openalex.org/subfields/1711","display_name":"Signal Processing"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}},{"id":"https://openalex.org/T10201","display_name":"Speech Recognition and Synthesis","score":0.9937999844551086,"subfield":{"id":"https://openalex.org/subfields/1702","display_name":"Artificial Intelligence"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}},{"id":"https://openalex.org/T10860","display_name":"Speech and Audio Processing","score":0.9886999726295471,"subfield":{"id":"https://openalex.org/subfields/1711","display_name":"Signal Processing"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}}],"keywords":[{"id":"https://openalex.org/keywords/speech-recognition","display_name":"Speech recognition","score":0.3866793215274811},{"id":"https://openalex.org/keywords/computer-science","display_name":"Computer science","score":0.30566781759262085}],"concepts":[{"id":"https://openalex.org/C28490314","wikidata":"https://www.wikidata.org/wiki/Q189436","display_name":"Speech recognition","level":1,"score":0.3866793215274811},{"id":"https://openalex.org/C41008148","wikidata":"https://www.wikidata.org/wiki/Q21198","display_name":"Computer science","level":0,"score":0.30566781759262085}],"mesh":[],"locations_count":1,"locations":[{"id":"doi:10.21437/eurospeech.1993-64","is_oa":false,"landing_page_url":"https://doi.org/10.21437/eurospeech.1993-64","pdf_url":null,"source":null,"license":null,"license_id":null,"version":"publishedVersion","is_accepted":true,"is_published":true,"raw_source_name":"3rd European Conference on Speech Communication and Technology (Eurospeech 1993)","raw_type":"proceedings-article"}],"best_oa_location":null,"sustainable_development_goals":[{"display_name":"Quality Education","score":0.9200000166893005,"id":"https://metadata.un.org/sdg/4"}],"awards":[],"funders":[],"has_content":{"grobid_xml":false,"pdf":false},"content_urls":null,"referenced_works_count":3,"referenced_works":["https://openalex.org/W107477321","https://openalex.org/W2148154194","https://openalex.org/W2907162034"],"related_works":["https://openalex.org/W4391375266","https://openalex.org/W2748952813","https://openalex.org/W2390279801","https://openalex.org/W2358668433","https://openalex.org/W4396701345","https://openalex.org/W2376932109","https://openalex.org/W2001405890","https://openalex.org/W4396696052","https://openalex.org/W4402327032","https://openalex.org/W2382290278"],"abstract_inverted_index":{"WITHINCLASSOPTIMIZATIONOFCEPSTRAFORSPEAKERRECOGNITIONJ.Thompson&J.S.MasonDepartmentofElectrical&ElectronicEngineering,UniversityCollegeofSwansea,SWANSEA,SA28PP,UKABSTRACTByidentifyinghighlysp":[0],"eakersp":[1,16,38,44,48,54,65,73,91],"eci":[2,17,39,45,49,51,55,66,74,92],"casp":[3],"ectsofcep-stralfeatures,thetaskofsp":[4],"eakerrecognitioncanb":[5],"ep":[6,11],"otentiallysimpli":[7],"ed.Theidenti":[8],"cationpro":[9],"cesscanb":[10],"erformedb":[12],"othphoneticallyandinthecepstraldomain.Thephoneticanalysisaimstodeterminethetemp":[13],"oralasp":[14],"ectsofutteranceswhichexhibitthehighestdegreeofsp":[15],"city,whilecepstralanalysisexaminesindividualcepstrawithinthesetemp":[18],"oraldivisions.Thispap":[19],"eraimstocom-plementworkthathasalreadyb":[20],"eenconductedonaphoneticbasis,byp":[21],"erforminganalysisup":[22],"onthein-dividualcepstralco":[23],"e\u000ecientswithinutterances.Keywords:Sp":[24],"eakerRecognition,Sp":[25],"eci-":[26],"city,Inter-sp":[27],"eakerVariation,tra-sp":[28],"eakaria-tion.1INTRODUCTIONInvestigationsintothephoneticasp":[29],"ectsofsp":[30],"eechsp":[31],"eakerrecognitionarerep":[32],"ortedbyEato":[33],"ck[1":[34],"],vandenHeuvel[2":[35],"]andBonastre[3].Intheseinvestiga-tions,ithasb":[36],"eenfoundthatsp":[37],"cityvariesnoticeablyb":[40],"etweendi":[41],"erentphoneticsubgroups,butremainsmoreconstantwithinthem.Thesephonemesubgroupscanb":[42],"ecrudelyrankedinthefollowingor-der:longvowels,nasals,shortfricatives,andplosives,wherelongvowelsexhibitthehighestdegreeofsp":[43],"city,andplosivesthelowest.Thispap":[46],"erexaminesthecepstralrepresentationsofutterances,inparticularthosesubgroupswhichex-hibitahighsp":[47],"citysuchaslongvowelsandnasals.Moresp":[50],"callytheaimistoidentifyindividualcepstralco":[52],"e\u000ecientswhichcontainthema-jorityofthesp":[53],"cinformation,andshowbyanalysisofcepstraldistributions,whsomesp":[56],"eak-ersp":[57],"erformsigni":[58],"cantlyworsethanothers.Severalb":[59],"ene":[60],"tsfromidentifyingthemoste":[61],"ectiveco":[62],"e\u000ecientsexist.Firstlytherecognitionsystemcanb":[63],"ereducedincomplexitybusingcepstrathatcon-tainahighdegreeofsp":[64],"city.Secondly,p":[67],"erhapsb":[68],"etteroverallp":[69],"erformancecanb":[70],"eachievedby`liftering'thefeaturevectorspriortoclassi":[71],"ca-tion[4],wherethecepstraweightingsarederiveddi-rectlyfromtheaverageoftheirsp":[72],"cityforagroupofsp":[75],"eakers.Exp":[76],"erimentssimilarnaturearerep":[77],"ortedonfrequencydomainfeaturesusingavariancemeasure[5":[78],"]andmultiplefeaturesusingdy-namicprogramming[6":[79],"][7].2CEPSTRALFEATURESThecepstralfeaturesconsideredherearegeneratedfromasp":[80],"eechdatabaseusing24channelDFT-simulatedmel":[81],"lter-bank[8":[82],"],witha95%overlapb":[83],"e-tweensuccessiveframes.This":[84],"lterbankiswidelyusedandleadstothestandardformofmelcepstra.Thesp":[85],"eechdatabaseconsistsof20sp":[86],"eakers(10male,10female),utteringthelettersofalphab":[87],"et(atoz)threetimes.Thesamplingrateofthedatais10KHz.Ofthe24p":[88],"ossiblestandardcepstralfeaturesonlyMFCCs1-14areusedinordertoremovethema":[89],"jor-ityofpitchinformation(whicmaexhibitamis-leadinglyhighsp":[90],"citybutisprobablyanunreliableparameterinpracticalsp":[93],"eakerrecognitionsystems[9":[94],"]).Alsothisreduceddimensionfeaturesetisinverseariance,zeromeanweighted,onaglobalp":[95],"o":[96],"oledbasis.":[97]},"counts_by_year":[],"updated_date":"2026-07-29T14:22:42.915294","created_date":"2025-10-10T00:00:00"}
