{"meta":{"query_hash":"ed141c5d9607","filters":{"venue":"American Speech"},"cohort_total":36,"direct_labels_cover":0,"predictions_cover":36,"exported":36,"export_cap":100000,"truncated":false,"label_status":"direct model label, unvalidated","prediction_status":"machine_predicted_unvalidated (Codex and Gemma teacher distillation)","score_status":"score_only:v0-immature-baseline","snapshot":{"source":"OpenAlex, pinned release, all 482 partitions","release":"2026-06-24","frame_built":"2026-07-12"},"permalink":"https://metacan.xera.ac/q/ed141c5d9607","api":"https://metacan.xera.ac/api/v1/cohort?venue=American+Speech"},"results":[{"id":"W1969441012","doi":"10.1215/00031283-2009-014","title":"REVISED PERCEPTIONS: CHANGING DIALECT PERCEPTIONS IN WISCONSIN AND MICHIGAN'S UPPER PENINSULA","year":2009,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":23,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Peninsula; Indexicality; Perception; Variety (cybernetics); Ethnic group; Geography; Variation (astronomy); Sociology; Economic geography; History; Linguistics; Anthropology; Psychology; Archaeology","score_opus":0.013192286435816535,"score_gpt":0.31857143749448824,"score_spread":0.3053791510586717,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W1969441012","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9585483,0.00006429477,0.00042039496,0.0079101585,0.00014691871,0.00023293987,0.000004821062,0.000079747566,0.032592427],"genre_scores_gemma":[0.98980457,0.00029929992,0.0055820667,0.0025621525,0.000326799,0.0000076629285,0.000007687326,0.000007308256,0.0014024479],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9988879,0.00021404073,0.00017244602,0.00024192965,0.00014689536,0.00033678985],"domain_scores_gemma":[0.9995429,0.00009530979,0.00007084365,0.00013198121,0.000054229225,0.00010475412],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00042927833,0.00009399118,0.0001923876,0.00024340083,0.00031212592,0.000051161664,0.00010286129,0.000048179867,0.000630651],"category_scores_gemma":[0.00057884795,0.00009890624,0.000040472372,0.00061422377,0.00040153027,0.00006841287,0.000019529416,0.00013319886,0.00007390336],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00008179847,0.00045596092,0.20212993,0.000011512616,0.00003524112,0.00013769015,0.3081103,0.000012664746,0.014324099,0.10284456,0.0032954141,0.36856082],"study_design_scores_gemma":[0.0011866585,0.00037701795,0.8487573,0.00008152506,0.000063228406,0.000051873565,0.09061191,0.00031216655,0.000046173507,0.007244829,0.05050762,0.0007596893],"about_ca_topic_score_codex":0.0021409746,"about_ca_topic_score_gemma":0.0017023872,"teacher_disagreement_score":0.64662737,"about_ca_system_score_codex":0.000059597292,"about_ca_system_score_gemma":0.00010532771,"threshold_uncertainty_score":0.6905186},"labels":[],"label_agreement":null},{"id":"W2027642902","doi":"10.1215/00031283-80-2-180","title":"THE USES AND MEANINGS OF THE FEMALE TITLE<i>MS.</i>","year":2005,"lang":"en","type":"article","venue":"American Speech","topic":"Gender Studies in Language","field":"Social Sciences","cited_by":50,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Psychology; Gender gap; Variation (astronomy); Medical education; Pedagogy; Medicine; Demographic economics","score_opus":0.013850086630987479,"score_gpt":0.29832190742116016,"score_spread":0.2844718207901727,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2027642902","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.15810591,0.00047628168,0.000005109041,0.003933137,0.00012412401,0.00006354656,0.0000014787705,0.00001642216,0.837274],"genre_scores_gemma":[0.9905082,0.00019491762,0.000546533,0.00037997545,0.000116053256,0.0000013802054,3.389268e-8,0.0000032304888,0.008249682],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9996402,0.00003836433,0.00004531507,0.000052766365,0.00012940113,0.00009390958],"domain_scores_gemma":[0.999728,0.00008446318,0.000044773835,0.000111819274,0.000016900714,0.000014018288],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00014923327,0.000027133621,0.000046871282,0.0000046562664,0.00021742508,0.000010775435,0.00015204832,0.000007309233,0.000104213774],"category_scores_gemma":[0.0001062,0.000015330354,0.000018013972,0.00012125818,0.00081156055,0.0000145781505,0.00005978818,0.00003678878,0.000033388867],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000036176618,0.00002218609,0.031130686,0.0000036683161,0.000044599397,0.0000012327617,0.05393452,8.029951e-7,0.000251754,0.028039776,0.1212981,0.76526904],"study_design_scores_gemma":[0.000022763736,0.000008000646,0.0057827253,0.000002297438,0.0000058732508,6.8536826e-7,0.016270189,9.990996e-7,0.00039562062,0.00015991603,0.9773159,0.00003504162],"about_ca_topic_score_codex":0.0015015741,"about_ca_topic_score_gemma":0.0019605488,"teacher_disagreement_score":0.85601777,"about_ca_system_score_codex":0.00001321221,"about_ca_system_score_gemma":0.00002049571,"threshold_uncertainty_score":0.29902288},"labels":[],"label_agreement":null},{"id":"W2053116124","doi":"10.1215/00031283-79-3-250","title":"REAL AND APPARENT TIME IN LANGUAGE CHANGE: LATE ADOPTION OF CHANGES IN MONTREAL ENGLISH","year":2004,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":117,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Language change; Variation (astronomy); Grading (engineering); Linguistics; Psychology","score_opus":0.014446154491451488,"score_gpt":0.29537334296094675,"score_spread":0.28092718846949527,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2053116124","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9828498,0.000023881046,0.0000027750584,0.00063325133,0.00007577572,0.00013349397,0.000006245434,0.000019974028,0.016254848],"genre_scores_gemma":[0.99872583,0.00030597838,0.00044501538,0.000121142955,0.00018338303,0.000011135318,0.000009183101,0.0000038553762,0.00019450158],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99946517,0.000069888614,0.00009876858,0.00011596757,0.000101847196,0.00014834534],"domain_scores_gemma":[0.99975955,0.000034336037,0.00007588776,0.000059323345,0.00003235572,0.000038563103],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0002577628,0.000046572153,0.00013701327,0.000102087484,0.000019827265,0.0000058951787,0.000054737844,0.000034505894,0.00004412395],"category_scores_gemma":[0.00016782705,0.000048100406,0.000009521368,0.00025796035,0.00015972756,0.000028165356,0.000018878292,0.00005359044,0.000008127307],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010771757,0.00033289532,0.15686798,0.000020338848,0.000015031331,0.0001654596,0.58236307,0.000027508384,0.0026178092,0.008859226,0.0000748332,0.24854814],"study_design_scores_gemma":[0.0014579087,0.00027838227,0.9604459,0.000051008912,0.00001261022,0.0000011054866,0.034178104,0.00004295825,0.00043959296,0.00096823665,0.0018948939,0.00022934363],"about_ca_topic_score_codex":0.20440362,"about_ca_topic_score_gemma":0.17869632,"teacher_disagreement_score":0.8035779,"about_ca_system_score_codex":0.00006768277,"about_ca_system_score_gemma":0.000033517288,"threshold_uncertainty_score":0.83629036},"labels":[],"label_agreement":null},{"id":"W2143503845","doi":"10.1215/00031283-2010-007","title":"THE ROLE OF SOCIAL FACTORS IN THE CANADIAN VOWEL SHIFT: EVIDENCE FROM TORONTO","year":2010,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":93,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Vowel; Ethnic group; Focus (optics); Sound change; Mid vowel; Variation (astronomy); Vowel length; Linguistics; Sociology; Geography; Psychology; Anthropology; Formant","score_opus":0.01657586381304114,"score_gpt":0.3150458259102862,"score_spread":0.298469962097245,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2143503845","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95040023,0.000040191142,0.0000029764074,0.0069874446,0.0004797646,0.00012287266,0.000014605549,0.000009920643,0.041941963],"genre_scores_gemma":[0.99899584,0.000014851289,0.00011843816,0.00041380152,0.00033814812,0.0000050455615,0.0000020868533,0.0000035658147,0.00010822056],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9991081,0.00020497692,0.00012205254,0.00010565823,0.00023282492,0.00022643746],"domain_scores_gemma":[0.99901783,0.0006218992,0.000098614866,0.00014816331,0.00004724171,0.00006623278],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00065894565,0.000051072388,0.00009313974,0.000014982252,0.0005012532,0.00004730287,0.00048070043,0.000050365605,0.0004337612],"category_scores_gemma":[0.0012717373,0.000032404332,0.000033524593,0.00014622345,0.00063636695,0.000046302564,0.000016025497,0.0001520188,0.00001279698],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012681388,0.000027436312,0.18624277,4.0879416e-7,0.000015819187,0.000005718036,0.6891986,1.5416471e-7,0.000691825,0.08928549,0.0012634307,0.033255633],"study_design_scores_gemma":[0.000039097835,0.000018019804,0.8235914,0.0000014508914,0.000006600076,1.311405e-7,0.04722087,0.0000034286843,0.00009980046,0.0042831674,0.12467242,0.000063607455],"about_ca_topic_score_codex":0.99684274,"about_ca_topic_score_gemma":0.9995293,"teacher_disagreement_score":0.6419778,"about_ca_system_score_codex":0.000112495174,"about_ca_system_score_gemma":0.0005442114,"threshold_uncertainty_score":0.47493806},"labels":[],"label_agreement":null},{"id":"W2154725923","doi":"10.1215/00031283-80-1-22","title":"THE NORTH AMERICAN REGIONAL VOCABULARY SURVEY: NEW VARIABLES AND METHODS IN THE STUDY OF NORTH AMERICAN ENGLISH","year":2005,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":68,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Vocabulary; Variation (astronomy); Geography; Nova scotia; American English; Regional variation; Lexical item; Linguistics; History; Regional science; Genealogy; Archaeology; Political science; Law","score_opus":0.034425505229728814,"score_gpt":0.3645259995402809,"score_spread":0.3301004943105521,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2154725923","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9943487,0.000037545957,0.00045154648,0.0022382268,0.00014582046,0.00042747907,0.000010148625,0.00003554783,0.0023049798],"genre_scores_gemma":[0.98452663,0.0002894286,0.012830166,0.0015373025,0.0005552966,0.00002071177,0.000010024621,0.000014819462,0.0002155974],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99520445,0.0031370015,0.00041186833,0.00034977501,0.00047808676,0.00041880886],"domain_scores_gemma":[0.9957132,0.0029934375,0.00052577467,0.0004122776,0.00021718159,0.00013812551],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0028867172,0.00016184489,0.00041937118,0.00009628693,0.00039643646,0.00006665258,0.0006345216,0.000016189095,0.000013756111],"category_scores_gemma":[0.00275685,0.000109746696,0.000044010485,0.0021891333,0.0021063806,0.000058982012,0.00008598853,0.00026255328,0.000003082374],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00004189863,0.00016283814,0.71782905,5.234061e-7,0.00003534178,0.0000043118575,0.025161268,0.00001633481,8.49577e-7,0.0004571425,0.0012518684,0.25503856],"study_design_scores_gemma":[0.00023122167,0.0003630458,0.88693184,0.0000010573583,0.000025008792,0.00000139444,0.029370412,0.000022913411,0.000001115048,0.000032599128,0.08289179,0.00012759648],"about_ca_topic_score_codex":0.29631448,"about_ca_topic_score_gemma":0.626948,"teacher_disagreement_score":0.33063355,"about_ca_system_score_codex":0.000065113505,"about_ca_system_score_gemma":0.00036250096,"threshold_uncertainty_score":0.77610475},"labels":[],"label_agreement":null},{"id":"W2171732113","doi":"10.1215/00031283-2726386","title":"A Weird (LANGUAGE) Tale: Variation and Change in the Adjectives of Strangeness","year":2014,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":19,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Adjective; Variation (astronomy); Linguistics; Grammar; Language change; Context (archaeology); Strangeness; History; Sociology; Noun; Philosophy; Archaeology","score_opus":0.019793404001440387,"score_gpt":0.3316941967168489,"score_spread":0.31190079271540855,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2171732113","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.98252505,0.000044571352,0.00082187116,0.001529201,0.000117887765,0.00018382187,0.0000038859616,0.000014353227,0.014759346],"genre_scores_gemma":[0.99852455,0.000029907185,0.00062058703,0.0005528176,0.00021053375,0.000011760896,0.0000016482235,0.0000027211004,0.000045486082],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.99931234,0.0002992223,0.00008081448,0.00009102353,0.000118435324,0.000098180164],"domain_scores_gemma":[0.9995506,0.00023771921,0.000081242055,0.000078542405,0.00003253786,0.000019382798],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0006323304,0.00003652637,0.000091189126,0.000044347355,0.000053258213,0.000010393039,0.00009282355,0.000021716924,0.000048150647],"category_scores_gemma":[0.00050110294,0.000028468305,0.000011967227,0.00024487014,0.00023355901,0.000029730434,0.000009937619,0.000041091283,0.0000026018372],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000021067237,0.000081450904,0.036350828,0.00000863324,0.0000113259475,0.000005558906,0.5801325,7.560015e-7,0.0012050043,0.07284907,0.000037080772,0.30929673],"study_design_scores_gemma":[0.00030590198,0.00014730646,0.95855093,0.00000782736,0.000015576281,0.0000022409606,0.033982728,0.00011559015,0.00013629682,0.0013674699,0.0052675665,0.00010054381],"about_ca_topic_score_codex":0.05693806,"about_ca_topic_score_gemma":0.007690692,"teacher_disagreement_score":0.92220014,"about_ca_system_score_codex":0.0000115192715,"about_ca_system_score_gemma":0.000025507907,"threshold_uncertainty_score":0.9493419},"labels":[],"label_agreement":null},{"id":"W2312271833","doi":"10.1215/00031283-2848989","title":"The Michigan Upper Peninsula English Vowel System in Finnish American Communities in Marquette County","year":2014,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Vowel; Obstruent; Dominance (genetics); Linguistics; Geography; Peninsula; History; Archaeology","score_opus":0.009139689905373548,"score_gpt":0.2711341489726681,"score_spread":0.2619944590672945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2312271833","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9045446,0.000025376778,0.00004301769,0.00215233,0.0007120807,0.00024142576,0.000022561035,0.000110692046,0.09214793],"genre_scores_gemma":[0.996808,0.000086944776,0.00061506196,0.0017133878,0.0002714922,0.000031865136,0.000015612573,0.000013384373,0.00044423682],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99751693,0.0011615945,0.0003385808,0.00018768439,0.00030548277,0.00048972835],"domain_scores_gemma":[0.99809617,0.0010414195,0.0002413518,0.0003781658,0.00015535657,0.000087531735],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0019366668,0.00013145173,0.00032260863,0.000115186325,0.00039787218,0.000103503124,0.0005532155,0.00004681858,0.0000418171],"category_scores_gemma":[0.0012313937,0.00011256454,0.00004750473,0.00065766973,0.0016266183,0.000051950057,0.00007828448,0.00031272013,0.000030772124],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00011094083,0.000158521,0.5820877,0.000018284354,0.000035061676,0.000041741278,0.17368963,0.000034440895,0.000039361865,0.19685201,0.005342281,0.041589998],"study_design_scores_gemma":[0.000503244,0.00012300421,0.09768867,0.00003568952,0.000010830752,0.0000037782063,0.26087034,0.00050807936,0.000014829478,0.00030532136,0.6396354,0.00030082802],"about_ca_topic_score_codex":0.21801017,"about_ca_topic_score_gemma":0.17183289,"teacher_disagreement_score":0.63429314,"about_ca_system_score_codex":0.0001748575,"about_ca_system_score_gemma":0.00015358864,"threshold_uncertainty_score":0.843279},"labels":[],"label_agreement":null},{"id":"W2317852315","doi":"10.1215/00031283-2322628","title":"Social and Linguistic Constraints on Relativizer Omission in Canadian English","year":2013,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":6,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Linguistics; Relative clause; Variety (cybernetics); Dependent clause; Animacy; Vernacular; Grammar; Head (geology); Sociology; Psychology; Computer science; Philosophy; Artificial intelligence","score_opus":0.020304784826196555,"score_gpt":0.31708917846608986,"score_spread":0.2967843936398933,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2317852315","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.41050357,0.0000037506504,0.000006407302,0.003087494,0.00032305662,0.00018186687,0.0000036995716,0.00003213219,0.58585805],"genre_scores_gemma":[0.9970737,0.0000072492303,0.00040402447,0.0011481356,0.00042419523,0.0000073212827,0.0000028808756,0.0000061017386,0.0009264034],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","domain_scores_codex":[0.99915874,0.00013671428,0.000115595256,0.00015936438,0.00011788521,0.00031170325],"domain_scores_gemma":[0.99939305,0.00016104765,0.000056224177,0.00005496351,0.00010550202,0.00022921742],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00029314836,0.000062578394,0.00012409323,0.00010867806,0.00021132648,0.000036957365,0.00007106345,0.000060265626,0.0011209901],"category_scores_gemma":[0.0033851091,0.000063468746,0.000015009598,0.00019555933,0.00062495575,0.000023728748,0.000010405986,0.00013582328,0.00009534298],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000018148281,0.0000805326,0.24606496,0.000004647639,0.000019474199,0.00010743408,0.1192491,8.8801556e-7,0.00006631234,0.3988634,0.009873019,0.22565208],"study_design_scores_gemma":[0.00091576454,0.0001669392,0.5049132,0.000028662558,0.0000148169,0.000002019216,0.01693808,0.000068505935,0.00002590743,0.014009104,0.4624431,0.0004738559],"about_ca_topic_score_codex":0.78514636,"about_ca_topic_score_gemma":0.35785133,"teacher_disagreement_score":0.58657014,"about_ca_system_score_codex":0.00015746108,"about_ca_system_score_gemma":0.00038549156,"threshold_uncertainty_score":0.9997921},"labels":[],"label_agreement":null},{"id":"W2460447947","doi":"10.1215/00031283-3633085","title":"Newspaper Dialectology: Harnessing the Power of the Mass Media to Study Canadian English","year":2016,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"Variation (astronomy); Newspaper; Dialectology; Vocabulary; American English; Alternation (linguistics); Mass media; Linguistics; Media studies; Geography; Sociology; History; Advertising","score_opus":0.01581522833933536,"score_gpt":0.28986082190909396,"score_spread":0.2740455935697586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2460447947","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.90998197,0.000006613832,0.000078115256,0.025269218,0.002054071,0.00037217804,0.00001186714,0.000027369064,0.062198605],"genre_scores_gemma":[0.99609965,0.00000264477,0.00024140553,0.0026036159,0.00031085967,0.000010727078,1.8726351e-7,0.0000072598436,0.0007236513],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.99877685,0.00042562815,0.00013845778,0.00015469978,0.00021976291,0.0002846192],"domain_scores_gemma":[0.99885887,0.00043947014,0.00009797532,0.00028281542,0.00017639181,0.00014448643],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005830277,0.00006473128,0.00013603963,0.000048089216,0.00031900388,0.000019293295,0.00044914908,0.000031306452,0.00073420664],"category_scores_gemma":[0.004696043,0.0000314753,0.000036940117,0.00044548113,0.00059602124,0.00002185802,0.0000402322,0.0000690895,0.000036236856],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000020723453,0.000083264196,0.62485045,5.3787596e-7,0.000068447946,0.000020906718,0.25373003,7.5777666e-7,0.0010344626,0.016313381,0.0073364987,0.096540555],"study_design_scores_gemma":[0.0004268451,0.00013824867,0.5567691,0.000009679475,0.00004696349,0.0000010630202,0.09290913,3.0988576e-7,0.000383665,0.001199489,0.34790593,0.00020956973],"about_ca_topic_score_codex":0.3662488,"about_ca_topic_score_gemma":0.61446786,"teacher_disagreement_score":0.34056944,"about_ca_system_score_codex":0.00009340872,"about_ca_system_score_gemma":0.0005099073,"threshold_uncertainty_score":0.80390483},"labels":[],"label_agreement":null},{"id":"W2546747998","doi":"10.1215/00031283-3701026","title":"Why Does Canadian English Use <i>try to</i> but British English Use <i>try And</i>? Let's Try and/to Figure It Out","year":2016,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":4,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"History; Linguistics; Philosophy","score_opus":0.015621130313769722,"score_gpt":0.26337026711753747,"score_spread":0.24774913680376776,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2546747998","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.94032156,0.000051828036,0.00022792033,0.031074125,0.0040445602,0.0007042345,0.0008013034,0.00037627455,0.022398164],"genre_scores_gemma":[0.90673524,0.00033262154,0.002787144,0.06612618,0.0025419283,0.000035754536,0.000015690805,0.000045786604,0.02137967],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9976886,0.00022695365,0.0002959637,0.000640916,0.0003546334,0.0007929142],"domain_scores_gemma":[0.9971949,0.0005655746,0.00011928075,0.0003417257,0.00052359654,0.0012549406],"candidate_categories":["metaresearch"],"consensus_categories":[],"category_scores_codex":[0.00043334483,0.00021179121,0.00037231296,0.00016929624,0.00043745383,0.00048066548,0.00032035867,0.00017223317,0.000611684],"category_scores_gemma":[0.009192065,0.00019007022,0.000060096736,0.00040316657,0.0006629579,0.00034459983,0.00011227249,0.00019631242,0.000049279774],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00006389694,0.00008804562,0.15198313,0.000017929811,0.00008194295,0.00053291564,0.062321264,0.0000010964712,0.0006403245,0.003422769,0.7131741,0.06767258],"study_design_scores_gemma":[0.0003319111,0.00008629312,0.010404868,0.000062135594,0.000025701225,0.000007952855,0.006273129,5.350771e-7,0.00009299899,0.00010381658,0.9822517,0.0003589774],"about_ca_topic_score_codex":0.7236761,"about_ca_topic_score_gemma":0.8038488,"teacher_disagreement_score":0.26907757,"about_ca_system_score_codex":0.00020094882,"about_ca_system_score_gemma":0.00047250712,"threshold_uncertainty_score":0.9991539},"labels":[],"label_agreement":null},{"id":"W2617712796","doi":"10.1215/00031283-4153186","title":"Is One Innovation Enough? Leaders, Covariation, and Language Change","year":2017,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":12,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Linguistics; Linguistic change; Scrutiny; Language change; Sociology; Possession (linguistics); Speech community; Psychology; Political science; Law","score_opus":0.10491769609340386,"score_gpt":0.3884767437273296,"score_spread":0.28355904763392575,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2617712796","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.7019009,0.000082887265,0.0015976713,0.11902673,0.0010288456,0.0005579504,0.000037240476,0.0001782546,0.17558949],"genre_scores_gemma":[0.9887861,0.000063635285,0.0032659252,0.004686245,0.00078917225,0.000009556236,0.00001034636,0.0000064344263,0.0023825734],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9993606,0.000051556955,0.00010968064,0.00015503235,0.00016323023,0.00015989818],"domain_scores_gemma":[0.99933964,0.000039746723,0.00022846546,0.00021052574,0.0001402135,0.00004140171],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00038132482,0.000053337826,0.00011034545,0.00006684761,0.000581528,0.00009250111,0.00015032901,0.000041352654,0.00040353977],"category_scores_gemma":[0.0010987786,0.000058031874,0.00001113403,0.00017008276,0.0005655344,0.000121529956,0.00003583722,0.00006298423,0.00006255822],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000021831245,0.00008408053,0.042968646,0.0000094215,0.000047990925,0.000018974071,0.18440266,3.4627064e-8,0.0015023396,0.4580947,0.006061976,0.30678737],"study_design_scores_gemma":[0.0010231213,0.00019680346,0.42560497,0.000024035367,0.000057742385,0.0000058193127,0.043535482,0.000042209547,0.0014326235,0.0054616765,0.52211475,0.000500742],"about_ca_topic_score_codex":0.03682115,"about_ca_topic_score_gemma":0.0012965216,"teacher_disagreement_score":0.5160528,"about_ca_system_score_codex":0.000032662672,"about_ca_system_score_gemma":0.00006158429,"threshold_uncertainty_score":0.96959275},"labels":[],"label_agreement":null},{"id":"W2619265518","doi":"10.1215/00031283-4153153","title":"Oral Histories as a Window to Sociolinguistic History and Language History: Exploring Earlier Ontario English with the Farm Work and Farm Life Since 1890 Oral History Collection","year":2016,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":26,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Victoria","funders":"","keywords":"Variation (astronomy); Oral history; Casual; Linguistics; History; Sociology; Interview; Sociolinguistics; Psychology; Anthropology; Law; Political science","score_opus":0.03444126782243698,"score_gpt":0.2564664985185454,"score_spread":0.22202523069610844,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2619265518","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.92477435,0.0016359679,0.00051257276,0.0040059527,0.00572117,0.0006633532,0.000005391003,0.0003073731,0.062373895],"genre_scores_gemma":[0.93979114,0.000049512808,0.0012288984,0.0025899275,0.0006955463,0.00007890426,0.0000016050586,0.000032380703,0.05553211],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.9983281,0.00025643688,0.00019605413,0.00044966117,0.00036154533,0.00040822226],"domain_scores_gemma":[0.998687,0.00030570186,0.00021753456,0.00024196651,0.00025119568,0.00029655916],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0005660417,0.00020703247,0.00033697067,0.000121646146,0.00032607865,0.00002343859,0.00019428256,0.00007409069,0.00044393845],"category_scores_gemma":[0.0020378416,0.00015155379,0.00004094609,0.00016675751,0.001678216,0.000075560914,0.000067666544,0.00022186985,0.00001612745],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":true,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0006086462,0.00009201853,0.034419887,0.000009977972,0.00010602578,0.000110610636,0.8514042,0.0000012643422,0.00024809304,0.022725994,0.058060948,0.032212332],"study_design_scores_gemma":[0.00045020934,0.00024587684,0.010002891,0.000015539934,0.00005636041,0.000004013533,0.011132341,4.0950533e-7,0.0000047941794,0.000028187196,0.977796,0.00026337267],"about_ca_topic_score_codex":0.22542895,"about_ca_topic_score_gemma":0.2878865,"teacher_disagreement_score":0.9197351,"about_ca_system_score_codex":0.00949068,"about_ca_system_score_gemma":0.0019222093,"threshold_uncertainty_score":0.9943117},"labels":[],"label_agreement":null},{"id":"W2808769745","doi":"10.1215/00031283-6926179","title":"Teaching Linguistics through Lexicography","year":2018,"lang":"en","type":"article","venue":"American Speech","topic":"Lexicography and Language Studies","field":"Arts and Humanities","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Citation; Icon; Conversation; History; Linguistics; Library science; Computer science; Philosophy","score_opus":0.023885673417659126,"score_gpt":0.28246792992665787,"score_spread":0.25858225650899874,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2808769745","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.06946299,0.00021054444,0.00010444122,0.00025873596,0.0010328658,0.000084539075,0.000023904573,0.00025693877,0.928565],"genre_scores_gemma":[0.9835345,0.000023662678,0.0037139938,0.0026347768,0.0069188797,0.000007405711,0.000006184753,0.000025090285,0.0031355042],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99905074,0.00004578306,0.00017108126,0.00023880914,0.00017413909,0.00031944772],"domain_scores_gemma":[0.99934685,0.00006584574,0.00010950252,0.00025356506,0.00017963092,0.00004458253],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00010853553,0.0001672296,0.00020355286,0.00009740322,0.00075003353,0.00010944185,0.00018238845,0.000016353968,0.00088127225],"category_scores_gemma":[0.00013067732,0.00013626194,0.0001378726,0.000109140965,0.0023547644,0.00006671963,0.000060565068,0.00019582568,0.00018571272],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00003057421,0.00014728896,0.0030533518,0.0000187402,0.0002252486,0.000028961122,0.18850999,6.368596e-8,0.000037938236,0.67170817,0.048371103,0.08786859],"study_design_scores_gemma":[0.000103907805,0.00033612567,0.00022083965,0.00001703618,0.000027075574,0.0000020096413,0.018798677,0.0000024851913,0.00019291762,0.005513588,0.9745619,0.00022339694],"about_ca_topic_score_codex":0.0076420642,"about_ca_topic_score_gemma":0.0018888285,"teacher_disagreement_score":0.92619085,"about_ca_system_score_codex":0.000009108306,"about_ca_system_score_gemma":0.000011512086,"threshold_uncertainty_score":0.99896616},"labels":[],"label_agreement":null},{"id":"W2809137360","doi":"10.1215/00031283-6926135","title":"New York City English in Film","year":2018,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":27,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"History; Icon; Citation; Linguistics; Library science; Computer science","score_opus":0.02829758963500316,"score_gpt":0.3210368167598412,"score_spread":0.292739227124838,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2809137360","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.16408147,0.000010491988,0.00077551213,0.0016629265,0.001478093,0.00011088098,0.0000026642533,0.0001130022,0.83176494],"genre_scores_gemma":[0.9802143,0.000006776815,0.007989986,0.0014801216,0.0022045956,0.0000018776834,0.0000016328128,0.0000048125107,0.008095902],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.9992806,0.00009365041,0.000106507054,0.00015011578,0.00014498751,0.00022416931],"domain_scores_gemma":[0.9995404,0.00008352295,0.00006252152,0.000118821525,0.00008848177,0.00010625919],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00030961377,0.00004786954,0.00010811142,0.000050968913,0.0000994018,0.000022176891,0.00016602063,0.000035918805,0.0028377622],"category_scores_gemma":[0.0014244851,0.000049796305,0.000021145263,0.00047478176,0.0003786128,0.00002539691,0.000022808177,0.00007495009,0.00020417766],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000054181906,0.00010713793,0.15636037,0.0000016213198,0.00001974131,0.00003992746,0.1355982,0.0000010325399,0.00010453899,0.14614499,0.28940576,0.2721625],"study_design_scores_gemma":[0.00019105685,0.00008325411,0.019959074,0.0000031051663,0.0000035826442,5.367401e-7,0.0059564407,0.000006684214,0.00012100573,0.0020451613,0.971525,0.0001050938],"about_ca_topic_score_codex":0.080217175,"about_ca_topic_score_gemma":0.021299848,"teacher_disagreement_score":0.8236691,"about_ca_system_score_codex":0.000058528505,"about_ca_system_score_gemma":0.00025011072,"threshold_uncertainty_score":0.99807376},"labels":[],"label_agreement":null},{"id":"W2889611509","doi":"10.1215/00031283-7251241","title":"Golly, Gosh, and Oh My God! What North American Dialects can Tell Us about Swear Words","year":2018,"lang":"en","type":"article","venue":"American Speech","topic":"Swearing, Euphemism, Multilingualism","field":"Social Sciences","cited_by":11,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Praise; Population; History; White (mutation); Sociology; Demography; Psychology; Social psychology","score_opus":0.014146741483473908,"score_gpt":0.31173229707894334,"score_spread":0.29758555559546945,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2889611509","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.97795194,0.0002264578,0.000016661666,0.0010568497,0.0005923877,0.00037685077,0.000014645333,0.00028920977,0.019474981],"genre_scores_gemma":[0.9892805,0.0015301211,0.0028854483,0.002318353,0.0018121368,0.00002153877,0.0000083956775,0.000075407756,0.0020680707],"study_design_codex":"design_other","study_design_gemma":"observational","domain_scores_codex":[0.9965085,0.0002546779,0.00038777926,0.0008624363,0.0008168485,0.0011698027],"domain_scores_gemma":[0.9977656,0.00025730347,0.00047567076,0.00063105964,0.00024331865,0.0006270203],"candidate_categories":["metaepi_narrow","sts"],"consensus_categories":[],"category_scores_codex":[0.00051944365,0.00037771516,0.0006343542,0.00017403043,0.00073876034,0.00051249785,0.00069080316,0.00008080894,0.00023820066],"category_scores_gemma":[0.0005676069,0.0003944383,0.000115962386,0.0019262042,0.006409293,0.00040448253,0.00023046424,0.00036690373,0.00020086388],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00007716068,0.00011105439,0.3254047,0.000011994436,0.00009056978,0.000095201016,0.06883076,0.0000016958681,0.0009517073,0.0001073328,0.001333009,0.6029848],"study_design_scores_gemma":[0.0009953874,0.0010198143,0.5607882,0.00014444183,0.00014969477,0.00004408675,0.08063737,0.00008850087,0.008770751,0.00017155151,0.34522906,0.0019611425],"about_ca_topic_score_codex":0.13344967,"about_ca_topic_score_gemma":0.15555023,"teacher_disagreement_score":0.6010237,"about_ca_system_score_codex":0.00030376323,"about_ca_system_score_gemma":0.00045038105,"threshold_uncertainty_score":0.99985075},"labels":[],"label_agreement":null},{"id":"W2945259153","doi":"10.1215/00031283-7587892","title":"Bag Across the Border","year":2019,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":31,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Ideology; Gentrification; Demographics; Sociocultural evolution; Distribution (mathematics); Geography; Interpersonal ties; Demographic economics; Sociology; Economic geography; Political science; Economic growth; Demography; Social science; Law; Anthropology; Economics","score_opus":0.012515321584955125,"score_gpt":0.36711628270793967,"score_spread":0.35460096112298456,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945259153","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5602239,0.000015250408,0.00012163907,0.011391415,0.0009950973,0.00017814008,0.000004545958,0.000069818074,0.4270002],"genre_scores_gemma":[0.97158617,0.000019925761,0.00064408715,0.0041763457,0.0003091775,0.0000038421886,0.0000013858129,0.0000050931408,0.02325397],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99932194,0.00009872932,0.00007244839,0.00011667198,0.00016508576,0.0002251478],"domain_scores_gemma":[0.99952096,0.00014101701,0.000060989914,0.00017512294,0.000059971728,0.000041928695],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00042732674,0.000040780797,0.00007950844,0.000008911416,0.00021283627,0.000035268145,0.00020985298,0.000021988011,0.002220559],"category_scores_gemma":[0.00032390328,0.000028698329,0.000028945504,0.00022826143,0.0003803549,0.000021726235,0.000033627715,0.00007439998,0.0013286622],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000034068595,0.00009678593,0.098165214,0.000003422201,0.0000571797,0.000020892208,0.0783501,0.000012052033,0.0005791193,0.33307952,0.027658109,0.46194354],"study_design_scores_gemma":[0.00008544929,0.000026788899,0.0110881245,7.3458176e-7,0.0000026547516,0.0000012182009,0.010972548,0.0000082384895,0.000026574182,0.00046671584,0.9772647,0.00005622396],"about_ca_topic_score_codex":0.013361284,"about_ca_topic_score_gemma":0.0014471369,"teacher_disagreement_score":0.9496066,"about_ca_system_score_codex":0.000025729756,"about_ca_system_score_gemma":0.000083371284,"threshold_uncertainty_score":0.9994489},"labels":[],"label_agreement":null},{"id":"W2945539495","doi":"10.1215/00031283-7603207","title":"Unlocking the Mystery of Dialect B","year":2019,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":32,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Raising (metalworking); Variety (cybernetics); Nike; Linguistics; History; Psychology; Computer science; Philosophy; Mathematics; Artificial intelligence","score_opus":0.014597910809429012,"score_gpt":0.30725463236367867,"score_spread":0.29265672155424965,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W2945539495","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.689127,0.0000141544315,0.0001448204,0.0018857031,0.0006091968,0.00013202577,0.0000015792982,0.000026228312,0.30805928],"genre_scores_gemma":[0.9960779,0.000011603925,0.00034126797,0.0007380288,0.0001874616,0.0000016692627,6.337251e-7,0.000003556148,0.0026378622],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994246,0.00013015888,0.00009233747,0.00008020973,0.00014787467,0.00012485209],"domain_scores_gemma":[0.9994495,0.00023116995,0.000107673455,0.00013508265,0.00005309672,0.000023476518],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00036525674,0.000032214364,0.00009708524,0.000023211593,0.00007307913,0.00000978313,0.00016087208,0.000015977957,0.0004746032],"category_scores_gemma":[0.000338393,0.000023217752,0.000031402415,0.00022515783,0.00027737606,0.000015018499,0.000021619855,0.000048813665,0.00013513947],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000055711043,0.000107007014,0.30417627,0.000015584897,0.00009946585,0.000015388903,0.060997732,0.000023754323,0.006435002,0.4079828,0.007780192,0.2123111],"study_design_scores_gemma":[0.00046941524,0.00029806775,0.054931093,0.0000227287,0.000041832587,0.00000645525,0.027771905,0.00007073356,0.0021097008,0.0036172385,0.91036254,0.00029829997],"about_ca_topic_score_codex":0.008403363,"about_ca_topic_score_gemma":0.00042260304,"teacher_disagreement_score":0.90258235,"about_ca_system_score_codex":0.00002086048,"about_ca_system_score_gemma":0.00009776959,"threshold_uncertainty_score":0.99819976},"labels":[],"label_agreement":null},{"id":"W3035976933","doi":"10.1215/00031283-8661833","title":"Interesting<i>Fellow</i>or Tough Old<i>Bird</i>?","year":2020,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":8,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Sociolinguistics; Variation (astronomy); Language change; Geography; Dialectology; Population; Sociology; Linguistics; Demography; History; Genealogy","score_opus":0.05470392004709201,"score_gpt":0.3453383602021966,"score_spread":0.2906344401551046,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3035976933","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.24631824,0.00004822486,0.0051030116,0.11026911,0.0025633327,0.00071002386,0.000023164213,0.0020231034,0.6329418],"genre_scores_gemma":[0.9592242,0.000031561303,0.0082737785,0.026174618,0.0012948598,0.000006315589,0.0000023120392,0.000022128766,0.004970186],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99885213,0.00014833108,0.00019207461,0.0002640897,0.000230347,0.00031303556],"domain_scores_gemma":[0.99918973,0.00020607341,0.00013290482,0.00013597158,0.000088979854,0.0002463325],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00021846018,0.00009562287,0.00019817363,0.000027464026,0.00018471788,0.00004747328,0.00030174133,0.00003911432,0.0015841691],"category_scores_gemma":[0.0028780361,0.00008554404,0.00005379633,0.0005057829,0.00044796095,0.000046098205,0.00006580113,0.00013801477,0.0008075554],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00055884785,0.00030897086,0.037899036,0.000037127695,0.00016403511,0.0010164503,0.17379804,0.000016647544,0.0031357883,0.17065942,0.26681346,0.34559217],"study_design_scores_gemma":[0.00028200296,0.00026270616,0.00074155704,0.0000067936644,0.00001968982,0.000005675863,0.0090245465,0.000066003704,0.00025556437,0.00047778705,0.98864466,0.00021303432],"about_ca_topic_score_codex":0.014737043,"about_ca_topic_score_gemma":0.0018111864,"teacher_disagreement_score":0.7218312,"about_ca_system_score_codex":0.000042646145,"about_ca_system_score_gemma":0.00020665415,"threshold_uncertainty_score":0.99997044},"labels":[],"label_agreement":null},{"id":"W3094887288","doi":"10.1215/00031283-8791763","title":"<i>Wait</i>, It’s a Discourse Marker","year":2020,"lang":"en","type":"article","venue":"American Speech","topic":"Language, Discourse, Communication Strategies","field":"Arts and Humanities","cited_by":5,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Linguistics; Verb; Variation (astronomy); Discourse marker; Meaning (existential); Function (biology); Psychology; Sociology; Philosophy","score_opus":0.04647680382892091,"score_gpt":0.301391317257425,"score_spread":0.25491451342850413,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3094887288","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.089315966,0.00019834243,0.00007512909,0.06247702,0.00016652897,0.00020342667,0.00007174007,0.00030244616,0.8471894],"genre_scores_gemma":[0.97789466,0.000062123334,0.00080629095,0.012663749,0.0006390649,0.000023997245,0.00003257944,0.0000385775,0.007838947],"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","domain_scores_codex":[0.99887174,0.00009476915,0.00023634071,0.00025766436,0.00025776573,0.0002817175],"domain_scores_gemma":[0.99903584,0.000080675556,0.00016000964,0.00050267036,0.00007601238,0.00014476955],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00009257601,0.00018446876,0.0002627457,0.000041633222,0.00020387172,0.00026517542,0.0005210612,0.0000168428,0.0060791415],"category_scores_gemma":[0.000043454453,0.00015347249,0.00010208542,0.0001147477,0.0011155946,0.00027388913,0.00013896267,0.00020653126,0.0010588926],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00010283674,0.0002293382,0.0012093586,0.000036560516,0.0002078691,0.00009033696,0.1409861,0.0000061028104,0.00042420998,0.33263624,0.43380976,0.09026128],"study_design_scores_gemma":[0.00024285713,0.00012845975,0.0003473239,0.000014972452,0.000042510685,0.0000054601433,0.17777102,0.000057516052,0.0001282983,0.0005823709,0.82034725,0.00033194476],"about_ca_topic_score_codex":0.0010720954,"about_ca_topic_score_gemma":0.0008300414,"teacher_disagreement_score":0.8885787,"about_ca_system_score_codex":0.000020279618,"about_ca_system_score_gemma":0.000070032365,"threshold_uncertainty_score":0.9997189},"labels":[],"label_agreement":null},{"id":"W3109451365","doi":"10.1215/00031283-8221002","title":"Diva Diction","year":2020,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":21,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"McGill University","funders":"","keywords":"Vowel; Mid vowel; Pronunciation; Period (music); Diction; Linguistics; American English; History; Formant; Geography; Art; Poetry; Philosophy","score_opus":0.031074270459744972,"score_gpt":0.32300944981076224,"score_spread":0.29193517935101726,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3109451365","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.09926824,0.00001603876,0.009844511,0.07768389,0.00091452186,0.0001790325,0.000006400048,0.00040804112,0.8116793],"genre_scores_gemma":[0.9901591,0.000019822248,0.0016658335,0.006688902,0.00075872766,0.0000019912202,0.000002434153,0.000003874715,0.00069933775],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.99951863,0.00006652039,0.000065288,0.00010448981,0.00012662944,0.000118444405],"domain_scores_gemma":[0.9997205,0.00003887622,0.00004758977,0.00004846823,0.000037100177,0.00010746219],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00009639396,0.000031464373,0.00006826086,0.000012971741,0.00012337141,0.000015606272,0.00008618074,0.000016595292,0.00069582806],"category_scores_gemma":[0.00063032587,0.00003189627,0.000021283196,0.00024251548,0.00018928366,0.000022329948,0.000015633816,0.000049938008,0.0003746251],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000059643895,0.000086054664,0.04531514,0.000004685065,0.00004506009,0.0000773117,0.0583092,0.000010310901,0.0019013368,0.36686102,0.08505814,0.4422721],"study_design_scores_gemma":[0.0000932047,0.00007533466,0.0041518644,6.6757553e-7,0.0000061188325,6.778776e-7,0.0035960362,0.00004277964,0.00010613643,0.0005226919,0.991334,0.00007045827],"about_ca_topic_score_codex":0.004867368,"about_ca_topic_score_gemma":0.00016926066,"teacher_disagreement_score":0.90627587,"about_ca_system_score_codex":0.000022168164,"about_ca_system_score_gemma":0.000057870137,"threshold_uncertainty_score":0.7618829},"labels":[],"label_agreement":null},{"id":"W3116226297","doi":"10.1215/00031283-8791781","title":"Dynamics of Short-<i>a</i>in Montreal and Quebec City English","year":2020,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Vowel; Linguistics; Artifact (error); Ethnic group; Psychology; Sociology; History; Geography; Anthropology","score_opus":0.01631406625981259,"score_gpt":0.2843554828187189,"score_spread":0.26804141655890634,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3116226297","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.91482925,0.00001062796,0.00028746857,0.0040495843,0.00011359153,0.00009530152,0.0000141779765,0.000033962137,0.08056602],"genre_scores_gemma":[0.9980243,0.000030188414,0.0008771073,0.00064416585,0.00015976345,0.0000014931942,0.00000450584,0.0000033747083,0.00025511227],"study_design_codex":"observational","study_design_gemma":"observational","domain_scores_codex":[0.9994256,0.0000766251,0.00012569419,0.00013059974,0.00011412441,0.00012737462],"domain_scores_gemma":[0.99964577,0.00008947091,0.000052658936,0.000055969267,0.00006613141,0.00009001919],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00015422833,0.00004450931,0.00015158772,0.00002456165,0.000034233777,0.000009868843,0.000087820095,0.000028835197,0.000098862314],"category_scores_gemma":[0.0008545978,0.00004696663,0.000018367564,0.00024415884,0.00036530732,0.000025716257,0.000027457327,0.00006830107,0.0000022946545],"study_design_candidate":"observational","study_design_consensus":"observational","about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0000507145,0.00007212518,0.75193864,0.000007952689,0.00001771245,0.000040137988,0.054984618,0.0000059037548,0.000087359345,0.051756117,0.00084581436,0.14019288],"study_design_scores_gemma":[0.0010602009,0.00044562877,0.86916894,0.000016859989,0.000057059675,0.0000014890434,0.077537395,0.0024085306,0.00024532998,0.0038824333,0.044586822,0.00058933126],"about_ca_topic_score_codex":0.33951253,"about_ca_topic_score_gemma":0.47970864,"teacher_disagreement_score":0.14019611,"about_ca_system_score_codex":0.000041336763,"about_ca_system_score_gemma":0.00008338753,"threshold_uncertainty_score":0.6648857},"labels":[],"label_agreement":null},{"id":"W3116595090","doi":"10.1215/00031283-8791772","title":"North Versus South","year":2020,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Ottawa","funders":"","keywords":"Spanish Civil War; Variation (astronomy); American English; Leverage (statistics); Linguistics; South carolina; History; Contextualization; Constraint (computer-aided design); Relation (database); Sociology; Political science; Statistics; Archaeology; Mathematics; Computer science; Philosophy; Interpretation (philosophy)","score_opus":0.05025621580387874,"score_gpt":0.3230608575838862,"score_spread":0.27280464178000746,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3116595090","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.48423398,0.000010015776,0.0010016423,0.021664448,0.0011037888,0.00015558171,0.000014755896,0.00028840956,0.49152738],"genre_scores_gemma":[0.99299556,0.000006966561,0.0019407516,0.0038919565,0.0007333913,0.0000018510817,0.000003751945,0.000005093694,0.00042069447],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.99942386,0.00006171414,0.00007233033,0.0001305091,0.00014957585,0.00016200419],"domain_scores_gemma":[0.9996174,0.00006189755,0.000059643822,0.00006712766,0.000045459437,0.00014850228],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00006860063,0.000042038573,0.00008873619,0.0000138352625,0.0001182199,0.000016505903,0.0001342774,0.000015762591,0.0005973725],"category_scores_gemma":[0.000928791,0.000042776377,0.000027463728,0.0003113573,0.00024368455,0.000017188475,0.000022033435,0.000062028565,0.0006033821],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00063458254,0.00011429313,0.15617345,0.000008521268,0.00015136336,0.00025653697,0.34515232,0.000023190421,0.00016912958,0.18143545,0.06040772,0.25547346],"study_design_scores_gemma":[0.0005252795,0.00023247048,0.007838412,6.687776e-7,0.000019503272,3.4583906e-7,0.017734224,0.000030726285,0.000041704254,0.000090563604,0.9733215,0.00016457742],"about_ca_topic_score_codex":0.002792727,"about_ca_topic_score_gemma":0.0013951164,"teacher_disagreement_score":0.9129138,"about_ca_system_score_codex":0.000021780359,"about_ca_system_score_gemma":0.000114814284,"threshold_uncertainty_score":0.7755458},"labels":[],"label_agreement":null},{"id":"W3143485133","doi":"10.1215/00031283-9116273","title":"Filipinos Front<i>Too</i>! A Sociophonetic Analysis of Toronto English /u/-Fronting","year":2021,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":7,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Tagalog; Linguistics; Vowel; History; Front (military); Ethnic group; Affect (linguistics); Psychology; Sociology; Geography; Anthropology; Philosophy","score_opus":0.011310578947850085,"score_gpt":0.30100702313233496,"score_spread":0.28969644418448487,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3143485133","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5764404,0.00065576425,0.0017160365,0.0011836688,0.0016793766,0.00018153976,0.00005373754,0.00015957163,0.41792992],"genre_scores_gemma":[0.98376054,0.00021920794,0.009868997,0.00081735465,0.0004516263,0.000006468102,0.000030069206,0.000009397636,0.004836358],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.99854785,0.00028040571,0.0002860268,0.00027714597,0.00032066062,0.0002879206],"domain_scores_gemma":[0.99864817,0.0002422385,0.00024321345,0.0002724361,0.00048854586,0.00010537936],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0004122972,0.000088370565,0.00041080586,0.000047449877,0.00016848181,0.000028233299,0.00019029404,0.000058138306,0.0052163065],"category_scores_gemma":[0.002371986,0.00009853079,0.0001919068,0.00056060584,0.000383004,0.000054812583,0.000054895256,0.00007525946,0.000020756068],"study_design_candidate":"qualitative","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000077743316,0.0008379439,0.21550626,0.000034648063,0.0042141248,0.000200671,0.46371332,0.00011411761,0.0071569793,0.038729265,0.049014825,0.22040011],"study_design_scores_gemma":[0.0009736694,0.00020267788,0.23917264,0.000021544533,0.002741839,0.0000025325119,0.22130063,0.0006479459,0.0023857236,0.0006366794,0.53104806,0.00086604885],"about_ca_topic_score_codex":0.08635316,"about_ca_topic_score_gemma":0.05552195,"teacher_disagreement_score":0.48203325,"about_ca_system_score_codex":0.00014829265,"about_ca_system_score_gemma":0.0003048271,"threshold_uncertainty_score":0.9956931},"labels":[],"label_agreement":null},{"id":"W3183133803","doi":"10.1215/00031283-9308384","title":"Regional Patterns in Prevelar Raising","year":2021,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Raising (metalworking); Geography; Distribution (mathematics); Physical geography; Climatology; Geology; Mathematics","score_opus":0.04081202377859353,"score_gpt":0.3486464287296638,"score_spread":0.30783440495107023,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W3183133803","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.8447335,0.00004841181,0.002096688,0.013790292,0.0005424548,0.00009662887,0.000005036321,0.000062962565,0.13862403],"genre_scores_gemma":[0.99220306,0.00006149576,0.0026187818,0.0025010796,0.00026466374,0.0000028876193,0.000006842154,0.0000047177637,0.002336501],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.99924314,0.00017548751,0.000104720544,0.00014739869,0.00015432024,0.00017492559],"domain_scores_gemma":[0.99963474,0.000101774756,0.000051986444,0.00009294615,0.000066968045,0.00005158604],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00024405358,0.000039345665,0.000097322794,0.000036967453,0.00007685649,0.00002023686,0.00008034992,0.000025707852,0.00073758344],"category_scores_gemma":[0.0005498387,0.000043831446,0.000024793713,0.00029963604,0.00013728083,0.000026071464,0.000022572869,0.0000749192,0.00004747677],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000022581355,0.00027300135,0.4991838,0.000008880501,0.000035026704,0.0011301923,0.034272958,0.000009247209,0.0013231828,0.17967695,0.008231703,0.27583247],"study_design_scores_gemma":[0.00048681584,0.000041922453,0.2112612,0.000030268993,0.0000107921005,0.000030845502,0.01444508,0.000016909496,0.00049684354,0.006514328,0.7663876,0.00027741384],"about_ca_topic_score_codex":0.010148964,"about_ca_topic_score_gemma":0.0064793383,"teacher_disagreement_score":0.7581559,"about_ca_system_score_codex":0.000059670852,"about_ca_system_score_gemma":0.00025756305,"threshold_uncertainty_score":0.99644256},"labels":[],"label_agreement":null},{"id":"W4220978711","doi":"10.1215/00031283-9766889","title":"Naturalistic Double Modals in North America","year":2022,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":9,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Modal verb; Sentence; Linguistics; Feature (linguistics); Modal; Naturalism; Natural language processing; Computer science; History; Artificial intelligence; Psychology; Geography; Verb","score_opus":0.02185511744994393,"score_gpt":0.31889253614783253,"score_spread":0.2970374186978886,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4220978711","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.78960365,0.000045929573,0.000101827856,0.0059959255,0.000996624,0.00031839727,0.000032820062,0.00013570205,0.20276915],"genre_scores_gemma":[0.9930787,0.0000187607,0.0010621949,0.00263556,0.0001543285,0.000045956192,0.00002572251,0.000008394713,0.0029704021],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.998792,0.00022476951,0.00015875092,0.000206377,0.0003181927,0.0002999193],"domain_scores_gemma":[0.99950147,0.00011227086,0.00012053377,0.00014574463,0.000042513762,0.00007748225],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.0002663972,0.00006733714,0.00016633248,0.000109664834,0.00034173264,0.00002002919,0.00027317787,0.000013720476,0.0016825566],"category_scores_gemma":[0.000244572,0.000074500145,0.000034775436,0.0011247528,0.00031251338,0.000030342875,0.00008720966,0.00020830691,0.00007993373],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00058713474,0.0009312227,0.26068437,0.000011915368,0.000093608636,0.0010718746,0.11996294,0.0030940163,0.00029891942,0.35042533,0.03849508,0.22434358],"study_design_scores_gemma":[0.0005237169,0.00013413554,0.017778808,8.598933e-7,0.000009304928,0.0000059206263,0.014938288,0.000119764416,0.000008334079,0.0008241054,0.9654328,0.00022396819],"about_ca_topic_score_codex":0.067324825,"about_ca_topic_score_gemma":0.010742957,"teacher_disagreement_score":0.9269377,"about_ca_system_score_codex":0.00021078781,"about_ca_system_score_gemma":0.00020287505,"threshold_uncertainty_score":0.99923},"labels":[],"label_agreement":null},{"id":"W4221038852","doi":"10.1215/00031283-9766922","title":"Second Dialect Acquisition “in Real Time”: Two Longitudinal Case Studies from YouTube","year":2022,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":2,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"Simon Fraser University","funders":"","keywords":"Salience (neuroscience); Tracking (education); History; Linguistics; Sociology; Psychology","score_opus":0.034290678188077406,"score_gpt":0.36002266291257606,"score_spread":0.32573198472449866,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4221038852","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.95821816,0.00004579948,0.0000129544505,0.00068406726,0.0005576659,0.0001375855,0.000086324326,0.000066136934,0.040191296],"genre_scores_gemma":[0.9961082,0.000033133136,0.0008234734,0.00055531313,0.000542144,0.0000317979,0.000031894888,0.000009222862,0.0018648746],"study_design_codex":"qualitative","study_design_gemma":"qualitative","domain_scores_codex":[0.9985157,0.00051329145,0.00018616102,0.0002770881,0.00024468062,0.00026303506],"domain_scores_gemma":[0.99925214,0.00034637708,0.00012802902,0.0001400383,0.00006632135,0.00006706604],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00060561363,0.00008530235,0.00022721576,0.00009192092,0.0005183796,0.000022181537,0.00013422804,0.000016927617,0.0068843737],"category_scores_gemma":[0.00027194046,0.000093628005,0.000044400877,0.00044063875,0.00033009722,0.000046437945,0.000104945146,0.00014475475,0.00009502614],"study_design_candidate":"qualitative","study_design_consensus":"qualitative","about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.0008586205,0.0013032312,0.3026113,0.00001835687,0.0008512692,0.067564614,0.38274857,0.0002856746,0.006585711,0.051614873,0.065923944,0.119633846],"study_design_scores_gemma":[0.0091562085,0.0031014243,0.12625208,0.000038553106,0.00055323465,0.0026033318,0.5664128,0.0007837305,0.0008697256,0.037126373,0.24981152,0.0032909997],"about_ca_topic_score_codex":0.07076747,"about_ca_topic_score_gemma":0.017147673,"teacher_disagreement_score":0.18388757,"about_ca_system_score_codex":0.00032565894,"about_ca_system_score_gemma":0.00014254538,"threshold_uncertainty_score":0.99402344},"labels":[],"label_agreement":null},{"id":"W4230998457","doi":"10.1215/00031283-7549976","title":"PRESIDENT’S MESSAGE","year":2019,"lang":"en","type":"article","venue":"American Speech","topic":"European Linguistics and Anthropology","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":false,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":false,"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science","score_opus":0.010783134484513405,"score_gpt":0.3240432566644507,"score_spread":0.3132601221799373,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4230998457","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.28590596,0.00003080229,0.000021837435,0.0009927963,0.00042477535,0.000070487135,0.0000017921299,0.00005514879,0.7124964],"genre_scores_gemma":[0.9662476,0.00011734472,0.0007414058,0.000546659,0.0003550326,8.053263e-7,0.000001087942,0.000010177704,0.03197987],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.99926144,0.00010493072,0.00008383542,0.00015215373,0.00016867551,0.0002289544],"domain_scores_gemma":[0.99959713,0.00007280538,0.00006398719,0.00015496711,0.000043031225,0.000068078356],"candidate_categories":["insufficient_payload"],"consensus_categories":["insufficient_payload"],"category_scores_codex":[0.00021616001,0.000051191397,0.00011081759,0.000028668741,0.00012545606,0.000029923202,0.00019958965,0.000015433927,0.0020715906],"category_scores_gemma":[0.0001613758,0.00004959528,0.00003294944,0.00016321297,0.0006477347,0.000027144526,0.00004007774,0.000070989525,0.0017267705],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000019122115,0.00011445826,0.16764581,0.000009904323,0.00004321609,0.0001391093,0.0076369816,0.0000041806898,0.0010454676,0.6864479,0.038002767,0.09889107],"study_design_scores_gemma":[0.00009489249,0.00008571942,0.008722732,0.0000055925834,0.0000062227105,0.000001218949,0.006230278,0.000005960306,0.00017776559,0.0013046953,0.9832303,0.000134617],"about_ca_topic_score_codex":0.009045152,"about_ca_topic_score_gemma":0.0005906617,"teacher_disagreement_score":0.94522756,"about_ca_system_score_codex":0.000022389402,"about_ca_system_score_gemma":0.00006154797,"threshold_uncertainty_score":0.9990505},"labels":[],"label_agreement":null},{"id":"W4286640707","doi":"10.1215/00031283-9940676","title":"Where Have All the Articles Gone? The Use of Zero Articles in Marmora and Lake, Ontario","year":2022,"lang":"en","type":"article","venue":"American Speech","topic":"Comparative and International Law Studies","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Zero (linguistics); History; Environmental science; Geography; Linguistics; Philosophy","score_opus":0.08384024049867037,"score_gpt":0.3180754849340773,"score_spread":0.23423524443540691,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4286640707","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9853584,0.00014287973,0.000006446597,0.009611592,0.000034369794,0.00013120219,0.000009367637,0.000006650533,0.0046990938],"genre_scores_gemma":[0.9957854,0.00009019138,0.00011740573,0.00073343335,0.000018348162,0.000023217615,7.646815e-7,0.0000030346778,0.0032281824],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991309,0.00022146378,0.00012739675,0.00010580757,0.00027601316,0.00013842357],"domain_scores_gemma":[0.9993722,0.00038147587,0.000082676816,0.000102418744,0.000040700128,0.000020569501],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00030605853,0.000052723553,0.000107815606,0.000022168959,0.00032658986,0.000041303076,0.00017328716,0.000004934371,0.0003973186],"category_scores_gemma":[0.00005690869,0.000033302327,0.000029375455,0.00019471269,0.00078016944,0.00008124419,0.000158399,0.00011754748,0.0000035851995],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000086269494,0.00015402236,0.7677187,0.000001794507,0.00009389803,0.000010471906,0.11159635,0.00012124267,0.00026425614,0.0350788,0.027640918,0.0572333],"study_design_scores_gemma":[0.00008445191,0.00006910686,0.29428953,0.0000032706957,0.000008717033,0.0000013200244,0.023078036,0.000045256744,0.000074840245,0.0016320336,0.6806581,0.000055351884],"about_ca_topic_score_codex":0.30313054,"about_ca_topic_score_gemma":0.9124613,"teacher_disagreement_score":0.65301716,"about_ca_system_score_codex":0.00006843921,"about_ca_system_score_gemma":0.00005295011,"threshold_uncertainty_score":0.70150995},"labels":[],"label_agreement":null},{"id":"W4296471838","doi":"10.1215/00031283-10104915","title":"Orderly Obsolescence: The Decline of /hw/ in Ontario","year":2022,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":3,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"PARRY; Obsolescence; Parallels; Scots; Vernacular; History; Immigration; Linguistics; Irish; Sound (geography); Sociology; Philosophy; Archaeology; Acoustics","score_opus":0.0228515118966771,"score_gpt":0.30946335748742826,"score_spread":0.28661184559075115,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4296471838","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.89650553,0.0000157057,0.000089751404,0.009905147,0.00054340804,0.000206394,0.0000041323706,0.000023542729,0.09270637],"genre_scores_gemma":[0.9940736,0.000005480696,0.00089301064,0.0019341845,0.000046526304,0.0000120995055,0.000002761118,0.0000033176755,0.0030289732],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9991029,0.0002298797,0.00015049164,0.00010462922,0.00026079273,0.00015128897],"domain_scores_gemma":[0.99954724,0.00013708702,0.00012085954,0.00012412184,0.00004369329,0.000026974778],"candidate_categories":["insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00067684223,0.000037849324,0.00011110051,0.000041025658,0.000194106,0.000007110582,0.00028413988,0.0000111347945,0.001949598],"category_scores_gemma":[0.00033168512,0.00003192667,0.000026432977,0.00044234045,0.00033865927,0.000012265987,0.00008850332,0.0001763785,0.000008520504],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00013166896,0.0004388245,0.600682,0.0000024769731,0.0000292403,0.0000737065,0.13302085,0.00009932761,0.00055206334,0.17403065,0.0072418004,0.08369739],"study_design_scores_gemma":[0.00043527674,0.00024608255,0.17779773,0.0000026103937,0.0000118089665,0.0000047836684,0.038589258,0.000034024684,0.000047362853,0.0052525974,0.7774406,0.00013788158],"about_ca_topic_score_codex":0.79315823,"about_ca_topic_score_gemma":0.7627458,"teacher_disagreement_score":0.77019876,"about_ca_system_score_codex":0.00016683283,"about_ca_system_score_gemma":0.00058268476,"threshold_uncertainty_score":0.99896276},"labels":[],"label_agreement":null},{"id":"W4312615804","doi":"10.1215/00031283-10096048","title":"It’s a Guy Thing","year":2022,"lang":"ru","type":"article","venue":"American Speech","topic":"Lexicography and Language Studies","field":"Arts and Humanities","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Alberta","funders":"","keywords":"Icon; Citation; History; History of the book; Art; Art history; Classics; Library science; Computer science","score_opus":0.01840523172098487,"score_gpt":0.250509379100885,"score_spread":0.23210414737990012,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4312615804","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.1418393,0.0052122357,0.000009449403,0.0115137715,0.0020574403,0.00036131774,0.00044680233,0.00023841017,0.83832127],"genre_scores_gemma":[0.95513564,0.0002937515,0.0001666665,0.013466499,0.001231106,0.00007115999,0.00002553336,0.00004899345,0.02956064],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.9979169,0.00021204283,0.00032775584,0.0004378951,0.0005313139,0.00057410146],"domain_scores_gemma":[0.9990302,0.00014068895,0.0002758812,0.00040505343,0.00005990632,0.00008828923],"candidate_categories":["metaepi_narrow","sts","insufficient_payload"],"consensus_categories":[],"category_scores_codex":[0.00031284717,0.00029207653,0.0004233386,0.00021201486,0.0025155363,0.00023562128,0.00042315797,0.00001659157,0.026548455],"category_scores_gemma":[0.000029606363,0.00028360315,0.0003131441,0.00033631476,0.0010760366,0.0001306155,0.0005366466,0.000601547,0.0002012999],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00009709438,0.00041356497,0.0013709088,0.000054726348,0.00063751824,0.0004752189,0.557084,0.000012687966,0.000036580564,0.03985959,0.28875783,0.11120026],"study_design_scores_gemma":[0.00018706659,0.00043310047,0.00018493614,0.000015486465,0.00007763736,0.000019678591,0.2845015,0.000013230092,0.000016081274,0.00025800048,0.71398365,0.00030963833],"about_ca_topic_score_codex":0.008793259,"about_ca_topic_score_gemma":0.0010170742,"teacher_disagreement_score":0.8132964,"about_ca_system_score_codex":0.000088326255,"about_ca_system_score_gemma":0.00006502472,"threshold_uncertainty_score":0.9999616},"labels":[],"label_agreement":null},{"id":"W4378189922","doi":"10.1215/00031283-10579455","title":"The Influence of English on Neologisms for Nonbinary Gender Identities and Sexual Orientations in Quebec French: Between Variation and Purism","year":2023,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistics, Language Diversity, and Identity","field":"Arts and Humanities","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"Université TÉLUQ","funders":"","keywords":"Neologism; Queer; Variety (cybernetics); Identity (music); Gender studies; Sociology; Variation (astronomy); Institution; Linguistics; Power (physics); History; Political science; Art; Aesthetics; Philosophy; Social science","score_opus":0.02769368386513524,"score_gpt":0.26160477262480214,"score_spread":0.2339110887596669,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378189922","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.99838394,0.000022502853,0.000011775941,0.000029775461,0.0007197864,0.00013561791,0.000106882006,0.000027330592,0.0005623861],"genre_scores_gemma":[0.9967914,0.00003762315,0.00006148053,0.00005299777,0.0012055399,0.000008387605,0.000028926934,0.000004942331,0.001808692],"study_design_codex":"qualitative","study_design_gemma":"observational","domain_scores_codex":[0.9994955,0.000029354575,0.0001307836,0.00012553045,0.000101293146,0.00011752773],"domain_scores_gemma":[0.9990954,0.0003575727,0.00009000833,0.00008686112,0.00034903252,0.000021119575],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00015116559,0.0000621461,0.0001114746,0.000095118856,0.0002214683,0.0000973867,0.00007525154,0.000016145494,0.000009449787],"category_scores_gemma":[0.001818637,0.000051895586,0.000014190043,0.00006719781,0.00052051124,0.00012308854,0.00005997172,0.000054624863,0.0000030505491],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00002804702,0.000049581344,0.31368414,0.00008641248,0.000088407614,0.000009014596,0.63186014,0.00004649791,0.000011697657,0.03880418,0.008820002,0.006511876],"study_design_scores_gemma":[0.00029640965,0.00019127777,0.86483943,0.000008984276,0.000043836433,5.7836832e-8,0.123277284,0.000033900087,0.000020759142,0.009099557,0.0020663452,0.00012215694],"about_ca_topic_score_codex":0.0780299,"about_ca_topic_score_gemma":0.027703844,"teacher_disagreement_score":0.55115527,"about_ca_system_score_codex":0.000017394515,"about_ca_system_score_gemma":0.000020202782,"threshold_uncertainty_score":0.99003804},"labels":[],"label_agreement":null},{"id":"W4378190398","doi":"10.1215/00031283-10579429","title":"Cultures and Complexities Concerning Place","year":2023,"lang":"en","type":"article","venue":"American Speech","topic":"Migration, Ethnicity, and Economy","field":"Social Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Epistemology; Sociology; Linguistics; Philosophy","score_opus":0.04406284468593775,"score_gpt":0.34030854368441377,"score_spread":0.296245698998476,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4378190398","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.9045573,0.00011875562,0.000018614513,0.0019264544,0.0001122443,0.00007795277,0.0000056452636,0.00024464817,0.09293842],"genre_scores_gemma":[0.99061817,0.00048089487,0.00036655267,0.00041766858,0.00033604665,0.0000063427033,0.0000074035524,0.000005692557,0.007761211],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994339,0.00006490186,0.000078751,0.00012578804,0.000108920176,0.0001877797],"domain_scores_gemma":[0.9996743,0.00011984692,0.00005412991,0.00006073813,0.000026592203,0.00006439177],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00023973444,0.000051967225,0.00010607267,0.00003485236,0.0003886338,0.0000630924,0.00008703026,0.000018169907,0.00013965591],"category_scores_gemma":[0.000077620425,0.000050739327,0.00001920359,0.00024903525,0.0006902042,0.00009918739,0.000026748898,0.000052334526,0.00013061555],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000016993014,0.000020840518,0.14126264,0.000019328942,0.000060410923,0.000013895727,0.412277,0.000048916838,0.00025699948,0.06399456,0.31674773,0.065280676],"study_design_scores_gemma":[0.00013397603,0.000048624563,0.051985573,0.000009972979,0.0000098681585,0.0000011981499,0.24326402,0.00019560503,0.00013363341,0.0028224527,0.7011756,0.00021953507],"about_ca_topic_score_codex":0.022106934,"about_ca_topic_score_gemma":0.014179563,"teacher_disagreement_score":0.38442782,"about_ca_system_score_codex":0.000023999055,"about_ca_system_score_gemma":0.000043543336,"threshold_uncertainty_score":0.9844049},"labels":[],"label_agreement":null},{"id":"W4404790465","doi":"10.1215/00031283-11466470","title":"Veteran Vowels: Early Western Canadian English in World War Oral Histories","year":2024,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":1,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"McGill University","funders":"","keywords":"History; Oral history; Linguistics; Geography; Archaeology; Philosophy","score_opus":0.02021930526267341,"score_gpt":0.30875526327894615,"score_spread":0.28853595801627274,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W4404790465","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.5344286,0.0003213879,0.000054824643,0.010395375,0.008260078,0.0003076714,0.000040188064,0.00039103755,0.4458008],"genre_scores_gemma":[0.98484534,0.000019189525,0.00055541395,0.001207422,0.0007186981,0.000009391753,0.0000057856555,0.000013714802,0.012625027],"study_design_codex":"qualitative","study_design_gemma":"not_applicable","domain_scores_codex":[0.99894536,0.00013102379,0.00015379311,0.00022552989,0.0001714798,0.0003727894],"domain_scores_gemma":[0.9994678,0.00010356638,0.00003127687,0.00012072596,0.00007143937,0.00020519622],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00036071418,0.000083320185,0.0001556139,0.00030287064,0.00011813465,0.00008286347,0.0001824876,0.000038349506,0.00046925026],"category_scores_gemma":[0.00039849817,0.000088480214,0.000035599263,0.0007891759,0.00032055555,0.000088579516,0.00001541603,0.0001695952,0.00015215995],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000027283826,0.000052242754,0.28334314,0.000016574555,0.000054615928,0.0015066017,0.43859795,0.0000035940327,0.00002868885,0.08780634,0.036801755,0.15176122],"study_design_scores_gemma":[0.00006557929,0.000036640806,0.011396572,0.0000104857245,0.000006804567,0.0000010590174,0.0030733424,0.000005631332,0.000004898328,0.00035989482,0.9849173,0.000121822704],"about_ca_topic_score_codex":0.9083636,"about_ca_topic_score_gemma":0.96184593,"teacher_disagreement_score":0.9481155,"about_ca_system_score_codex":0.00053376,"about_ca_system_score_gemma":0.0006666632,"threshold_uncertainty_score":0.5137961},"labels":[],"label_agreement":null},{"id":"W7106006481","doi":"10.1215/00031283-11466518","title":"Globalization, Localization, and the Preservation of French in North America (including Creole Relevancies)","year":2025,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Globalization; French; Immigration; Creole language; Order (exchange)","score_opus":0.014163237455987066,"score_gpt":0.30386821265679426,"score_spread":0.2897049752008072,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7106006481","genre_codex":"empirical","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":"empirical","domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.40506503,0.0006247094,0.2847615,0.03338742,0.0010223496,0.0016570424,0.000034295666,0.0001811437,0.27326652],"genre_scores_gemma":[0.9958558,0.00028741968,0.001269948,0.001674715,0.000046540637,0.000009902658,0.00001191381,0.000003327163,0.0008404006],"study_design_codex":"observational","study_design_gemma":"not_applicable","domain_scores_codex":[0.9990346,0.00026596757,0.00025002137,0.0001424505,0.00017206011,0.0001349384],"domain_scores_gemma":[0.9991279,0.00029462657,0.00018187064,0.00013426683,0.0002364653,0.000024835685],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00044527857,0.000055470176,0.00016651554,0.000086930144,0.00017811466,0.000021732525,0.00015453123,0.00002620436,0.000043289143],"category_scores_gemma":[0.0031169765,0.00004622837,0.00001826934,0.0016282246,0.0009663657,0.000063713,0.000054021246,0.00005081653,0.0000025697157],"study_design_candidate":"observational","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":true,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.00005154945,0.00006182173,0.6309521,0.000018682296,0.000028509527,0.0000017542848,0.02080956,0.0007895962,0.000017605855,0.32262558,0.0063567557,0.018286444],"study_design_scores_gemma":[0.0021130843,0.00007430018,0.2627685,0.000053084514,0.000065303066,7.160521e-7,0.013406335,0.010147399,0.00008817559,0.024542453,0.6864752,0.00026542714],"about_ca_topic_score_codex":0.0428369,"about_ca_topic_score_gemma":0.01816977,"teacher_disagreement_score":0.68011844,"about_ca_system_score_codex":0.00005867451,"about_ca_system_score_gemma":0.00018234196,"threshold_uncertainty_score":0.9997461},"labels":[],"label_agreement":null},{"id":"W7106020807","doi":"10.1215/00031283-12250394","title":"Editors’ Notes","year":2025,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":false,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"","funders":"","keywords":"Centennial; Theme (computing); Miscellany; American studies; Style (visual arts); Microform","score_opus":0.014742262375275553,"score_gpt":0.34265305012923414,"score_spread":0.3279107877539586,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7106020807","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.021453504,0.000017394215,0.0030269972,0.025398323,0.0074825087,0.00011309087,0.0000029899584,0.00016650393,0.9423387],"genre_scores_gemma":[0.98163813,0.00001762318,0.0032997932,0.0031449301,0.0030835764,0.00000474561,0.0000018449352,0.0000029892524,0.0088063935],"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"not_applicable","domain_scores_codex":[0.9994907,0.00006967306,0.00007578329,0.00010694941,0.00010763916,0.00014928466],"domain_scores_gemma":[0.9995324,0.00022102626,0.000039620478,0.000095914904,0.00007015096,0.000040856805],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00020916296,0.000036559388,0.00008273054,0.00005170499,0.00016289143,0.000020607647,0.00012370401,0.000026056417,0.00036839923],"category_scores_gemma":[0.001551587,0.00003646393,0.000024348032,0.0003547201,0.00033014527,0.000016647453,0.000021383728,0.000052648764,0.00014166172],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000012594538,0.00006650157,0.02546764,0.000002676146,0.000030919917,0.00001624695,0.0041005732,0.0000011398351,0.00029621663,0.52914715,0.2323603,0.20849803],"study_design_scores_gemma":[0.000062479485,0.000014458923,0.003744882,0.0000022414606,0.0000065404843,1.8157834e-7,0.0011913177,0.000003249118,0.00017647722,0.0032698691,0.9914809,0.00004741479],"about_ca_topic_score_codex":0.01406752,"about_ca_topic_score_gemma":0.0007870554,"teacher_disagreement_score":0.9601846,"about_ca_system_score_codex":0.000041007206,"about_ca_system_score_gemma":0.00018662916,"threshold_uncertainty_score":0.99249786},"labels":[],"label_agreement":null},{"id":"W7106022641","doi":"10.1215/00031283-12250406","title":"<i>American Speech</i> : The Next 100 Years","year":2025,"lang":"en","type":"article","venue":"American Speech","topic":"Linguistic Variation and Morphology","field":"Social Sciences","cited_by":0,"is_retracted":false,"has_abstract":true,"route_ca_aff":true,"route_ca_fund":false,"route_ca_venue":false,"route_about_ca":true,"ca_institutions":"University of Toronto","funders":"","keywords":"Feeling; Speech community; American English; English language; Plain language","score_opus":0.022253641920094044,"score_gpt":0.33468920614445635,"score_spread":0.3124355642243623,"validation_status":"score_only:v0-immature-baseline","prediction":{"id":"W7106022641","genre_codex":"other","genre_gemma":"empirical","domain_codex":null,"domain_gemma":null,"model_version":"codex-gemma-dda1882f352a","genre_candidate":"empirical","genre_consensus":null,"domain_candidate":null,"domain_consensus":null,"prediction_status":"machine_predicted_unvalidated","genre_scores_codex":[0.14674172,0.0000661187,0.00066296896,0.033307888,0.0015282073,0.00035410406,0.000008836443,0.0002350562,0.8170951],"genre_scores_gemma":[0.9638951,0.00014577825,0.0032277193,0.015044425,0.00049681537,0.000014787554,0.0000039121824,0.000011618885,0.017159844],"study_design_codex":"design_other","study_design_gemma":"not_applicable","domain_scores_codex":[0.9986458,0.0002689373,0.00018599762,0.00025283537,0.00027554692,0.0003709223],"domain_scores_gemma":[0.9989869,0.00034100312,0.00015288328,0.0003345425,0.000097111515,0.00008755391],"candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00056036975,0.000099778146,0.0002106346,0.00008308494,0.00034161835,0.00009155047,0.00048642175,0.000033988177,0.0004550162],"category_scores_gemma":[0.0008162189,0.00008315688,0.000076214674,0.0012396099,0.0016109641,0.000041607436,0.000080833684,0.00018021057,0.0003208823],"study_design_candidate":"not_applicable","study_design_consensus":null,"about_ca_topic_candidate":true,"about_ca_topic_consensus":false,"about_ca_system_candidate":false,"about_ca_system_consensus":false,"study_design_scores_codex":[0.000024154366,0.00007008064,0.006953853,0.0000019549946,0.000056378667,0.00004820725,0.00478764,0.0000022117636,0.00022292926,0.14012247,0.0748277,0.7728824],"study_design_scores_gemma":[0.000121849924,0.000050309172,0.010299221,0.0000045141373,0.000025378818,0.0000025325173,0.00816592,0.000009973997,0.000085857726,0.0023925109,0.9787213,0.00012062474],"about_ca_topic_score_codex":0.04948828,"about_ca_topic_score_gemma":0.0041647777,"teacher_disagreement_score":0.9038936,"about_ca_system_score_codex":0.00009861555,"about_ca_system_score_gemma":0.00038535998,"threshold_uncertainty_score":0.95684123},"labels":[],"label_agreement":null}]}