{"id":"W1597484020","doi":"10.82308/15144","title":"The conceptualization and development of a high-stakes video listening test within an AUA framework in a military context","year":2012,"lang":"en","type":"article","venue":"eScholarship@McGill (McGill)","topic":"Hearing Impairment and Communication","field":"Psychology","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"McGill University","funders":"","keywords":"Conceptualization; Context (archaeology); Active listening; Test (biology); Psychology; Computer science; Communication; History; Artificial intelligence","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":true,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.043706,0.0008586649,0.0006378581,0.007999402,0.006460223,0.01567633,0.005633964,0.008761458,0.002318579],"category_scores_gemma":[0.05505238,0.0008038576,0.0008643104,0.003480889,0.05092347,0.0208321,0.007932067,0.007565675,0.0005098328],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.01549476,"about_ca_system_score_gemma":0.01529511,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01219518,"about_ca_topic_score_gemma":0.009518377,"domain_scores_codex":[0.934689,0.05202457,0.002528296,0.002884452,0.006170418,0.001703179],"domain_scores_gemma":[0.932,0.04899542,0.004630727,0.003187038,0.007793867,0.00339282],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"observational","study_design_scores_codex":[0.0000342248,0.0003403458,0.005540332,0.000434133,0.00001119597,0.001137214,0.16188,0.0008500624,0.001347555,0.785772,0.0009859534,0.04166706],"study_design_scores_gemma":[0.00009718885,0.001002269,0.01569145,0.005049512,0.00006045543,0.003467323,0.3342846,0.01907326,0.002969774,0.4700625,0.1479461,0.000295644],"study_design_candidate":"observational","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.2755622,0.004416762,0.4164584,0.0558835,0.0006907328,0.004642009,0.0002164159,0.0003770399,0.2417529],"genre_scores_gemma":[0.8458674,0.000722779,0.1456355,0.002131785,0.00005984589,0.002703045,0.00007240898,0.00006733764,0.002739925],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.043706,"threshold_uncertainty_score":0.2311422,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04383227410290567,"score_gpt":0.2932648193070564,"score_spread":0.2494325452041507,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}