{"id":"W2579409202","doi":"10.21700/ijcis.2016.119","title":"Automatic Diacritics Restoration for Dialectal Arabic Text","year":2016,"lang":"en","type":"article","venue":"International Journal of Computing and Information Sciences","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":false,"ca_institutions":"","funders":"","keywords":"Arabic; Linguistics; Natural language processing; Artificial intelligence; Computer science; Philosophy","routes":{"ca_aff":false,"ca_fund":false,"ca_venue":true,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.000994292,0.00005888324,0.00008524315,0.0003106866,0.0001258302,0.0004905954,0.0006901675,0.0000264003,0.000002053219],"category_scores_gemma":[0.0008293909,0.00003555465,0.00003773719,0.0001485055,0.00009431999,0.005569918,0.00008408923,0.00004679311,0.00000204297],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00004275952,"about_ca_system_score_gemma":0.0001158555,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.000001766442,"about_ca_topic_score_gemma":2.975851e-7,"domain_scores_codex":[0.9989414,0.00002410738,0.0004182168,0.0000637337,0.0004550481,0.00009753952],"domain_scores_gemma":[0.9983656,0.000351234,0.0004551684,0.0000516285,0.0007380468,0.00003826842],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.000005417393,0.000009316777,0.0003336108,0.0000097103,0.000009760938,9.531036e-7,0.001004986,0.00002733773,0.0007500576,0.1054738,0.0005658644,0.8918092],"study_design_scores_gemma":[0.002254259,0.00123014,0.005071513,0.001353657,0.00001775513,0.0010562,0.0003529746,0.7217699,0.0282083,0.2233342,0.01477015,0.0005809678],"study_design_candidate":"design_other","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.07052427,0.000131952,0.9244246,0.003961939,0.0007161943,0.00004806824,0.000001333377,0.00005109988,0.0001405351],"genre_scores_gemma":[0.7630242,0.00001866469,0.2365407,0.0003021823,0.0001085453,6.306345e-7,3.14086e-7,9.494397e-7,0.000003850434],"genre_candidate":"methods","genre_consensus":null,"teacher_disagreement_score":0.8912282,"threshold_uncertainty_score":0.4730822,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.01376920173908248,"score_gpt":0.3163899634614012,"score_spread":0.3026207617223188,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}