{"id":"W4416157662","doi":"10.48550/arxiv.2511.06497","title":"Rethinking what Matters: Effective and Robust Multilingual Realignment for Low-Resource Languages","year":2025,"lang":"","type":"preprint","venue":"ArXiv.org","topic":"ICT in Developing Communities","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"Natural Sciences and Engineering Research Council of Canada; Alliance de recherche numérique du Canada; Grand Équipement National De Calcul Intensif; Compute Canada","keywords":"Word (group theory); Transfer (computing); Empirical research; Multilingualism; Language model; Word identification; Data collection","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005588632,0.001970272,0.001356074,0.001502996,0.001533903,0.002886067,0.002256203,0.00107948,0.005510727],"category_scores_gemma":[0.02906005,0.0009545535,0.001036613,0.001943769,0.001630327,0.007230925,0.005138861,0.002215779,0.007366847],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005254614,"about_ca_system_score_gemma":0.002330651,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003890855,"about_ca_topic_score_gemma":0.005958298,"domain_scores_codex":[0.994397,0.003005482,0.0003733591,0.001392049,0.0005188759,0.0003131775],"domain_scores_gemma":[0.9894053,0.004387138,0.0006059345,0.003570407,0.001654958,0.000376274],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001995721,0.0006804802,0.02588865,0.001322697,0.0003187522,0.000726499,0.006145367,0.03032459,0.1054718,0.006954427,0.01970362,0.8004674],"study_design_scores_gemma":[0.0007606097,0.001785016,0.03168373,0.0006393582,0.0006821119,0.002019873,0.01847718,0.4087258,0.2994922,0.08178157,0.1531506,0.0008019739],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.5543759,0.001617823,0.4100187,0.00171154,0.0004388273,0.0004503046,0.003887242,0.01880318,0.008696391],"genre_scores_gemma":[0.7063149,0.0004846108,0.2732578,0.0006882804,0.000108744,0.0005371277,0.009988673,0.005121394,0.003498366],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.005588632,"threshold_uncertainty_score":0.02955586,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04868336307078377,"score_gpt":0.3027608409348006,"score_spread":0.2540774778640169,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}