{"id":"W26623053","doi":"10.3892/mco.2015.580","title":"Generating Coherent Extracts of Single Documents Using Latent Semantic Analysis","year":2003,"lang":"en","type":"article","venue":"Molecular and Clinical Oncology","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"ca_institutions":"","funders":"University of Toronto; University of Pennsylvania","keywords":"Latent semantic analysis; Computer science; Natural language processing; Coherence (philosophical gambling strategy); Artificial intelligence; Information retrieval; Similarity (geometry); Identification (biology); Semantic similarity; Semantics (computer science); Probabilistic latent semantic analysis; Vector space model; Topic model; Statistics; Mathematics; Programming language","routes":{"ca_aff":false,"ca_fund":true,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":true},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0007323169,0.0001012519,0.000404108,0.00009612935,0.00004717805,0.00004316695,0.0002375984,0.0001537269,0.00001318732],"category_scores_gemma":[0.0003142348,0.00008496641,0.0001539752,0.0003832289,0.00007819179,0.00009639321,0.000137,0.0001628572,8.183551e-7],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00004006658,"about_ca_system_score_gemma":0.00008110629,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00003629999,"about_ca_topic_score_gemma":0.0000156612,"domain_scores_codex":[0.9984458,0.0003595922,0.0005308704,0.0003283172,0.0001547347,0.0001806534],"domain_scores_gemma":[0.9991204,0.0001610774,0.0002718437,0.0002683456,0.00008880336,0.0000895609],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"bench_or_experimental","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.00003829281,0.002415109,0.04999947,0.0001002494,0.001856812,0.0008842329,0.0003889501,0.001387324,0.6210775,0.03775078,0.00008383671,0.2840174],"study_design_scores_gemma":[0.00518574,0.006395265,0.006686507,0.0002660621,0.004391502,0.0003275307,0.00008912936,0.4349619,0.4726268,0.06423373,0.00280388,0.002031918],"study_design_candidate":"bench_or_experimental","study_design_consensus":"bench_or_experimental","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.4761303,0.001211944,0.5222098,0.00008023566,0.00008382711,0.00006613375,2.511301e-7,0.00003423393,0.0001832585],"genre_scores_gemma":[0.6626456,0.00001268326,0.3370157,0.0003007077,0.000007221882,0.000001425038,8.636046e-7,0.000003585561,0.00001218288],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.4335746,"threshold_uncertainty_score":0.346483,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04939897924259892,"score_gpt":0.3915064841222118,"score_spread":0.3421075048796129,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}