{"id":"W4417273487","doi":"10.2139/ssrn.5835542","title":"Event-Driven Alpha from Large Language Models: A Theoretical Framework for Contrastive Financial Representation Learning","year":2025,"lang":"","type":"preprint","venue":"SSRN Electronic Journal","topic":"Stock Market Forecasting Methods","field":"Decision Sciences","cited_by":0,"is_retracted":false,"has_abstract":false,"ca_institutions":"Canadian Imperial Bank of Commerce (Canada)","funders":"","keywords":"Interpretation (philosophy); Representation (politics); Financial market; Focus (optics); Alpha (finance); Event (particle physics); Modern portfolio theory; Downside risk","routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false,"invisible_to_affiliation_only":false},"retraction":null,"screen":null,"direct_labels":[],"prediction":{"model_version":"codex-gemma-dda1882f352a","candidate_categories":["metaresearch","metaepi_narrow","sts","scholarly_communication","research_integrity"],"consensus_categories":["metaresearch","research_integrity"],"category_scores_codex":[0.03462706,0.001167233,0.002185397,0.0009397474,0.001866016,0.001121289,0.002943061,0.001775128,0.0008988627],"category_scores_gemma":[0.1004301,0.001070939,0.001976863,0.001303453,0.0005396162,0.0006313641,0.001558232,0.01961043,0.00004263978],"about_ca_system_candidate":true,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002839686,"about_ca_system_score_gemma":0.01733335,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0001572173,"about_ca_topic_score_gemma":0.0003483141,"domain_scores_codex":[0.9776148,0.006776671,0.003245084,0.002725923,0.002863786,0.006773746],"domain_scores_gemma":[0.9635592,0.02984653,0.003109246,0.001275683,0.001743927,0.0004654061],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.003083307,0.0002371372,0.001095008,0.00002132203,0.000777347,0.00001982317,0.005698027,0.01940827,0.00004963988,0.7160385,0.0001452687,0.2534263],"study_design_scores_gemma":[0.002098623,0.0005620281,0.0003247061,0.0008182665,0.0004978125,0.00006953433,0.01227286,0.2910293,0.00008431514,0.6912472,0.000343772,0.0006515976],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.04952333,0.006377889,0.935098,0.00110775,0.003990622,0.001875538,0.0004807245,0.00009480134,0.00145136],"genre_scores_gemma":[0.9029507,0.003282601,0.08657025,0.0002608578,0.003106578,0.0002493604,0.0001136443,0.0001113604,0.003354638],"genre_candidate":"methods","genre_consensus":null,"teacher_disagreement_score":0.8534274,"threshold_uncertainty_score":0.9999157,"prediction_status":"machine_predicted_unvalidated"},"machine_scores":{"provisional":true,"baseline":true,"maturity_gate_passed":false,"score_opus":0.04087777380982899,"score_gpt":0.4059147191940084,"score_spread":0.3650369453841794,"validation_status":"score_only:v0-immature-baseline","note":"Baseline scores from an immature model (maturity gate not passed). Scores rank; they never assert a category."}}