Compare commits
5
Commits
| Author | SHA1 | Date | |
|---|---|---|---|
|
|
84d20216f1 | ||
|
|
cdfd3105aa | ||
|
|
2b3ff82e9c | ||
|
|
bae78d4fbb | ||
|
|
1b5f4d1e76 |
+1
-1
@@ -1,5 +1,5 @@
|
|||||||
# TradeAC custom-qlib-code snapshot (auto-generated)
|
# TradeAC custom-qlib-code snapshot (auto-generated)
|
||||||
# parent repo HEAD : 22a57a81c62aa2d99f9362a50b52fb4c0d82d1c2
|
# parent repo HEAD : cdfd3105aabf69c817cd9acee24072096124dcc2
|
||||||
# tac-qlib/tac_qlib/contrib
|
# tac-qlib/tac_qlib/contrib
|
||||||
# tac-qlib/tac_qlib/data
|
# tac-qlib/tac_qlib/data
|
||||||
# per-file hashes (git hash-object):
|
# per-file hashes (git hash-object):
|
||||||
|
|||||||
@@ -1,6 +1,6 @@
|
|||||||
# QUEUE-me
|
# QUEUE-ly
|
||||||
# HMM regime overlay (Q10): identical TopkDropout selection, entry gated on sp_hmm_p_regime1 >= 0.5.
|
# Fractional-Kelly sizing (Q06): same selection, sizing ∝ score magnitude, capped at 0.5× equal-weight notional.
|
||||||
# Acceptance: net_max_drawdown < 7.69% AND net_IR >= 0.21. No regime column enters feature_fields.
|
# Acceptance: net_IR > 0.21 AND net_ann > +2.13% with total_cost not higher than reference.
|
||||||
|
|
||||||
|
|
||||||
|
|
||||||
@@ -49,7 +49,7 @@ qlib_init:
|
|||||||
module_path: qlib.workflow.expm
|
module_path: qlib.workflow.expm
|
||||||
kwargs:
|
kwargs:
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
||||||
default_exp_name: "tac-rd-q10-regime"
|
default_exp_name: "tac-rd-q06-kelly"
|
||||||
|
|
||||||
task:
|
task:
|
||||||
model:
|
model:
|
||||||
@@ -120,15 +120,15 @@ task:
|
|||||||
kwargs:
|
kwargs:
|
||||||
config:
|
config:
|
||||||
strategy:
|
strategy:
|
||||||
class: RegimeGateDropoutStrategy
|
class: FractionalKellyDropoutStrategy
|
||||||
module_path: tac_qlib.contrib.strategy.regime_gate
|
module_path: tac_qlib.contrib.strategy.kelly_dropout
|
||||||
kwargs:
|
kwargs:
|
||||||
signal: "<PRED>"
|
signal: "<PRED>"
|
||||||
topk: 10
|
topk: 10
|
||||||
n_drop: 1
|
n_drop: 1
|
||||||
only_tradable: true
|
only_tradable: true
|
||||||
risk_degree: 0.95
|
risk_degree: 0.95
|
||||||
regime_threshold: 0.5
|
cap_frac: 0.5
|
||||||
backtest:
|
backtest:
|
||||||
start_time: 2026-01-04
|
start_time: 2026-01-04
|
||||||
end_time: 2026-08-10
|
end_time: 2026-08-10
|
||||||
+8
-9
@@ -1,13 +1,14 @@
|
|||||||
# -----------------------------------------------------------------------------
|
# -----------------------------------------------------------------------------
|
||||||
# EXP 20 - R1: 2-seed ensemble (seeds 42,7), TopkDropout baseline.
|
# EXP 15 - Strategy A (baseline): parallel reference model + TopkDropout.
|
||||||
#
|
#
|
||||||
# Runtime cut: 2 seeds instead of 5. Everything else identical to the reference
|
# Model = RankICEnsembleLGBModel with PARALLEL=5 (thread-pool seed training,
|
||||||
# (test 2026-01-04..2026-08-10, SPY, costs 5bp/15bp). Measures whether the
|
# rank_ensemble.py) - the speedup means the 5-seed 3000-round ensemble trains in
|
||||||
# 2-seed ensemble keeps the reference quality at ~2/5 the training time.
|
# ~1/5th the wall time of the serial reference. Strategy = TopkDropout topk=10
|
||||||
|
# n_drop=2 risk_degree=0.95 (the reference's recorded strategy).
|
||||||
#
|
#
|
||||||
# Run:
|
# Run:
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp20-risk-limit-improve/r1_2seed.yaml \
|
# rd_run_workflow config_path=experiments/workflows/exp15-kelly-size/a_topk_baseline.yaml \
|
||||||
# experiment_name=tac-rd-risk-limit
|
# experiment_name=tac-rd-kelly-size
|
||||||
# -----------------------------------------------------------------------------
|
# -----------------------------------------------------------------------------
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
{%- set LAKE = TAC_LAKE_DIR %}
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
||||||
@@ -41,7 +42,7 @@ qlib_init:
|
|||||||
module_path: qlib.workflow.expm
|
module_path: qlib.workflow.expm
|
||||||
kwargs:
|
kwargs:
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
default_exp_name: "tac-rd-kelly-size"
|
||||||
|
|
||||||
task:
|
task:
|
||||||
model:
|
model:
|
||||||
@@ -62,8 +63,6 @@ task:
|
|||||||
reg_alpha: 0.1
|
reg_alpha: 0.1
|
||||||
reg_lambda: 1.0
|
reg_lambda: 1.0
|
||||||
seeds: "42,7,2026,99,123"
|
seeds: "42,7,2026,99,123"
|
||||||
weight_mode: rolling_ic
|
|
||||||
rolling_ic_window: 21
|
|
||||||
parallel: 5
|
parallel: 5
|
||||||
|
|
||||||
dataset:
|
dataset:
|
||||||
+20
-12
@@ -1,13 +1,18 @@
|
|||||||
# -----------------------------------------------------------------------------
|
# -----------------------------------------------------------------------------
|
||||||
# EXP 20 - R1: 2-seed ensemble (seeds 42,7), TopkDropout baseline.
|
# EXP 15 - Strategy B: parallel reference model + KellyWeightStrategy.
|
||||||
#
|
#
|
||||||
# Runtime cut: 2 seeds instead of 5. Everything else identical to the reference
|
# Model = RankICEnsembleLGBModel PARALLEL=5 (same as A). Strategy =
|
||||||
# (test 2026-01-04..2026-08-10, SPY, costs 5bp/15bp). Measures whether the
|
# KellyWeightStrategy (tac_qlib.contrib.strategy.kelly_weight): applies the
|
||||||
# 2-seed ensemble keeps the reference quality at ~2/5 the training time.
|
# Kelly criterion to position SIZING -
|
||||||
|
# f* = kelly_fraction * mu_i / var_i (Gaussian Kelly, mu_i > 0)
|
||||||
|
# with per-name mu/var estimated from the rolling signal history (no
|
||||||
|
# lookahead), floored/capped and normalized to the risk_degree leverage budget.
|
||||||
|
# This is the piece the reference's equal-weight TopkDropout never tunes: it
|
||||||
|
# weights names by edge/risk instead of equal-weight top-k.
|
||||||
#
|
#
|
||||||
# Run:
|
# Run:
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp20-risk-limit-improve/r1_2seed.yaml \
|
# rd_run_workflow config_path=experiments/workflows/exp15-kelly-size/b_kelly_weight.yaml \
|
||||||
# experiment_name=tac-rd-risk-limit
|
# experiment_name=tac-rd-kelly-size
|
||||||
# -----------------------------------------------------------------------------
|
# -----------------------------------------------------------------------------
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
{%- set LAKE = TAC_LAKE_DIR %}
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
||||||
@@ -41,7 +46,7 @@ qlib_init:
|
|||||||
module_path: qlib.workflow.expm
|
module_path: qlib.workflow.expm
|
||||||
kwargs:
|
kwargs:
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
default_exp_name: "tac-rd-kelly-size"
|
||||||
|
|
||||||
task:
|
task:
|
||||||
model:
|
model:
|
||||||
@@ -112,15 +117,18 @@ task:
|
|||||||
kwargs:
|
kwargs:
|
||||||
config:
|
config:
|
||||||
strategy:
|
strategy:
|
||||||
class: MomentumGateTopk
|
class: KellyWeightStrategy
|
||||||
module_path: tac_qlib.contrib.strategy.momentum_gate
|
module_path: tac_qlib.contrib.strategy.kelly_weight
|
||||||
kwargs:
|
kwargs:
|
||||||
signal: "<PRED>"
|
signal: "<PRED>"
|
||||||
topk: 10
|
topk: 10
|
||||||
n_drop: 2
|
lookback: 20
|
||||||
min_momentum: 0.0
|
min_obs: 10
|
||||||
only_tradable: true
|
kelly_fraction: 0.5
|
||||||
|
max_weight: 0.15
|
||||||
|
min_weight: 0.0
|
||||||
risk_degree: 0.95
|
risk_degree: 0.95
|
||||||
|
max_turnover: 0.30
|
||||||
backtest:
|
backtest:
|
||||||
start_time: 2026-01-04
|
start_time: 2026-01-04
|
||||||
end_time: 2026-08-10
|
end_time: 2026-08-10
|
||||||
@@ -1,141 +0,0 @@
|
|||||||
# -----------------------------------------------------------------------------
|
|
||||||
# EXP 18 - Risk-limit control: reference model + TopkDropout baseline (A).
|
|
||||||
#
|
|
||||||
# Signal/model identical to the reference (tac-rd-rank-ensemble-isolated,
|
|
||||||
# run 0cea66d9...): RankICEnsembleLGBModel (parallel, 5 seeds) on the 50-ETF
|
|
||||||
# SP-5d panel, test 2026-01-04..2026-08-10. This workflow reproduces the
|
|
||||||
# unconstrained TopkDropout baseline net-of-cost so the risk-limited variant
|
|
||||||
# (same pred, liquidity/size/concentration caps) can be compared 1:1.
|
|
||||||
#
|
|
||||||
# The risk_limits spec itself is applied via rd_backtest / rd_strategy_targets
|
|
||||||
# (tool-level param, not a YAML key); this run records the unconstrained
|
|
||||||
# baseline that the limit A/B is measured against.
|
|
||||||
#
|
|
||||||
# Run:
|
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp18-risk-limit/a_baseline.yaml \
|
|
||||||
# experiment_name=tac-rd-risk-limit
|
|
||||||
# -----------------------------------------------------------------------------
|
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
|
||||||
{%- set SP_FIELDS = "sp_ret,sp_jump_ratio,sp_jump_flag,sp_jump_tail,sp_max_move,sp_rv1,sp_rv5,sp_rv22,sp_vol_ratio_5_22,sp_vol_ratio_1_22,sp_trend_slope_5,sp_trend_slope_20,sp_trend_slope_60,sp_logp,sp_hurst_exponent,sp_sig_level1_lead,sp_sig_level1_lag,sp_sig_level2_lead_lag,sp_sig_level2_lag_lead" %}
|
|
||||||
|
|
||||||
qlib_init:
|
|
||||||
provider_uri: "{{ LAKE }}"
|
|
||||||
region: us
|
|
||||||
expression_cache: null
|
|
||||||
dataset_cache: null
|
|
||||||
|
|
||||||
calendar_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeCalendarProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
instrument_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeInstrumentProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
markets: {}
|
|
||||||
feature_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeFeatureProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
|
|
||||||
exp_manager:
|
|
||||||
class: MLflowExpManager
|
|
||||||
module_path: qlib.workflow.expm
|
|
||||||
kwargs:
|
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
|
||||||
|
|
||||||
task:
|
|
||||||
model:
|
|
||||||
class: RankICEnsembleLGBModel
|
|
||||||
module_path: tac_qlib.contrib.model.rank_ensemble
|
|
||||||
kwargs:
|
|
||||||
loss: mse
|
|
||||||
learning_rate: 0.02
|
|
||||||
num_leaves: 31
|
|
||||||
n_estimators: 3000
|
|
||||||
num_boost_round: 3000
|
|
||||||
early_stopping_rounds: 200
|
|
||||||
min_data_in_leaf: 20
|
|
||||||
lambda_l2: 0.5
|
|
||||||
colsample_bytree: 0.8
|
|
||||||
subsample: 0.8
|
|
||||||
subsample_freq: 1
|
|
||||||
reg_alpha: 0.1
|
|
||||||
reg_lambda: 1.0
|
|
||||||
seeds: "42,7,2026,99,123"
|
|
||||||
parallel: 5
|
|
||||||
|
|
||||||
dataset:
|
|
||||||
class: DatasetH
|
|
||||||
module_path: qlib.data.dataset
|
|
||||||
kwargs:
|
|
||||||
handler:
|
|
||||||
class: TACHandler
|
|
||||||
module_path: tac_qlib.contrib.data.handler
|
|
||||||
kwargs:
|
|
||||||
instruments: "{{ UNIVERSE }}"
|
|
||||||
start_time: 2015-01-03
|
|
||||||
end_time: 2026-08-14
|
|
||||||
fit_start_time: 2016-01-04
|
|
||||||
fit_end_time: 2025-09-01
|
|
||||||
freq: day
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
label: "Ref($close,-6)/Ref($close,-1)-1"
|
|
||||||
feature_fields: "$open,$high,$low,$close,$vwap,$volume,{{ SP_FIELDS }}"
|
|
||||||
infer_processors:
|
|
||||||
- class: DropAllNaN
|
|
||||||
kwargs: {}
|
|
||||||
- class: ProcessInf
|
|
||||||
kwargs: {}
|
|
||||||
- class: CSRankNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: ZScoreNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: Fillna
|
|
||||||
kwargs: {}
|
|
||||||
segments:
|
|
||||||
train: [2016-01-04, 2025-09-01]
|
|
||||||
valid: [2025-09-03, 2026-01-03]
|
|
||||||
test: [2026-01-04, 2026-08-10]
|
|
||||||
|
|
||||||
record:
|
|
||||||
- class: SignalRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs: {}
|
|
||||||
- class: SigAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
ana_long_short: true
|
|
||||||
ann_scaler: 252
|
|
||||||
- class: PortAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
config:
|
|
||||||
strategy:
|
|
||||||
class: TopkDropoutStrategy
|
|
||||||
module_path: qlib.contrib.strategy
|
|
||||||
kwargs:
|
|
||||||
signal: "<PRED>"
|
|
||||||
topk: 10
|
|
||||||
n_drop: 2
|
|
||||||
only_tradable: true
|
|
||||||
risk_degree: 0.95
|
|
||||||
backtest:
|
|
||||||
start_time: 2026-01-04
|
|
||||||
end_time: 2026-08-10
|
|
||||||
account: 1000000
|
|
||||||
benchmark: SPY
|
|
||||||
exchange_kwargs:
|
|
||||||
codes: "{{ UNIVERSE }}"
|
|
||||||
deal_price: $close
|
|
||||||
freq: day
|
|
||||||
open_cost: 0.0005
|
|
||||||
close_cost: 0.0015
|
|
||||||
min_cost: 5.0
|
|
||||||
risk_analysis_freq: 1d
|
|
||||||
@@ -1,135 +0,0 @@
|
|||||||
# -----------------------------------------------------------------------------
|
|
||||||
# EXP 20 - R1: 2-seed ensemble (seeds 42,7), TopkDropout baseline.
|
|
||||||
#
|
|
||||||
# Runtime cut: 2 seeds instead of 5. Everything else identical to the reference
|
|
||||||
# (test 2026-01-04..2026-08-10, SPY, costs 5bp/15bp). Measures whether the
|
|
||||||
# 2-seed ensemble keeps the reference quality at ~2/5 the training time.
|
|
||||||
#
|
|
||||||
# Run:
|
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp20-risk-limit-improve/r1_2seed.yaml \
|
|
||||||
# experiment_name=tac-rd-risk-limit
|
|
||||||
# -----------------------------------------------------------------------------
|
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
|
||||||
{%- set SP_FIELDS = "sp_ret,sp_jump_ratio,sp_jump_flag,sp_jump_tail,sp_max_move,sp_rv1,sp_rv5,sp_rv22,sp_vol_ratio_5_22,sp_vol_ratio_1_22,sp_trend_slope_5,sp_trend_slope_20,sp_trend_slope_60,sp_logp,sp_hurst_exponent,sp_sig_level1_lead,sp_sig_level1_lag,sp_sig_level2_lead_lag,sp_sig_level2_lag_lead" %}
|
|
||||||
|
|
||||||
qlib_init:
|
|
||||||
provider_uri: "{{ LAKE }}"
|
|
||||||
region: us
|
|
||||||
expression_cache: null
|
|
||||||
dataset_cache: null
|
|
||||||
|
|
||||||
calendar_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeCalendarProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
instrument_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeInstrumentProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
markets: {}
|
|
||||||
feature_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeFeatureProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
|
|
||||||
exp_manager:
|
|
||||||
class: MLflowExpManager
|
|
||||||
module_path: qlib.workflow.expm
|
|
||||||
kwargs:
|
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
|
||||||
|
|
||||||
task:
|
|
||||||
model:
|
|
||||||
class: RankICEnsembleLGBModel
|
|
||||||
module_path: tac_qlib.contrib.model.rank_ensemble
|
|
||||||
kwargs:
|
|
||||||
loss: mse
|
|
||||||
learning_rate: 0.02
|
|
||||||
num_leaves: 31
|
|
||||||
n_estimators: 3000
|
|
||||||
num_boost_round: 3000
|
|
||||||
early_stopping_rounds: 200
|
|
||||||
min_data_in_leaf: 20
|
|
||||||
lambda_l2: 0.5
|
|
||||||
colsample_bytree: 0.8
|
|
||||||
subsample: 0.8
|
|
||||||
subsample_freq: 1
|
|
||||||
reg_alpha: 0.1
|
|
||||||
reg_lambda: 1.0
|
|
||||||
seeds: "42,7,2026,99,123"
|
|
||||||
parallel: 5
|
|
||||||
|
|
||||||
dataset:
|
|
||||||
class: DatasetH
|
|
||||||
module_path: qlib.data.dataset
|
|
||||||
kwargs:
|
|
||||||
handler:
|
|
||||||
class: TACHandler
|
|
||||||
module_path: tac_qlib.contrib.data.handler
|
|
||||||
kwargs:
|
|
||||||
instruments: "{{ UNIVERSE }}"
|
|
||||||
start_time: 2015-01-03
|
|
||||||
end_time: 2026-08-14
|
|
||||||
fit_start_time: 2016-01-04
|
|
||||||
fit_end_time: 2025-09-01
|
|
||||||
freq: day
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
label: "Ref($close,-6)/Ref($close,-1)-1"
|
|
||||||
feature_fields: "$open,$high,$low,$close,$vwap,$volume,{{ SP_FIELDS }}"
|
|
||||||
infer_processors:
|
|
||||||
- class: DropAllNaN
|
|
||||||
kwargs: {}
|
|
||||||
- class: ProcessInf
|
|
||||||
kwargs: {}
|
|
||||||
- class: CSRankNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: ZScoreNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: Fillna
|
|
||||||
kwargs: {}
|
|
||||||
segments:
|
|
||||||
train: [2016-01-04, 2025-09-01]
|
|
||||||
valid: [2025-09-03, 2026-01-03]
|
|
||||||
test: [2026-01-04, 2026-08-10]
|
|
||||||
|
|
||||||
record:
|
|
||||||
- class: SignalRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs: {}
|
|
||||||
- class: SigAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
ana_long_short: true
|
|
||||||
ann_scaler: 252
|
|
||||||
- class: PortAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
config:
|
|
||||||
strategy:
|
|
||||||
class: TopkDropoutStrategy
|
|
||||||
module_path: qlib.contrib.strategy
|
|
||||||
kwargs:
|
|
||||||
signal: "<PRED>"
|
|
||||||
topk: 10
|
|
||||||
n_drop: 2
|
|
||||||
only_tradable: true
|
|
||||||
risk_degree: 0.95
|
|
||||||
backtest:
|
|
||||||
start_time: 2026-01-04
|
|
||||||
end_time: 2026-08-10
|
|
||||||
account: 1000000
|
|
||||||
benchmark: SPY
|
|
||||||
exchange_kwargs:
|
|
||||||
codes: "{{ UNIVERSE }}"
|
|
||||||
deal_price: $close
|
|
||||||
freq: day
|
|
||||||
open_cost: 0.0005
|
|
||||||
close_cost: 0.0015
|
|
||||||
min_cost: 5.0
|
|
||||||
risk_analysis_freq: 1d
|
|
||||||
@@ -1,135 +0,0 @@
|
|||||||
# -----------------------------------------------------------------------------
|
|
||||||
# EXP 20 - R1: 2-seed ensemble (seeds 42,7), TopkDropout baseline.
|
|
||||||
#
|
|
||||||
# Runtime cut: 2 seeds instead of 5. Everything else identical to the reference
|
|
||||||
# (test 2026-01-04..2026-08-10, SPY, costs 5bp/15bp). Measures whether the
|
|
||||||
# 2-seed ensemble keeps the reference quality at ~2/5 the training time.
|
|
||||||
#
|
|
||||||
# Run:
|
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp20-risk-limit-improve/r1_2seed.yaml \
|
|
||||||
# experiment_name=tac-rd-risk-limit
|
|
||||||
# -----------------------------------------------------------------------------
|
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
|
||||||
{%- set SP_FIELDS = "sp_ret,sp_jump_ratio,sp_jump_flag,sp_jump_tail,sp_max_move,sp_rv1,sp_rv5,sp_rv22,sp_vol_ratio_5_22,sp_vol_ratio_1_22,sp_trend_slope_5,sp_trend_slope_20,sp_trend_slope_60,sp_logp,sp_hurst_exponent,sp_sig_level1_lead,sp_sig_level1_lag,sp_sig_level2_lead_lag,sp_sig_level2_lag_lead" %}
|
|
||||||
|
|
||||||
qlib_init:
|
|
||||||
provider_uri: "{{ LAKE }}"
|
|
||||||
region: us
|
|
||||||
expression_cache: null
|
|
||||||
dataset_cache: null
|
|
||||||
|
|
||||||
calendar_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeCalendarProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
instrument_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeInstrumentProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
markets: {}
|
|
||||||
feature_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeFeatureProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
|
|
||||||
exp_manager:
|
|
||||||
class: MLflowExpManager
|
|
||||||
module_path: qlib.workflow.expm
|
|
||||||
kwargs:
|
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
|
||||||
|
|
||||||
task:
|
|
||||||
model:
|
|
||||||
class: RankICEnsembleLGBModel
|
|
||||||
module_path: tac_qlib.contrib.model.rank_ensemble
|
|
||||||
kwargs:
|
|
||||||
loss: mse
|
|
||||||
learning_rate: 0.02
|
|
||||||
num_leaves: 31
|
|
||||||
n_estimators: 3000
|
|
||||||
num_boost_round: 3000
|
|
||||||
early_stopping_rounds: 200
|
|
||||||
min_data_in_leaf: 20
|
|
||||||
lambda_l2: 0.5
|
|
||||||
colsample_bytree: 0.8
|
|
||||||
subsample: 0.8
|
|
||||||
subsample_freq: 1
|
|
||||||
reg_alpha: 0.1
|
|
||||||
reg_lambda: 1.0
|
|
||||||
seeds: "42,7,2026,99,123"
|
|
||||||
parallel: 1
|
|
||||||
|
|
||||||
dataset:
|
|
||||||
class: DatasetH
|
|
||||||
module_path: qlib.data.dataset
|
|
||||||
kwargs:
|
|
||||||
handler:
|
|
||||||
class: TACHandler
|
|
||||||
module_path: tac_qlib.contrib.data.handler
|
|
||||||
kwargs:
|
|
||||||
instruments: "{{ UNIVERSE }}"
|
|
||||||
start_time: 2015-01-03
|
|
||||||
end_time: 2026-08-14
|
|
||||||
fit_start_time: 2016-01-04
|
|
||||||
fit_end_time: 2025-09-01
|
|
||||||
freq: day
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
label: "Ref($close,-6)/Ref($close,-1)-1"
|
|
||||||
feature_fields: "$open,$high,$low,$close,$vwap,$volume,{{ SP_FIELDS }}"
|
|
||||||
infer_processors:
|
|
||||||
- class: DropAllNaN
|
|
||||||
kwargs: {}
|
|
||||||
- class: ProcessInf
|
|
||||||
kwargs: {}
|
|
||||||
- class: CSRankNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: ZScoreNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: Fillna
|
|
||||||
kwargs: {}
|
|
||||||
segments:
|
|
||||||
train: [2016-01-04, 2025-09-01]
|
|
||||||
valid: [2025-09-03, 2026-01-03]
|
|
||||||
test: [2026-01-04, 2026-08-10]
|
|
||||||
|
|
||||||
record:
|
|
||||||
- class: SignalRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs: {}
|
|
||||||
- class: SigAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
ana_long_short: true
|
|
||||||
ann_scaler: 252
|
|
||||||
- class: PortAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
config:
|
|
||||||
strategy:
|
|
||||||
class: TopkDropoutStrategy
|
|
||||||
module_path: qlib.contrib.strategy
|
|
||||||
kwargs:
|
|
||||||
signal: "<PRED>"
|
|
||||||
topk: 10
|
|
||||||
n_drop: 2
|
|
||||||
only_tradable: true
|
|
||||||
risk_degree: 0.95
|
|
||||||
backtest:
|
|
||||||
start_time: 2026-01-04
|
|
||||||
end_time: 2026-08-10
|
|
||||||
account: 1000000
|
|
||||||
benchmark: SPY
|
|
||||||
exchange_kwargs:
|
|
||||||
codes: "{{ UNIVERSE }}"
|
|
||||||
deal_price: $close
|
|
||||||
freq: day
|
|
||||||
open_cost: 0.0005
|
|
||||||
close_cost: 0.0015
|
|
||||||
min_cost: 5.0
|
|
||||||
risk_analysis_freq: 1d
|
|
||||||
@@ -1,135 +0,0 @@
|
|||||||
# -----------------------------------------------------------------------------
|
|
||||||
# EXP 20 - R1: 2-seed ensemble (seeds 42,7), TopkDropout baseline.
|
|
||||||
#
|
|
||||||
# Runtime cut: 2 seeds instead of 5. Everything else identical to the reference
|
|
||||||
# (test 2026-01-04..2026-08-10, SPY, costs 5bp/15bp). Measures whether the
|
|
||||||
# 2-seed ensemble keeps the reference quality at ~2/5 the training time.
|
|
||||||
#
|
|
||||||
# Run:
|
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp20-risk-limit-improve/r1_2seed.yaml \
|
|
||||||
# experiment_name=tac-rd-risk-limit
|
|
||||||
# -----------------------------------------------------------------------------
|
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
|
||||||
{%- set SP_FIELDS = "sp_ret,sp_jump_ratio,sp_jump_flag,sp_jump_tail,sp_max_move,sp_rv1,sp_rv5,sp_rv22,sp_vol_ratio_5_22,sp_vol_ratio_1_22,sp_trend_slope_5,sp_trend_slope_20,sp_trend_slope_60,sp_logp,sp_hurst_exponent,sp_sig_level1_lead,sp_sig_level1_lag,sp_sig_level2_lead_lag,sp_sig_level2_lag_lead" %}
|
|
||||||
|
|
||||||
qlib_init:
|
|
||||||
provider_uri: "{{ LAKE }}"
|
|
||||||
region: us
|
|
||||||
expression_cache: null
|
|
||||||
dataset_cache: null
|
|
||||||
|
|
||||||
calendar_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeCalendarProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
instrument_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeInstrumentProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
markets: {}
|
|
||||||
feature_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeFeatureProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
|
|
||||||
exp_manager:
|
|
||||||
class: MLflowExpManager
|
|
||||||
module_path: qlib.workflow.expm
|
|
||||||
kwargs:
|
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
|
||||||
|
|
||||||
task:
|
|
||||||
model:
|
|
||||||
class: RankICEnsembleLGBModel
|
|
||||||
module_path: tac_qlib.contrib.model.rank_ensemble
|
|
||||||
kwargs:
|
|
||||||
loss: mse
|
|
||||||
learning_rate: 0.02
|
|
||||||
num_leaves: 31
|
|
||||||
n_estimators: 3000
|
|
||||||
num_boost_round: 3000
|
|
||||||
early_stopping_rounds: 200
|
|
||||||
min_data_in_leaf: 20
|
|
||||||
lambda_l2: 0.5
|
|
||||||
colsample_bytree: 0.8
|
|
||||||
subsample: 0.8
|
|
||||||
subsample_freq: 1
|
|
||||||
reg_alpha: 0.1
|
|
||||||
reg_lambda: 1.0
|
|
||||||
seeds: "42,7"
|
|
||||||
parallel: 2
|
|
||||||
|
|
||||||
dataset:
|
|
||||||
class: DatasetH
|
|
||||||
module_path: qlib.data.dataset
|
|
||||||
kwargs:
|
|
||||||
handler:
|
|
||||||
class: TACHandler
|
|
||||||
module_path: tac_qlib.contrib.data.handler
|
|
||||||
kwargs:
|
|
||||||
instruments: "{{ UNIVERSE }}"
|
|
||||||
start_time: 2015-01-03
|
|
||||||
end_time: 2026-08-14
|
|
||||||
fit_start_time: 2016-01-04
|
|
||||||
fit_end_time: 2025-09-01
|
|
||||||
freq: day
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
label: "Ref($close,-6)/Ref($close,-1)-1"
|
|
||||||
feature_fields: "$open,$high,$low,$close,$vwap,$volume,{{ SP_FIELDS }}"
|
|
||||||
infer_processors:
|
|
||||||
- class: DropAllNaN
|
|
||||||
kwargs: {}
|
|
||||||
- class: ProcessInf
|
|
||||||
kwargs: {}
|
|
||||||
- class: CSRankNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: ZScoreNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: Fillna
|
|
||||||
kwargs: {}
|
|
||||||
segments:
|
|
||||||
train: [2016-01-04, 2025-09-01]
|
|
||||||
valid: [2025-09-03, 2026-01-03]
|
|
||||||
test: [2026-01-04, 2026-08-10]
|
|
||||||
|
|
||||||
record:
|
|
||||||
- class: SignalRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs: {}
|
|
||||||
- class: SigAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
ana_long_short: true
|
|
||||||
ann_scaler: 252
|
|
||||||
- class: PortAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
config:
|
|
||||||
strategy:
|
|
||||||
class: TopkDropoutStrategy
|
|
||||||
module_path: qlib.contrib.strategy
|
|
||||||
kwargs:
|
|
||||||
signal: "<PRED>"
|
|
||||||
topk: 10
|
|
||||||
n_drop: 2
|
|
||||||
only_tradable: true
|
|
||||||
risk_degree: 0.95
|
|
||||||
backtest:
|
|
||||||
start_time: 2026-01-04
|
|
||||||
end_time: 2026-08-10
|
|
||||||
account: 1000000
|
|
||||||
benchmark: SPY
|
|
||||||
exchange_kwargs:
|
|
||||||
codes: "{{ UNIVERSE }}"
|
|
||||||
deal_price: $close
|
|
||||||
freq: day
|
|
||||||
open_cost: 0.0005
|
|
||||||
close_cost: 0.0015
|
|
||||||
min_cost: 5.0
|
|
||||||
risk_analysis_freq: 1d
|
|
||||||
@@ -1,138 +0,0 @@
|
|||||||
# -----------------------------------------------------------------------------
|
|
||||||
# EXP 20 - R1: 2-seed ensemble (seeds 42,7), TopkDropout baseline.
|
|
||||||
#
|
|
||||||
# Runtime cut: 2 seeds instead of 5. Everything else identical to the reference
|
|
||||||
# (test 2026-01-04..2026-08-10, SPY, costs 5bp/15bp). Measures whether the
|
|
||||||
# 2-seed ensemble keeps the reference quality at ~2/5 the training time.
|
|
||||||
#
|
|
||||||
# Run:
|
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp20-risk-limit-improve/r1_2seed.yaml \
|
|
||||||
# experiment_name=tac-rd-risk-limit
|
|
||||||
# -----------------------------------------------------------------------------
|
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
|
||||||
{%- set SP_FIELDS = "sp_ret,sp_jump_ratio,sp_jump_flag,sp_jump_tail,sp_max_move,sp_rv1,sp_rv5,sp_rv22,sp_vol_ratio_5_22,sp_vol_ratio_1_22,sp_trend_slope_5,sp_trend_slope_20,sp_trend_slope_60,sp_logp,sp_hurst_exponent,sp_sig_level1_lead,sp_sig_level1_lag,sp_sig_level2_lead_lag,sp_sig_level2_lag_lead" %}
|
|
||||||
|
|
||||||
qlib_init:
|
|
||||||
provider_uri: "{{ LAKE }}"
|
|
||||||
region: us
|
|
||||||
expression_cache: null
|
|
||||||
dataset_cache: null
|
|
||||||
|
|
||||||
calendar_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeCalendarProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
instrument_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeInstrumentProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
markets: {}
|
|
||||||
feature_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeFeatureProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
|
|
||||||
exp_manager:
|
|
||||||
class: MLflowExpManager
|
|
||||||
module_path: qlib.workflow.expm
|
|
||||||
kwargs:
|
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
|
||||||
|
|
||||||
task:
|
|
||||||
model:
|
|
||||||
class: RankICEnsembleLGBModel
|
|
||||||
module_path: tac_qlib.contrib.model.rank_ensemble
|
|
||||||
kwargs:
|
|
||||||
loss: mse
|
|
||||||
learning_rate: 0.02
|
|
||||||
num_leaves: 31
|
|
||||||
n_estimators: 3000
|
|
||||||
num_boost_round: 3000
|
|
||||||
early_stopping_rounds: 200
|
|
||||||
min_data_in_leaf: 20
|
|
||||||
lambda_l2: 0.5
|
|
||||||
colsample_bytree: 0.8
|
|
||||||
subsample: 0.8
|
|
||||||
subsample_freq: 1
|
|
||||||
reg_alpha: 0.1
|
|
||||||
reg_lambda: 1.0
|
|
||||||
seeds: "42,7,2026,99,123"
|
|
||||||
parallel: 5
|
|
||||||
|
|
||||||
dataset:
|
|
||||||
class: DatasetH
|
|
||||||
module_path: qlib.data.dataset
|
|
||||||
kwargs:
|
|
||||||
handler:
|
|
||||||
class: TACHandler
|
|
||||||
module_path: tac_qlib.contrib.data.handler
|
|
||||||
kwargs:
|
|
||||||
instruments: "{{ UNIVERSE }}"
|
|
||||||
start_time: 2015-01-03
|
|
||||||
end_time: 2026-08-14
|
|
||||||
fit_start_time: 2016-01-04
|
|
||||||
fit_end_time: 2025-09-01
|
|
||||||
freq: day
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
label: "Ref($close,-6)/Ref($close,-1)-1"
|
|
||||||
feature_fields: "$open,$high,$low,$close,$vwap,$volume,{{ SP_FIELDS }}"
|
|
||||||
infer_processors:
|
|
||||||
- class: DropAllNaN
|
|
||||||
kwargs: {}
|
|
||||||
- class: ProcessInf
|
|
||||||
kwargs: {}
|
|
||||||
- class: CSRankNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: ZScoreNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: Fillna
|
|
||||||
kwargs: {}
|
|
||||||
segments:
|
|
||||||
train: [2016-01-04, 2025-09-01]
|
|
||||||
valid: [2025-09-03, 2026-01-03]
|
|
||||||
test: [2026-01-04, 2026-08-10]
|
|
||||||
|
|
||||||
record:
|
|
||||||
- class: SignalRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs: {}
|
|
||||||
- class: SigAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
ana_long_short: true
|
|
||||||
ann_scaler: 252
|
|
||||||
- class: PortAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
config:
|
|
||||||
strategy:
|
|
||||||
class: HmmRiskTopk
|
|
||||||
module_path: tac_qlib.contrib.strategy.hmm_risk
|
|
||||||
kwargs:
|
|
||||||
signal: "<PRED>"
|
|
||||||
topk: 10
|
|
||||||
n_drop: 2
|
|
||||||
hmm_pause_pct: 0.70
|
|
||||||
drawdown_pause_pct: 8.0
|
|
||||||
liquidity_floor_adv: 5000000
|
|
||||||
only_tradable: true
|
|
||||||
risk_degree: 0.95
|
|
||||||
backtest:
|
|
||||||
start_time: 2026-01-04
|
|
||||||
end_time: 2026-08-10
|
|
||||||
account: 1000000
|
|
||||||
benchmark: SPY
|
|
||||||
exchange_kwargs:
|
|
||||||
codes: "{{ UNIVERSE }}"
|
|
||||||
deal_price: $close
|
|
||||||
freq: day
|
|
||||||
open_cost: 0.0005
|
|
||||||
close_cost: 0.0015
|
|
||||||
min_cost: 5.0
|
|
||||||
risk_analysis_freq: 1d
|
|
||||||
@@ -1,135 +0,0 @@
|
|||||||
# -----------------------------------------------------------------------------
|
|
||||||
# EXP 20 - R1: 2-seed ensemble (seeds 42,7), TopkDropout baseline.
|
|
||||||
#
|
|
||||||
# Runtime cut: 2 seeds instead of 5. Everything else identical to the reference
|
|
||||||
# (test 2026-01-04..2026-08-10, SPY, costs 5bp/15bp). Measures whether the
|
|
||||||
# 2-seed ensemble keeps the reference quality at ~2/5 the training time.
|
|
||||||
#
|
|
||||||
# Run:
|
|
||||||
# rd_run_workflow config_path=experiments/workflows/exp20-risk-limit-improve/r1_2seed.yaml \
|
|
||||||
# experiment_name=tac-rd-risk-limit
|
|
||||||
# -----------------------------------------------------------------------------
|
|
||||||
{%- set LAKE = TAC_LAKE_DIR %}
|
|
||||||
{%- set UNIVERSE = "SPY,QQQ,DIA,IWM,MDY,VTI,VOO,VEA,VWO,VT,EFA,EEM,TLT,IEF,SHY,AGG,BND,LQD,HYG,JNK,EMB,GLD,SLV,USO,UNG,DBA,DBC,XLK,XLF,XLE,XLV,XLI,XLY,XLP,XLU,XLB,XLRE,ARKK,SMH,SOXX,IBB,XBI,ITA,XAR,ICLN,TAN,FDN,IGV,ESPO,REM" %}
|
|
||||||
{%- set SP_FIELDS = "sp_ret,sp_jump_ratio,sp_jump_flag,sp_jump_tail,sp_max_move,sp_rv1,sp_rv5,sp_rv22,sp_vol_ratio_5_22,sp_vol_ratio_1_22,sp_trend_slope_5,sp_trend_slope_20,sp_trend_slope_60,sp_logp,sp_hurst_exponent,sp_sig_level1_lead,sp_sig_level1_lag,sp_sig_level2_lead_lag,sp_sig_level2_lag_lead,sma_3,ema_3" %}
|
|
||||||
|
|
||||||
qlib_init:
|
|
||||||
provider_uri: "{{ LAKE }}"
|
|
||||||
region: us
|
|
||||||
expression_cache: null
|
|
||||||
dataset_cache: null
|
|
||||||
|
|
||||||
calendar_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeCalendarProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
instrument_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeInstrumentProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
markets: {}
|
|
||||||
feature_provider:
|
|
||||||
class: tac_qlib.data.providers.LakeFeatureProvider
|
|
||||||
kwargs:
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
|
|
||||||
exp_manager:
|
|
||||||
class: MLflowExpManager
|
|
||||||
module_path: qlib.workflow.expm
|
|
||||||
kwargs:
|
|
||||||
uri: "sqlite:///{{ LAKE }}/mlruns.db"
|
|
||||||
default_exp_name: "tac-rd-risk-limit"
|
|
||||||
|
|
||||||
task:
|
|
||||||
model:
|
|
||||||
class: RankICEnsembleLGBModel
|
|
||||||
module_path: tac_qlib.contrib.model.rank_ensemble
|
|
||||||
kwargs:
|
|
||||||
loss: mse
|
|
||||||
learning_rate: 0.02
|
|
||||||
num_leaves: 31
|
|
||||||
n_estimators: 3000
|
|
||||||
num_boost_round: 3000
|
|
||||||
early_stopping_rounds: 200
|
|
||||||
min_data_in_leaf: 20
|
|
||||||
lambda_l2: 0.5
|
|
||||||
colsample_bytree: 0.8
|
|
||||||
subsample: 0.8
|
|
||||||
subsample_freq: 1
|
|
||||||
reg_alpha: 0.1
|
|
||||||
reg_lambda: 1.0
|
|
||||||
seeds: "42,7,2026,99,123"
|
|
||||||
parallel: 5
|
|
||||||
|
|
||||||
dataset:
|
|
||||||
class: DatasetH
|
|
||||||
module_path: qlib.data.dataset
|
|
||||||
kwargs:
|
|
||||||
handler:
|
|
||||||
class: TACHandler
|
|
||||||
module_path: tac_qlib.contrib.data.handler
|
|
||||||
kwargs:
|
|
||||||
instruments: "{{ UNIVERSE }}"
|
|
||||||
start_time: 2015-01-03
|
|
||||||
end_time: 2026-08-14
|
|
||||||
fit_start_time: 2016-01-04
|
|
||||||
fit_end_time: 2025-09-01
|
|
||||||
freq: day
|
|
||||||
lake_root: "{{ LAKE }}"
|
|
||||||
market: US
|
|
||||||
label: "Ref($close,-6)/Ref($close,-1)-1"
|
|
||||||
feature_fields: "$open,$high,$low,$close,$vwap,$volume,{{ SP_FIELDS }}"
|
|
||||||
infer_processors:
|
|
||||||
- class: DropAllNaN
|
|
||||||
kwargs: {}
|
|
||||||
- class: ProcessInf
|
|
||||||
kwargs: {}
|
|
||||||
- class: CSRankNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: ZScoreNorm
|
|
||||||
kwargs: {}
|
|
||||||
- class: Fillna
|
|
||||||
kwargs: {}
|
|
||||||
segments:
|
|
||||||
train: [2016-01-04, 2025-09-01]
|
|
||||||
valid: [2025-09-03, 2026-01-03]
|
|
||||||
test: [2026-01-04, 2026-08-10]
|
|
||||||
|
|
||||||
record:
|
|
||||||
- class: SignalRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs: {}
|
|
||||||
- class: SigAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
ana_long_short: true
|
|
||||||
ann_scaler: 252
|
|
||||||
- class: PortAnaRecord
|
|
||||||
module_path: qlib.workflow.record_temp
|
|
||||||
kwargs:
|
|
||||||
config:
|
|
||||||
strategy:
|
|
||||||
class: TopkDropoutStrategy
|
|
||||||
module_path: qlib.contrib.strategy
|
|
||||||
kwargs:
|
|
||||||
signal: "<PRED>"
|
|
||||||
topk: 10
|
|
||||||
n_drop: 2
|
|
||||||
only_tradable: true
|
|
||||||
risk_degree: 0.95
|
|
||||||
backtest:
|
|
||||||
start_time: 2026-01-04
|
|
||||||
end_time: 2026-08-10
|
|
||||||
account: 1000000
|
|
||||||
benchmark: SPY
|
|
||||||
exchange_kwargs:
|
|
||||||
codes: "{{ UNIVERSE }}"
|
|
||||||
deal_price: $close
|
|
||||||
freq: day
|
|
||||||
open_cost: 0.0005
|
|
||||||
close_cost: 0.0015
|
|
||||||
min_cost: 5.0
|
|
||||||
risk_analysis_freq: 1d
|
|
||||||
Reference in New Issue
Block a user