Joosep Pata
add cms model v2.4.0
d303317
checkpoint_freq: 1
comet: true
comet_name: particleflow-pt
comet_offline: false
comet_step_freq: 10
conv_type: attention
data_dir: /scratch/persistent/joosep/tensorflow_datasets
dataset: cms
dtype: bfloat16
enabled_test_datasets:
- cms_pf_qcd
finetune: null
gpu_batch_multiplier: 5
gpus: 1
load: null
lr: 0.0001
lr_schedule: cosinedecay
lr_schedule_config:
onecycle:
pct_start: 0.3
make_plots: null
model:
attention:
activation: relu
attention_type: flash
conv_type: attention
dropout_conv_id_ff: 0.0
dropout_conv_id_mha: 0.0
dropout_conv_reg_ff: 0.0
dropout_conv_reg_mha: 0.0
dropout_ff: 0.0
head_dim: 16
num_convs: 3
num_heads: 16
use_pre_layernorm: true
cos_phi_mode: linear
energy_mode: direct-elemtype-split
eta_mode: linear
gnn_lsh:
activation: elu
bin_size: 320
conv_type: gnn_lsh
distance_dim: 128
dropout_ff: 0.0
embedding_dim: 512
ffn_dist_hidden_dim: 128
ffn_dist_num_layers: 2
layernorm: true
max_num_bins: 200
num_convs: 8
num_node_messages: 2
width: 512
input_encoding: split
learned_representation_mode: last
mamba:
activation: elu
conv_type: mamba
d_conv: 4
d_state: 32
dropout_ff: 0.0
embedding_dim: 1024
expand: 2
num_convs: 4
width: 1024
pt_mode: direct-elemtype-split
sin_phi_mode: linear
trainable: all
ntest: 1000
ntrain: null
num_epochs: 10
num_workers: 8
nvalid: null
patience: 20
prefetch_factor: 50
raytune:
asha:
brackets: 1
grace_period: 10
max_t: 200
reduction_factor: 4
default_metric: val_loss
default_mode: min
hyperband:
max_t: 200
reduction_factor: 4
hyperopt:
n_random_steps: 10
local_dir: null
nevergrad:
n_random_steps: 10
sched: asha
search_alg: hyperopt
save_attention: false
sort_data: true
start_epoch: 1
test: null
test_dataset:
cms_pf_qcd:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_qcd_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_single_ele:
splits:
- 1
version: 2.5.0
cms_pf_ttbar:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ttbar_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ztt:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ztt_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
train: true
train_dataset:
cms:
physical_nopu:
batch_size: 8
samples:
cms_pf_qcd_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ttbar_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ztt_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
physical_pu:
batch_size: 1
samples:
cms_pf_qcd:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ttbar:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ztt:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
val_freq: null
valid_dataset:
cms:
physical_nopu:
batch_size: 8
samples:
cms_pf_qcd_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ttbar_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ztt_nopu:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
physical_pu:
batch_size: 1
samples:
cms_pf_qcd:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ttbar:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0
cms_pf_ztt:
splits:
- 1
- 2
- 3
- 4
- 5
- 6
- 7
- 8
- 9
- 10
version: 2.5.0