chang-m-yun commited on
Commit
990ab17
·
verified ·
1 Parent(s): 722a17f

Upload folder using huggingface_hub

Browse files
This view is limited to 50 files because it contains too many changes.   See raw diff
Files changed (50) hide show
  1. README.md +120 -0
  2. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.args.json +42 -0
  3. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.batch_loss.tsv +0 -0
  4. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.bias_formatting.stdout.txt +1 -0
  5. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_data_params.tsv +3 -0
  6. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_formatting.stdout.txt +1 -0
  7. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_model_params.tsv +9 -0
  8. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt +1 -0
  9. fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.epoch_loss.csv +12 -0
  10. fold_0/model.bias_scaled.fold_0.ENCSR385AMY.h5 +3 -0
  11. fold_0/model.bias_scaled.fold_0.ENCSR385AMY.tar +3 -0
  12. fold_0/model.chrombpnet.fold_0.ENCSR385AMY.h5 +3 -0
  13. fold_0/model.chrombpnet.fold_0.ENCSR385AMY.tar +3 -0
  14. fold_0/model.chrombpnet_nobias.fold_0.ENCSR385AMY.h5 +3 -0
  15. fold_0/model.chrombpnet_nobias.fold_0.ENCSR385AMY.tar +3 -0
  16. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.args.json +50 -0
  17. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.batch_loss.tsv +0 -0
  18. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.bias_formatting.stdout.txt +1 -0
  19. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_data_params.tsv +3 -0
  20. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_formatting.stdout.txt +1 -0
  21. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_model_params.tsv +9 -0
  22. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt +1 -0
  23. fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.epoch_loss.csv +15 -0
  24. fold_1/model.bias_scaled.fold_1.ENCSR385AMY.h5 +3 -0
  25. fold_1/model.bias_scaled.fold_1.ENCSR385AMY.tar +3 -0
  26. fold_1/model.chrombpnet.fold_1.ENCSR385AMY.h5 +3 -0
  27. fold_1/model.chrombpnet.fold_1.ENCSR385AMY.tar +3 -0
  28. fold_1/model.chrombpnet_nobias.fold_1.ENCSR385AMY.h5 +3 -0
  29. fold_1/model.chrombpnet_nobias.fold_1.ENCSR385AMY.tar +3 -0
  30. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.args.json +50 -0
  31. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.batch_loss.tsv +0 -0
  32. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.bias_formatting.stdout.txt +1 -0
  33. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_data_params.tsv +3 -0
  34. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_formatting.stdout.txt +1 -0
  35. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_model_params.tsv +9 -0
  36. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt +1 -0
  37. fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.epoch_loss.csv +14 -0
  38. fold_2/model.bias_scaled.fold_2.ENCSR385AMY.h5 +3 -0
  39. fold_2/model.bias_scaled.fold_2.ENCSR385AMY.tar +3 -0
  40. fold_2/model.chrombpnet.fold_2.ENCSR385AMY.h5 +3 -0
  41. fold_2/model.chrombpnet.fold_2.ENCSR385AMY.tar +3 -0
  42. fold_2/model.chrombpnet_nobias.fold_2.ENCSR385AMY.h5 +3 -0
  43. fold_2/model.chrombpnet_nobias.fold_2.ENCSR385AMY.tar +3 -0
  44. fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.args.json +50 -0
  45. fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.batch_loss.tsv +0 -0
  46. fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.bias_formatting.stdout.txt +1 -0
  47. fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_data_params.tsv +3 -0
  48. fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_formatting.stdout.txt +1 -0
  49. fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_model_params.tsv +9 -0
  50. fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt +1 -0
README.md ADDED
@@ -0,0 +1,120 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ ---
2
+ license: mit
3
+ library_name: chrombpnet
4
+ tags:
5
+ - encode
6
+ - chrombpnet
7
+ - chromatin-accessibility
8
+ - DNASE
9
+ - limb
10
+ - hg38
11
+ ---
12
+ # ENCODE ChromBPNet Atlas
13
+ As part of the ENCODE 4 Project, we trained ChromBPNet models on 1,512 ENCODE DNAse-seq and ATAC-seq across 408 biosamples. Here, we provide all models for open-source use.
14
+
15
+ For more information about the models, see:
16
+ - Main ENCODE 4 Paper
17
+ - [A unified lexicon of predictive DNA sequence motifs from ENCODE transcription factor binding and chromatin accessibility assays](https://doi.org/10.5281/zenodo.17123347) (Deshpande et al., Zenodo 2025)
18
+ - [ChromBPNet: bias factorized, base-resolution deep learning models of chromatin accessibility reveal cis-regulatory sequence syntax, transcription factor footprints and regulatory variants](https://doi.org/10.1101/2024.12.25.630221) (Pampari et al., bioRxiv 2024)
19
+
20
+ ## ChromBPNet model: DNASE in hindlimb muscle (ENCSR385AMY)
21
+ - Model: ChromBPNet
22
+ - Assay: DNASE-seq
23
+ - Experiment: [ENCSR385AMY](https://www.encodeproject.org/experiments/ENCSR385AMY/)
24
+ - Model annotation: [ENCSR003PKQ](https://www.encodeproject.org/annotations/ENCSR003PKQ/)
25
+ - Biosample: hindlimb muscle (Full name: Homo sapiens hindlimb muscle tissue male embryo (120 days))
26
+ - Cell slim(s): None
27
+ - Organ slim(s): musculature-of-body,limb
28
+ - Developmental slim(s): mesoderm
29
+ - System slim(s): musculature
30
+ - Assembly: hg38
31
+
32
+ ## Directory structure
33
+ - `fold_0`: Model of 5-fold cross-validation: Fold 0
34
+ - `model.chrombpnet.fold_0.encid.h5`: full chrombpnet model that combines both bias and corrected model in .h5 format
35
+ - `model.chrombpnet_nobias.fold_0.encid.h5`: bias-corrected accessibility model in .h5 format (Use for all biological discovery)
36
+ - `model.bias_scaled.fold_0.encid.h5`: bias model in .h5 format
37
+ - `model.chrombpnet.fold_0.encid.tar`: full chrombpnet model that combines both bias and corrected model in SavedModel format. After being untarred, it results in a directory named "chrombpnet".
38
+ - `model.chrombpnet_nobias.fold_0.encid.tar`: bias-corrected accessibility model in SavedModel format (Use for all biological discovery). After being untarred, it results in a directory named "chrombpnet_wo_bias".
39
+ - `model.bias_scaled.fold_0.encid.tar`: bias model in SavedModel format. After being untarred, it results in a directory named "bias_model_scaled".
40
+ - `logs.models.fold_0.encid`: folder containing log files for training models
41
+ - `fold_1`: Model of 5-fold coss-validation: Fold 1
42
+ - `fold_2`: Model of 5-fold cross-validation: Fold 2
43
+ - `fold_3`: Model of 5-fold cross-validation: Fold 3
44
+ - `fold_4`: Model of 5-fold cross-validation: Fold 4
45
+
46
+ # Instructions
47
+ ## 1. Pseudocode for loading models in .h5 format
48
+
49
+ (1) Use the code in python after appropriately defining `model_in_h5_format` and `inputs`. \
50
+ (2) `inputs` is a one hot encoded sequence of shape (N,2114,4). Here N corresponds to the
51
+ number of tested sequences, 2114 is the input sequence length and 4 corresponds to [A,C,G,T].
52
+
53
+ ```python
54
+ import tensorflow as tf
55
+ from tensorflow.keras.utils import get_custom_objects
56
+ from tensorflow.keras.models import load_model
57
+
58
+ custom_objects={"tf": tf}
59
+ get_custom_objects().update(custom_objects)
60
+
61
+ model=load_model(model_in_h5_format,compile=False)
62
+ outputs = model(inputs)
63
+ ```
64
+
65
+ The list `outputs` consists of two elements. The first element has a shape of (N, 1000) and
66
+ contains logit predictions for a 1000-base-pair output. The second element, with a shape of
67
+ (N, 1), contains logcount predictions. To transform these predictions into per-base signals,
68
+ follow the provided pseudo code lines below.
69
+
70
+ ```python
71
+ import numpy as np
72
+
73
+ def softmax(x, temp=1):
74
+ norm_x = x - np.mean(x,axis=1, keepdims=True)
75
+ return np.exp(temp*norm_x)/np.sum(np.exp(temp*norm_x), axis=1, keepdims=True)
76
+
77
+ predictions = softmax(outputs[0]) * (np.exp(outputs[1])-1)
78
+ ```
79
+
80
+ ## 2. Pseudocode for loading models in .tar format
81
+
82
+ (1) First untar the directory as follows `tar -xvf model.tar`. \
83
+ (2) Use the code below in python after appropriately defining `model_dir_untared` and `inputs`. \
84
+ (3) `inputs` is a one hot encoded sequence of shape (N,2114,4). Here N corresponds to the number
85
+ of tested sequences, 2114 is the input sequence length and 4 corresponds to ACGT.
86
+
87
+ Reference: https://www.tensorflow.org/api_docs/python/tf/saved_model/load
88
+
89
+ ```python
90
+ import tensorflow as tf
91
+
92
+ model = tf.saved_model.load('model_dir_untared')
93
+ outputs = model.signatures['serving_default'](**{'sequence':inputs.astype('float32')})
94
+ ```
95
+
96
+ The variable `outputs` represents a dictionary containing two key-value pairs. The first key
97
+ is `logits_profile_predictions`, holding a value with a shape of (N, 1000). This value corresponds
98
+ to logit predictions for a 1000-base-pair output. The second key, named `logcount_predictions``,
99
+ is associated with a value of shape (N, 1), representing logcount predictions. To transform these
100
+ predictions into per-base signals, utilize the provided pseudo code lines mentioned below.
101
+
102
+ ```python
103
+ import numpy as np
104
+ def softmax(x, temp=1):
105
+ norm_x = x - np.mean(x,axis=1, keepdims=True)
106
+ return np.exp(temp*norm_x)/np.sum(np.exp(temp*norm_x), axis=1, keepdims=True)
107
+
108
+ predictions = softmax(outputs["logits_profile_predictions"]) * (np.exp(outputs["logcount_predictions"])-1)
109
+ ```
110
+
111
+ ## Docker image to load and use the models
112
+ - https://hub.docker.com/r/kundajelab/chrombpnet-atlas/ (tag:v1)
113
+
114
+ ## Code for ChromBPNet
115
+ - https://github.com/kundajelab/chrombpnet/
116
+
117
+ # License & citation
118
+ External data users may freely download, analyze and publish results based on any ENCODE data without restrictions.
119
+
120
+ Released under the [ENCODE data-use policy](https://www.encodeproject.org/about/data-use-policy/). Please cite the ENCODE Project Consortium and the model software: [ChromBPNet](https://github.com/kundajelab/chrombpnet) (Pampari et al., bioRxiv 2024).
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.args.json ADDED
@@ -0,0 +1,42 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "cmd": "pipeline",
3
+ "genome": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_no_alt_analysis_set_GCA_000001405.15.fasta",
4
+ "chrom_sizes": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_EBV.chrom.sizes.tsv",
5
+ "bigwig": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/DNASE/ENCSR385AMY/preprocessing/bigWigs/ENCSR385AMY.bigWig",
6
+ "output_dir": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY//fold0/",
7
+ "data_type": "DNASE",
8
+ "peaks": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY//fold0/auxiliary/filtered.peaks.bed",
9
+ "nonpeaks": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY//fold0/auxiliary/filtered.nonpeaks.bed",
10
+ "chr_fold_path": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_0.json",
11
+ "outlier_threshold": 0.9999,
12
+ "ATAC_ref_path": null,
13
+ "DNASE_ref_path": null,
14
+ "num_samples": 10000,
15
+ "inputlen": 2114,
16
+ "outputlen": 1000,
17
+ "seed": 1234,
18
+ "epochs": 50,
19
+ "early_stop": 5,
20
+ "learning_rate": 0.001,
21
+ "trackables": [
22
+ "logcount_predictions_loss",
23
+ "loss",
24
+ "logits_profile_predictions_loss",
25
+ "val_logcount_predictions_loss",
26
+ "val_loss",
27
+ "val_logits_profile_predictions_loss"
28
+ ],
29
+ "architecture_from_file": "/home/groups/akundaje/ziwei75/anaconda3/envs/chrombpnet/lib/python3.8/site-packages/chrombpnet/training/models/chrombpnet_with_bias_model.py",
30
+ "file_prefix": null,
31
+ "html_prefix": "./",
32
+ "bias_model_path": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_bias_model/filtered_negatives_models/ENCSR385AMY/models/bias.h5",
33
+ "negative_sampling_ratio": 0.1,
34
+ "filters": 512,
35
+ "n_dilation_layers": 8,
36
+ "max_jitter": 500,
37
+ "batch_size": 64,
38
+ "output_prefix": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY//fold0/models/chrombpnet",
39
+ "chr": "chr8",
40
+ "pwm_width": 24,
41
+ "params": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY//fold0/logs/chrombpnet_model_params.tsv"
42
+ }
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.batch_loss.tsv ADDED
The diff for this file is too large to render. See raw diff
 
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY/fold0/models/bias_model_scaled.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_0/new_model_format/bias_model_scaled.tar with get_new_tf_model_format.py
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_data_params.tsv ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ counts_sum_min_thresh 0.0
2
+ counts_sum_max_thresh 5132.92
3
+ trainings_pts_post_thresh 171261
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY/fold0/models/chrombpnet.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_0/new_model_format/chrombpnet.tar with get_new_tf_model_format.py
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_model_params.tsv ADDED
@@ -0,0 +1,9 @@
 
 
 
 
 
 
 
 
 
 
1
+ counts_loss_weight 12.2
2
+ filters 512
3
+ n_dil_layers 8
4
+ bias_model_path /oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY//fold0/models/bias_model_scaled.h5
5
+ inputlen 2114
6
+ outputlen 1000
7
+ max_jitter 500
8
+ chr_fold_path /oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_0.json
9
+ negative_sampling_ratio 0.1
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_model/bias_filtered/ENCSR385AMY/fold0/models/chrombpnet_nobias.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_0/new_model_format/chrombpnet_nobias.tar with get_new_tf_model_format.py
fold_0/logs.models.fold_0.ENCSR385AMY/logfile.modelling.fold_0.ENCSR385AMY.epoch_loss.csv ADDED
@@ -0,0 +1,12 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ epoch,logcount_predictions_loss,logits_profile_predictions_loss,loss,val_logcount_predictions_loss,val_logits_profile_predictions_loss,val_loss
2
+ 0,1.9770828485488892,522.44189453125,546.5613403320312,0.878165602684021,507.1616516113281,517.8751831054688
3
+ 1,0.8748293519020081,495.7141418457031,506.3877258300781,0.8365785479545593,494.6370849609375,504.8434143066406
4
+ 2,0.7963923811912537,488.0716552734375,497.7874755859375,0.7952279448509216,495.7333984375,505.435302734375
5
+ 3,0.7414178848266602,483.43304443359375,492.4776611328125,0.6943086385726929,493.35467529296875,501.82525634765625
6
+ 4,0.7072145342826843,479.7239685058594,488.3514404296875,0.6661702990531921,493.9805603027344,502.1078186035156
7
+ 5,0.6854123473167419,476.93798828125,485.2994689941406,0.699362576007843,490.8526611328125,499.38507080078125
8
+ 6,0.6538832783699036,474.883056640625,482.8609619140625,0.6603286862373352,494.6324768066406,502.688720703125
9
+ 7,0.6392768621444702,472.5515441894531,480.3503112792969,0.7141312956809998,491.0947265625,499.80706787109375
10
+ 8,0.626732349395752,470.98797607421875,478.634033203125,0.6582152843475342,495.2186584472656,503.2485656738281
11
+ 9,0.6128236651420593,469.64739990234375,477.1233215332031,0.6511006355285645,499.6353454589844,507.57879638671875
12
+ 10,0.5974838137626648,468.0060119628906,475.29583740234375,0.6157989501953125,502.32843017578125,509.84100341796875
fold_0/model.bias_scaled.fold_0.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0ae27f76e5091e7c278af748a2124c76346ddf0bb165c95e7a451cb4f6094be8
3
+ size 2691928
fold_0/model.bias_scaled.fold_0.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:81db4ee3e3fa5b5be725dbc90440517ba8f598afea84a2ee7664c0dfe0bfd490
3
+ size 1198080
fold_0/model.chrombpnet.fold_0.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:b323c74ba6355fa98766f0695d90b4a75fe263cb588535ddc09563cd64d972c7
3
+ size 77538952
fold_0/model.chrombpnet.fold_0.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:268d0b12a78c1834d7b71bd50b2a56ea8dd22eb1a74a1644d7f325d55fbc2f7b
3
+ size 27525120
fold_0/model.chrombpnet_nobias.fold_0.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:4791a15cdd32113d43f0af0db2e58252028f8fc1bf0285ac9d0a508298a7d151
3
+ size 25582648
fold_0/model.chrombpnet_nobias.fold_0.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:1a829c2c39cc0062507d9d8650df8211c146ff340fd15c246422e71051ca0d54
3
+ size 26060800
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.args.json ADDED
@@ -0,0 +1,50 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "cmd": "pipeline",
3
+ "genome": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_no_alt_analysis_set_GCA_000001405.15.fasta",
4
+ "chrom_sizes": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_EBV.chrom.sizes.tsv",
5
+ "input_bam_file": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/DNASE/ENCSR385AMY/preprocessing/bigWigs/ENCSR385AMY.bigWig",
6
+ "input_fragment_file": null,
7
+ "input_tagalign_file": null,
8
+ "output_dir": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1",
9
+ "data_type": "DNASE",
10
+ "peaks": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/auxiliary/filtered.peaks.bed",
11
+ "nonpeaks": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/auxiliary/filtered.nonpeaks.bed",
12
+ "chr_fold_path": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_1.json",
13
+ "outlier_threshold": 0.9999,
14
+ "ATAC_ref_path": null,
15
+ "DNASE_ref_path": null,
16
+ "num_samples": 10000,
17
+ "inputlen": 2114,
18
+ "outputlen": 1000,
19
+ "seed": 1234,
20
+ "epochs": 50,
21
+ "early_stop": 5,
22
+ "learning_rate": 0.001,
23
+ "trackables": [
24
+ "logcount_predictions_loss",
25
+ "loss",
26
+ "logits_profile_predictions_loss",
27
+ "val_logcount_predictions_loss",
28
+ "val_loss",
29
+ "val_logits_profile_predictions_loss"
30
+ ],
31
+ "architecture_from_file": "/home/users/vhecht/chrombpnet/chrombpnet/chrombpnet/training/models/chrombpnet_with_bias_model.py",
32
+ "file_prefix": null,
33
+ "html_prefix": "./",
34
+ "bsort": false,
35
+ "tmpdir": null,
36
+ "no_st": false,
37
+ "bias_model_path": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_bias_model/filtered_negatives_models/ENCSR385AMY/models/bias.h5",
38
+ "negative_sampling_ratio": 0.1,
39
+ "filters": 512,
40
+ "n_dilation_layers": 8,
41
+ "max_jitter": 500,
42
+ "batch_size": 64,
43
+ "output_prefix": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/models/chrombpnet",
44
+ "bigwig": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/auxiliary/data_unstranded.bw",
45
+ "plus_shift": null,
46
+ "minus_shift": null,
47
+ "chr": "chr12",
48
+ "pwm_width": 24,
49
+ "params": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/logs/chrombpnet_model_params.tsv"
50
+ }
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.batch_loss.tsv ADDED
The diff for this file is too large to render. See raw diff
 
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/models/bias_model_scaled.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_1/new_model_format/bias_model_scaled.tar with get_new_tf_model_format.py
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_data_params.tsv ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ counts_sum_min_thresh 0.0
2
+ counts_sum_max_thresh 5143.72
3
+ trainings_pts_post_thresh 173979
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/models/chrombpnet.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_1/new_model_format/chrombpnet.tar with get_new_tf_model_format.py
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_model_params.tsv ADDED
@@ -0,0 +1,9 @@
 
 
 
 
 
 
 
 
 
 
1
+ counts_loss_weight 12.2
2
+ filters 512
3
+ n_dil_layers 8
4
+ bias_model_path /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/models/bias_model_scaled.h5
5
+ inputlen 2114
6
+ outputlen 1000
7
+ max_jitter 500
8
+ chr_fold_path /oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_1.json
9
+ negative_sampling_ratio 0.1
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_1/models/chrombpnet_nobias.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_1/new_model_format/chrombpnet_nobias.tar with get_new_tf_model_format.py
fold_1/logs.models.fold_1.ENCSR385AMY/logfile.modelling.fold_1.ENCSR385AMY.epoch_loss.csv ADDED
@@ -0,0 +1,15 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ epoch,logcount_predictions_loss,logits_profile_predictions_loss,loss,val_logcount_predictions_loss,val_logits_profile_predictions_loss,val_loss
2
+ 0,2.424912214279175,529.6612548828125,559.24462890625,0.9572994709014893,602.5901489257812,614.2691650390625
3
+ 1,0.9195420145988464,498.1009521484375,509.3187255859375,0.8413581848144531,579.21533203125,589.4797973632812
4
+ 2,0.8402743339538574,488.75689697265625,499.00860595703125,0.7647204995155334,576.4014282226562,585.73095703125
5
+ 3,0.7942594289779663,483.0146484375,492.7039489746094,0.7088714838027954,566.425048828125,575.0730590820312
6
+ 4,0.7513942718505859,479.6485595703125,488.8154296875,0.6680809259414673,567.3887329101562,575.5392456054688
7
+ 5,0.7172863483428955,476.3273620605469,485.07843017578125,0.646342933177948,566.6307373046875,574.515869140625
8
+ 6,0.6939284801483154,472.8807067871094,481.3465576171875,0.6559173464775085,569.8682861328125,577.8707275390625
9
+ 7,0.6707106828689575,470.9489440917969,479.1319885253906,0.6465588808059692,567.7783813476562,575.6661987304688
10
+ 8,0.6497904658317566,468.7471618652344,476.6751403808594,0.6884491443634033,565.8653564453125,574.2647705078125
11
+ 9,0.638763427734375,466.5169982910156,474.31011962890625,0.6580843329429626,569.990966796875,578.0197143554688
12
+ 10,0.6197393536567688,465.3634948730469,472.9242858886719,0.6305816173553467,570.5250244140625,578.2179565429688
13
+ 11,0.6102444529533386,463.60302734375,471.0477600097656,0.7974913120269775,573.4833984375,583.2124633789062
14
+ 12,0.5966919660568237,462.4775695800781,469.7580871582031,0.6064395308494568,568.424072265625,575.82275390625
15
+ 13,0.590790867805481,460.1824951171875,467.390869140625,0.6153243780136108,570.3412475585938,577.8485107421875
fold_1/model.bias_scaled.fold_1.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:6e1b7c91434c16ce25406b1c134ef9a2428b83222806eb9f37f98b8eb1ef03a0
3
+ size 2691928
fold_1/model.bias_scaled.fold_1.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0668e8cb1932dbc4918c86494b7cec4bd69fab72513d7a2cb239516b6c36b85a
3
+ size 1198080
fold_1/model.chrombpnet.fold_1.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:4bbd35af3ab0259db2452d2b8e6626cc0bec9053e493168b4b0166719a80a30f
3
+ size 77538840
fold_1/model.chrombpnet.fold_1.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2ad05ee8f5cbd864e07461beb728673eea7db3b5f44e37db34ac891b7801b5db
3
+ size 27525120
fold_1/model.chrombpnet_nobias.fold_1.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:590fe49e7fa86985110edd60c7dfe6aecbada2561e4b4a32d5a1178ec6849f41
3
+ size 25582648
fold_1/model.chrombpnet_nobias.fold_1.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:bd9b7a16df3f10d04922820be12dfb2de5618542a8fed64b9e34084e1808da71
3
+ size 26060800
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.args.json ADDED
@@ -0,0 +1,50 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "cmd": "pipeline",
3
+ "genome": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_no_alt_analysis_set_GCA_000001405.15.fasta",
4
+ "chrom_sizes": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_EBV.chrom.sizes.tsv",
5
+ "input_bam_file": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/DNASE/ENCSR385AMY/preprocessing/bigWigs/ENCSR385AMY.bigWig",
6
+ "input_fragment_file": null,
7
+ "input_tagalign_file": null,
8
+ "output_dir": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2",
9
+ "data_type": "DNASE",
10
+ "peaks": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/auxiliary/filtered.peaks.bed",
11
+ "nonpeaks": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/auxiliary/filtered.nonpeaks.bed",
12
+ "chr_fold_path": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_2.json",
13
+ "outlier_threshold": 0.9999,
14
+ "ATAC_ref_path": null,
15
+ "DNASE_ref_path": null,
16
+ "num_samples": 10000,
17
+ "inputlen": 2114,
18
+ "outputlen": 1000,
19
+ "seed": 1234,
20
+ "epochs": 50,
21
+ "early_stop": 5,
22
+ "learning_rate": 0.001,
23
+ "trackables": [
24
+ "logcount_predictions_loss",
25
+ "loss",
26
+ "logits_profile_predictions_loss",
27
+ "val_logcount_predictions_loss",
28
+ "val_loss",
29
+ "val_logits_profile_predictions_loss"
30
+ ],
31
+ "architecture_from_file": "/home/users/vhecht/chrombpnet/chrombpnet/chrombpnet/training/models/chrombpnet_with_bias_model.py",
32
+ "file_prefix": null,
33
+ "html_prefix": "./",
34
+ "bsort": false,
35
+ "tmpdir": null,
36
+ "no_st": false,
37
+ "bias_model_path": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_bias_model/filtered_negatives_models/ENCSR385AMY/models/bias.h5",
38
+ "negative_sampling_ratio": 0.1,
39
+ "filters": 512,
40
+ "n_dilation_layers": 8,
41
+ "max_jitter": 500,
42
+ "batch_size": 64,
43
+ "output_prefix": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/models/chrombpnet",
44
+ "bigwig": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/auxiliary/data_unstranded.bw",
45
+ "plus_shift": null,
46
+ "minus_shift": null,
47
+ "chr": "chr22",
48
+ "pwm_width": 24,
49
+ "params": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/logs/chrombpnet_model_params.tsv"
50
+ }
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.batch_loss.tsv ADDED
The diff for this file is too large to render. See raw diff
 
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/models/bias_model_scaled.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_2/new_model_format/bias_model_scaled.tar with get_new_tf_model_format.py
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_data_params.tsv ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ counts_sum_min_thresh 0.0
2
+ counts_sum_max_thresh 5126.81
3
+ trainings_pts_post_thresh 178753
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/models/chrombpnet.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_2/new_model_format/chrombpnet.tar with get_new_tf_model_format.py
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_model_params.tsv ADDED
@@ -0,0 +1,9 @@
 
 
 
 
 
 
 
 
 
 
1
+ counts_loss_weight 12.2
2
+ filters 512
3
+ n_dil_layers 8
4
+ bias_model_path /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/models/bias_model_scaled.h5
5
+ inputlen 2114
6
+ outputlen 1000
7
+ max_jitter 500
8
+ chr_fold_path /oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_2.json
9
+ negative_sampling_ratio 0.1
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_2/models/chrombpnet_nobias.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_2/new_model_format/chrombpnet_nobias.tar with get_new_tf_model_format.py
fold_2/logs.models.fold_2.ENCSR385AMY/logfile.modelling.fold_2.ENCSR385AMY.epoch_loss.csv ADDED
@@ -0,0 +1,14 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ epoch,logcount_predictions_loss,logits_profile_predictions_loss,loss,val_logcount_predictions_loss,val_logits_profile_predictions_loss,val_loss
2
+ 0,2.771857500076294,536.7540283203125,570.571044921875,0.9925419092178345,533.5272216796875,545.6359252929688
3
+ 1,0.9409024119377136,502.71185302734375,514.1902465820312,0.812254786491394,519.439208984375,529.348876953125
4
+ 2,0.8516588807106018,493.65301513671875,504.0433654785156,0.8800976276397705,517.45556640625,528.192626953125
5
+ 3,0.7830417156219482,488.1264953613281,497.6792297363281,1.090691089630127,515.3241577148438,528.630615234375
6
+ 4,0.7456042170524597,483.96575927734375,493.06182861328125,0.7167993783950806,504.13275146484375,512.8778076171875
7
+ 5,0.7080942392349243,480.1345520019531,488.7730712890625,0.7253450751304626,503.6178283691406,512.4674072265625
8
+ 6,0.6795583367347717,478.2273864746094,486.5172424316406,0.7077139019966125,507.39202880859375,516.026123046875
9
+ 7,0.6550350785255432,474.9714050292969,482.9627685546875,0.7205138802528381,502.78802490234375,511.5780334472656
10
+ 8,0.6312917470932007,473.41387939453125,481.115478515625,0.6696923971176147,506.3115234375,514.481689453125
11
+ 9,0.6226451992988586,470.5510559082031,478.1474609375,0.6414923071861267,504.35382080078125,512.1802368164062
12
+ 10,0.6068309545516968,469.2301940917969,476.63287353515625,0.6340993642807007,505.0211181640625,512.7572021484375
13
+ 11,0.5914859175682068,467.615234375,474.8310241699219,0.6398987770080566,506.36077880859375,514.1672973632812
14
+ 12,0.5783110857009888,465.83984375,472.8951110839844,0.6854179501533508,506.67529296875,515.0374755859375
fold_2/model.bias_scaled.fold_2.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:455ff2f8cabfe7d734b473857bfccc3a765a25eb8659ca03d6084d1d5622d9ca
3
+ size 2691928
fold_2/model.bias_scaled.fold_2.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7f4b97d0cd2e444a1f12e97b9da6e4db2b4a88fb675ab2c5550ce3dd5f9457a8
3
+ size 1198080
fold_2/model.chrombpnet.fold_2.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:125751b1a8478345c73d935d3d3cbb5e1aa8aaa09d026dd29dcda2b86a1a261c
3
+ size 77538840
fold_2/model.chrombpnet.fold_2.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0be2ccceeb5f5573a47e30efaab7239fbecae9ea1d64ba43bdede86e952b881a
3
+ size 27525120
fold_2/model.chrombpnet_nobias.fold_2.ENCSR385AMY.h5 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:10a5c6281c1c277dbedb37f9fc54f365ffda3e925e254bb07751dbf6d1a44f67
3
+ size 25582648
fold_2/model.chrombpnet_nobias.fold_2.ENCSR385AMY.tar ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:82be39ad5c702144940f5ee0f11891fd0f93233dc37452e6a98057d351e93eef
3
+ size 26060800
fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.args.json ADDED
@@ -0,0 +1,50 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "cmd": "pipeline",
3
+ "genome": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_no_alt_analysis_set_GCA_000001405.15.fasta",
4
+ "chrom_sizes": "/oak/stanford/groups/akundaje/ziwei75/atac_seq_pipeline/hg38/GRCh38_EBV.chrom.sizes.tsv",
5
+ "input_bam_file": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/DNASE/ENCSR385AMY/preprocessing/bigWigs/ENCSR385AMY.bigWig",
6
+ "input_fragment_file": null,
7
+ "input_tagalign_file": null,
8
+ "output_dir": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3",
9
+ "data_type": "DNASE",
10
+ "peaks": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/auxiliary/filtered.peaks.bed",
11
+ "nonpeaks": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/auxiliary/filtered.nonpeaks.bed",
12
+ "chr_fold_path": "/oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_3.json",
13
+ "outlier_threshold": 0.9999,
14
+ "ATAC_ref_path": null,
15
+ "DNASE_ref_path": null,
16
+ "num_samples": 10000,
17
+ "inputlen": 2114,
18
+ "outputlen": 1000,
19
+ "seed": 1234,
20
+ "epochs": 50,
21
+ "early_stop": 5,
22
+ "learning_rate": 0.001,
23
+ "trackables": [
24
+ "logcount_predictions_loss",
25
+ "loss",
26
+ "logits_profile_predictions_loss",
27
+ "val_logcount_predictions_loss",
28
+ "val_loss",
29
+ "val_logits_profile_predictions_loss"
30
+ ],
31
+ "architecture_from_file": "/home/users/vhecht/chrombpnet/chrombpnet/chrombpnet/training/models/chrombpnet_with_bias_model.py",
32
+ "file_prefix": null,
33
+ "html_prefix": "./",
34
+ "bsort": false,
35
+ "tmpdir": null,
36
+ "no_st": false,
37
+ "bias_model_path": "/oak/stanford/groups/akundaje/ziwei75/chromatin_atlas_bias/DNase_bias_model/filtered_negatives_models/ENCSR385AMY/models/bias.h5",
38
+ "negative_sampling_ratio": 0.1,
39
+ "filters": 512,
40
+ "n_dilation_layers": 8,
41
+ "max_jitter": 500,
42
+ "batch_size": 64,
43
+ "output_prefix": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/models/chrombpnet",
44
+ "bigwig": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/auxiliary/data_unstranded.bw",
45
+ "plus_shift": null,
46
+ "minus_shift": null,
47
+ "chr": "chr6",
48
+ "pwm_width": 24,
49
+ "params": "/oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/logs/chrombpnet_model_params.tsv"
50
+ }
fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.batch_loss.tsv ADDED
The diff for this file is too large to render. See raw diff
 
fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/models/bias_model_scaled.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_3/new_model_format/bias_model_scaled.tar with get_new_tf_model_format.py
fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_data_params.tsv ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ counts_sum_min_thresh 0.0
2
+ counts_sum_max_thresh 5257.42
3
+ trainings_pts_post_thresh 171787
fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/models/chrombpnet.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_3/new_model_format/chrombpnet.tar with get_new_tf_model_format.py
fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_model_params.tsv ADDED
@@ -0,0 +1,9 @@
 
 
 
 
 
 
 
 
 
 
1
+ counts_loss_weight 12.2
2
+ filters 512
3
+ n_dil_layers 8
4
+ bias_model_path /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/models/bias_model_scaled.h5
5
+ inputlen 2114
6
+ outputlen 1000
7
+ max_jitter 500
8
+ chr_fold_path /oak/stanford/groups/akundaje/projects/chromatin-atlas-2022/splits/fold_3.json
9
+ negative_sampling_ratio 0.1
fold_3/logs.models.fold_3.ENCSR385AMY/logfile.modelling.fold_3.ENCSR385AMY.chrombpnet_no_bias_formatting.stdout.txt ADDED
@@ -0,0 +1 @@
 
 
1
+ Converting /oak/stanford/groups/akundaje/vhecht/chromatin_atlas_bias_corrected/DNASE_model/ENCSR385AMY/fold_3/models/chrombpnet_nobias.h5 to /oak/stanford/groups/akundaje/vhecht/chromatin-atlas-2022/DNASE/ENCSR385AMY/fold_3/new_model_format/chrombpnet_nobias.tar with get_new_tf_model_format.py