; oligo-analysis  -v 1 -sort -i $RSAT/public_html/tmp/www-data/2026/06/17/tmp_sequence_2026-06-17.141420_XRUSe9.fasta.purged -format fasta -lth occ_sig 0 -uth rank 50 -return occ,proba,rank -2str -noov -quick_if_possible -seqtype dna -markov 2 -pseudo 0.01 -l 8 -o $RSAT/public_html/tmp/www-data/2026/06/17/oligo-analysis_2026-06-17.141420_y6Waxm_8nt.tab
; Citation: van Helden et al. (1998). J Mol Biol 281(5), 827-42. 
; Program version              	1.169
; Quick counting mode          
; Detection of over-represented words (right-tail test)
; Oligomer length              	8
; Input file                   	$RSAT/public_html/tmp/www-data/2026/06/17/tmp_sequence_2026-06-17.141420_XRUSe9.fasta.purged
; Input format                 	fasta
; Output file                  	$RSAT/public_html/tmp/www-data/2026/06/17/oligo-analysis_2026-06-17.141420_y6Waxm_8nt.tab
; Discard overlapping matches
; Counted on both strands
; 	grouped by pairs of reverse complements
; Background model             	Markov
; Background estimation method 	Markov model estimated from input sequences
; Markov chain order           	2
; Pseudo-frequency             	0.01
; Pseudo-frequency per oligo   	3.03988326848249e-07
; Sequence type                	DNA
; Nb of sequences              	1
; Sum of sequence lengths      	6935
; discarded residues           	NA (quick mode)	 (other letters than ACGT)
; discarded occurrences        	NA (quick mode)	 (contain discarded residues)
; nb possible positions        	NA (quick mode)
; total oligo occurrences      	6206
; total overlapping occurrences	16
; total non overlapping occ    	6190
; alphabet size                	4
; nb possible oligomers        	32896
; oligomers tested for significance	5121
; Sequences:
;	2-UTseq835-UTL128-PhiB002	6935
;
; column headers
;	1	seq            	oligomer sequence
;	2	id             	oligomer identifier
;	3	exp_freq       	expected relative frequency
;	4	occ            	observed occurrences
;	5	exp_occ        	expected occurrences
;	6	occ_P          	occurrence probability (binomial)
;	7	occ_E          	E-value for occurrences (binomial)
;	8	occ_sig        	occurrence significance (binomial)
;	9	rank           	rank
;	10	ovl_occ        	number of overlapping occurrences (discarded from the count)
;	11	forbocc        	forbidden positions (to avoid self-overlap)
#seq	id	exp_freq	occ	exp_occ	occ_P	occ_E	occ_sig	rank	ovl_occ	forbocc
gtcaacac	gtcaacac|gtgttgac	0.0000261471057	6	0.16	2.2e-08	1.1e-04	3.95	1	0	42
aagtcaac	aagtcaac|gttgactt	0.0000276250767	6	0.17	3e-08	1.6e-04	3.81	2	0	42
agtcaaca	agtcaaca|tgttgact	0.0000329716707	6	0.20	8.5e-08	4.4e-04	3.36	3	0	42
attgtcag	attgtcag|ctgacaat	0.0000349121697	5	0.22	3.3e-06	1.7e-02	1.77	4	0	35
aaagtcaa	aaagtcaa|ttgacttt	0.0000435166557	5	0.27	9.6e-06	4.9e-02	1.31	5	0	35
gaatacga	gaatacga|tcgtattc	0.0000231222003	4	0.14	1.6e-05	8.1e-02	1.09	6	0	28
cccgccgc	cccgccgc|gcggcggg	0.0000095527643	3	0.06	3.3e-05	1.7e-01	0.77	7	0	21
aacgattc	aacgattc|gaatcgtt	0.0000137474052	3	0.09	9.7e-05	5.0e-01	0.30	8	0	21
cgttttaa	cgttttaa|ttaaaacg	0.0000411004617	4	0.26	0.00014	7.4e-01	0.13	9	0	28
acgattca	acgattca|tgaatcgt	0.0000165203853	3	0.10	0.00017	8.5e-01	0.07	10	0	21
; Host name	rsat
; Job started	2026-06-17.141422
; Job done	2026-06-17.141423
; Seconds	1.26
;	user	1.26
;	system	0.06
;	cuser	0.1
;	csystem	0.01
