diff --git "a/lr4e-4_total_batch_size5120_seq_len128/log2.txt" "b/lr4e-4_total_batch_size5120_seq_len128/log2.txt" --- "a/lr4e-4_total_batch_size5120_seq_len128/log2.txt" +++ "b/lr4e-4_total_batch_size5120_seq_len128/log2.txt" @@ -1,1043 +1,5203 @@ -max_steps: 1000 -0 val loss 11.6870 -0 val perplexity 119016.5000 -0 train 11.683437 (lr=1.3986e-06) (hash(x)=6154740) -1 train 11.714834 (lr=2.7972e-06) (hash(x)=6605789) -2 train 11.724317 (lr=4.1958e-06) (hash(x)=6696707) -3 train 11.702266 (lr=5.5944e-06) (hash(x)=5113283) -4 train 11.619727 (lr=6.9930e-06) (hash(x)=5607069) -5 train 11.654514 (lr=8.3916e-06) (hash(x)=7007748) -6 train 11.585394 (lr=9.7902e-06) (hash(x)=7643299) -7 train 11.596355 (lr=1.1189e-05) (hash(x)=5534142) -8 train 11.551446 (lr=1.2587e-05) (hash(x)=6539293) -9 train 11.563625 (lr=1.3986e-05) (hash(x)=5974274) -10 train 11.548813 (lr=1.5385e-05) (hash(x)=7199085) -11 train 11.482031 (lr=1.6783e-05) (hash(x)=6763231) -12 train 11.380842 (lr=1.8182e-05) (hash(x)=6720398) -13 train 11.360128 (lr=1.9580e-05) (hash(x)=6706513) -14 train 11.364397 (lr=2.0979e-05) (hash(x)=5727925) -15 train 11.290984 (lr=2.2378e-05) (hash(x)=6328012) -16 train 11.216109 (lr=2.3776e-05) (hash(x)=5812722) -17 train 11.163255 (lr=2.5175e-05) (hash(x)=6183012) -18 train 11.194344 (lr=2.6573e-05) (hash(x)=5952539) -19 train 11.130493 (lr=2.7972e-05) (hash(x)=7052161) -20 train 10.930115 (lr=2.9371e-05) (hash(x)=5810871) -21 train 10.979438 (lr=3.0769e-05) (hash(x)=6586466) -22 train 10.851637 (lr=3.2168e-05) (hash(x)=5061781) -23 train 10.863820 (lr=3.3566e-05) (hash(x)=5556098) -24 train 10.915709 (lr=3.4965e-05) (hash(x)=7268535) -25 train 10.771225 (lr=3.6364e-05) (hash(x)=8155572) -26 train 10.657117 (lr=3.7762e-05) (hash(x)=5019590) -27 train 10.572108 (lr=3.9161e-05) (hash(x)=4696943) -28 train 10.519890 (lr=4.0559e-05) (hash(x)=5984010) -29 train 10.677322 (lr=4.1958e-05) (hash(x)=7649389) -30 train 10.422065 (lr=4.3357e-05) (hash(x)=6332817) -31 train 10.346922 (lr=4.4755e-05) (hash(x)=6219138) -32 train 10.494532 (lr=4.6154e-05) (hash(x)=6895681) -33 train 10.375053 (lr=4.7552e-05) (hash(x)=7075794) -34 train 10.271528 (lr=4.8951e-05) (hash(x)=5024163) -35 train 10.334936 (lr=5.0350e-05) (hash(x)=6049878) -36 train 10.403502 (lr=5.1748e-05) (hash(x)=4774943) -37 train 10.398000 (lr=5.3147e-05) (hash(x)=6385927) -38 train 10.220635 (lr=5.4545e-05) (hash(x)=5070217) -39 train 10.221786 (lr=5.5944e-05) (hash(x)=5491972) -40 train 10.286844 (lr=5.7343e-05) (hash(x)=5695751) -41 train 10.220202 (lr=5.8741e-05) (hash(x)=6061420) -42 train 10.138086 (lr=6.0140e-05) (hash(x)=7836710) -43 train 10.250729 (lr=6.1538e-05) (hash(x)=8075458) -44 train 10.070316 (lr=6.2937e-05) (hash(x)=7190584) -45 train 10.125028 (lr=6.4336e-05) (hash(x)=7238769) -46 train 10.279922 (lr=6.5734e-05) (hash(x)=4931205) -47 train 10.091259 (lr=6.7133e-05) (hash(x)=5917741) -48 train 10.034421 (lr=6.8531e-05) (hash(x)=5530785) -49 train 10.084161 (lr=6.9930e-05) (hash(x)=5430605) -50 val loss 10.0612 -50 val perplexity 23416.7500 -50 train 10.069185 (lr=7.1329e-05) (hash(x)=6056774) -51 train 9.960612 (lr=7.2727e-05) (hash(x)=5750403) -52 train 9.968672 (lr=7.4126e-05) (hash(x)=6187497) -53 train 9.958236 (lr=7.5524e-05) (hash(x)=4792640) -54 train 10.064965 (lr=7.6923e-05) (hash(x)=6868430) -55 train 9.912956 (lr=7.8322e-05) (hash(x)=6112458) -56 train 9.852873 (lr=7.9720e-05) (hash(x)=5523601) -57 train 9.787319 (lr=8.1119e-05) (hash(x)=5801122) -58 train 9.881976 (lr=8.2517e-05) (hash(x)=6826295) -59 train 9.995114 (lr=8.3916e-05) (hash(x)=5806490) -60 train 9.822445 (lr=8.5315e-05) (hash(x)=6223100) -61 train 9.759766 (lr=8.6713e-05) (hash(x)=5754140) -62 train 9.739307 (lr=8.8112e-05) (hash(x)=6675367) -63 train 9.644345 (lr=8.9510e-05) (hash(x)=6568379) -64 train 9.501768 (lr=9.0909e-05) (hash(x)=5648233) -65 train 9.672204 (lr=9.2308e-05) (hash(x)=7134932) -66 train 9.627056 (lr=9.3706e-05) (hash(x)=5242536) -67 train 9.665130 (lr=9.5105e-05) (hash(x)=7212403) -68 train 9.505008 (lr=9.6503e-05) (hash(x)=5342967) -69 train 9.610952 (lr=9.7902e-05) (hash(x)=7720703) -70 train 9.525063 (lr=9.9301e-05) (hash(x)=7351337) -71 train 9.558569 (lr=1.0070e-04) (hash(x)=7386082) -72 train 9.406364 (lr=1.0210e-04) (hash(x)=6931318) -73 train 9.404094 (lr=1.0350e-04) (hash(x)=6848513) -74 train 9.399150 (lr=1.0490e-04) (hash(x)=7714274) -75 train 9.452235 (lr=1.0629e-04) (hash(x)=6359743) -76 train 9.264580 (lr=1.0769e-04) (hash(x)=6801650) -77 train 9.226777 (lr=1.0909e-04) (hash(x)=7536538) -78 train 9.317209 (lr=1.1049e-04) (hash(x)=6957410) -79 train 9.313169 (lr=1.1189e-04) (hash(x)=6151674) -80 train 9.268032 (lr=1.1329e-04) (hash(x)=6033844) -81 train 9.132256 (lr=1.1469e-04) (hash(x)=6317099) -82 train 9.070349 (lr=1.1608e-04) (hash(x)=5357844) -83 train 9.122319 (lr=1.1748e-04) (hash(x)=7168704) -84 train 9.251307 (lr=1.1888e-04) (hash(x)=7157938) -85 train 9.034517 (lr=1.2028e-04) (hash(x)=6736092) -86 train 8.936291 (lr=1.2168e-04) (hash(x)=5383012) -87 train 8.991364 (lr=1.2308e-04) (hash(x)=7859060) -88 train 8.970185 (lr=1.2448e-04) (hash(x)=5869288) -89 train 8.944607 (lr=1.2587e-04) (hash(x)=6958223) -90 train 8.808808 (lr=1.2727e-04) (hash(x)=6731344) -91 train 8.763073 (lr=1.2867e-04) (hash(x)=6976239) -92 train 8.726782 (lr=1.3007e-04) (hash(x)=7769441) -93 train 8.755550 (lr=1.3147e-04) (hash(x)=6956845) -94 train 8.856450 (lr=1.3287e-04) (hash(x)=6638037) -95 train 8.550882 (lr=1.3427e-04) (hash(x)=6391506) -96 train 8.845491 (lr=1.3566e-04) (hash(x)=7319928) -97 train 8.793879 (lr=1.3706e-04) (hash(x)=6113701) -98 train 8.449331 (lr=1.3846e-04) (hash(x)=6246453) -99 train 8.519605 (lr=1.3986e-04) (hash(x)=5245999) -100 val loss 8.5686 -100 val perplexity 5263.9146 -100 train 8.578829 (lr=1.4126e-04) (hash(x)=6227256) -101 train 8.509744 (lr=1.4266e-04) (hash(x)=5774189) -102 train 8.398865 (lr=1.4406e-04) (hash(x)=5588220) -103 train 8.480408 (lr=1.4545e-04) (hash(x)=5985675) -104 train 8.346600 (lr=1.4685e-04) (hash(x)=6259214) -105 train 8.453425 (lr=1.4825e-04) (hash(x)=6451776) -106 train 8.139659 (lr=1.4965e-04) (hash(x)=4940406) -107 train 8.260358 (lr=1.5105e-04) (hash(x)=4355733) -108 train 8.303938 (lr=1.5245e-04) (hash(x)=6186065) -109 train 8.337807 (lr=1.5385e-04) (hash(x)=6669188) -110 train 9.373795 (lr=1.5524e-04) (hash(x)=7118997) -111 train 8.518595 (lr=1.5664e-04) (hash(x)=6984772) -112 train 8.155628 (lr=1.5804e-04) (hash(x)=6472877) -113 train 8.311499 (lr=1.5944e-04) (hash(x)=7440580) -114 train 8.334013 (lr=1.6084e-04) (hash(x)=6378428) -115 train 8.339608 (lr=1.6224e-04) (hash(x)=6244096) -116 train 8.340956 (lr=1.6364e-04) (hash(x)=9100578) -117 train 8.612723 (lr=1.6503e-04) (hash(x)=8381560) -118 train 8.327225 (lr=1.6643e-04) (hash(x)=6066072) -119 train 8.059075 (lr=1.6783e-04) (hash(x)=6603717) -120 train 7.984727 (lr=1.6923e-04) (hash(x)=5138438) -121 train 7.869165 (lr=1.7063e-04) (hash(x)=6437516) -122 train 8.058741 (lr=1.7203e-04) (hash(x)=5743807) -123 train 7.859820 (lr=1.7343e-04) (hash(x)=5669522) -124 train 8.065757 (lr=1.7483e-04) (hash(x)=5963623) -125 train 7.815839 (lr=1.7622e-04) (hash(x)=5505889) -126 train 8.219624 (lr=1.7762e-04) (hash(x)=7607512) -127 train 7.993526 (lr=1.7902e-04) (hash(x)=6190579) -128 train 8.177776 (lr=1.8042e-04) (hash(x)=7444834) -129 train 8.087092 (lr=1.8182e-04) (hash(x)=5153196) -130 train 7.981976 (lr=1.8322e-04) (hash(x)=6495326) -131 train 8.088462 (lr=1.8462e-04) (hash(x)=4912983) -132 train 8.271357 (lr=1.8601e-04) (hash(x)=7094046) -133 train 7.855439 (lr=1.8741e-04) (hash(x)=5109125) -134 train 7.815718 (lr=1.8881e-04) (hash(x)=7119136) -135 train 7.861666 (lr=1.9021e-04) (hash(x)=7276303) -136 train 8.110420 (lr=1.9161e-04) (hash(x)=9147123) -137 train 8.102427 (lr=1.9301e-04) (hash(x)=6645716) -138 train 8.029046 (lr=1.9441e-04) (hash(x)=6347875) -139 train 7.806607 (lr=1.9580e-04) (hash(x)=5479947) -140 train 7.927036 (lr=1.9720e-04) (hash(x)=5832595) -141 train 8.217505 (lr=1.9860e-04) (hash(x)=7850360) -142 train 8.590468 (lr=2.0000e-04) (hash(x)=8499683) -143 train 8.117359 (lr=2.0140e-04) (hash(x)=6630378) -144 train 8.050580 (lr=2.0280e-04) (hash(x)=5340846) -145 train 7.878900 (lr=2.0420e-04) (hash(x)=6295358) -146 train 7.917219 (lr=2.0559e-04) (hash(x)=5974729) -147 train 8.450937 (lr=2.0699e-04) (hash(x)=8097010) -148 train 8.198786 (lr=2.0839e-04) (hash(x)=6681250) -149 train 7.617529 (lr=2.0979e-04) (hash(x)=3958188) -150 val loss 7.8662 -150 val perplexity 2607.7197 -150 train 7.818198 (lr=2.1119e-04) (hash(x)=5975539) -151 train 9.711157 (lr=2.1259e-04) (hash(x)=9989744) -152 train 8.821021 (lr=2.1399e-04) (hash(x)=6748487) -153 train 7.778756 (lr=2.1538e-04) (hash(x)=6343832) -154 train 7.855833 (lr=2.1678e-04) (hash(x)=5688949) -155 train 8.195263 (lr=2.1818e-04) (hash(x)=6694932) -156 train 7.828443 (lr=2.1958e-04) (hash(x)=5872835) -157 train 7.897624 (lr=2.2098e-04) (hash(x)=6429137) -158 train 8.120209 (lr=2.2238e-04) (hash(x)=6999390) -159 train 7.968429 (lr=2.2378e-04) (hash(x)=6317894) -160 train 7.690858 (lr=2.2517e-04) (hash(x)=5813657) -161 train 7.793367 (lr=2.2657e-04) (hash(x)=5165595) -162 train 7.734653 (lr=2.2797e-04) (hash(x)=5682633) -163 train 7.942009 (lr=2.2937e-04) (hash(x)=5046327) -164 train 8.024662 (lr=2.3077e-04) (hash(x)=6195808) -165 train 7.881701 (lr=2.3217e-04) (hash(x)=7060363) -166 train 7.873632 (lr=2.3357e-04) (hash(x)=5837126) -167 train 7.847794 (lr=2.3497e-04) (hash(x)=7947703) -168 train 7.484069 (lr=2.3636e-04) (hash(x)=5277865) -169 train 7.740170 (lr=2.3776e-04) (hash(x)=5971240) -170 train 7.682817 (lr=2.3916e-04) (hash(x)=6590408) -171 train 7.842447 (lr=2.4056e-04) (hash(x)=5916068) -172 train 7.779253 (lr=2.4196e-04) (hash(x)=6069727) -173 train 7.809871 (lr=2.4336e-04) (hash(x)=7426277) -174 train 7.972969 (lr=2.4476e-04) (hash(x)=5356513) -175 train 7.737049 (lr=2.4615e-04) (hash(x)=5777498) -176 train 8.004381 (lr=2.4755e-04) (hash(x)=5989756) -177 train 7.739745 (lr=2.4895e-04) (hash(x)=6403331) -178 train 7.780223 (lr=2.5035e-04) (hash(x)=7650667) -179 train 7.965009 (lr=2.5175e-04) (hash(x)=5753092) -180 train 7.714940 (lr=2.5315e-04) (hash(x)=5751457) -181 train 7.589414 (lr=2.5455e-04) (hash(x)=5981486) -182 train 7.635160 (lr=2.5594e-04) (hash(x)=5918229) -183 train 7.883895 (lr=2.5734e-04) (hash(x)=7723226) -184 train 7.962802 (lr=2.5874e-04) (hash(x)=6111721) -185 train 7.578727 (lr=2.6014e-04) (hash(x)=6069558) -186 train 7.837216 (lr=2.6154e-04) (hash(x)=5626336) -187 train 8.072463 (lr=2.6294e-04) (hash(x)=7007016) -188 train 7.572512 (lr=2.6434e-04) (hash(x)=5883013) -189 train 7.457197 (lr=2.6573e-04) (hash(x)=5027958) -190 train 7.813180 (lr=2.6713e-04) (hash(x)=5474948) -191 train 7.631199 (lr=2.6853e-04) (hash(x)=6491229) -192 train 7.726456 (lr=2.6993e-04) (hash(x)=5518341) -193 train 7.834788 (lr=2.7133e-04) (hash(x)=7036653) -194 train 7.817052 (lr=2.7273e-04) (hash(x)=5527616) -195 train 8.100171 (lr=2.7413e-04) (hash(x)=6390752) -196 train 7.767852 (lr=2.7552e-04) (hash(x)=6032557) -197 train 8.478538 (lr=2.7692e-04) (hash(x)=7758167) -198 train 8.345364 (lr=2.7832e-04) (hash(x)=7968070) -199 train 7.871463 (lr=2.7972e-04) (hash(x)=6019389) -200 val loss 7.7821 -200 val perplexity 2397.2014 -200 train 7.883311 (lr=2.8112e-04) (hash(x)=7276744) -201 train 7.909128 (lr=2.8252e-04) (hash(x)=6503191) -202 train 7.744545 (lr=2.8392e-04) (hash(x)=6006880) -203 train 7.957067 (lr=2.8531e-04) (hash(x)=7662067) -204 train 7.719955 (lr=2.8671e-04) (hash(x)=6297345) -205 train 7.953323 (lr=2.8811e-04) (hash(x)=7901992) -206 train 8.076226 (lr=2.8951e-04) (hash(x)=6579950) -207 train 7.601086 (lr=2.9091e-04) (hash(x)=4648609) -208 train 7.848987 (lr=2.9231e-04) (hash(x)=6903216) -209 train 7.687524 (lr=2.9371e-04) (hash(x)=5897288) -210 train 7.758507 (lr=2.9510e-04) (hash(x)=7300160) -211 train 7.608027 (lr=2.9650e-04) (hash(x)=4725966) -212 train 7.430196 (lr=2.9790e-04) (hash(x)=5060060) -213 train 7.959236 (lr=2.9930e-04) (hash(x)=6243442) -214 train 7.658198 (lr=3.0070e-04) (hash(x)=5893816) -215 train 7.678847 (lr=3.0210e-04) (hash(x)=5558355) -216 train 7.831263 (lr=3.0350e-04) (hash(x)=4378747) -217 train 7.689771 (lr=3.0490e-04) (hash(x)=5400407) -218 train 7.959946 (lr=3.0629e-04) (hash(x)=6900554) -219 train 7.831897 (lr=3.0769e-04) (hash(x)=6524933) -220 train 7.763245 (lr=3.0909e-04) (hash(x)=5409944) -221 train 7.847879 (lr=3.1049e-04) (hash(x)=5889724) -222 train 7.771236 (lr=3.1189e-04) (hash(x)=4970496) -223 train 7.565698 (lr=3.1329e-04) (hash(x)=6369326) -224 train 7.664714 (lr=3.1469e-04) (hash(x)=6563975) -225 train 7.846850 (lr=3.1608e-04) (hash(x)=5911906) -226 train 7.939465 (lr=3.1748e-04) (hash(x)=6111462) -227 train 7.945687 (lr=3.1888e-04) (hash(x)=6022625) -228 train 7.745234 (lr=3.2028e-04) (hash(x)=5884663) -229 train 7.843312 (lr=3.2168e-04) (hash(x)=6584810) -230 train 7.801966 (lr=3.2308e-04) (hash(x)=5289998) -231 train 7.510679 (lr=3.2448e-04) (hash(x)=4906853) -232 train 7.579675 (lr=3.2587e-04) (hash(x)=5862071) -233 train 7.219170 (lr=3.2727e-04) (hash(x)=4564564) -234 train 7.591276 (lr=3.2867e-04) (hash(x)=5723971) -235 train 7.708779 (lr=3.3007e-04) (hash(x)=5589269) -236 train 7.665421 (lr=3.3147e-04) (hash(x)=5683368) -237 train 7.913085 (lr=3.3287e-04) (hash(x)=6969180) -238 train 7.796215 (lr=3.3427e-04) (hash(x)=5341421) -239 train 7.567434 (lr=3.3566e-04) (hash(x)=6437376) -240 train 8.076303 (lr=3.3706e-04) (hash(x)=7527007) -241 train 7.919109 (lr=3.3846e-04) (hash(x)=7116258) -242 train 7.877451 (lr=3.3986e-04) (hash(x)=6463930) -243 train 7.907961 (lr=3.4126e-04) (hash(x)=7304892) -244 train 8.047992 (lr=3.4266e-04) (hash(x)=5881808) -245 train 7.837342 (lr=3.4406e-04) (hash(x)=7366721) -246 train 7.760063 (lr=3.4545e-04) (hash(x)=6457521) -247 train 8.144821 (lr=3.4685e-04) (hash(x)=7084093) -248 train 7.808409 (lr=3.4825e-04) (hash(x)=6015300) -249 train 8.056373 (lr=3.4965e-04) (hash(x)=7197652) -250 val loss 7.7773 -250 val perplexity 2385.8423 -250 train 8.080597 (lr=3.5105e-04) (hash(x)=4830828) -251 train 7.805977 (lr=3.5245e-04) (hash(x)=5908178) -252 train 7.789102 (lr=3.5385e-04) (hash(x)=6299092) -253 train 7.798859 (lr=3.5524e-04) (hash(x)=6349974) -254 train 7.802348 (lr=3.5664e-04) (hash(x)=5641494) -255 train 7.781238 (lr=3.5804e-04) (hash(x)=7048804) -256 train 8.006168 (lr=3.5944e-04) (hash(x)=7035244) -257 train 7.726725 (lr=3.6084e-04) (hash(x)=7909039) -258 train 7.948016 (lr=3.6224e-04) (hash(x)=5470939) -259 train 7.773250 (lr=3.6364e-04) (hash(x)=6085549) -260 train 8.079906 (lr=3.6503e-04) (hash(x)=5882649) -261 train 7.917882 (lr=3.6643e-04) (hash(x)=7463181) -262 train 8.210739 (lr=3.6783e-04) (hash(x)=5215574) -263 train 7.656866 (lr=3.6923e-04) (hash(x)=5752594) -264 train 7.814951 (lr=3.7063e-04) (hash(x)=6416619) -265 train 7.814353 (lr=3.7203e-04) (hash(x)=6114601) -266 train 7.759092 (lr=3.7343e-04) (hash(x)=5646171) -267 train 7.688155 (lr=3.7483e-04) (hash(x)=7662769) -268 train 7.963375 (lr=3.7622e-04) (hash(x)=6394376) -269 train 7.553571 (lr=3.7762e-04) (hash(x)=7485666) -270 train 7.640589 (lr=3.7902e-04) (hash(x)=6636823) -271 train 7.575523 (lr=3.8042e-04) (hash(x)=6393520) -272 train 7.766035 (lr=3.8182e-04) (hash(x)=5235659) -273 train 7.842965 (lr=3.8322e-04) (hash(x)=5586291) -274 train 8.059671 (lr=3.8462e-04) (hash(x)=7034674) -275 train 8.154101 (lr=3.8601e-04) (hash(x)=5942867) -276 train 7.854487 (lr=3.8741e-04) (hash(x)=7812344) -277 train 8.128255 (lr=3.8881e-04) (hash(x)=7574027) -278 train 8.080032 (lr=3.9021e-04) (hash(x)=5960216) -279 train 7.866420 (lr=3.9161e-04) (hash(x)=6793550) -280 train 7.837950 (lr=3.9301e-04) (hash(x)=6879965) -281 train 7.543940 (lr=3.9441e-04) (hash(x)=6264789) -282 train 7.698717 (lr=3.9580e-04) (hash(x)=5175356) -283 train 7.926527 (lr=3.9720e-04) (hash(x)=7105976) -284 train 7.934340 (lr=3.9860e-04) (hash(x)=7286386) -285 train 8.124363 (lr=4.0000e-04) (hash(x)=7244279) -286 train 7.606567 (lr=4.0000e-04) (hash(x)=4877247) -287 train 7.569118 (lr=4.0000e-04) (hash(x)=6581348) -288 train 7.564267 (lr=3.9999e-04) (hash(x)=6397208) -289 train 7.678329 (lr=3.9998e-04) (hash(x)=6543825) -290 train 7.451445 (lr=3.9997e-04) (hash(x)=6088993) -291 train 7.783363 (lr=3.9996e-04) (hash(x)=5555598) -292 train 7.851149 (lr=3.9994e-04) (hash(x)=6924240) -293 train 8.049478 (lr=3.9991e-04) (hash(x)=7254232) -294 train 7.581804 (lr=3.9989e-04) (hash(x)=7689743) -295 train 7.550489 (lr=3.9986e-04) (hash(x)=6235837) -296 train 7.804015 (lr=3.9983e-04) (hash(x)=5216883) -297 train 7.435771 (lr=3.9979e-04) (hash(x)=4738550) -298 train 7.616967 (lr=3.9975e-04) (hash(x)=5126918) -299 train 7.774863 (lr=3.9971e-04) (hash(x)=5591770) -300 val loss 7.7129 -300 val perplexity 2237.1272 -300 train 8.668016 (lr=3.9966e-04) (hash(x)=8001573) -301 train 8.370278 (lr=3.9961e-04) (hash(x)=8012491) -302 train 8.050730 (lr=3.9955e-04) (hash(x)=3997428) -303 train 7.358860 (lr=3.9950e-04) (hash(x)=7083021) -304 train 7.944624 (lr=3.9944e-04) (hash(x)=7177371) -305 train 7.770327 (lr=3.9937e-04) (hash(x)=6456022) -306 train 7.872246 (lr=3.9930e-04) (hash(x)=5238861) -307 train 7.862182 (lr=3.9923e-04) (hash(x)=7281348) -308 train 7.997207 (lr=3.9916e-04) (hash(x)=6848441) -309 train 8.047302 (lr=3.9908e-04) (hash(x)=7824794) -310 train 8.186137 (lr=3.9900e-04) (hash(x)=7440208) -311 train 7.972701 (lr=3.9891e-04) (hash(x)=6155821) -312 train 7.693457 (lr=3.9882e-04) (hash(x)=5350461) -313 train 7.940413 (lr=3.9873e-04) (hash(x)=7807564) -314 train 7.728491 (lr=3.9864e-04) (hash(x)=5992165) -315 train 7.688345 (lr=3.9854e-04) (hash(x)=5736241) -316 train 7.666516 (lr=3.9843e-04) (hash(x)=6536253) -317 train 7.865399 (lr=3.9833e-04) (hash(x)=5688415) -318 train 7.497298 (lr=3.9822e-04) (hash(x)=5839705) -319 train 7.677273 (lr=3.9811e-04) (hash(x)=5657123) -320 train 7.717196 (lr=3.9799e-04) (hash(x)=6354364) -321 train 7.809872 (lr=3.9787e-04) (hash(x)=6102985) -322 train 8.071723 (lr=3.9775e-04) (hash(x)=7161967) -323 train 8.007324 (lr=3.9762e-04) (hash(x)=6452095) -324 train 7.631840 (lr=3.9749e-04) (hash(x)=5547396) -325 train 7.940474 (lr=3.9736e-04) (hash(x)=5848921) -326 train 7.720636 (lr=3.9722e-04) (hash(x)=4890294) -327 train 7.432765 (lr=3.9708e-04) (hash(x)=5312267) -328 train 7.513749 (lr=3.9694e-04) (hash(x)=5348929) -329 train 7.471559 (lr=3.9679e-04) (hash(x)=5025297) -330 train 7.505089 (lr=3.9664e-04) (hash(x)=5670052) -331 train 7.356798 (lr=3.9648e-04) (hash(x)=5434493) -332 train 8.054082 (lr=3.9633e-04) (hash(x)=9122453) -333 train 7.766166 (lr=3.9616e-04) (hash(x)=5197474) -334 train 7.603662 (lr=3.9600e-04) (hash(x)=5550786) -335 train 7.666357 (lr=3.9583e-04) (hash(x)=6830813) -336 train 7.490107 (lr=3.9566e-04) (hash(x)=5258435) -337 train 7.524698 (lr=3.9549e-04) (hash(x)=5242548) -338 train 7.362160 (lr=3.9531e-04) (hash(x)=3754886) -339 train 7.215922 (lr=3.9513e-04) (hash(x)=4752771) -340 train 7.464192 (lr=3.9494e-04) (hash(x)=5926875) -341 train 8.004081 (lr=3.9475e-04) (hash(x)=6691249) -342 train 7.723438 (lr=3.9456e-04) (hash(x)=7514623) -343 train 7.812165 (lr=3.9437e-04) (hash(x)=6424933) -344 train 7.747590 (lr=3.9417e-04) (hash(x)=6347757) -345 train 7.708298 (lr=3.9397e-04) (hash(x)=6960120) -346 train 7.726888 (lr=3.9376e-04) (hash(x)=7597197) -347 train 7.654065 (lr=3.9356e-04) (hash(x)=5786517) -348 train 7.732175 (lr=3.9334e-04) (hash(x)=6424363) -349 train 7.754531 (lr=3.9313e-04) (hash(x)=7229689) -350 val loss 7.6790 -350 val perplexity 2162.5024 -350 train 7.852203 (lr=3.9291e-04) (hash(x)=6772284) -351 train 7.939065 (lr=3.9269e-04) (hash(x)=6680023) -352 train 7.837459 (lr=3.9246e-04) (hash(x)=6218520) -353 train 7.896614 (lr=3.9223e-04) (hash(x)=7688393) -354 train 7.720498 (lr=3.9200e-04) (hash(x)=6327977) -355 train 7.606984 (lr=3.9177e-04) (hash(x)=6474729) -356 train 7.730735 (lr=3.9153e-04) (hash(x)=6657266) -357 train 7.744442 (lr=3.9129e-04) (hash(x)=5946650) -358 train 7.644190 (lr=3.9104e-04) (hash(x)=6734878) -359 train 7.726983 (lr=3.9079e-04) (hash(x)=7523279) -360 train 7.573905 (lr=3.9054e-04) (hash(x)=6179478) -361 train 7.626390 (lr=3.9029e-04) (hash(x)=5531605) -362 train 7.613953 (lr=3.9003e-04) (hash(x)=5844980) -363 train 7.281720 (lr=3.8977e-04) (hash(x)=5508768) -364 train 7.134599 (lr=3.8950e-04) (hash(x)=6337190) -365 train 8.774722 (lr=3.8923e-04) (hash(x)=6415982) -366 train 7.787969 (lr=3.8896e-04) (hash(x)=6828959) -367 train 7.746159 (lr=3.8869e-04) (hash(x)=6635925) -368 train 7.471364 (lr=3.8841e-04) (hash(x)=6358540) -369 train 7.853199 (lr=3.8813e-04) (hash(x)=5923706) -370 train 7.516103 (lr=3.8785e-04) (hash(x)=4843320) -371 train 7.742843 (lr=3.8756e-04) (hash(x)=6663801) -372 train 7.572280 (lr=3.8727e-04) (hash(x)=5583287) -373 train 7.525884 (lr=3.8697e-04) (hash(x)=6113540) -374 train 7.653235 (lr=3.8667e-04) (hash(x)=6389936) -375 train 7.604330 (lr=3.8637e-04) (hash(x)=5869441) -376 train 7.666893 (lr=3.8607e-04) (hash(x)=5148617) -377 train 7.929458 (lr=3.8576e-04) (hash(x)=6913454) -378 train 7.606462 (lr=3.8545e-04) (hash(x)=5356233) -379 train 7.587533 (lr=3.8514e-04) (hash(x)=6698878) -380 train 7.561080 (lr=3.8482e-04) (hash(x)=6280665) -381 train 7.779469 (lr=3.8450e-04) (hash(x)=6094899) -382 train 7.309037 (lr=3.8418e-04) (hash(x)=6249746) -383 train 7.636660 (lr=3.8385e-04) (hash(x)=6349004) -384 train 7.577253 (lr=3.8352e-04) (hash(x)=7013468) -385 train 7.548975 (lr=3.8319e-04) (hash(x)=6045933) -386 train 7.871435 (lr=3.8286e-04) (hash(x)=6571735) -387 train 7.498989 (lr=3.8252e-04) (hash(x)=4986137) -388 train 7.167762 (lr=3.8217e-04) (hash(x)=5244912) -389 train 7.498528 (lr=3.8183e-04) (hash(x)=4798229) -390 train 7.577496 (lr=3.8148e-04) (hash(x)=5815783) -391 train 7.346290 (lr=3.8113e-04) (hash(x)=6008454) -392 train 7.551310 (lr=3.8077e-04) (hash(x)=6407333) -393 train 7.436758 (lr=3.8042e-04) (hash(x)=5938362) -394 train 7.289711 (lr=3.8006e-04) (hash(x)=6077124) -395 train 7.600750 (lr=3.7969e-04) (hash(x)=6550770) -396 train 7.420817 (lr=3.7933e-04) (hash(x)=6181528) -397 train 7.688013 (lr=3.7896e-04) (hash(x)=7055344) -398 train 7.801223 (lr=3.7858e-04) (hash(x)=7689348) -399 train 7.725225 (lr=3.7821e-04) (hash(x)=7682741) -400 val loss 7.6094 -400 val perplexity 2017.1411 -400 train 7.511962 (lr=3.7783e-04) (hash(x)=5189401) -401 train 7.518741 (lr=3.7744e-04) (hash(x)=7372095) -402 train 7.668220 (lr=3.7706e-04) (hash(x)=6838219) -403 train 8.134772 (lr=3.7667e-04) (hash(x)=7892158) -404 train 7.775495 (lr=3.7628e-04) (hash(x)=5677077) -405 train 7.401995 (lr=3.7588e-04) (hash(x)=5929740) -406 train 7.558940 (lr=3.7549e-04) (hash(x)=5285055) -407 train 7.650312 (lr=3.7509e-04) (hash(x)=7794028) -408 train 7.726353 (lr=3.7468e-04) (hash(x)=6640802) -409 train 7.822291 (lr=3.7428e-04) (hash(x)=6839923) -410 train 7.277285 (lr=3.7387e-04) (hash(x)=4020997) -411 train 7.458729 (lr=3.7345e-04) (hash(x)=7093523) -412 train 7.784377 (lr=3.7304e-04) (hash(x)=7370045) -413 train 7.507353 (lr=3.7262e-04) (hash(x)=5341312) -414 train 7.512074 (lr=3.7220e-04) (hash(x)=5589312) -415 train 7.986567 (lr=3.7177e-04) (hash(x)=6166062) -416 train 7.887094 (lr=3.7135e-04) (hash(x)=8354149) -417 train 7.509351 (lr=3.7092e-04) (hash(x)=5992164) -418 train 7.547449 (lr=3.7048e-04) (hash(x)=4748657) -419 train 7.541067 (lr=3.7005e-04) (hash(x)=6645781) -420 train 7.770974 (lr=3.6961e-04) (hash(x)=6689147) -421 train 7.642218 (lr=3.6917e-04) (hash(x)=6207724) -422 train 7.557092 (lr=3.6872e-04) (hash(x)=6232863) -423 train 7.306893 (lr=3.6828e-04) (hash(x)=4811192) -424 train 7.341501 (lr=3.6782e-04) (hash(x)=6211761) -425 train 7.530269 (lr=3.6737e-04) (hash(x)=6906335) -426 train 7.448237 (lr=3.6692e-04) (hash(x)=4922365) -427 train 7.491903 (lr=3.6646e-04) (hash(x)=5970866) -428 train 7.444016 (lr=3.6599e-04) (hash(x)=6417833) -429 train 7.543698 (lr=3.6553e-04) (hash(x)=5238015) -430 train 7.432292 (lr=3.6506e-04) (hash(x)=5375451) -431 train 7.556702 (lr=3.6459e-04) (hash(x)=7822680) -432 train 7.571167 (lr=3.6412e-04) (hash(x)=6173200) -433 train 7.307553 (lr=3.6364e-04) (hash(x)=5478019) -434 train 7.529016 (lr=3.6316e-04) (hash(x)=5597859) -435 train 7.444973 (lr=3.6268e-04) (hash(x)=6317317) -436 train 7.626946 (lr=3.6220e-04) (hash(x)=7120467) -437 train 7.724288 (lr=3.6171e-04) (hash(x)=5832344) -438 train 7.608081 (lr=3.6122e-04) (hash(x)=6813802) -439 train 7.667047 (lr=3.6073e-04) (hash(x)=6431409) -440 train 7.744073 (lr=3.6023e-04) (hash(x)=5886861) -441 train 7.410055 (lr=3.5974e-04) (hash(x)=5900130) -442 train 7.698512 (lr=3.5924e-04) (hash(x)=6061563) -443 train 7.647983 (lr=3.5873e-04) (hash(x)=6653337) -444 train 7.555354 (lr=3.5823e-04) (hash(x)=6575992) -445 train 7.568292 (lr=3.5772e-04) (hash(x)=6696784) -446 train 7.717435 (lr=3.5721e-04) (hash(x)=6242161) -447 train 7.539690 (lr=3.5669e-04) (hash(x)=5323032) -448 train 7.336827 (lr=3.5618e-04) (hash(x)=6439367) -449 train 7.581370 (lr=3.5566e-04) (hash(x)=5372100) -450 val loss 7.5729 -450 val perplexity 1944.7410 -450 train 7.643305 (lr=3.5514e-04) (hash(x)=6344849) -451 train 7.298000 (lr=3.5461e-04) (hash(x)=5125339) -452 train 7.260176 (lr=3.5408e-04) (hash(x)=4876156) -453 train 7.400402 (lr=3.5355e-04) (hash(x)=5664463) -454 train 7.294392 (lr=3.5302e-04) (hash(x)=6010042) -455 train 7.540538 (lr=3.5249e-04) (hash(x)=7363286) -456 train 8.173691 (lr=3.5195e-04) (hash(x)=4873268) -457 train 7.884413 (lr=3.5141e-04) (hash(x)=7444468) -458 train 7.905367 (lr=3.5087e-04) (hash(x)=6698134) -459 train 7.703500 (lr=3.5032e-04) (hash(x)=7670050) -460 train 7.708179 (lr=3.4977e-04) (hash(x)=6280866) -461 train 7.541913 (lr=3.4922e-04) (hash(x)=8189420) -462 train 7.666384 (lr=3.4867e-04) (hash(x)=6878292) -463 train 7.350158 (lr=3.4812e-04) (hash(x)=5616075) -464 train 7.556487 (lr=3.4756e-04) (hash(x)=5839939) -465 train 7.437470 (lr=3.4700e-04) (hash(x)=6360924) -466 train 7.913586 (lr=3.4644e-04) (hash(x)=6466475) -467 train 7.657649 (lr=3.4587e-04) (hash(x)=6593764) -468 train 8.092769 (lr=3.4530e-04) (hash(x)=7389698) -469 train 7.748988 (lr=3.4473e-04) (hash(x)=6764377) -470 train 7.509034 (lr=3.4416e-04) (hash(x)=7293570) -471 train 7.544131 (lr=3.4359e-04) (hash(x)=5672608) -472 train 7.408926 (lr=3.4301e-04) (hash(x)=5587975) -473 train 7.305445 (lr=3.4243e-04) (hash(x)=6985256) -474 train 7.932063 (lr=3.4185e-04) (hash(x)=6549231) -475 train 7.584744 (lr=3.4127e-04) (hash(x)=6828653) -476 train 7.410558 (lr=3.4068e-04) (hash(x)=5478970) -477 train 7.449081 (lr=3.4009e-04) (hash(x)=5199845) -478 train 7.520986 (lr=3.3950e-04) (hash(x)=6482073) -479 train 7.604218 (lr=3.3891e-04) (hash(x)=6839867) -480 train 7.603945 (lr=3.3831e-04) (hash(x)=7073454) -481 train 7.904164 (lr=3.3771e-04) (hash(x)=6928712) -482 train 7.573225 (lr=3.3711e-04) (hash(x)=5962982) -483 train 7.646687 (lr=3.3651e-04) (hash(x)=8426992) -484 train 7.570432 (lr=3.3590e-04) (hash(x)=6422394) -485 train 7.563037 (lr=3.3530e-04) (hash(x)=5330355) -486 train 7.035078 (lr=3.3469e-04) (hash(x)=4111491) -487 train 7.035528 (lr=3.3408e-04) (hash(x)=4916928) -488 train 7.207923 (lr=3.3346e-04) (hash(x)=6610244) -489 train 7.738671 (lr=3.3285e-04) (hash(x)=7287928) -490 train 7.586032 (lr=3.3223e-04) (hash(x)=6525083) -491 train 7.549232 (lr=3.3161e-04) (hash(x)=6484050) -492 train 7.689845 (lr=3.3099e-04) (hash(x)=6573424) -493 train 7.540273 (lr=3.3036e-04) (hash(x)=5062850) -494 train 7.517158 (lr=3.2973e-04) (hash(x)=6215995) -495 train 7.628617 (lr=3.2911e-04) (hash(x)=8353379) -496 train 7.506550 (lr=3.2847e-04) (hash(x)=6344823) -497 train 7.503284 (lr=3.2784e-04) (hash(x)=6805358) -498 train 7.402885 (lr=3.2721e-04) (hash(x)=6018070) -499 train 7.670665 (lr=3.2657e-04) (hash(x)=6552510) -500 val loss 7.5385 -500 val perplexity 1878.9983 -500 train 7.472359 (lr=3.2593e-04) (hash(x)=7218108) -501 train 7.663666 (lr=3.2529e-04) (hash(x)=5105545) -502 train 7.650087 (lr=3.2464e-04) (hash(x)=5997860) -503 train 7.435174 (lr=3.2400e-04) (hash(x)=4838871) -504 train 7.471886 (lr=3.2335e-04) (hash(x)=7105888) -505 train 7.803447 (lr=3.2270e-04) (hash(x)=5945286) -506 train 6.972367 (lr=3.2205e-04) (hash(x)=4381720) -507 train 6.651798 (lr=3.2140e-04) (hash(x)=3021697) -508 train 7.056805 (lr=3.2074e-04) (hash(x)=6826973) -509 train 7.653838 (lr=3.2008e-04) (hash(x)=6443070) -510 train 7.341478 (lr=3.1943e-04) (hash(x)=5520637) -511 train 7.597086 (lr=3.1876e-04) (hash(x)=6795665) -512 train 7.391040 (lr=3.1810e-04) (hash(x)=6063891) -513 train 7.892898 (lr=3.1744e-04) (hash(x)=7683068) -514 train 7.155515 (lr=3.1677e-04) (hash(x)=5526694) -515 train 7.179341 (lr=3.1610e-04) (hash(x)=5486935) -516 train 7.567066 (lr=3.1543e-04) (hash(x)=6002252) -517 train 7.534157 (lr=3.1476e-04) (hash(x)=6115232) -518 train 7.810001 (lr=3.1408e-04) (hash(x)=9039136) -519 train 7.467688 (lr=3.1341e-04) (hash(x)=6678038) -520 train 7.619728 (lr=3.1273e-04) (hash(x)=5383926) -521 train 7.672162 (lr=3.1205e-04) (hash(x)=5747629) -522 train 7.557688 (lr=3.1137e-04) (hash(x)=7090040) -523 train 7.751288 (lr=3.1069e-04) (hash(x)=6657714) -524 train 7.519429 (lr=3.1000e-04) (hash(x)=5414034) -525 train 7.539899 (lr=3.0931e-04) (hash(x)=6190642) -526 train 7.662959 (lr=3.0862e-04) (hash(x)=7431177) -527 train 7.527627 (lr=3.0793e-04) (hash(x)=6112215) -528 train 7.646883 (lr=3.0724e-04) (hash(x)=6785874) -529 train 7.409170 (lr=3.0655e-04) (hash(x)=7314989) -530 train 7.627702 (lr=3.0585e-04) (hash(x)=6222614) -531 train 7.712043 (lr=3.0516e-04) (hash(x)=8353143) -532 train 7.794209 (lr=3.0446e-04) (hash(x)=6752498) -533 train 7.523862 (lr=3.0376e-04) (hash(x)=5912570) -534 train 7.683999 (lr=3.0306e-04) (hash(x)=5621785) -535 train 7.583555 (lr=3.0235e-04) (hash(x)=5915361) -536 train 7.408762 (lr=3.0165e-04) (hash(x)=6853672) -537 train 7.357994 (lr=3.0094e-04) (hash(x)=6369494) -538 train 7.551110 (lr=3.0023e-04) (hash(x)=6039652) -539 train 7.680251 (lr=2.9952e-04) (hash(x)=6254885) -540 train 7.523490 (lr=2.9881e-04) (hash(x)=5829860) -541 train 7.459058 (lr=2.9810e-04) (hash(x)=5425120) -542 train 7.214892 (lr=2.9738e-04) (hash(x)=5209746) -543 train 7.830770 (lr=2.9667e-04) (hash(x)=5771588) -544 train 7.543005 (lr=2.9595e-04) (hash(x)=8298790) -545 train 7.579824 (lr=2.9523e-04) (hash(x)=6763967) -546 train 7.572336 (lr=2.9451e-04) (hash(x)=4882397) -547 train 7.336516 (lr=2.9379e-04) (hash(x)=5561507) -548 train 7.339019 (lr=2.9307e-04) (hash(x)=5762755) -549 train 7.460457 (lr=2.9234e-04) (hash(x)=7081151) -550 val loss 7.4936 -550 val perplexity 1796.4297 -550 train 7.356226 (lr=2.9162e-04) (hash(x)=5047520) -551 train 7.495140 (lr=2.9089e-04) (hash(x)=5688829) -552 train 7.373357 (lr=2.9016e-04) (hash(x)=7008925) -553 train 7.523464 (lr=2.8943e-04) (hash(x)=7821554) -554 train 7.538155 (lr=2.8870e-04) (hash(x)=6469035) -555 train 7.351624 (lr=2.8797e-04) (hash(x)=5371951) -556 train 7.655009 (lr=2.8723e-04) (hash(x)=6436050) -557 train 7.574912 (lr=2.8650e-04) (hash(x)=8120149) -558 train 7.821643 (lr=2.8576e-04) (hash(x)=6189935) -559 train 7.515394 (lr=2.8502e-04) (hash(x)=5443305) -560 train 7.617784 (lr=2.8428e-04) (hash(x)=6294312) -561 train 7.236239 (lr=2.8354e-04) (hash(x)=7066030) -562 train 7.650414 (lr=2.8280e-04) (hash(x)=7744989) -563 train 7.790557 (lr=2.8206e-04) (hash(x)=5016757) -564 train 7.587486 (lr=2.8132e-04) (hash(x)=8001284) -565 train 7.524496 (lr=2.8057e-04) (hash(x)=6296163) -566 train 7.660974 (lr=2.7982e-04) (hash(x)=6027688) -567 train 7.569398 (lr=2.7908e-04) (hash(x)=6901933) -568 train 7.437927 (lr=2.7833e-04) (hash(x)=5124305) -569 train 7.547812 (lr=2.7758e-04) (hash(x)=8056633) -570 train 7.617347 (lr=2.7683e-04) (hash(x)=6677566) -571 train 7.645174 (lr=2.7607e-04) (hash(x)=6019085) -572 train 7.398640 (lr=2.7532e-04) (hash(x)=5924495) -573 train 7.496363 (lr=2.7457e-04) (hash(x)=7003893) -574 train 7.461383 (lr=2.7381e-04) (hash(x)=5563075) -575 train 7.270102 (lr=2.7306e-04) (hash(x)=4294425) -576 train 7.459836 (lr=2.7230e-04) (hash(x)=5677870) -577 train 7.345787 (lr=2.7154e-04) (hash(x)=6568540) -578 train 7.579265 (lr=2.7078e-04) (hash(x)=7473394) -579 train 7.606874 (lr=2.7002e-04) (hash(x)=6095229) -580 train 7.610292 (lr=2.6926e-04) (hash(x)=6946168) -581 train 7.416889 (lr=2.6850e-04) (hash(x)=5512150) -582 train 7.065933 (lr=2.6773e-04) (hash(x)=4380268) -583 train 7.130316 (lr=2.6697e-04) (hash(x)=4363941) -584 train 7.292055 (lr=2.6620e-04) (hash(x)=4836009) -585 train 7.551021 (lr=2.6544e-04) (hash(x)=6457783) -586 train 7.422904 (lr=2.6467e-04) (hash(x)=6378784) -587 train 7.336568 (lr=2.6390e-04) (hash(x)=4693798) -588 train 7.671266 (lr=2.6314e-04) (hash(x)=9193310) -589 train 8.412244 (lr=2.6237e-04) (hash(x)=8759806) -590 train 8.784805 (lr=2.6160e-04) (hash(x)=9176140) -591 train 7.467399 (lr=2.6083e-04) (hash(x)=5486335) -592 train 7.676648 (lr=2.6005e-04) (hash(x)=7868436) -593 train 7.575103 (lr=2.5928e-04) (hash(x)=6964993) -594 train 7.791656 (lr=2.5851e-04) (hash(x)=6741270) -595 train 8.048973 (lr=2.5773e-04) (hash(x)=7907450) -596 train 7.905000 (lr=2.5696e-04) (hash(x)=6443452) -597 train 7.502116 (lr=2.5618e-04) (hash(x)=5236642) -598 train 7.662479 (lr=2.5541e-04) (hash(x)=5604795) -599 train 7.613657 (lr=2.5463e-04) (hash(x)=7295165) -600 val loss 7.5941 -600 val perplexity 1986.3958 -600 train 7.598695 (lr=2.5385e-04) (hash(x)=6034549) -601 train 7.626319 (lr=2.5307e-04) (hash(x)=4913602) -602 train 7.583829 (lr=2.5230e-04) (hash(x)=6271294) -603 train 7.656682 (lr=2.5152e-04) (hash(x)=6814026) -604 train 7.544633 (lr=2.5074e-04) (hash(x)=6848040) -605 train 7.775752 (lr=2.4996e-04) (hash(x)=7042051) -606 train 7.718560 (lr=2.4917e-04) (hash(x)=6481002) -607 train 7.384568 (lr=2.4839e-04) (hash(x)=6267424) -608 train 7.670825 (lr=2.4761e-04) (hash(x)=7665306) -609 train 7.614493 (lr=2.4683e-04) (hash(x)=5614727) -610 train 7.607079 (lr=2.4604e-04) (hash(x)=7039197) -611 train 7.670383 (lr=2.4526e-04) (hash(x)=8086437) -612 train 7.734481 (lr=2.4448e-04) (hash(x)=6993846) -613 train 7.768439 (lr=2.4369e-04) (hash(x)=5900143) -614 train 7.634926 (lr=2.4291e-04) (hash(x)=4841318) -615 train 7.609076 (lr=2.4212e-04) (hash(x)=5270452) -616 train 7.519147 (lr=2.4133e-04) (hash(x)=5955026) -617 train 7.589430 (lr=2.4055e-04) (hash(x)=8617707) -618 train 7.421934 (lr=2.3976e-04) (hash(x)=5159401) -619 train 7.433770 (lr=2.3897e-04) (hash(x)=6420820) -620 train 7.409893 (lr=2.3818e-04) (hash(x)=6628863) -621 train 7.219546 (lr=2.3740e-04) (hash(x)=6215930) -622 train 7.556403 (lr=2.3661e-04) (hash(x)=5248645) -623 train 7.451512 (lr=2.3582e-04) (hash(x)=6305297) -624 train 8.019797 (lr=2.3503e-04) (hash(x)=6989704) -625 train 7.530154 (lr=2.3424e-04) (hash(x)=4026682) -626 train 7.409697 (lr=2.3345e-04) (hash(x)=6367688) -627 train 7.472428 (lr=2.3266e-04) (hash(x)=7889849) -628 train 7.642259 (lr=2.3187e-04) (hash(x)=7334719) -629 train 7.769471 (lr=2.3108e-04) (hash(x)=6141721) -630 train 8.547319 (lr=2.3029e-04) (hash(x)=5056572) -631 train 7.972347 (lr=2.2950e-04) (hash(x)=6040077) -632 train 8.060642 (lr=2.2871e-04) (hash(x)=6866651) -633 train 8.293114 (lr=2.2792e-04) (hash(x)=6155449) -634 train 7.873876 (lr=2.2713e-04) (hash(x)=6690174) -635 train 7.580360 (lr=2.2633e-04) (hash(x)=5652497) -636 train 7.625008 (lr=2.2554e-04) (hash(x)=5931050) -637 train 7.704155 (lr=2.2475e-04) (hash(x)=6961314) -638 train 7.488312 (lr=2.2396e-04) (hash(x)=5973636) -639 train 7.489496 (lr=2.2317e-04) (hash(x)=7130251) -640 train 7.263834 (lr=2.2238e-04) (hash(x)=7503390) -641 train 7.272105 (lr=2.2158e-04) (hash(x)=5927461) -642 train 7.510231 (lr=2.2079e-04) (hash(x)=6196741) -643 train 7.378845 (lr=2.2000e-04) (hash(x)=6610177) -644 train 7.352052 (lr=2.1921e-04) (hash(x)=6635147) -645 train 7.656536 (lr=2.1842e-04) (hash(x)=7277580) -646 train 7.282983 (lr=2.1762e-04) (hash(x)=5050080) -647 train 7.363386 (lr=2.1683e-04) (hash(x)=6508350) -648 train 7.595304 (lr=2.1604e-04) (hash(x)=5276338) -649 train 7.743566 (lr=2.1525e-04) (hash(x)=6536034) -650 val loss 7.4993 -650 val perplexity 1806.7867 -650 train 7.540031 (lr=2.1446e-04) (hash(x)=6944772) -651 train 7.556524 (lr=2.1367e-04) (hash(x)=6994983) -652 train 7.513436 (lr=2.1287e-04) (hash(x)=7172017) -653 train 7.572359 (lr=2.1208e-04) (hash(x)=8700721) -654 train 7.562931 (lr=2.1129e-04) (hash(x)=6774360) -655 train 7.394775 (lr=2.1050e-04) (hash(x)=5859576) -656 train 7.122480 (lr=2.0971e-04) (hash(x)=5899275) -657 train 7.059240 (lr=2.0892e-04) (hash(x)=5264962) -658 train 6.860204 (lr=2.0813e-04) (hash(x)=5679861) -659 train 7.304075 (lr=2.0734e-04) (hash(x)=5487065) -660 train 7.066036 (lr=2.0655e-04) (hash(x)=4239476) -661 train 7.380940 (lr=2.0576e-04) (hash(x)=5731624) -662 train 7.522326 (lr=2.0497e-04) (hash(x)=5883465) -663 train 7.277260 (lr=2.0418e-04) (hash(x)=4892065) -664 train 7.116191 (lr=2.0339e-04) (hash(x)=5858782) -665 train 7.418376 (lr=2.0260e-04) (hash(x)=5489496) -666 train 7.080954 (lr=2.0182e-04) (hash(x)=4485195) -667 train 7.119433 (lr=2.0103e-04) (hash(x)=4933674) -668 train 7.377406 (lr=2.0024e-04) (hash(x)=5746292) -669 train 7.380608 (lr=1.9945e-04) (hash(x)=7021003) -670 train 7.578075 (lr=1.9867e-04) (hash(x)=5876710) -671 train 7.622410 (lr=1.9788e-04) (hash(x)=7317289) -672 train 7.732048 (lr=1.9709e-04) (hash(x)=5598226) -673 train 7.517103 (lr=1.9631e-04) (hash(x)=7869305) -674 train 7.592114 (lr=1.9552e-04) (hash(x)=6611408) -675 train 7.565008 (lr=1.9474e-04) (hash(x)=6811522) -676 train 7.554715 (lr=1.9396e-04) (hash(x)=6704714) -677 train 7.417961 (lr=1.9317e-04) (hash(x)=6601423) -678 train 7.628778 (lr=1.9239e-04) (hash(x)=6726071) -679 train 7.073271 (lr=1.9161e-04) (hash(x)=5510218) -680 train 7.881392 (lr=1.9083e-04) (hash(x)=7950952) -681 train 7.560656 (lr=1.9004e-04) (hash(x)=7180298) -682 train 7.378563 (lr=1.8926e-04) (hash(x)=6068813) -683 train 7.513981 (lr=1.8848e-04) (hash(x)=7304235) -684 train 7.530277 (lr=1.8770e-04) (hash(x)=7441806) -685 train 7.653310 (lr=1.8693e-04) (hash(x)=8111920) -686 train 8.254439 (lr=1.8615e-04) (hash(x)=6222783) -687 train 7.511406 (lr=1.8537e-04) (hash(x)=6752265) -688 train 7.330371 (lr=1.8459e-04) (hash(x)=6147634) -689 train 7.367846 (lr=1.8382e-04) (hash(x)=6788720) -690 train 7.493330 (lr=1.8304e-04) (hash(x)=6413518) -691 train 7.387491 (lr=1.8227e-04) (hash(x)=5994476) -692 train 7.327689 (lr=1.8149e-04) (hash(x)=5462082) -693 train 7.312479 (lr=1.8072e-04) (hash(x)=5862533) -694 train 7.480178 (lr=1.7995e-04) (hash(x)=7132796) -695 train 7.392669 (lr=1.7917e-04) (hash(x)=6530867) -696 train 7.427022 (lr=1.7840e-04) (hash(x)=7932207) -697 train 7.381664 (lr=1.7763e-04) (hash(x)=5914734) -698 train 7.570195 (lr=1.7686e-04) (hash(x)=6361594) -699 train 7.522934 (lr=1.7610e-04) (hash(x)=5746260) -700 val loss 7.4286 -700 val perplexity 1683.3726 -700 train 7.453722 (lr=1.7533e-04) (hash(x)=7065743) -701 train 7.390742 (lr=1.7456e-04) (hash(x)=6856097) -702 train 7.463803 (lr=1.7380e-04) (hash(x)=7588648) -703 train 7.390601 (lr=1.7303e-04) (hash(x)=5790078) -704 train 7.385214 (lr=1.7227e-04) (hash(x)=6069422) -705 train 7.735109 (lr=1.7150e-04) (hash(x)=5805002) -706 train 7.402860 (lr=1.7074e-04) (hash(x)=5344711) -707 train 7.274633 (lr=1.6998e-04) (hash(x)=6430135) -708 train 7.365939 (lr=1.6922e-04) (hash(x)=6317763) -709 train 7.301821 (lr=1.6846e-04) (hash(x)=6156715) -710 train 7.646048 (lr=1.6770e-04) (hash(x)=6321494) -711 train 7.687255 (lr=1.6694e-04) (hash(x)=7614023) -712 train 7.692809 (lr=1.6619e-04) (hash(x)=6740380) -713 train 7.522471 (lr=1.6543e-04) (hash(x)=4861744) -714 train 7.374860 (lr=1.6468e-04) (hash(x)=6542179) -715 train 7.430492 (lr=1.6393e-04) (hash(x)=5244861) -716 train 7.552281 (lr=1.6317e-04) (hash(x)=7306636) -717 train 7.379325 (lr=1.6242e-04) (hash(x)=7163697) -718 train 7.581982 (lr=1.6167e-04) (hash(x)=6421642) -719 train 7.362127 (lr=1.6092e-04) (hash(x)=5245146) -720 train 7.442839 (lr=1.6018e-04) (hash(x)=6046027) -721 train 7.225340 (lr=1.5943e-04) (hash(x)=6153866) -722 train 7.422318 (lr=1.5868e-04) (hash(x)=5827481) -723 train 7.333190 (lr=1.5794e-04) (hash(x)=6415565) -724 train 7.498818 (lr=1.5720e-04) (hash(x)=6409570) -725 train 7.945497 (lr=1.5646e-04) (hash(x)=7835853) -726 train 7.059450 (lr=1.5572e-04) (hash(x)=4827589) -727 train 7.184995 (lr=1.5498e-04) (hash(x)=5786972) -728 train 7.603070 (lr=1.5424e-04) (hash(x)=6736612) -729 train 7.464903 (lr=1.5350e-04) (hash(x)=6733560) -730 train 7.637909 (lr=1.5277e-04) (hash(x)=4877208) -731 train 7.113995 (lr=1.5203e-04) (hash(x)=6131703) -732 train 7.215689 (lr=1.5130e-04) (hash(x)=6533769) -733 train 7.212533 (lr=1.5057e-04) (hash(x)=6001331) -734 train 7.784758 (lr=1.4984e-04) (hash(x)=10602643) -735 train 7.712900 (lr=1.4911e-04) (hash(x)=6346459) -736 train 7.368340 (lr=1.4838e-04) (hash(x)=6728215) -737 train 7.657687 (lr=1.4766e-04) (hash(x)=8943770) -738 train 7.495002 (lr=1.4693e-04) (hash(x)=7141912) -739 train 7.455205 (lr=1.4621e-04) (hash(x)=6504131) -740 train 7.310729 (lr=1.4549e-04) (hash(x)=6461667) -741 train 7.536925 (lr=1.4477e-04) (hash(x)=5869339) -742 train 7.941216 (lr=1.4405e-04) (hash(x)=7948065) -743 train 7.168941 (lr=1.4333e-04) (hash(x)=5209234) -744 train 7.300319 (lr=1.4262e-04) (hash(x)=6372244) -745 train 7.478199 (lr=1.4190e-04) (hash(x)=7678937) -746 train 7.382864 (lr=1.4119e-04) (hash(x)=6519438) -747 train 7.456671 (lr=1.4048e-04) (hash(x)=6163272) -748 train 7.490024 (lr=1.3977e-04) (hash(x)=7025209) -749 train 7.333843 (lr=1.3906e-04) (hash(x)=5774172) -750 val loss 7.3960 -750 val perplexity 1629.4705 -750 train 7.503047 (lr=1.3835e-04) (hash(x)=5327301) -751 train 7.291592 (lr=1.3765e-04) (hash(x)=5676768) -752 train 6.999124 (lr=1.3694e-04) (hash(x)=4823304) -753 train 7.102332 (lr=1.3624e-04) (hash(x)=6443895) -754 train 7.922614 (lr=1.3554e-04) (hash(x)=8032890) -755 train 7.967155 (lr=1.3484e-04) (hash(x)=6090561) -756 train 7.320749 (lr=1.3415e-04) (hash(x)=4514262) -757 train 7.754841 (lr=1.3345e-04) (hash(x)=7788628) -758 train 7.693462 (lr=1.3276e-04) (hash(x)=7239456) -759 train 7.744527 (lr=1.3207e-04) (hash(x)=5863092) -760 train 7.414846 (lr=1.3138e-04) (hash(x)=5572628) -761 train 7.481321 (lr=1.3069e-04) (hash(x)=7141354) -762 train 7.679090 (lr=1.3000e-04) (hash(x)=7272638) -763 train 7.548131 (lr=1.2931e-04) (hash(x)=7201312) -764 train 7.542058 (lr=1.2863e-04) (hash(x)=6015388) -765 train 7.633627 (lr=1.2795e-04) (hash(x)=6414278) -766 train 7.389520 (lr=1.2727e-04) (hash(x)=7137466) -767 train 7.992478 (lr=1.2659e-04) (hash(x)=7019489) -768 train 7.542756 (lr=1.2592e-04) (hash(x)=7233453) -769 train 7.249399 (lr=1.2524e-04) (hash(x)=7914626) -770 train 7.352845 (lr=1.2457e-04) (hash(x)=5764080) -771 train 7.396888 (lr=1.2390e-04) (hash(x)=6225608) -772 train 7.448322 (lr=1.2323e-04) (hash(x)=8097255) -773 train 7.464288 (lr=1.2256e-04) (hash(x)=5998078) -774 train 7.598824 (lr=1.2190e-04) (hash(x)=5416254) -775 train 7.305884 (lr=1.2124e-04) (hash(x)=5483019) -776 train 7.075982 (lr=1.2057e-04) (hash(x)=4702208) -777 train 7.426601 (lr=1.1992e-04) (hash(x)=5911642) -778 train 7.462656 (lr=1.1926e-04) (hash(x)=6132487) -779 train 7.602303 (lr=1.1860e-04) (hash(x)=5903258) -780 train 7.524452 (lr=1.1795e-04) (hash(x)=7915382) -781 train 7.529359 (lr=1.1730e-04) (hash(x)=5632006) -782 train 7.327503 (lr=1.1665e-04) (hash(x)=6518211) -783 train 7.411739 (lr=1.1600e-04) (hash(x)=5968716) -784 train 7.534944 (lr=1.1536e-04) (hash(x)=7344525) -785 train 7.436255 (lr=1.1471e-04) (hash(x)=6401968) -786 train 7.360393 (lr=1.1407e-04) (hash(x)=6276127) -787 train 7.465710 (lr=1.1343e-04) (hash(x)=5778017) -788 train 7.023957 (lr=1.1279e-04) (hash(x)=5387306) -789 train 6.992826 (lr=1.1216e-04) (hash(x)=5772567) -790 train 7.221425 (lr=1.1153e-04) (hash(x)=6383748) -791 train 7.493791 (lr=1.1089e-04) (hash(x)=7780194) -792 train 7.362870 (lr=1.1027e-04) (hash(x)=7119030) -793 train 7.496387 (lr=1.0964e-04) (hash(x)=6424771) -794 train 7.302902 (lr=1.0901e-04) (hash(x)=6540151) -795 train 7.295705 (lr=1.0839e-04) (hash(x)=6140998) -796 train 7.391833 (lr=1.0777e-04) (hash(x)=6208271) -797 train 7.261333 (lr=1.0715e-04) (hash(x)=7859566) -798 train 7.814253 (lr=1.0654e-04) (hash(x)=7064477) -799 train 7.040195 (lr=1.0592e-04) (hash(x)=3784321) -800 val loss 7.3842 -800 val perplexity 1610.3484 -800 train 7.389319 (lr=1.0531e-04) (hash(x)=4472758) -801 train 7.322342 (lr=1.0470e-04) (hash(x)=5557891) -802 train 7.455420 (lr=1.0410e-04) (hash(x)=7969325) -803 train 7.474804 (lr=1.0349e-04) (hash(x)=5860821) -804 train 7.479715 (lr=1.0289e-04) (hash(x)=6750848) -805 train 7.432245 (lr=1.0229e-04) (hash(x)=5674826) -806 train 7.225125 (lr=1.0169e-04) (hash(x)=5529163) -807 train 7.610019 (lr=1.0109e-04) (hash(x)=7774109) -808 train 7.402578 (lr=1.0050e-04) (hash(x)=6762509) -809 train 7.511266 (lr=9.9910e-05) (hash(x)=5466449) -810 train 7.366781 (lr=9.9321e-05) (hash(x)=7046935) -811 train 8.244202 (lr=9.8735e-05) (hash(x)=7997664) -812 train 7.527016 (lr=9.8151e-05) (hash(x)=6882594) -813 train 7.705925 (lr=9.7569e-05) (hash(x)=7006517) -814 train 7.564469 (lr=9.6990e-05) (hash(x)=6832530) -815 train 7.388801 (lr=9.6413e-05) (hash(x)=6576749) -816 train 7.497406 (lr=9.5838e-05) (hash(x)=8276629) -817 train 7.316368 (lr=9.5266e-05) (hash(x)=6896198) -818 train 7.508495 (lr=9.4696e-05) (hash(x)=5829252) -819 train 7.657996 (lr=9.4129e-05) (hash(x)=7266655) -820 train 7.383558 (lr=9.3564e-05) (hash(x)=6015975) -821 train 7.364309 (lr=9.3001e-05) (hash(x)=5696696) -822 train 7.410369 (lr=9.2441e-05) (hash(x)=5411666) -823 train 7.629680 (lr=9.1884e-05) (hash(x)=7072404) -824 train 7.315626 (lr=9.1328e-05) (hash(x)=4910095) -825 train 7.424473 (lr=9.0776e-05) (hash(x)=6590657) -826 train 7.353024 (lr=9.0226e-05) (hash(x)=7665574) -827 train 7.405118 (lr=8.9678e-05) (hash(x)=6626459) -828 train 7.428163 (lr=8.9133e-05) (hash(x)=6731971) -829 train 7.423668 (lr=8.8591e-05) (hash(x)=5644977) -830 train 7.382772 (lr=8.8051e-05) (hash(x)=6238654) -831 train 7.264947 (lr=8.7513e-05) (hash(x)=6556025) -832 train 7.580822 (lr=8.6978e-05) (hash(x)=7175563) -833 train 7.459033 (lr=8.6446e-05) (hash(x)=6407492) -834 train 7.285389 (lr=8.5916e-05) (hash(x)=4934335) -835 train 7.300157 (lr=8.5389e-05) (hash(x)=5942129) -836 train 7.337029 (lr=8.4865e-05) (hash(x)=5995643) -837 train 7.404269 (lr=8.4343e-05) (hash(x)=6012163) -838 train 7.545713 (lr=8.3824e-05) (hash(x)=8464831) -839 train 7.774149 (lr=8.3307e-05) (hash(x)=7325027) -840 train 7.636600 (lr=8.2793e-05) (hash(x)=6785865) -841 train 7.550599 (lr=8.2282e-05) (hash(x)=4425520) -842 train 7.334558 (lr=8.1773e-05) (hash(x)=5388267) -843 train 7.595325 (lr=8.1267e-05) (hash(x)=7322467) -844 train 7.362306 (lr=8.0764e-05) (hash(x)=6681766) -845 train 7.458833 (lr=8.0263e-05) (hash(x)=7482800) -846 train 7.374401 (lr=7.9765e-05) (hash(x)=5554493) -847 train 7.599457 (lr=7.9270e-05) (hash(x)=6373412) -848 train 7.021658 (lr=7.8778e-05) (hash(x)=5117517) -849 train 7.258224 (lr=7.8288e-05) (hash(x)=6981426) -850 val loss 7.3779 -850 val perplexity 1600.2367 -850 train 7.218916 (lr=7.7801e-05) (hash(x)=6886188) -851 train 7.424624 (lr=7.7317e-05) (hash(x)=7332255) -852 train 7.436696 (lr=7.6836e-05) (hash(x)=6172042) -853 train 7.589216 (lr=7.6357e-05) (hash(x)=5930894) -854 train 7.635510 (lr=7.5881e-05) (hash(x)=7448958) -855 train 7.429552 (lr=7.5408e-05) (hash(x)=5262868) -856 train 7.302075 (lr=7.4938e-05) (hash(x)=5558427) -857 train 7.190957 (lr=7.4470e-05) (hash(x)=5585769) -858 train 7.074392 (lr=7.4005e-05) (hash(x)=5838081) -859 train 7.253884 (lr=7.3544e-05) (hash(x)=5688247) -860 train 7.357229 (lr=7.3085e-05) (hash(x)=5162020) -861 train 7.555761 (lr=7.2628e-05) (hash(x)=7462079) -862 train 7.404220 (lr=7.2175e-05) (hash(x)=6516108) -863 train 7.197952 (lr=7.1725e-05) (hash(x)=8055563) -864 train 7.409320 (lr=7.1277e-05) (hash(x)=6271901) -865 train 7.435908 (lr=7.0832e-05) (hash(x)=6221701) -866 train 7.431785 (lr=7.0390e-05) (hash(x)=5772861) -867 train 7.454789 (lr=6.9952e-05) (hash(x)=5352405) -868 train 7.438743 (lr=6.9516e-05) (hash(x)=6111630) -869 train 7.469532 (lr=6.9082e-05) (hash(x)=6730666) -870 train 7.646123 (lr=6.8652e-05) (hash(x)=8698315) -871 train 7.321314 (lr=6.8225e-05) (hash(x)=5932790) -872 train 7.698706 (lr=6.7801e-05) (hash(x)=8099892) -873 train 7.362211 (lr=6.7379e-05) (hash(x)=6223114) -874 train 7.465238 (lr=6.6961e-05) (hash(x)=5798363) -875 train 7.291901 (lr=6.6545e-05) (hash(x)=6249312) -876 train 7.439408 (lr=6.6133e-05) (hash(x)=6929692) -877 train 7.550559 (lr=6.5723e-05) (hash(x)=7242827) -878 train 7.301716 (lr=6.5317e-05) (hash(x)=6332123) -879 train 7.210484 (lr=6.4913e-05) (hash(x)=5680154) -880 train 7.595014 (lr=6.4513e-05) (hash(x)=6352331) -881 train 7.403541 (lr=6.4115e-05) (hash(x)=6332332) -882 train 7.452321 (lr=6.3721e-05) (hash(x)=6180674) -883 train 7.404085 (lr=6.3329e-05) (hash(x)=8156280) -884 train 7.426672 (lr=6.2941e-05) (hash(x)=5874247) -885 train 7.371298 (lr=6.2556e-05) (hash(x)=6659781) -886 train 7.493792 (lr=6.2173e-05) (hash(x)=5780147) -887 train 7.313876 (lr=6.1794e-05) (hash(x)=5914217) -888 train 7.157481 (lr=6.1418e-05) (hash(x)=5411762) -889 train 7.472383 (lr=6.1045e-05) (hash(x)=6714110) -890 train 7.149372 (lr=6.0675e-05) (hash(x)=5999685) -891 train 7.332294 (lr=6.0308e-05) (hash(x)=7120215) -892 train 7.582332 (lr=5.9944e-05) (hash(x)=6546587) -893 train 7.346202 (lr=5.9583e-05) (hash(x)=6593413) -894 train 7.456257 (lr=5.9225e-05) (hash(x)=6369927) -895 train 7.350650 (lr=5.8871e-05) (hash(x)=6424128) -896 train 7.335809 (lr=5.8519e-05) (hash(x)=5736158) -897 train 7.159677 (lr=5.8171e-05) (hash(x)=5903043) -898 train 7.245536 (lr=5.7826e-05) (hash(x)=4419128) -899 train 7.069800 (lr=5.7484e-05) (hash(x)=4390027) -900 val loss 7.3551 -900 val perplexity 1564.1499 -900 train 7.308222 (lr=5.7145e-05) (hash(x)=6728135) -901 train 7.435823 (lr=5.6809e-05) (hash(x)=6945760) -902 train 7.396671 (lr=5.6476e-05) (hash(x)=6081534) -903 train 7.469185 (lr=5.6147e-05) (hash(x)=7804089) -904 train 7.526792 (lr=5.5821e-05) (hash(x)=6225832) -905 train 7.565935 (lr=5.5497e-05) (hash(x)=6273417) -906 train 7.607038 (lr=5.5178e-05) (hash(x)=7775633) -907 train 7.564097 (lr=5.4861e-05) (hash(x)=7130267) -908 train 7.345338 (lr=5.4547e-05) (hash(x)=6554076) -909 train 7.456250 (lr=5.4237e-05) (hash(x)=6140697) -910 train 7.271626 (lr=5.3930e-05) (hash(x)=6128181) -911 train 7.385910 (lr=5.3626e-05) (hash(x)=6490149) -912 train 6.928421 (lr=5.3325e-05) (hash(x)=5422426) -913 train 7.370100 (lr=5.3028e-05) (hash(x)=4847733) -914 train 7.373724 (lr=5.2734e-05) (hash(x)=6349747) -915 train 7.668162 (lr=5.2443e-05) (hash(x)=9276378) -916 train 7.275504 (lr=5.2155e-05) (hash(x)=6869524) -917 train 7.300123 (lr=5.1871e-05) (hash(x)=7082053) -918 train 7.355165 (lr=5.1589e-05) (hash(x)=7799351) -919 train 7.519971 (lr=5.1311e-05) (hash(x)=6283540) -920 train 7.304411 (lr=5.1037e-05) (hash(x)=6660860) -921 train 7.242377 (lr=5.0765e-05) (hash(x)=5200956) -922 train 7.342465 (lr=5.0497e-05) (hash(x)=5653241) -923 train 7.494324 (lr=5.0232e-05) (hash(x)=5781680) -924 train 7.461452 (lr=4.9971e-05) (hash(x)=6146666) -925 train 7.314250 (lr=4.9712e-05) (hash(x)=8652907) -926 train 7.114859 (lr=4.9457e-05) (hash(x)=5311546) -927 train 7.249785 (lr=4.9206e-05) (hash(x)=5728337) -928 train 7.229902 (lr=4.8957e-05) (hash(x)=5413240) -929 train 7.349687 (lr=4.8712e-05) (hash(x)=5832401) -930 train 7.109508 (lr=4.8470e-05) (hash(x)=7420230) -931 train 7.447092 (lr=4.8232e-05) (hash(x)=6278202) -932 train 7.297539 (lr=4.7997e-05) (hash(x)=6189873) -933 train 7.422616 (lr=4.7765e-05) (hash(x)=7403868) -934 train 7.268861 (lr=4.7537e-05) (hash(x)=6867459) -935 train 7.485469 (lr=4.7312e-05) (hash(x)=7400272) -936 train 7.508608 (lr=4.7090e-05) (hash(x)=7936275) -937 train 7.229595 (lr=4.6871e-05) (hash(x)=5943413) -938 train 7.445310 (lr=4.6656e-05) (hash(x)=5442166) -939 train 7.338399 (lr=4.6445e-05) (hash(x)=5795128) -940 train 7.523449 (lr=4.6236e-05) (hash(x)=6831680) -941 train 7.247817 (lr=4.6031e-05) (hash(x)=5519980) -942 train 7.243730 (lr=4.5830e-05) (hash(x)=6765461) -943 train 7.199632 (lr=4.5631e-05) (hash(x)=6473432) -944 train 7.178991 (lr=4.5437e-05) (hash(x)=6490395) -945 train 7.411891 (lr=4.5245e-05) (hash(x)=5484669) -946 train 7.390506 (lr=4.5057e-05) (hash(x)=7278970) -947 train 7.377321 (lr=4.4872e-05) (hash(x)=5201502) -948 train 7.628700 (lr=4.4691e-05) (hash(x)=7277915) -949 train 7.284154 (lr=4.4513e-05) (hash(x)=7227057) -950 val loss 7.3435 -950 val perplexity 1546.1652 -950 train 7.572911 (lr=4.4338e-05) (hash(x)=8178996) -951 train 7.713424 (lr=4.4167e-05) (hash(x)=6922646) -952 train 7.215410 (lr=4.4000e-05) (hash(x)=5941772) -953 train 7.174014 (lr=4.3835e-05) (hash(x)=6056673) -954 train 7.589973 (lr=4.3674e-05) (hash(x)=6912178) -955 train 7.107974 (lr=4.3517e-05) (hash(x)=5484641) -956 train 7.452599 (lr=4.3363e-05) (hash(x)=5121117) -957 train 7.482172 (lr=4.3212e-05) (hash(x)=7100865) -958 train 7.160964 (lr=4.3065e-05) (hash(x)=5334280) -959 train 7.272405 (lr=4.2921e-05) (hash(x)=7277141) -960 train 7.261897 (lr=4.2781e-05) (hash(x)=6401960) -961 train 7.311950 (lr=4.2644e-05) (hash(x)=6407603) -962 train 7.450472 (lr=4.2510e-05) (hash(x)=5786089) -963 train 7.306540 (lr=4.2380e-05) (hash(x)=5604614) -964 train 7.339895 (lr=4.2253e-05) (hash(x)=5454287) -965 train 7.479698 (lr=4.2130e-05) (hash(x)=6200612) -966 train 7.418107 (lr=4.2010e-05) (hash(x)=5691759) -967 train 7.381812 (lr=4.1894e-05) (hash(x)=7443124) -968 train 7.400855 (lr=4.1781e-05) (hash(x)=6711366) -969 train 7.318673 (lr=4.1672e-05) (hash(x)=6728764) -970 train 7.627553 (lr=4.1566e-05) (hash(x)=6870240) -971 train 7.339466 (lr=4.1463e-05) (hash(x)=5601968) -972 train 7.080988 (lr=4.1364e-05) (hash(x)=5311504) -973 train 7.358827 (lr=4.1269e-05) (hash(x)=6877761) -974 train 7.431441 (lr=4.1177e-05) (hash(x)=6457033) -975 train 7.382578 (lr=4.1088e-05) (hash(x)=5632735) -976 train 7.576962 (lr=4.1003e-05) (hash(x)=6024880) -977 train 7.265551 (lr=4.0921e-05) (hash(x)=5658845) -978 train 7.156719 (lr=4.0843e-05) (hash(x)=6089461) -979 train 7.201392 (lr=4.0768e-05) (hash(x)=5594810) -980 train 7.007952 (lr=4.0697e-05) (hash(x)=4667026) -981 train 6.872203 (lr=4.0629e-05) (hash(x)=4738523) -982 train 7.193507 (lr=4.0564e-05) (hash(x)=7451876) -983 train 7.457832 (lr=4.0503e-05) (hash(x)=5881015) -984 train 7.443588 (lr=4.0446e-05) (hash(x)=6255031) -985 train 7.263505 (lr=4.0392e-05) (hash(x)=5263325) -986 train 7.258874 (lr=4.0341e-05) (hash(x)=6654585) -987 train 7.319571 (lr=4.0294e-05) (hash(x)=6865361) -988 train 7.416778 (lr=4.0251e-05) (hash(x)=6978986) -989 train 7.520934 (lr=4.0211e-05) (hash(x)=5644619) -990 train 7.341804 (lr=4.0174e-05) (hash(x)=6551023) -991 train 7.284855 (lr=4.0141e-05) (hash(x)=5809411) -992 train 7.375855 (lr=4.0112e-05) (hash(x)=5501837) -993 train 7.392744 (lr=4.0085e-05) (hash(x)=7044374) -994 train 7.355122 (lr=4.0063e-05) (hash(x)=6526130) -995 train 7.469051 (lr=4.0044e-05) (hash(x)=6293025) -996 train 7.465250 (lr=4.0028e-05) (hash(x)=6295768) -997 train 7.271845 (lr=4.0016e-05) (hash(x)=5189023) -998 train 7.086488 (lr=4.0007e-05) (hash(x)=6255467) -999 val loss 7.3384 -999 val perplexity 1538.2334 -999 train 7.425029 (lr=4.0002e-05) (hash(x)=7232964) +max_steps: 5000 +0 val loss 11.6964 +0 val perplexity 120134.1641 +0 train 11.681432 (lr=2.7972e-07) (hash(x)=22886834) +1 train 11.705652 (lr=5.5944e-07) (hash(x)=26375038) +2 train 11.700677 (lr=8.3916e-07) (hash(x)=30598777) +3 train 11.719830 (lr=1.1189e-06) (hash(x)=27234506) +4 train 11.680297 (lr=1.3986e-06) (hash(x)=27767880) +5 train 11.678657 (lr=1.6783e-06) (hash(x)=23702020) +6 train 11.669114 (lr=1.9580e-06) (hash(x)=31986844) +7 train 11.689692 (lr=2.2378e-06) (hash(x)=20782690) +8 train 11.681740 (lr=2.5175e-06) (hash(x)=25201599) +9 train 11.681633 (lr=2.7972e-06) (hash(x)=23094976) +10 train 11.697703 (lr=3.0769e-06) (hash(x)=23841096) +11 train 11.633875 (lr=3.3566e-06) (hash(x)=26532095) +12 train 11.626414 (lr=3.6364e-06) (hash(x)=24432298) +13 train 11.600323 (lr=3.9161e-06) (hash(x)=27151649) +14 train 11.642756 (lr=4.1958e-06) (hash(x)=24596846) +15 train 11.636794 (lr=4.4755e-06) (hash(x)=23890908) +16 train 11.594416 (lr=4.7552e-06) (hash(x)=28913955) +17 train 11.587169 (lr=5.0350e-06) (hash(x)=25588236) +18 train 11.604854 (lr=5.3147e-06) (hash(x)=23770034) +19 train 11.568233 (lr=5.5944e-06) (hash(x)=24011372) +20 train 11.543827 (lr=5.8741e-06) (hash(x)=25441898) +21 train 11.529476 (lr=6.1538e-06) (hash(x)=28375581) +22 train 11.519797 (lr=6.4336e-06) (hash(x)=24046679) +23 train 11.544165 (lr=6.7133e-06) (hash(x)=24611628) +24 train 11.557258 (lr=6.9930e-06) (hash(x)=26169030) +25 train 11.488771 (lr=7.2727e-06) (hash(x)=30298407) +26 train 11.406871 (lr=7.5524e-06) (hash(x)=23711112) +27 train 11.440053 (lr=7.8322e-06) (hash(x)=19245352) +28 train 11.378906 (lr=8.1119e-06) (hash(x)=21529136) +29 train 11.432513 (lr=8.3916e-06) (hash(x)=28936608) +30 train 11.362607 (lr=8.6713e-06) (hash(x)=24339013) +31 train 11.315014 (lr=8.9510e-06) (hash(x)=25767553) +32 train 11.361699 (lr=9.2308e-06) (hash(x)=26439905) +33 train 11.279980 (lr=9.5105e-06) (hash(x)=31093473) +34 train 11.227470 (lr=9.7902e-06) (hash(x)=25450374) +35 train 11.256378 (lr=1.0070e-05) (hash(x)=24809873) +36 train 11.261599 (lr=1.0350e-05) (hash(x)=23253252) +37 train 11.240863 (lr=1.0629e-05) (hash(x)=27852919) +38 train 11.143288 (lr=1.0909e-05) (hash(x)=23327497) +39 train 11.129188 (lr=1.1189e-05) (hash(x)=22512166) +40 train 11.142281 (lr=1.1469e-05) (hash(x)=22859419) +41 train 11.106739 (lr=1.1748e-05) (hash(x)=27620338) +42 train 11.056373 (lr=1.2028e-05) (hash(x)=26397837) +43 train 11.071579 (lr=1.2308e-05) (hash(x)=28092148) +44 train 10.930977 (lr=1.2587e-05) (hash(x)=24662703) +45 train 10.918567 (lr=1.2867e-05) (hash(x)=27938767) +46 train 11.046695 (lr=1.3147e-05) (hash(x)=26037988) +47 train 10.892894 (lr=1.3427e-05) (hash(x)=24732833) +48 train 10.874605 (lr=1.3706e-05) (hash(x)=25259526) +49 train 10.896080 (lr=1.3986e-05) (hash(x)=23200230) +50 val loss 10.8452 +50 val perplexity 51287.5547 +50 train 10.866057 (lr=1.4266e-05) (hash(x)=26721357) +51 train 10.803464 (lr=1.4545e-05) (hash(x)=22694718) +52 train 10.776980 (lr=1.4825e-05) (hash(x)=28066766) +53 train 10.781597 (lr=1.5105e-05) (hash(x)=23125151) +54 train 10.790044 (lr=1.5385e-05) (hash(x)=27193725) +55 train 10.705384 (lr=1.5664e-05) (hash(x)=25129410) +56 train 10.651207 (lr=1.5944e-05) (hash(x)=24263988) +57 train 10.641276 (lr=1.6224e-05) (hash(x)=23059154) +58 train 10.681503 (lr=1.6503e-05) (hash(x)=26063864) +59 train 10.750127 (lr=1.6783e-05) (hash(x)=27858570) +60 train 10.613048 (lr=1.7063e-05) (hash(x)=23874620) +61 train 10.583058 (lr=1.7343e-05) (hash(x)=22402617) +62 train 10.536903 (lr=1.7622e-05) (hash(x)=23600822) +63 train 10.469747 (lr=1.7902e-05) (hash(x)=26582391) +64 train 10.379783 (lr=1.8182e-05) (hash(x)=23225283) +65 train 10.482368 (lr=1.8462e-05) (hash(x)=26075451) +66 train 10.508182 (lr=1.8741e-05) (hash(x)=24723419) +67 train 10.500177 (lr=1.9021e-05) (hash(x)=27279806) +68 train 10.338507 (lr=1.9301e-05) (hash(x)=25870391) +69 train 10.449234 (lr=1.9580e-05) (hash(x)=26188136) +70 train 10.387484 (lr=1.9860e-05) (hash(x)=30373443) +71 train 10.461452 (lr=2.0140e-05) (hash(x)=26472336) +72 train 10.332763 (lr=2.0420e-05) (hash(x)=26651572) +73 train 10.355941 (lr=2.0699e-05) (hash(x)=26376212) +74 train 10.363896 (lr=2.0979e-05) (hash(x)=26733350) +75 train 10.423232 (lr=2.1259e-05) (hash(x)=28301589) +76 train 10.276484 (lr=2.1538e-05) (hash(x)=27599559) +77 train 10.249597 (lr=2.1818e-05) (hash(x)=28035221) +78 train 10.363973 (lr=2.2098e-05) (hash(x)=25016783) +79 train 10.352807 (lr=2.2378e-05) (hash(x)=27654289) +80 train 10.392247 (lr=2.2657e-05) (hash(x)=24597558) +81 train 10.274711 (lr=2.2937e-05) (hash(x)=21560904) +82 train 10.266309 (lr=2.3217e-05) (hash(x)=21983837) +83 train 10.319811 (lr=2.3497e-05) (hash(x)=24995715) +84 train 10.361818 (lr=2.3776e-05) (hash(x)=29876413) +85 train 10.290359 (lr=2.4056e-05) (hash(x)=23792508) +86 train 10.222163 (lr=2.4336e-05) (hash(x)=25509120) +87 train 10.259757 (lr=2.4615e-05) (hash(x)=26559876) +88 train 10.269846 (lr=2.4895e-05) (hash(x)=23569647) +89 train 10.277059 (lr=2.5175e-05) (hash(x)=25758852) +90 train 10.193969 (lr=2.5455e-05) (hash(x)=25706298) +91 train 10.098770 (lr=2.5734e-05) (hash(x)=28364895) +92 train 10.187699 (lr=2.6014e-05) (hash(x)=25304663) +93 train 10.186853 (lr=2.6294e-05) (hash(x)=25269299) +94 train 10.342330 (lr=2.6573e-05) (hash(x)=25870566) +95 train 10.164028 (lr=2.6853e-05) (hash(x)=21770329) +96 train 10.335703 (lr=2.7133e-05) (hash(x)=27595900) +97 train 10.333036 (lr=2.7413e-05) (hash(x)=24785397) +98 train 10.074593 (lr=2.7692e-05) (hash(x)=21521480) +99 train 10.119580 (lr=2.7972e-05) (hash(x)=24628606) +100 val loss 10.1578 +100 val perplexity 25791.5098 +100 train 10.140892 (lr=2.8252e-05) (hash(x)=24670150) +101 train 10.168761 (lr=2.8531e-05) (hash(x)=23181910) +102 train 10.094563 (lr=2.8811e-05) (hash(x)=22714991) +103 train 10.272529 (lr=2.9091e-05) (hash(x)=22723459) +104 train 10.105750 (lr=2.9371e-05) (hash(x)=21524316) +105 train 10.145549 (lr=2.9650e-05) (hash(x)=25506632) +106 train 9.951752 (lr=2.9930e-05) (hash(x)=21675672) +107 train 10.076271 (lr=3.0210e-05) (hash(x)=22897919) +108 train 10.118456 (lr=3.0490e-05) (hash(x)=23321631) +109 train 10.074455 (lr=3.0769e-05) (hash(x)=26546719) +110 train 10.668172 (lr=3.1049e-05) (hash(x)=31962348) +111 train 10.086867 (lr=3.1329e-05) (hash(x)=30338342) +112 train 9.984250 (lr=3.1608e-05) (hash(x)=23724471) +113 train 10.087273 (lr=3.1888e-05) (hash(x)=29175888) +114 train 10.128128 (lr=3.2168e-05) (hash(x)=23256716) +115 train 10.078434 (lr=3.2448e-05) (hash(x)=27063280) +116 train 10.060738 (lr=3.2727e-05) (hash(x)=31057659) +117 train 10.180396 (lr=3.3007e-05) (hash(x)=32915097) +118 train 10.071523 (lr=3.3287e-05) (hash(x)=28842717) +119 train 9.997933 (lr=3.3566e-05) (hash(x)=25678059) +120 train 9.989501 (lr=3.3846e-05) (hash(x)=21593510) +121 train 9.930816 (lr=3.4126e-05) (hash(x)=20083773) +122 train 10.021068 (lr=3.4406e-05) (hash(x)=23002820) +123 train 9.899536 (lr=3.4685e-05) (hash(x)=21853028) +124 train 9.945949 (lr=3.4965e-05) (hash(x)=26985625) +125 train 9.817840 (lr=3.5245e-05) (hash(x)=21808483) +126 train 9.982982 (lr=3.5524e-05) (hash(x)=28873251) +127 train 9.900296 (lr=3.5804e-05) (hash(x)=26109335) +128 train 9.941215 (lr=3.6084e-05) (hash(x)=26334674) +129 train 9.964335 (lr=3.6364e-05) (hash(x)=24916754) +130 train 9.924269 (lr=3.6643e-05) (hash(x)=25449624) +131 train 10.025691 (lr=3.6923e-05) (hash(x)=25334848) +132 train 9.969893 (lr=3.7203e-05) (hash(x)=27484863) +133 train 9.835284 (lr=3.7483e-05) (hash(x)=24917705) +134 train 9.707572 (lr=3.7762e-05) (hash(x)=25143449) +135 train 9.752698 (lr=3.8042e-05) (hash(x)=25044885) +136 train 9.895902 (lr=3.8322e-05) (hash(x)=27821028) +137 train 9.942389 (lr=3.8601e-05) (hash(x)=28747022) +138 train 9.887458 (lr=3.8881e-05) (hash(x)=27182888) +139 train 9.752761 (lr=3.9161e-05) (hash(x)=23678349) +140 train 9.929805 (lr=3.9441e-05) (hash(x)=23593235) +141 train 9.959308 (lr=3.9720e-05) (hash(x)=28529813) +142 train 10.132763 (lr=4.0000e-05) (hash(x)=32074661) +143 train 9.810905 (lr=4.0280e-05) (hash(x)=28870690) +144 train 9.833708 (lr=4.0559e-05) (hash(x)=27307705) +145 train 9.799855 (lr=4.0839e-05) (hash(x)=25044834) +146 train 9.844878 (lr=4.1119e-05) (hash(x)=23712023) +147 train 9.963409 (lr=4.1399e-05) (hash(x)=32982615) +148 train 9.835284 (lr=4.1678e-05) (hash(x)=30113660) +149 train 9.742824 (lr=4.1958e-05) (hash(x)=20970960) +150 val loss 9.7075 +150 val perplexity 16439.6582 +150 train 9.750391 (lr=4.2238e-05) (hash(x)=23132684) +151 train 10.371020 (lr=4.2517e-05) (hash(x)=35279941) +152 train 10.016452 (lr=4.2797e-05) (hash(x)=31227444) +153 train 9.604296 (lr=4.3077e-05) (hash(x)=25529472) +154 train 9.571414 (lr=4.3357e-05) (hash(x)=24350409) +155 train 9.894919 (lr=4.3636e-05) (hash(x)=26400041) +156 train 9.675529 (lr=4.3916e-05) (hash(x)=25262621) +157 train 9.660502 (lr=4.4196e-05) (hash(x)=24656138) +158 train 9.677133 (lr=4.4476e-05) (hash(x)=26803414) +159 train 9.583070 (lr=4.4755e-05) (hash(x)=25015923) +160 train 9.488092 (lr=4.5035e-05) (hash(x)=23581172) +161 train 9.511117 (lr=4.5315e-05) (hash(x)=22924885) +162 train 9.535074 (lr=4.5594e-05) (hash(x)=23414296) +163 train 9.607943 (lr=4.5874e-05) (hash(x)=24853586) +164 train 9.759442 (lr=4.6154e-05) (hash(x)=25000130) +165 train 9.614139 (lr=4.6434e-05) (hash(x)=27004780) +166 train 9.538065 (lr=4.6713e-05) (hash(x)=26148573) +167 train 9.558585 (lr=4.6993e-05) (hash(x)=26740855) +168 train 9.359779 (lr=4.7273e-05) (hash(x)=20965419) +169 train 9.494944 (lr=4.7552e-05) (hash(x)=23950114) +170 train 9.351913 (lr=4.7832e-05) (hash(x)=24951982) +171 train 9.475811 (lr=4.8112e-05) (hash(x)=24584116) +172 train 9.507161 (lr=4.8392e-05) (hash(x)=24378759) +173 train 9.410069 (lr=4.8671e-05) (hash(x)=25718516) +174 train 9.603662 (lr=4.8951e-05) (hash(x)=28424396) +175 train 9.474504 (lr=4.9231e-05) (hash(x)=22262151) +176 train 9.564561 (lr=4.9510e-05) (hash(x)=26438412) +177 train 9.392962 (lr=4.9790e-05) (hash(x)=23025303) +178 train 9.299335 (lr=5.0070e-05) (hash(x)=24190770) +179 train 9.334599 (lr=5.0350e-05) (hash(x)=26627860) +180 train 9.413591 (lr=5.0629e-05) (hash(x)=23663439) +181 train 9.294341 (lr=5.0909e-05) (hash(x)=23700532) +182 train 9.337937 (lr=5.1189e-05) (hash(x)=23075676) +183 train 9.396532 (lr=5.1469e-05) (hash(x)=26621834) +184 train 9.341347 (lr=5.1748e-05) (hash(x)=29426269) +185 train 9.228568 (lr=5.2028e-05) (hash(x)=21821465) +186 train 9.291506 (lr=5.2308e-05) (hash(x)=26506130) +187 train 9.422121 (lr=5.2587e-05) (hash(x)=26930630) +188 train 9.206786 (lr=5.2867e-05) (hash(x)=22993793) +189 train 9.193172 (lr=5.3147e-05) (hash(x)=19557946) +190 train 9.281096 (lr=5.3427e-05) (hash(x)=23572891) +191 train 9.135388 (lr=5.3706e-05) (hash(x)=23234741) +192 train 9.230703 (lr=5.3986e-05) (hash(x)=25547951) +193 train 9.149540 (lr=5.4266e-05) (hash(x)=26713563) +194 train 9.225824 (lr=5.4545e-05) (hash(x)=25913622) +195 train 9.319949 (lr=5.4825e-05) (hash(x)=28102443) +196 train 9.203295 (lr=5.5105e-05) (hash(x)=23093351) +197 train 9.540465 (lr=5.5385e-05) (hash(x)=31689122) +198 train 9.436080 (lr=5.5664e-05) (hash(x)=32252517) +199 train 9.184306 (lr=5.5944e-05) (hash(x)=25470563) +200 val loss 9.1570 +200 val perplexity 9480.1797 +200 train 9.318839 (lr=5.6224e-05) (hash(x)=25597614) +201 train 9.189975 (lr=5.6503e-05) (hash(x)=23757479) +202 train 9.006488 (lr=5.6783e-05) (hash(x)=25140048) +203 train 9.163328 (lr=5.7063e-05) (hash(x)=28282861) +204 train 9.018702 (lr=5.7343e-05) (hash(x)=24754885) +205 train 9.202818 (lr=5.7622e-05) (hash(x)=26731964) +206 train 9.277732 (lr=5.7902e-05) (hash(x)=26660561) +207 train 8.982871 (lr=5.8182e-05) (hash(x)=21799102) +208 train 9.114472 (lr=5.8462e-05) (hash(x)=27025986) +209 train 8.972461 (lr=5.8741e-05) (hash(x)=24672077) +210 train 9.026059 (lr=5.9021e-05) (hash(x)=25322984) +211 train 8.968822 (lr=5.9301e-05) (hash(x)=23471769) +212 train 8.720599 (lr=5.9580e-05) (hash(x)=20766491) +213 train 9.170511 (lr=5.9860e-05) (hash(x)=24058931) +214 train 8.949053 (lr=6.0140e-05) (hash(x)=23872843) +215 train 8.973722 (lr=6.0420e-05) (hash(x)=23275479) +216 train 9.073886 (lr=6.0699e-05) (hash(x)=24914695) +217 train 8.954366 (lr=6.0979e-05) (hash(x)=24364396) +218 train 8.915129 (lr=6.1259e-05) (hash(x)=27986474) +219 train 9.036798 (lr=6.1538e-05) (hash(x)=24473581) +220 train 8.919971 (lr=6.1818e-05) (hash(x)=22974689) +221 train 8.978792 (lr=6.2098e-05) (hash(x)=23774644) +222 train 8.958130 (lr=6.2378e-05) (hash(x)=23245327) +223 train 8.812987 (lr=6.2657e-05) (hash(x)=22091862) +224 train 8.812299 (lr=6.2937e-05) (hash(x)=24362839) +225 train 8.803389 (lr=6.3217e-05) (hash(x)=25482303) +226 train 8.938492 (lr=6.3497e-05) (hash(x)=24911853) +227 train 8.939918 (lr=6.3776e-05) (hash(x)=26018202) +228 train 8.758079 (lr=6.4056e-05) (hash(x)=26124495) +229 train 8.712245 (lr=6.4336e-05) (hash(x)=24560096) +230 train 8.802936 (lr=6.4615e-05) (hash(x)=24695331) +231 train 8.811640 (lr=6.4895e-05) (hash(x)=17430373) +232 train 8.723782 (lr=6.5175e-05) (hash(x)=21813345) +233 train 8.369038 (lr=6.5455e-05) (hash(x)=20098681) +234 train 8.537743 (lr=6.5734e-05) (hash(x)=25095928) +235 train 8.561785 (lr=6.6014e-05) (hash(x)=24078083) +236 train 8.585090 (lr=6.6294e-05) (hash(x)=22901505) +237 train 8.795768 (lr=6.6573e-05) (hash(x)=26595592) +238 train 8.699968 (lr=6.6853e-05) (hash(x)=27663196) +239 train 8.629430 (lr=6.7133e-05) (hash(x)=22954861) +240 train 8.847414 (lr=6.7413e-05) (hash(x)=30159234) +241 train 8.700428 (lr=6.7692e-05) (hash(x)=29294271) +242 train 8.706714 (lr=6.7972e-05) (hash(x)=23728322) +243 train 8.735248 (lr=6.8252e-05) (hash(x)=28695016) +244 train 8.771428 (lr=6.8531e-05) (hash(x)=26702728) +245 train 8.632565 (lr=6.8811e-05) (hash(x)=27100115) +246 train 8.526557 (lr=6.9091e-05) (hash(x)=25082752) +247 train 8.877078 (lr=6.9371e-05) (hash(x)=26671799) +248 train 8.749573 (lr=6.9650e-05) (hash(x)=23718946) +249 train 8.768171 (lr=6.9930e-05) (hash(x)=28137394) +250 val loss 8.5759 +250 val perplexity 5302.2915 +250 train 8.717565 (lr=7.0210e-05) (hash(x)=23893495) +251 train 8.589409 (lr=7.0490e-05) (hash(x)=23166092) +252 train 8.503679 (lr=7.0769e-05) (hash(x)=25907665) +253 train 8.421417 (lr=7.1049e-05) (hash(x)=25885986) +254 train 8.485579 (lr=7.1329e-05) (hash(x)=25262712) +255 train 8.489021 (lr=7.1608e-05) (hash(x)=26924723) +256 train 8.602015 (lr=7.1888e-05) (hash(x)=28744736) +257 train 8.432463 (lr=7.2168e-05) (hash(x)=26140590) +258 train 8.593373 (lr=7.2448e-05) (hash(x)=25780449) +259 train 8.510340 (lr=7.2727e-05) (hash(x)=25138659) +260 train 8.578063 (lr=7.3007e-05) (hash(x)=27244046) +261 train 8.643265 (lr=7.3287e-05) (hash(x)=27224685) +262 train 8.676275 (lr=7.3566e-05) (hash(x)=28274477) +263 train 8.338901 (lr=7.3846e-05) (hash(x)=23557495) +264 train 8.433718 (lr=7.4126e-05) (hash(x)=24680596) +265 train 8.298604 (lr=7.4406e-05) (hash(x)=23928957) +266 train 8.300453 (lr=7.4685e-05) (hash(x)=23761390) +267 train 8.349640 (lr=7.4965e-05) (hash(x)=25288123) +268 train 8.412529 (lr=7.5245e-05) (hash(x)=28705502) +269 train 8.214411 (lr=7.5524e-05) (hash(x)=23246294) +270 train 8.127835 (lr=7.5804e-05) (hash(x)=28639079) +271 train 8.057528 (lr=7.6084e-05) (hash(x)=27804380) +272 train 8.333439 (lr=7.6364e-05) (hash(x)=24172235) +273 train 8.405916 (lr=7.6643e-05) (hash(x)=23089140) +274 train 8.565055 (lr=7.6923e-05) (hash(x)=27163701) +275 train 8.619143 (lr=7.7203e-05) (hash(x)=26993263) +276 train 8.403809 (lr=7.7483e-05) (hash(x)=28224233) +277 train 8.593162 (lr=7.7762e-05) (hash(x)=27397203) +278 train 8.509372 (lr=7.8042e-05) (hash(x)=27923882) +279 train 8.360802 (lr=7.8322e-05) (hash(x)=26654908) +280 train 8.366153 (lr=7.8601e-05) (hash(x)=24213147) +281 train 8.150020 (lr=7.8881e-05) (hash(x)=21965022) +282 train 8.248713 (lr=7.9161e-05) (hash(x)=25465685) +283 train 8.320612 (lr=7.9441e-05) (hash(x)=27347722) +284 train 8.332123 (lr=7.9720e-05) (hash(x)=26732050) +285 train 8.441467 (lr=8.0000e-05) (hash(x)=28314127) +286 train 8.240745 (lr=8.0280e-05) (hash(x)=21471186) +287 train 8.162442 (lr=8.0559e-05) (hash(x)=23627518) +288 train 8.208360 (lr=8.0839e-05) (hash(x)=20870353) +289 train 8.200301 (lr=8.1119e-05) (hash(x)=25024764) +290 train 8.008876 (lr=8.1399e-05) (hash(x)=20683822) +291 train 8.278331 (lr=8.1678e-05) (hash(x)=21768671) +292 train 8.245654 (lr=8.1958e-05) (hash(x)=25557309) +293 train 8.355617 (lr=8.2238e-05) (hash(x)=25076667) +294 train 8.103814 (lr=8.2517e-05) (hash(x)=23765822) +295 train 8.093103 (lr=8.2797e-05) (hash(x)=21889990) +296 train 8.184050 (lr=8.3077e-05) (hash(x)=26339893) +297 train 8.030908 (lr=8.3357e-05) (hash(x)=20932794) +298 train 8.073474 (lr=8.3636e-05) (hash(x)=21750070) +299 train 8.191063 (lr=8.3916e-05) (hash(x)=23665838) +300 val loss 8.1502 +300 val perplexity 3464.2300 +300 train 8.797359 (lr=8.4196e-05) (hash(x)=32888061) +301 train 8.600707 (lr=8.4476e-05) (hash(x)=30223582) +302 train 8.397466 (lr=8.4755e-05) (hash(x)=26908418) +303 train 7.817451 (lr=8.5035e-05) (hash(x)=22528001) +304 train 8.355705 (lr=8.5315e-05) (hash(x)=27452187) +305 train 8.123291 (lr=8.5594e-05) (hash(x)=25181641) +306 train 8.218148 (lr=8.5874e-05) (hash(x)=25546593) +307 train 8.188229 (lr=8.6154e-05) (hash(x)=22487328) +308 train 8.276194 (lr=8.6434e-05) (hash(x)=27804274) +309 train 8.356766 (lr=8.6713e-05) (hash(x)=26544630) +310 train 8.395320 (lr=8.6993e-05) (hash(x)=27738934) +311 train 8.263254 (lr=8.7273e-05) (hash(x)=29248942) +312 train 8.077911 (lr=8.7552e-05) (hash(x)=25103452) +313 train 8.242172 (lr=8.7832e-05) (hash(x)=25052066) +314 train 8.067119 (lr=8.8112e-05) (hash(x)=24481302) +315 train 8.017956 (lr=8.8392e-05) (hash(x)=23543273) +316 train 8.053043 (lr=8.8671e-05) (hash(x)=25608244) +317 train 8.204776 (lr=8.8951e-05) (hash(x)=27451288) +318 train 7.840873 (lr=8.9231e-05) (hash(x)=22806491) +319 train 8.049987 (lr=8.9510e-05) (hash(x)=25533417) +320 train 8.005333 (lr=8.9790e-05) (hash(x)=24557997) +321 train 8.030066 (lr=9.0070e-05) (hash(x)=24432899) +322 train 8.191503 (lr=9.0350e-05) (hash(x)=27583287) +323 train 8.107607 (lr=9.0629e-05) (hash(x)=25552036) +324 train 7.921436 (lr=9.0909e-05) (hash(x)=24201868) +325 train 8.185282 (lr=9.1189e-05) (hash(x)=28149782) +326 train 8.054921 (lr=9.1469e-05) (hash(x)=25529698) +327 train 7.699329 (lr=9.1748e-05) (hash(x)=20612533) +328 train 7.767572 (lr=9.2028e-05) (hash(x)=20699000) +329 train 7.737216 (lr=9.2308e-05) (hash(x)=19774173) +330 train 7.765160 (lr=9.2587e-05) (hash(x)=21681646) +331 train 7.655166 (lr=9.2867e-05) (hash(x)=20216795) +332 train 8.276549 (lr=9.3147e-05) (hash(x)=27697998) +333 train 8.051928 (lr=9.3427e-05) (hash(x)=25896435) +334 train 7.938391 (lr=9.3706e-05) (hash(x)=21585310) +335 train 7.980289 (lr=9.3986e-05) (hash(x)=24677740) +336 train 7.874303 (lr=9.4266e-05) (hash(x)=22027900) +337 train 7.878704 (lr=9.4545e-05) (hash(x)=21835643) +338 train 7.780991 (lr=9.4825e-05) (hash(x)=20558462) +339 train 7.743780 (lr=9.5105e-05) (hash(x)=16722715) +340 train 7.813516 (lr=9.5385e-05) (hash(x)=20972655) +341 train 8.183291 (lr=9.5664e-05) (hash(x)=26303975) +342 train 8.032054 (lr=9.5944e-05) (hash(x)=22938170) +343 train 7.998839 (lr=9.6224e-05) (hash(x)=25347203) +344 train 7.975152 (lr=9.6503e-05) (hash(x)=27398686) +345 train 7.963434 (lr=9.6783e-05) (hash(x)=25973417) +346 train 7.907532 (lr=9.7063e-05) (hash(x)=26918389) +347 train 7.897779 (lr=9.7343e-05) (hash(x)=24272489) +348 train 7.931806 (lr=9.7622e-05) (hash(x)=25593714) +349 train 7.984360 (lr=9.7902e-05) (hash(x)=29260846) +350 val loss 7.9140 +350 val perplexity 2735.2266 +350 train 8.105070 (lr=9.8182e-05) (hash(x)=27951602) +351 train 8.170658 (lr=9.8462e-05) (hash(x)=28922363) +352 train 8.050817 (lr=9.8741e-05) (hash(x)=27210734) +353 train 8.051019 (lr=9.9021e-05) (hash(x)=26322572) +354 train 7.954082 (lr=9.9301e-05) (hash(x)=27084665) +355 train 7.862874 (lr=9.9580e-05) (hash(x)=25510798) +356 train 7.957891 (lr=9.9860e-05) (hash(x)=24970921) +357 train 8.009904 (lr=1.0014e-04) (hash(x)=24138948) +358 train 7.899603 (lr=1.0042e-04) (hash(x)=24790211) +359 train 7.933965 (lr=1.0070e-04) (hash(x)=25631397) +360 train 7.828710 (lr=1.0098e-04) (hash(x)=23226625) +361 train 7.858880 (lr=1.0126e-04) (hash(x)=24001903) +362 train 7.858968 (lr=1.0154e-04) (hash(x)=24587948) +363 train 7.575414 (lr=1.0182e-04) (hash(x)=21333676) +364 train 7.457738 (lr=1.0210e-04) (hash(x)=23673779) +365 train 8.741650 (lr=1.0238e-04) (hash(x)=30770484) +366 train 7.965117 (lr=1.0266e-04) (hash(x)=26564899) +367 train 7.940494 (lr=1.0294e-04) (hash(x)=26237983) +368 train 7.710056 (lr=1.0322e-04) (hash(x)=23764356) +369 train 8.036211 (lr=1.0350e-04) (hash(x)=26205744) +370 train 7.774320 (lr=1.0378e-04) (hash(x)=19208770) +371 train 7.977690 (lr=1.0406e-04) (hash(x)=25976502) +372 train 7.774491 (lr=1.0434e-04) (hash(x)=23983933) +373 train 7.774879 (lr=1.0462e-04) (hash(x)=24080636) +374 train 7.891135 (lr=1.0490e-04) (hash(x)=24404047) +375 train 7.868119 (lr=1.0517e-04) (hash(x)=24742645) +376 train 7.881885 (lr=1.0545e-04) (hash(x)=24159600) +377 train 8.084444 (lr=1.0573e-04) (hash(x)=28677257) +378 train 7.876456 (lr=1.0601e-04) (hash(x)=25604111) +379 train 7.857339 (lr=1.0629e-04) (hash(x)=27086333) +380 train 7.799113 (lr=1.0657e-04) (hash(x)=25188207) +381 train 7.977731 (lr=1.0685e-04) (hash(x)=27855233) +382 train 7.613178 (lr=1.0713e-04) (hash(x)=19470039) +383 train 7.882783 (lr=1.0741e-04) (hash(x)=26157660) +384 train 7.801929 (lr=1.0769e-04) (hash(x)=25291570) +385 train 7.766172 (lr=1.0797e-04) (hash(x)=25046062) +386 train 8.043158 (lr=1.0825e-04) (hash(x)=27020337) +387 train 7.717350 (lr=1.0853e-04) (hash(x)=23616370) +388 train 7.372201 (lr=1.0881e-04) (hash(x)=19113218) +389 train 7.725282 (lr=1.0909e-04) (hash(x)=24302232) +390 train 7.764499 (lr=1.0937e-04) (hash(x)=22188949) +391 train 7.576656 (lr=1.0965e-04) (hash(x)=22582169) +392 train 7.771549 (lr=1.0993e-04) (hash(x)=24700570) +393 train 7.693165 (lr=1.1021e-04) (hash(x)=22773833) +394 train 7.566705 (lr=1.1049e-04) (hash(x)=21875928) +395 train 7.791702 (lr=1.1077e-04) (hash(x)=26233189) +396 train 7.672849 (lr=1.1105e-04) (hash(x)=24321467) +397 train 7.898304 (lr=1.1133e-04) (hash(x)=26431507) +398 train 7.999419 (lr=1.1161e-04) (hash(x)=28690877) +399 train 7.941890 (lr=1.1189e-04) (hash(x)=26431960) +400 val loss 7.8147 +400 val perplexity 2476.8091 +400 train 7.786003 (lr=1.1217e-04) (hash(x)=24580300) +401 train 7.785663 (lr=1.1245e-04) (hash(x)=25112360) +402 train 7.856658 (lr=1.1273e-04) (hash(x)=27597243) +403 train 8.324168 (lr=1.1301e-04) (hash(x)=30707498) +404 train 7.959845 (lr=1.1329e-04) (hash(x)=28485465) +405 train 7.633141 (lr=1.1357e-04) (hash(x)=22586447) +406 train 7.784090 (lr=1.1385e-04) (hash(x)=23175270) +407 train 7.835039 (lr=1.1413e-04) (hash(x)=25716176) +408 train 7.929630 (lr=1.1441e-04) (hash(x)=26861373) +409 train 8.064385 (lr=1.1469e-04) (hash(x)=25118971) +410 train 7.553987 (lr=1.1497e-04) (hash(x)=19829066) +411 train 7.687543 (lr=1.1524e-04) (hash(x)=26256420) +412 train 7.923865 (lr=1.1552e-04) (hash(x)=27796153) +413 train 7.711610 (lr=1.1580e-04) (hash(x)=22633318) +414 train 7.717796 (lr=1.1608e-04) (hash(x)=22589383) +415 train 8.111142 (lr=1.1636e-04) (hash(x)=28019788) +416 train 8.052124 (lr=1.1664e-04) (hash(x)=28970440) +417 train 7.696166 (lr=1.1692e-04) (hash(x)=27396089) +418 train 7.716078 (lr=1.1720e-04) (hash(x)=21183513) +419 train 7.752925 (lr=1.1748e-04) (hash(x)=23510110) +420 train 7.952095 (lr=1.1776e-04) (hash(x)=28833467) +421 train 7.827655 (lr=1.1804e-04) (hash(x)=23646926) +422 train 7.770056 (lr=1.1832e-04) (hash(x)=24697272) +423 train 7.584944 (lr=1.1860e-04) (hash(x)=20382963) +424 train 7.587346 (lr=1.1888e-04) (hash(x)=23467595) +425 train 7.739859 (lr=1.1916e-04) (hash(x)=24304768) +426 train 7.669890 (lr=1.1944e-04) (hash(x)=21392328) +427 train 7.681346 (lr=1.1972e-04) (hash(x)=25339466) +428 train 7.628826 (lr=1.2000e-04) (hash(x)=22092542) +429 train 7.742391 (lr=1.2028e-04) (hash(x)=22088696) +430 train 7.670621 (lr=1.2056e-04) (hash(x)=22184471) +431 train 7.766905 (lr=1.2084e-04) (hash(x)=24489647) +432 train 7.782338 (lr=1.2112e-04) (hash(x)=26794132) +433 train 7.541985 (lr=1.2140e-04) (hash(x)=22940357) +434 train 7.752074 (lr=1.2168e-04) (hash(x)=23719522) +435 train 7.611131 (lr=1.2196e-04) (hash(x)=22927699) +436 train 7.771802 (lr=1.2224e-04) (hash(x)=26068576) +437 train 7.900320 (lr=1.2252e-04) (hash(x)=27631132) +438 train 7.771894 (lr=1.2280e-04) (hash(x)=26739991) +439 train 7.848542 (lr=1.2308e-04) (hash(x)=25128502) +440 train 7.879604 (lr=1.2336e-04) (hash(x)=25657260) +441 train 7.640905 (lr=1.2364e-04) (hash(x)=23576982) +442 train 7.867248 (lr=1.2392e-04) (hash(x)=27117886) +443 train 7.836057 (lr=1.2420e-04) (hash(x)=25808969) +444 train 7.743042 (lr=1.2448e-04) (hash(x)=24738238) +445 train 7.750248 (lr=1.2476e-04) (hash(x)=23429962) +446 train 7.892906 (lr=1.2503e-04) (hash(x)=25075165) +447 train 7.759227 (lr=1.2531e-04) (hash(x)=25231390) +448 train 7.527074 (lr=1.2559e-04) (hash(x)=22055054) +449 train 7.769775 (lr=1.2587e-04) (hash(x)=25395441) +450 val loss 7.7544 +450 val perplexity 2331.8193 +450 train 7.788908 (lr=1.2615e-04) (hash(x)=25863209) +451 train 7.471836 (lr=1.2643e-04) (hash(x)=21154388) +452 train 7.458272 (lr=1.2671e-04) (hash(x)=21600876) +453 train 7.600598 (lr=1.2699e-04) (hash(x)=24278611) +454 train 7.471626 (lr=1.2727e-04) (hash(x)=23221720) +455 train 7.654756 (lr=1.2755e-04) (hash(x)=22708977) +456 train 8.245894 (lr=1.2783e-04) (hash(x)=23637758) +457 train 7.989209 (lr=1.2811e-04) (hash(x)=28228490) +458 train 8.024860 (lr=1.2839e-04) (hash(x)=28638071) +459 train 7.840843 (lr=1.2867e-04) (hash(x)=27258353) +460 train 7.880686 (lr=1.2895e-04) (hash(x)=26604728) +461 train 7.719940 (lr=1.2923e-04) (hash(x)=23252199) +462 train 7.844332 (lr=1.2951e-04) (hash(x)=26441427) +463 train 7.583513 (lr=1.2979e-04) (hash(x)=24364920) +464 train 7.738811 (lr=1.3007e-04) (hash(x)=25623792) +465 train 7.632152 (lr=1.3035e-04) (hash(x)=23283905) +466 train 8.007241 (lr=1.3063e-04) (hash(x)=26025267) +467 train 7.808859 (lr=1.3091e-04) (hash(x)=27243972) +468 train 8.191301 (lr=1.3119e-04) (hash(x)=30449945) +469 train 7.920693 (lr=1.3147e-04) (hash(x)=28113043) +470 train 7.708316 (lr=1.3175e-04) (hash(x)=25182521) +471 train 7.745363 (lr=1.3203e-04) (hash(x)=24932925) +472 train 7.551541 (lr=1.3231e-04) (hash(x)=20353098) +473 train 7.492368 (lr=1.3259e-04) (hash(x)=19001259) +474 train 8.010329 (lr=1.3287e-04) (hash(x)=27585685) +475 train 7.748430 (lr=1.3315e-04) (hash(x)=26371091) +476 train 7.568470 (lr=1.3343e-04) (hash(x)=24891798) +477 train 7.604988 (lr=1.3371e-04) (hash(x)=24258817) +478 train 7.667824 (lr=1.3399e-04) (hash(x)=24330263) +479 train 7.776863 (lr=1.3427e-04) (hash(x)=26913684) +480 train 7.774129 (lr=1.3455e-04) (hash(x)=26338455) +481 train 8.035324 (lr=1.3483e-04) (hash(x)=27753043) +482 train 7.766191 (lr=1.3510e-04) (hash(x)=26123289) +483 train 7.851757 (lr=1.3538e-04) (hash(x)=29239611) +484 train 7.773692 (lr=1.3566e-04) (hash(x)=26553003) +485 train 7.734964 (lr=1.3594e-04) (hash(x)=22984557) +486 train 7.229377 (lr=1.3622e-04) (hash(x)=16947491) +487 train 7.254314 (lr=1.3650e-04) (hash(x)=18017792) +488 train 7.411728 (lr=1.3678e-04) (hash(x)=19918608) +489 train 7.909356 (lr=1.3706e-04) (hash(x)=23374526) +490 train 7.797341 (lr=1.3734e-04) (hash(x)=25009505) +491 train 7.727072 (lr=1.3762e-04) (hash(x)=27574089) +492 train 7.907964 (lr=1.3790e-04) (hash(x)=24122664) +493 train 7.751007 (lr=1.3818e-04) (hash(x)=26154906) +494 train 7.736349 (lr=1.3846e-04) (hash(x)=25192767) +495 train 7.878122 (lr=1.3874e-04) (hash(x)=28613882) +496 train 7.734479 (lr=1.3902e-04) (hash(x)=23547219) +497 train 7.729660 (lr=1.3930e-04) (hash(x)=25272182) +498 train 7.597702 (lr=1.3958e-04) (hash(x)=24992761) +499 train 7.870467 (lr=1.3986e-04) (hash(x)=26981914) +500 val loss 7.7307 +500 val perplexity 2277.2988 +500 train 7.629415 (lr=1.4014e-04) (hash(x)=22051933) +501 train 7.827107 (lr=1.4042e-04) (hash(x)=24232348) +502 train 7.793712 (lr=1.4070e-04) (hash(x)=23158331) +503 train 7.630573 (lr=1.4098e-04) (hash(x)=22652243) +504 train 7.680202 (lr=1.4126e-04) (hash(x)=23805602) +505 train 7.976804 (lr=1.4154e-04) (hash(x)=25411991) +506 train 7.298094 (lr=1.4182e-04) (hash(x)=18827215) +507 train 6.947015 (lr=1.4210e-04) (hash(x)=15446025) +508 train 7.288272 (lr=1.4238e-04) (hash(x)=20516263) +509 train 7.817081 (lr=1.4266e-04) (hash(x)=27846176) +510 train 7.519442 (lr=1.4294e-04) (hash(x)=23342449) +511 train 7.783105 (lr=1.4322e-04) (hash(x)=27194521) +512 train 7.581825 (lr=1.4350e-04) (hash(x)=23008284) +513 train 8.055969 (lr=1.4378e-04) (hash(x)=29430001) +514 train 7.358160 (lr=1.4406e-04) (hash(x)=22579319) +515 train 7.412921 (lr=1.4434e-04) (hash(x)=25264518) +516 train 7.762818 (lr=1.4462e-04) (hash(x)=25359075) +517 train 7.759334 (lr=1.4490e-04) (hash(x)=25568956) +518 train 8.035318 (lr=1.4517e-04) (hash(x)=32004108) +519 train 7.667718 (lr=1.4545e-04) (hash(x)=24936836) +520 train 7.834546 (lr=1.4573e-04) (hash(x)=27263338) +521 train 7.899292 (lr=1.4601e-04) (hash(x)=27452099) +522 train 7.757212 (lr=1.4629e-04) (hash(x)=25965406) +523 train 7.962968 (lr=1.4657e-04) (hash(x)=28197282) +524 train 7.670166 (lr=1.4685e-04) (hash(x)=22466209) +525 train 7.729378 (lr=1.4713e-04) (hash(x)=22931889) +526 train 7.870152 (lr=1.4741e-04) (hash(x)=26903920) +527 train 7.738748 (lr=1.4769e-04) (hash(x)=24765578) +528 train 7.827399 (lr=1.4797e-04) (hash(x)=27811359) +529 train 7.620152 (lr=1.4825e-04) (hash(x)=25078649) +530 train 7.818067 (lr=1.4853e-04) (hash(x)=25572416) +531 train 7.894573 (lr=1.4881e-04) (hash(x)=27448185) +532 train 7.974401 (lr=1.4909e-04) (hash(x)=25923719) +533 train 7.710668 (lr=1.4937e-04) (hash(x)=24804856) +534 train 7.800461 (lr=1.4965e-04) (hash(x)=23207829) +535 train 7.741437 (lr=1.4993e-04) (hash(x)=23107416) +536 train 7.600116 (lr=1.5021e-04) (hash(x)=26739531) +537 train 7.560553 (lr=1.5049e-04) (hash(x)=24960796) +538 train 7.761447 (lr=1.5077e-04) (hash(x)=24667802) +539 train 7.865750 (lr=1.5105e-04) (hash(x)=26755138) +540 train 7.745243 (lr=1.5133e-04) (hash(x)=25537132) +541 train 7.635506 (lr=1.5161e-04) (hash(x)=24542526) +542 train 7.423744 (lr=1.5189e-04) (hash(x)=21296355) +543 train 8.012912 (lr=1.5217e-04) (hash(x)=29314255) +544 train 7.704752 (lr=1.5245e-04) (hash(x)=26001799) +545 train 7.772429 (lr=1.5273e-04) (hash(x)=27347755) +546 train 7.744110 (lr=1.5301e-04) (hash(x)=25107798) +547 train 7.559268 (lr=1.5329e-04) (hash(x)=22112669) +548 train 7.571847 (lr=1.5357e-04) (hash(x)=21897967) +549 train 7.699791 (lr=1.5385e-04) (hash(x)=25161929) +550 val loss 7.7041 +550 val perplexity 2217.3101 +550 train 7.633374 (lr=1.5413e-04) (hash(x)=27465106) +551 train 7.689460 (lr=1.5441e-04) (hash(x)=24013079) +552 train 7.600994 (lr=1.5469e-04) (hash(x)=23142015) +553 train 7.707086 (lr=1.5497e-04) (hash(x)=26768629) +554 train 7.712857 (lr=1.5524e-04) (hash(x)=26393383) +555 train 7.548347 (lr=1.5552e-04) (hash(x)=22537194) +556 train 7.806934 (lr=1.5580e-04) (hash(x)=24046036) +557 train 7.755718 (lr=1.5608e-04) (hash(x)=24974360) +558 train 7.993939 (lr=1.5636e-04) (hash(x)=28379928) +559 train 7.723049 (lr=1.5664e-04) (hash(x)=25322001) +560 train 7.809935 (lr=1.5692e-04) (hash(x)=26622031) +561 train 7.471033 (lr=1.5720e-04) (hash(x)=20562247) +562 train 7.864546 (lr=1.5748e-04) (hash(x)=27381885) +563 train 7.987198 (lr=1.5776e-04) (hash(x)=27028126) +564 train 7.831239 (lr=1.5804e-04) (hash(x)=28882928) +565 train 7.752853 (lr=1.5832e-04) (hash(x)=25666355) +566 train 7.857646 (lr=1.5860e-04) (hash(x)=24330810) +567 train 7.771210 (lr=1.5888e-04) (hash(x)=26690440) +568 train 7.642743 (lr=1.5916e-04) (hash(x)=22923592) +569 train 7.763253 (lr=1.5944e-04) (hash(x)=27348418) +570 train 7.818096 (lr=1.5972e-04) (hash(x)=28849848) +571 train 7.828838 (lr=1.6000e-04) (hash(x)=26967331) +572 train 7.606356 (lr=1.6028e-04) (hash(x)=22831467) +573 train 7.683434 (lr=1.6056e-04) (hash(x)=24765121) +574 train 7.653260 (lr=1.6084e-04) (hash(x)=24331857) +575 train 7.481162 (lr=1.6112e-04) (hash(x)=22598512) +576 train 7.676166 (lr=1.6140e-04) (hash(x)=25149353) +577 train 7.523717 (lr=1.6168e-04) (hash(x)=23725598) +578 train 7.781949 (lr=1.6196e-04) (hash(x)=26449557) +579 train 7.789368 (lr=1.6224e-04) (hash(x)=24697985) +580 train 7.788190 (lr=1.6252e-04) (hash(x)=26923059) +581 train 7.634191 (lr=1.6280e-04) (hash(x)=25201962) +582 train 7.302167 (lr=1.6308e-04) (hash(x)=20931520) +583 train 7.370741 (lr=1.6336e-04) (hash(x)=18473911) +584 train 7.478966 (lr=1.6364e-04) (hash(x)=21306267) +585 train 7.744124 (lr=1.6392e-04) (hash(x)=25982840) +586 train 7.641984 (lr=1.6420e-04) (hash(x)=25364874) +587 train 7.544707 (lr=1.6448e-04) (hash(x)=23172124) +588 train 7.852381 (lr=1.6476e-04) (hash(x)=27876897) +589 train 8.477080 (lr=1.6503e-04) (hash(x)=34646114) +590 train 8.794193 (lr=1.6531e-04) (hash(x)=35153576) +591 train 7.715353 (lr=1.6559e-04) (hash(x)=22322442) +592 train 7.959844 (lr=1.6587e-04) (hash(x)=27907331) +593 train 7.875203 (lr=1.6615e-04) (hash(x)=26211794) +594 train 8.028293 (lr=1.6643e-04) (hash(x)=29291512) +595 train 8.249352 (lr=1.6671e-04) (hash(x)=29659121) +596 train 8.071311 (lr=1.6699e-04) (hash(x)=29674399) +597 train 7.706496 (lr=1.6727e-04) (hash(x)=23538306) +598 train 7.820385 (lr=1.6755e-04) (hash(x)=21991524) +599 train 7.749043 (lr=1.6783e-04) (hash(x)=26324153) +600 val loss 7.7073 +600 val perplexity 2224.5654 +600 train 7.690235 (lr=1.6811e-04) (hash(x)=23712082) +601 train 7.756430 (lr=1.6839e-04) (hash(x)=24910403) +602 train 7.699323 (lr=1.6867e-04) (hash(x)=26737205) +603 train 7.844785 (lr=1.6895e-04) (hash(x)=26939970) +604 train 7.710744 (lr=1.6923e-04) (hash(x)=27651943) +605 train 7.997003 (lr=1.6951e-04) (hash(x)=27515446) +606 train 7.915620 (lr=1.6979e-04) (hash(x)=26753129) +607 train 7.609352 (lr=1.7007e-04) (hash(x)=23446058) +608 train 7.900656 (lr=1.7035e-04) (hash(x)=27587849) +609 train 7.769150 (lr=1.7063e-04) (hash(x)=25308253) +610 train 7.784833 (lr=1.7091e-04) (hash(x)=26615098) +611 train 7.832840 (lr=1.7119e-04) (hash(x)=29981801) +612 train 7.888113 (lr=1.7147e-04) (hash(x)=29592345) +613 train 7.932767 (lr=1.7175e-04) (hash(x)=23470413) +614 train 7.824748 (lr=1.7203e-04) (hash(x)=24742370) +615 train 7.828477 (lr=1.7231e-04) (hash(x)=24843741) +616 train 7.719766 (lr=1.7259e-04) (hash(x)=25192548) +617 train 7.808021 (lr=1.7287e-04) (hash(x)=27176996) +618 train 7.624393 (lr=1.7315e-04) (hash(x)=23964552) +619 train 7.631208 (lr=1.7343e-04) (hash(x)=22855363) +620 train 7.605110 (lr=1.7371e-04) (hash(x)=26332996) +621 train 7.418732 (lr=1.7399e-04) (hash(x)=22960957) +622 train 7.775790 (lr=1.7427e-04) (hash(x)=22752597) +623 train 7.645708 (lr=1.7455e-04) (hash(x)=23197102) +624 train 8.167784 (lr=1.7483e-04) (hash(x)=27383319) +625 train 7.705582 (lr=1.7510e-04) (hash(x)=26132276) +626 train 7.574229 (lr=1.7538e-04) (hash(x)=19810497) +627 train 7.659574 (lr=1.7566e-04) (hash(x)=25704919) +628 train 7.857672 (lr=1.7594e-04) (hash(x)=27174264) +629 train 7.989839 (lr=1.7622e-04) (hash(x)=22280814) +630 train 8.702450 (lr=1.7650e-04) (hash(x)=24279448) +631 train 8.163171 (lr=1.7678e-04) (hash(x)=23054940) +632 train 8.244329 (lr=1.7706e-04) (hash(x)=23942400) +633 train 8.433205 (lr=1.7734e-04) (hash(x)=24712416) +634 train 8.062888 (lr=1.7762e-04) (hash(x)=24564658) +635 train 7.801651 (lr=1.7790e-04) (hash(x)=24909904) +636 train 7.779956 (lr=1.7818e-04) (hash(x)=23049534) +637 train 7.871587 (lr=1.7846e-04) (hash(x)=24321591) +638 train 7.649572 (lr=1.7874e-04) (hash(x)=26153298) +639 train 7.673213 (lr=1.7902e-04) (hash(x)=27140757) +640 train 7.420553 (lr=1.7930e-04) (hash(x)=25115907) +641 train 7.464143 (lr=1.7958e-04) (hash(x)=26563770) +642 train 7.707849 (lr=1.7986e-04) (hash(x)=28089252) +643 train 7.598185 (lr=1.8014e-04) (hash(x)=25993110) +644 train 7.552217 (lr=1.8042e-04) (hash(x)=25150008) +645 train 7.849532 (lr=1.8070e-04) (hash(x)=28520222) +646 train 7.531822 (lr=1.8098e-04) (hash(x)=21349943) +647 train 7.551662 (lr=1.8126e-04) (hash(x)=25149419) +648 train 7.767348 (lr=1.8154e-04) (hash(x)=25730641) +649 train 7.927057 (lr=1.8182e-04) (hash(x)=26112813) +650 val loss 7.6845 +650 val perplexity 2174.4495 +650 train 7.741924 (lr=1.8210e-04) (hash(x)=25907805) +651 train 7.796501 (lr=1.8238e-04) (hash(x)=27623643) +652 train 7.729521 (lr=1.8266e-04) (hash(x)=26484959) +653 train 7.806695 (lr=1.8294e-04) (hash(x)=29199854) +654 train 7.770873 (lr=1.8322e-04) (hash(x)=28369628) +655 train 7.581839 (lr=1.8350e-04) (hash(x)=24727764) +656 train 7.325067 (lr=1.8378e-04) (hash(x)=22610673) +657 train 7.272020 (lr=1.8406e-04) (hash(x)=22667179) +658 train 7.025464 (lr=1.8434e-04) (hash(x)=18477300) +659 train 7.490548 (lr=1.8462e-04) (hash(x)=23155773) +660 train 7.265726 (lr=1.8490e-04) (hash(x)=19461032) +661 train 7.628860 (lr=1.8517e-04) (hash(x)=23453788) +662 train 7.745768 (lr=1.8545e-04) (hash(x)=24543466) +663 train 7.489150 (lr=1.8573e-04) (hash(x)=21935931) +664 train 7.315831 (lr=1.8601e-04) (hash(x)=19910292) +665 train 7.645569 (lr=1.8629e-04) (hash(x)=24481079) +666 train 7.307378 (lr=1.8657e-04) (hash(x)=18922411) +667 train 7.349472 (lr=1.8685e-04) (hash(x)=20054917) +668 train 7.581896 (lr=1.8713e-04) (hash(x)=24850470) +669 train 7.628011 (lr=1.8741e-04) (hash(x)=25907741) +670 train 7.791395 (lr=1.8769e-04) (hash(x)=26873522) +671 train 7.811131 (lr=1.8797e-04) (hash(x)=27606073) +672 train 7.947141 (lr=1.8825e-04) (hash(x)=26209645) +673 train 7.746366 (lr=1.8853e-04) (hash(x)=25202001) +674 train 7.796944 (lr=1.8881e-04) (hash(x)=25569462) +675 train 7.751663 (lr=1.8909e-04) (hash(x)=26534487) +676 train 7.762586 (lr=1.8937e-04) (hash(x)=26455057) +677 train 7.604237 (lr=1.8965e-04) (hash(x)=24095850) +678 train 7.842903 (lr=1.8993e-04) (hash(x)=25287752) +679 train 7.253073 (lr=1.9021e-04) (hash(x)=22450341) +680 train 8.053784 (lr=1.9049e-04) (hash(x)=29004853) +681 train 7.720787 (lr=1.9077e-04) (hash(x)=27993763) +682 train 7.527099 (lr=1.9105e-04) (hash(x)=26382658) +683 train 7.701692 (lr=1.9133e-04) (hash(x)=25013073) +684 train 7.762369 (lr=1.9161e-04) (hash(x)=30595809) +685 train 7.817664 (lr=1.9189e-04) (hash(x)=30934371) +686 train 8.430835 (lr=1.9217e-04) (hash(x)=33060834) +687 train 7.732685 (lr=1.9245e-04) (hash(x)=25945859) +688 train 7.565568 (lr=1.9273e-04) (hash(x)=23375678) +689 train 7.601855 (lr=1.9301e-04) (hash(x)=25218689) +690 train 7.726300 (lr=1.9329e-04) (hash(x)=28127397) +691 train 7.645743 (lr=1.9357e-04) (hash(x)=24418091) +692 train 7.559189 (lr=1.9385e-04) (hash(x)=22761099) +693 train 7.514260 (lr=1.9413e-04) (hash(x)=24615466) +694 train 7.666876 (lr=1.9441e-04) (hash(x)=24699240) +695 train 7.567143 (lr=1.9469e-04) (hash(x)=21613707) +696 train 7.633594 (lr=1.9497e-04) (hash(x)=24977554) +697 train 7.594537 (lr=1.9524e-04) (hash(x)=24348175) +698 train 7.771253 (lr=1.9552e-04) (hash(x)=25102767) +699 train 7.728400 (lr=1.9580e-04) (hash(x)=26386157) +700 val loss 7.6312 +700 val perplexity 2061.5151 +700 train 7.638048 (lr=1.9608e-04) (hash(x)=26423460) +701 train 7.604745 (lr=1.9636e-04) (hash(x)=25771047) +702 train 7.670741 (lr=1.9664e-04) (hash(x)=26999875) +703 train 7.594495 (lr=1.9692e-04) (hash(x)=24396519) +704 train 7.583020 (lr=1.9720e-04) (hash(x)=22588122) +705 train 7.947483 (lr=1.9748e-04) (hash(x)=25142399) +706 train 7.595601 (lr=1.9776e-04) (hash(x)=20440214) +707 train 7.490608 (lr=1.9804e-04) (hash(x)=23265507) +708 train 7.575449 (lr=1.9832e-04) (hash(x)=24563470) +709 train 7.514136 (lr=1.9860e-04) (hash(x)=22514858) +710 train 7.843362 (lr=1.9888e-04) (hash(x)=26691212) +711 train 7.895192 (lr=1.9916e-04) (hash(x)=29138828) +712 train 7.856945 (lr=1.9944e-04) (hash(x)=28028528) +713 train 7.712360 (lr=1.9972e-04) (hash(x)=20531210) +714 train 7.563757 (lr=2.0000e-04) (hash(x)=25075352) +715 train 7.655571 (lr=2.0028e-04) (hash(x)=24265353) +716 train 7.752625 (lr=2.0056e-04) (hash(x)=24635726) +717 train 7.571253 (lr=2.0084e-04) (hash(x)=24999726) +718 train 7.795953 (lr=2.0112e-04) (hash(x)=27412910) +719 train 7.572258 (lr=2.0140e-04) (hash(x)=24685515) +720 train 7.646062 (lr=2.0168e-04) (hash(x)=23780329) +721 train 7.434561 (lr=2.0196e-04) (hash(x)=25071701) +722 train 7.577852 (lr=2.0224e-04) (hash(x)=23767130) +723 train 7.572559 (lr=2.0252e-04) (hash(x)=24876269) +724 train 7.703425 (lr=2.0280e-04) (hash(x)=26405773) +725 train 8.104126 (lr=2.0308e-04) (hash(x)=31733180) +726 train 7.267221 (lr=2.0336e-04) (hash(x)=21337509) +727 train 7.367560 (lr=2.0364e-04) (hash(x)=22825749) +728 train 7.782383 (lr=2.0392e-04) (hash(x)=28638695) +729 train 7.662183 (lr=2.0420e-04) (hash(x)=26393943) +730 train 7.801141 (lr=2.0448e-04) (hash(x)=27563583) +731 train 7.310507 (lr=2.0476e-04) (hash(x)=21239652) +732 train 7.459405 (lr=2.0503e-04) (hash(x)=23986428) +733 train 7.422882 (lr=2.0531e-04) (hash(x)=24943881) +734 train 7.970260 (lr=2.0559e-04) (hash(x)=29691448) +735 train 7.895206 (lr=2.0587e-04) (hash(x)=28767869) +736 train 7.580738 (lr=2.0615e-04) (hash(x)=23628188) +737 train 7.839057 (lr=2.0643e-04) (hash(x)=29341482) +738 train 7.680745 (lr=2.0671e-04) (hash(x)=30336570) +739 train 7.624516 (lr=2.0699e-04) (hash(x)=25614301) +740 train 7.535782 (lr=2.0727e-04) (hash(x)=24160500) +741 train 7.736092 (lr=2.0755e-04) (hash(x)=26030058) +742 train 8.071903 (lr=2.0783e-04) (hash(x)=29243936) +743 train 7.408698 (lr=2.0811e-04) (hash(x)=21159060) +744 train 7.486439 (lr=2.0839e-04) (hash(x)=23701853) +745 train 7.663116 (lr=2.0867e-04) (hash(x)=24629937) +746 train 7.574273 (lr=2.0895e-04) (hash(x)=25110108) +747 train 7.643520 (lr=2.0923e-04) (hash(x)=26751788) +748 train 7.669766 (lr=2.0951e-04) (hash(x)=26430427) +749 train 7.519864 (lr=2.0979e-04) (hash(x)=26012353) +750 val loss 7.5737 +750 val perplexity 1946.3228 +750 train 7.644215 (lr=2.1007e-04) (hash(x)=22735910) +751 train 7.469612 (lr=2.1035e-04) (hash(x)=25045397) +752 train 7.186100 (lr=2.1063e-04) (hash(x)=21554427) +753 train 7.280351 (lr=2.1091e-04) (hash(x)=23751143) +754 train 8.065994 (lr=2.1119e-04) (hash(x)=28602273) +755 train 8.089591 (lr=2.1147e-04) (hash(x)=29989709) +756 train 7.497457 (lr=2.1175e-04) (hash(x)=22331648) +757 train 7.902699 (lr=2.1203e-04) (hash(x)=31017246) +758 train 7.826366 (lr=2.1231e-04) (hash(x)=29709045) +759 train 7.938167 (lr=2.1259e-04) (hash(x)=25560928) +760 train 7.647357 (lr=2.1287e-04) (hash(x)=25075464) +761 train 7.686697 (lr=2.1315e-04) (hash(x)=27352253) +762 train 7.850749 (lr=2.1343e-04) (hash(x)=28187891) +763 train 7.726692 (lr=2.1371e-04) (hash(x)=26062687) +764 train 7.719066 (lr=2.1399e-04) (hash(x)=27427811) +765 train 7.812685 (lr=2.1427e-04) (hash(x)=27614522) +766 train 7.562927 (lr=2.1455e-04) (hash(x)=26129544) +767 train 8.117548 (lr=2.1483e-04) (hash(x)=28959222) +768 train 7.698925 (lr=2.1510e-04) (hash(x)=26860067) +769 train 7.443318 (lr=2.1538e-04) (hash(x)=25122598) +770 train 7.562387 (lr=2.1566e-04) (hash(x)=25245030) +771 train 7.571133 (lr=2.1594e-04) (hash(x)=25434884) +772 train 7.622682 (lr=2.1622e-04) (hash(x)=27732790) +773 train 7.616994 (lr=2.1650e-04) (hash(x)=27824438) +774 train 7.750951 (lr=2.1678e-04) (hash(x)=27201953) +775 train 7.482411 (lr=2.1706e-04) (hash(x)=21248405) +776 train 7.272460 (lr=2.1734e-04) (hash(x)=22805934) +777 train 7.575542 (lr=2.1762e-04) (hash(x)=26482588) +778 train 7.590247 (lr=2.1790e-04) (hash(x)=24153691) +779 train 7.714342 (lr=2.1818e-04) (hash(x)=25044192) +780 train 7.667870 (lr=2.1846e-04) (hash(x)=25910078) +781 train 7.657702 (lr=2.1874e-04) (hash(x)=28645524) +782 train 7.506832 (lr=2.1902e-04) (hash(x)=24368498) +783 train 7.580148 (lr=2.1930e-04) (hash(x)=25830182) +784 train 7.665555 (lr=2.1958e-04) (hash(x)=29181807) +785 train 7.592044 (lr=2.1986e-04) (hash(x)=25585137) +786 train 7.509233 (lr=2.2014e-04) (hash(x)=24798246) +787 train 7.592251 (lr=2.2042e-04) (hash(x)=26621419) +788 train 7.148273 (lr=2.2070e-04) (hash(x)=21446891) +789 train 7.116399 (lr=2.2098e-04) (hash(x)=22165286) +790 train 7.347281 (lr=2.2126e-04) (hash(x)=23477219) +791 train 7.660460 (lr=2.2154e-04) (hash(x)=25173113) +792 train 7.527387 (lr=2.2182e-04) (hash(x)=25853788) +793 train 7.622820 (lr=2.2210e-04) (hash(x)=27267091) +794 train 7.467865 (lr=2.2238e-04) (hash(x)=23743694) +795 train 7.468395 (lr=2.2266e-04) (hash(x)=24400133) +796 train 7.543504 (lr=2.2294e-04) (hash(x)=23663639) +797 train 7.431534 (lr=2.2322e-04) (hash(x)=23103223) +798 train 7.940204 (lr=2.2350e-04) (hash(x)=28748411) +799 train 7.205155 (lr=2.2378e-04) (hash(x)=23486277) +800 val loss 7.5189 +800 val perplexity 1842.5638 +800 train 7.492054 (lr=2.2406e-04) (hash(x)=25678518) +801 train 7.417660 (lr=2.2434e-04) (hash(x)=23421286) +802 train 7.593326 (lr=2.2462e-04) (hash(x)=26054104) +803 train 7.588770 (lr=2.2490e-04) (hash(x)=25978130) +804 train 7.600864 (lr=2.2517e-04) (hash(x)=26006525) +805 train 7.576558 (lr=2.2545e-04) (hash(x)=25769432) +806 train 7.427661 (lr=2.2573e-04) (hash(x)=22430795) +807 train 7.748378 (lr=2.2601e-04) (hash(x)=28916006) +808 train 7.559361 (lr=2.2629e-04) (hash(x)=25166800) +809 train 7.698528 (lr=2.2657e-04) (hash(x)=24226056) +810 train 7.493913 (lr=2.2685e-04) (hash(x)=25057934) +811 train 8.344107 (lr=2.2713e-04) (hash(x)=33195100) +812 train 7.698192 (lr=2.2741e-04) (hash(x)=26312888) +813 train 7.870038 (lr=2.2769e-04) (hash(x)=27730410) +814 train 7.714164 (lr=2.2797e-04) (hash(x)=27372474) +815 train 7.549402 (lr=2.2825e-04) (hash(x)=25556929) +816 train 7.640759 (lr=2.2853e-04) (hash(x)=26909985) +817 train 7.465608 (lr=2.2881e-04) (hash(x)=25991247) +818 train 7.632067 (lr=2.2909e-04) (hash(x)=27438141) +819 train 7.760474 (lr=2.2937e-04) (hash(x)=29536986) +820 train 7.525220 (lr=2.2965e-04) (hash(x)=24478391) +821 train 7.494245 (lr=2.2993e-04) (hash(x)=26125216) +822 train 7.514154 (lr=2.3021e-04) (hash(x)=26422130) +823 train 7.771907 (lr=2.3049e-04) (hash(x)=29648798) +824 train 7.453516 (lr=2.3077e-04) (hash(x)=21247770) +825 train 7.561367 (lr=2.3105e-04) (hash(x)=23195388) +826 train 7.486989 (lr=2.3133e-04) (hash(x)=25796725) +827 train 7.538743 (lr=2.3161e-04) (hash(x)=23124767) +828 train 7.577192 (lr=2.3189e-04) (hash(x)=25233464) +829 train 7.535623 (lr=2.3217e-04) (hash(x)=25713275) +830 train 7.454314 (lr=2.3245e-04) (hash(x)=25550167) +831 train 7.361808 (lr=2.3273e-04) (hash(x)=24976217) +832 train 7.669777 (lr=2.3301e-04) (hash(x)=28536827) +833 train 7.608041 (lr=2.3329e-04) (hash(x)=27500801) +834 train 7.409499 (lr=2.3357e-04) (hash(x)=25545765) +835 train 7.453441 (lr=2.3385e-04) (hash(x)=23632825) +836 train 7.490743 (lr=2.3413e-04) (hash(x)=25708009) +837 train 7.516370 (lr=2.3441e-04) (hash(x)=24456276) +838 train 7.666122 (lr=2.3469e-04) (hash(x)=29189855) +839 train 7.884171 (lr=2.3497e-04) (hash(x)=31019606) +840 train 7.761130 (lr=2.3524e-04) (hash(x)=26328013) +841 train 7.683393 (lr=2.3552e-04) (hash(x)=25027904) +842 train 7.486340 (lr=2.3580e-04) (hash(x)=23734189) +843 train 7.701155 (lr=2.3608e-04) (hash(x)=28236580) +844 train 7.486937 (lr=2.3636e-04) (hash(x)=26509780) +845 train 7.557217 (lr=2.3664e-04) (hash(x)=25386473) +846 train 7.490911 (lr=2.3692e-04) (hash(x)=24052671) +847 train 7.708571 (lr=2.3720e-04) (hash(x)=28269421) +848 train 7.117858 (lr=2.3748e-04) (hash(x)=22251724) +849 train 7.375565 (lr=2.3776e-04) (hash(x)=24308447) +850 val loss 7.4968 +850 val perplexity 1802.2134 +850 train 7.333034 (lr=2.3804e-04) (hash(x)=24242830) +851 train 7.548427 (lr=2.3832e-04) (hash(x)=25563279) +852 train 7.568655 (lr=2.3860e-04) (hash(x)=26354481) +853 train 7.684388 (lr=2.3888e-04) (hash(x)=26152637) +854 train 7.718719 (lr=2.3916e-04) (hash(x)=28051025) +855 train 7.543172 (lr=2.3944e-04) (hash(x)=24865358) +856 train 7.460557 (lr=2.3972e-04) (hash(x)=24288911) +857 train 7.312505 (lr=2.4000e-04) (hash(x)=22230964) +858 train 7.211596 (lr=2.4028e-04) (hash(x)=21303832) +859 train 7.375171 (lr=2.4056e-04) (hash(x)=22155546) +860 train 7.476053 (lr=2.4084e-04) (hash(x)=25296428) +861 train 7.665496 (lr=2.4112e-04) (hash(x)=29142319) +862 train 7.443094 (lr=2.4140e-04) (hash(x)=25545430) +863 train 7.263494 (lr=2.4168e-04) (hash(x)=26984272) +864 train 7.473186 (lr=2.4196e-04) (hash(x)=25429005) +865 train 7.563750 (lr=2.4224e-04) (hash(x)=27077032) +866 train 7.554494 (lr=2.4252e-04) (hash(x)=26494424) +867 train 7.577733 (lr=2.4280e-04) (hash(x)=23193673) +868 train 7.540765 (lr=2.4308e-04) (hash(x)=25075134) +869 train 7.571756 (lr=2.4336e-04) (hash(x)=27112558) +870 train 7.730793 (lr=2.4364e-04) (hash(x)=27436608) +871 train 7.431173 (lr=2.4392e-04) (hash(x)=24544116) +872 train 7.775933 (lr=2.4420e-04) (hash(x)=31632686) +873 train 7.459625 (lr=2.4448e-04) (hash(x)=25890184) +874 train 7.620160 (lr=2.4476e-04) (hash(x)=22887555) +875 train 7.442400 (lr=2.4503e-04) (hash(x)=24547533) +876 train 7.561583 (lr=2.4531e-04) (hash(x)=26553496) +877 train 7.651526 (lr=2.4559e-04) (hash(x)=27467688) +878 train 7.409323 (lr=2.4587e-04) (hash(x)=24766934) +879 train 7.291982 (lr=2.4615e-04) (hash(x)=22059850) +880 train 7.717697 (lr=2.4643e-04) (hash(x)=22871702) +881 train 7.534711 (lr=2.4671e-04) (hash(x)=23893130) +882 train 7.562534 (lr=2.4699e-04) (hash(x)=25125691) +883 train 7.513532 (lr=2.4727e-04) (hash(x)=25994573) +884 train 7.537247 (lr=2.4755e-04) (hash(x)=26076345) +885 train 7.468202 (lr=2.4783e-04) (hash(x)=26577783) +886 train 7.587338 (lr=2.4811e-04) (hash(x)=27395225) +887 train 7.427608 (lr=2.4839e-04) (hash(x)=23926632) +888 train 7.278594 (lr=2.4867e-04) (hash(x)=21737239) +889 train 7.585377 (lr=2.4895e-04) (hash(x)=23574207) +890 train 7.219833 (lr=2.4923e-04) (hash(x)=24365231) +891 train 7.387833 (lr=2.4951e-04) (hash(x)=27111369) +892 train 7.662264 (lr=2.4979e-04) (hash(x)=27290015) +893 train 7.458275 (lr=2.5007e-04) (hash(x)=23979820) +894 train 7.547031 (lr=2.5035e-04) (hash(x)=26450121) +895 train 7.451131 (lr=2.5063e-04) (hash(x)=27025333) +896 train 7.450174 (lr=2.5091e-04) (hash(x)=23624605) +897 train 7.339319 (lr=2.5119e-04) (hash(x)=22846386) +898 train 7.376901 (lr=2.5147e-04) (hash(x)=22970561) +899 train 7.203673 (lr=2.5175e-04) (hash(x)=16908068) +900 val loss 7.4605 +900 val perplexity 1738.0703 +900 train 7.411653 (lr=2.5203e-04) (hash(x)=24661446) +901 train 7.558141 (lr=2.5231e-04) (hash(x)=25664727) +902 train 7.486881 (lr=2.5259e-04) (hash(x)=25667011) +903 train 7.562444 (lr=2.5287e-04) (hash(x)=29120407) +904 train 7.594120 (lr=2.5315e-04) (hash(x)=23385735) +905 train 7.614103 (lr=2.5343e-04) (hash(x)=25564213) +906 train 7.677373 (lr=2.5371e-04) (hash(x)=25413898) +907 train 7.609928 (lr=2.5399e-04) (hash(x)=27092710) +908 train 7.362422 (lr=2.5427e-04) (hash(x)=25789923) +909 train 7.445940 (lr=2.5455e-04) (hash(x)=28533197) +910 train 7.305092 (lr=2.5483e-04) (hash(x)=22982996) +911 train 7.369134 (lr=2.5510e-04) (hash(x)=23827393) +912 train 7.084000 (lr=2.5538e-04) (hash(x)=21242640) +913 train 7.527719 (lr=2.5566e-04) (hash(x)=24154233) +914 train 7.513446 (lr=2.5594e-04) (hash(x)=24331967) +915 train 7.783154 (lr=2.5622e-04) (hash(x)=32812727) +916 train 7.428472 (lr=2.5650e-04) (hash(x)=23572994) +917 train 7.425394 (lr=2.5678e-04) (hash(x)=26305435) +918 train 7.521004 (lr=2.5706e-04) (hash(x)=26268355) +919 train 7.617249 (lr=2.5734e-04) (hash(x)=27230027) +920 train 7.446352 (lr=2.5762e-04) (hash(x)=23885377) +921 train 7.402520 (lr=2.5790e-04) (hash(x)=23532437) +922 train 7.473314 (lr=2.5818e-04) (hash(x)=25577034) +923 train 7.594168 (lr=2.5846e-04) (hash(x)=25703381) +924 train 7.589828 (lr=2.5874e-04) (hash(x)=27113866) +925 train 7.427747 (lr=2.5902e-04) (hash(x)=26961429) +926 train 7.222391 (lr=2.5930e-04) (hash(x)=21355372) +927 train 7.376823 (lr=2.5958e-04) (hash(x)=24968260) +928 train 7.332779 (lr=2.5986e-04) (hash(x)=25357517) +929 train 7.444406 (lr=2.6014e-04) (hash(x)=24854265) +930 train 7.166138 (lr=2.6042e-04) (hash(x)=21102770) +931 train 7.530703 (lr=2.6070e-04) (hash(x)=25676468) +932 train 7.397392 (lr=2.6098e-04) (hash(x)=22809869) +933 train 7.515044 (lr=2.6126e-04) (hash(x)=25503865) +934 train 7.375136 (lr=2.6154e-04) (hash(x)=24853995) +935 train 7.547159 (lr=2.6182e-04) (hash(x)=27544803) +936 train 7.582397 (lr=2.6210e-04) (hash(x)=25981933) +937 train 7.348733 (lr=2.6238e-04) (hash(x)=24658683) +938 train 7.489519 (lr=2.6266e-04) (hash(x)=23855201) +939 train 7.430785 (lr=2.6294e-04) (hash(x)=24331407) +940 train 7.648281 (lr=2.6322e-04) (hash(x)=29265551) +941 train 7.332585 (lr=2.6350e-04) (hash(x)=21892556) +942 train 7.378923 (lr=2.6378e-04) (hash(x)=27183405) +943 train 7.341776 (lr=2.6406e-04) (hash(x)=26540663) +944 train 7.296903 (lr=2.6434e-04) (hash(x)=25718393) +945 train 7.501007 (lr=2.6462e-04) (hash(x)=26819462) +946 train 7.492872 (lr=2.6490e-04) (hash(x)=27427540) +947 train 7.453744 (lr=2.6517e-04) (hash(x)=25532657) +948 train 7.694570 (lr=2.6545e-04) (hash(x)=27641372) +949 train 7.422721 (lr=2.6573e-04) (hash(x)=26515570) +950 val loss 7.4455 +950 val perplexity 1712.1124 +950 train 7.656774 (lr=2.6601e-04) (hash(x)=26911957) +951 train 7.800310 (lr=2.6629e-04) (hash(x)=25856625) +952 train 7.335590 (lr=2.6657e-04) (hash(x)=25219129) +953 train 7.278200 (lr=2.6685e-04) (hash(x)=25260471) +954 train 7.663245 (lr=2.6713e-04) (hash(x)=29373370) +955 train 7.208820 (lr=2.6741e-04) (hash(x)=23437426) +956 train 7.560237 (lr=2.6769e-04) (hash(x)=23769521) +957 train 7.603613 (lr=2.6797e-04) (hash(x)=25961833) +958 train 7.273198 (lr=2.6825e-04) (hash(x)=23582666) +959 train 7.385067 (lr=2.6853e-04) (hash(x)=23164356) +960 train 7.344376 (lr=2.6881e-04) (hash(x)=24443114) +961 train 7.398649 (lr=2.6909e-04) (hash(x)=25052665) +962 train 7.538443 (lr=2.6937e-04) (hash(x)=27802272) +963 train 7.405740 (lr=2.6965e-04) (hash(x)=25957896) +964 train 7.408214 (lr=2.6993e-04) (hash(x)=26737251) +965 train 7.537184 (lr=2.7021e-04) (hash(x)=24723263) +966 train 7.507994 (lr=2.7049e-04) (hash(x)=24707011) +967 train 7.432172 (lr=2.7077e-04) (hash(x)=25646282) +968 train 7.487988 (lr=2.7105e-04) (hash(x)=27544665) +969 train 7.374023 (lr=2.7133e-04) (hash(x)=25851993) +970 train 7.704267 (lr=2.7161e-04) (hash(x)=29059700) +971 train 7.413632 (lr=2.7189e-04) (hash(x)=21513584) +972 train 7.130844 (lr=2.7217e-04) (hash(x)=23151267) +973 train 7.430782 (lr=2.7245e-04) (hash(x)=26017176) +974 train 7.488590 (lr=2.7273e-04) (hash(x)=26979518) +975 train 7.476188 (lr=2.7301e-04) (hash(x)=23843233) +976 train 7.652481 (lr=2.7329e-04) (hash(x)=24193010) +977 train 7.402660 (lr=2.7357e-04) (hash(x)=21476847) +978 train 7.282546 (lr=2.7385e-04) (hash(x)=21366504) +979 train 7.294516 (lr=2.7413e-04) (hash(x)=23226697) +980 train 7.121054 (lr=2.7441e-04) (hash(x)=19961773) +981 train 6.978329 (lr=2.7469e-04) (hash(x)=19772969) +982 train 7.268858 (lr=2.7497e-04) (hash(x)=23110142) +983 train 7.549109 (lr=2.7524e-04) (hash(x)=24506028) +984 train 7.485195 (lr=2.7552e-04) (hash(x)=25480731) +985 train 7.318305 (lr=2.7580e-04) (hash(x)=21077417) +986 train 7.340670 (lr=2.7608e-04) (hash(x)=23686713) +987 train 7.377912 (lr=2.7636e-04) (hash(x)=26024321) +988 train 7.465696 (lr=2.7664e-04) (hash(x)=27424109) +989 train 7.539336 (lr=2.7692e-04) (hash(x)=27786174) +990 train 7.443260 (lr=2.7720e-04) (hash(x)=25232502) +991 train 7.416818 (lr=2.7748e-04) (hash(x)=22781277) +992 train 7.517578 (lr=2.7776e-04) (hash(x)=26184527) +993 train 7.474454 (lr=2.7804e-04) (hash(x)=24459895) +994 train 7.441172 (lr=2.7832e-04) (hash(x)=25244624) +995 train 7.544560 (lr=2.7860e-04) (hash(x)=24451843) +996 train 7.545395 (lr=2.7888e-04) (hash(x)=22129897) +997 train 7.340460 (lr=2.7916e-04) (hash(x)=21116390) +998 train 7.210990 (lr=2.7944e-04) (hash(x)=20650070) +999 train 7.525536 (lr=2.7972e-04) (hash(x)=24948650) +1000 val loss 7.4128 +1000 val perplexity 1657.0664 +1000 train 7.384187 (lr=2.8000e-04) (hash(x)=25444553) +1001 train 7.424169 (lr=2.8028e-04) (hash(x)=25617781) +1002 train 7.307617 (lr=2.8056e-04) (hash(x)=23862434) +1003 train 7.334592 (lr=2.8084e-04) (hash(x)=25559534) +1004 train 7.550586 (lr=2.8112e-04) (hash(x)=26577585) +1005 train 7.372856 (lr=2.8140e-04) (hash(x)=25546274) +1006 train 7.461088 (lr=2.8168e-04) (hash(x)=26284202) +1007 train 7.488636 (lr=2.8196e-04) (hash(x)=26373991) +1008 train 7.357298 (lr=2.8224e-04) (hash(x)=24612851) +1009 train 7.465871 (lr=2.8252e-04) (hash(x)=26410662) +1010 train 7.494233 (lr=2.8280e-04) (hash(x)=23824841) +1011 train 7.785019 (lr=2.8308e-04) (hash(x)=27756673) +1012 train 7.508904 (lr=2.8336e-04) (hash(x)=25427447) +1013 train 7.558198 (lr=2.8364e-04) (hash(x)=23661686) +1014 train 7.465250 (lr=2.8392e-04) (hash(x)=25129504) +1015 train 7.362031 (lr=2.8420e-04) (hash(x)=23402396) +1016 train 7.597111 (lr=2.8448e-04) (hash(x)=26145557) +1017 train 7.356005 (lr=2.8476e-04) (hash(x)=26547918) +1018 train 7.488058 (lr=2.8503e-04) (hash(x)=26653070) +1019 train 7.480762 (lr=2.8531e-04) (hash(x)=28250354) +1020 train 7.469599 (lr=2.8559e-04) (hash(x)=25437401) +1021 train 7.521523 (lr=2.8587e-04) (hash(x)=25598861) +1022 train 7.358792 (lr=2.8615e-04) (hash(x)=24572587) +1023 train 7.435652 (lr=2.8643e-04) (hash(x)=21562133) +1024 train 7.824697 (lr=2.8671e-04) (hash(x)=27718876) +1025 train 7.383985 (lr=2.8699e-04) (hash(x)=25312391) +1026 train 7.474351 (lr=2.8727e-04) (hash(x)=26579535) +1027 train 7.714063 (lr=2.8755e-04) (hash(x)=27253861) +1028 train 7.556354 (lr=2.8783e-04) (hash(x)=28451867) +1029 train 7.293367 (lr=2.8811e-04) (hash(x)=26795921) +1030 train 7.433921 (lr=2.8839e-04) (hash(x)=26147208) +1031 train 7.504082 (lr=2.8867e-04) (hash(x)=25009210) +1032 train 7.386035 (lr=2.8895e-04) (hash(x)=25213771) +1033 train 7.302167 (lr=2.8923e-04) (hash(x)=26254538) +1034 train 7.370673 (lr=2.8951e-04) (hash(x)=25565614) +1035 train 7.200353 (lr=2.8979e-04) (hash(x)=23052577) +1036 train 8.019770 (lr=2.9007e-04) (hash(x)=29630613) +1037 train 7.280005 (lr=2.9035e-04) (hash(x)=23667224) +1038 train 7.404008 (lr=2.9063e-04) (hash(x)=25740670) +1039 train 7.467009 (lr=2.9091e-04) (hash(x)=27161811) +1040 train 7.612668 (lr=2.9119e-04) (hash(x)=26385663) +1041 train 7.267751 (lr=2.9147e-04) (hash(x)=26313522) +1042 train 7.369490 (lr=2.9175e-04) (hash(x)=26814686) +1043 train 7.543861 (lr=2.9203e-04) (hash(x)=27302459) +1044 train 7.861290 (lr=2.9231e-04) (hash(x)=26758132) +1045 train 7.454199 (lr=2.9259e-04) (hash(x)=26837963) +1046 train 7.254780 (lr=2.9287e-04) (hash(x)=22089547) +1047 train 7.693677 (lr=2.9315e-04) (hash(x)=29515100) +1048 train 7.425275 (lr=2.9343e-04) (hash(x)=25471442) +1049 train 7.386746 (lr=2.9371e-04) (hash(x)=26674478) +1050 val loss 7.3973 +1050 val perplexity 1631.5153 +1050 train 7.365238 (lr=2.9399e-04) (hash(x)=25373386) +1051 train 7.385129 (lr=2.9427e-04) (hash(x)=25318001) +1052 train 7.560556 (lr=2.9455e-04) (hash(x)=27255021) +1053 train 7.474577 (lr=2.9483e-04) (hash(x)=25174043) +1054 train 7.215147 (lr=2.9510e-04) (hash(x)=23857597) +1055 train 7.470221 (lr=2.9538e-04) (hash(x)=25305929) +1056 train 7.528258 (lr=2.9566e-04) (hash(x)=27009246) +1057 train 7.596395 (lr=2.9594e-04) (hash(x)=26175477) +1058 train 7.287915 (lr=2.9622e-04) (hash(x)=22700025) +1059 train 7.362848 (lr=2.9650e-04) (hash(x)=24339043) +1060 train 7.312262 (lr=2.9678e-04) (hash(x)=22139900) +1061 train 8.017066 (lr=2.9706e-04) (hash(x)=25412772) +1062 train 7.488575 (lr=2.9734e-04) (hash(x)=27020849) +1063 train 7.392455 (lr=2.9762e-04) (hash(x)=26808543) +1064 train 7.235596 (lr=2.9790e-04) (hash(x)=23061527) +1065 train 7.400952 (lr=2.9818e-04) (hash(x)=24738650) +1066 train 7.246457 (lr=2.9846e-04) (hash(x)=24036715) +1067 train 7.298908 (lr=2.9874e-04) (hash(x)=25763991) +1068 train 7.575664 (lr=2.9902e-04) (hash(x)=27393753) +1069 train 7.697719 (lr=2.9930e-04) (hash(x)=28182190) +1070 train 7.340035 (lr=2.9958e-04) (hash(x)=23358569) +1071 train 7.498947 (lr=2.9986e-04) (hash(x)=25669509) +1072 train 7.617255 (lr=3.0014e-04) (hash(x)=29139024) +1073 train 7.543153 (lr=3.0042e-04) (hash(x)=25616522) +1074 train 7.443698 (lr=3.0070e-04) (hash(x)=25695789) +1075 train 7.612895 (lr=3.0098e-04) (hash(x)=27676869) +1076 train 7.468233 (lr=3.0126e-04) (hash(x)=25952695) +1077 train 7.454184 (lr=3.0154e-04) (hash(x)=26316170) +1078 train 7.344853 (lr=3.0182e-04) (hash(x)=24081867) +1079 train 7.559426 (lr=3.0210e-04) (hash(x)=28435805) +1080 train 7.437314 (lr=3.0238e-04) (hash(x)=23375063) +1081 train 7.424577 (lr=3.0266e-04) (hash(x)=26869022) +1082 train 7.327545 (lr=3.0294e-04) (hash(x)=25793007) +1083 train 7.303977 (lr=3.0322e-04) (hash(x)=23455211) +1084 train 7.307920 (lr=3.0350e-04) (hash(x)=20441501) +1085 train 7.698088 (lr=3.0378e-04) (hash(x)=29321187) +1086 train 7.861182 (lr=3.0406e-04) (hash(x)=32627505) +1087 train 7.355141 (lr=3.0434e-04) (hash(x)=26482758) +1088 train 7.259872 (lr=3.0462e-04) (hash(x)=21431511) +1089 train 7.489113 (lr=3.0490e-04) (hash(x)=26046639) +1090 train 7.502962 (lr=3.0517e-04) (hash(x)=27464841) +1091 train 7.389236 (lr=3.0545e-04) (hash(x)=27068280) +1092 train 7.265472 (lr=3.0573e-04) (hash(x)=23119133) +1093 train 7.454473 (lr=3.0601e-04) (hash(x)=26782091) +1094 train 7.480035 (lr=3.0629e-04) (hash(x)=26265326) +1095 train 7.525990 (lr=3.0657e-04) (hash(x)=24929929) +1096 train 7.377866 (lr=3.0685e-04) (hash(x)=23158628) +1097 train 7.465065 (lr=3.0713e-04) (hash(x)=25950541) +1098 train 7.264986 (lr=3.0741e-04) (hash(x)=22093912) +1099 train 7.394191 (lr=3.0769e-04) (hash(x)=25373676) +1100 val loss 7.4087 +1100 val perplexity 1650.2693 +1100 train 7.154834 (lr=3.0797e-04) (hash(x)=18986670) +1101 train 7.618595 (lr=3.0825e-04) (hash(x)=27283187) +1102 train 7.393076 (lr=3.0853e-04) (hash(x)=25474743) +1103 train 7.327985 (lr=3.0881e-04) (hash(x)=25043037) +1104 train 7.231018 (lr=3.0909e-04) (hash(x)=23156261) +1105 train 7.689672 (lr=3.0937e-04) (hash(x)=27027534) +1106 train 7.273616 (lr=3.0965e-04) (hash(x)=22733630) +1107 train 7.523715 (lr=3.0993e-04) (hash(x)=27906976) +1108 train 7.480759 (lr=3.1021e-04) (hash(x)=27848655) +1109 train 7.416396 (lr=3.1049e-04) (hash(x)=23889709) +1110 train 8.465189 (lr=3.1077e-04) (hash(x)=33189918) +1111 train 8.017624 (lr=3.1105e-04) (hash(x)=29151257) +1112 train 7.669467 (lr=3.1133e-04) (hash(x)=23881512) +1113 train 7.432558 (lr=3.1161e-04) (hash(x)=21307974) +1114 train 7.491772 (lr=3.1189e-04) (hash(x)=25264524) +1115 train 7.639293 (lr=3.1217e-04) (hash(x)=26405613) +1116 train 7.315568 (lr=3.1245e-04) (hash(x)=21918678) +1117 train 7.318353 (lr=3.1273e-04) (hash(x)=24233887) +1118 train 6.917146 (lr=3.1301e-04) (hash(x)=19490509) +1119 train 7.498358 (lr=3.1329e-04) (hash(x)=26400365) +1120 train 7.728421 (lr=3.1357e-04) (hash(x)=28572086) +1121 train 7.203462 (lr=3.1385e-04) (hash(x)=22293114) +1122 train 7.342314 (lr=3.1413e-04) (hash(x)=26845479) +1123 train 7.273404 (lr=3.1441e-04) (hash(x)=23971905) +1124 train 7.460256 (lr=3.1469e-04) (hash(x)=25639959) +1125 train 7.056110 (lr=3.1497e-04) (hash(x)=20076502) +1126 train 7.466146 (lr=3.1524e-04) (hash(x)=25089255) +1127 train 7.339288 (lr=3.1552e-04) (hash(x)=24098812) +1128 train 7.360893 (lr=3.1580e-04) (hash(x)=23493707) +1129 train 7.582637 (lr=3.1608e-04) (hash(x)=27610410) +1130 train 7.340979 (lr=3.1636e-04) (hash(x)=24540186) +1131 train 7.622218 (lr=3.1664e-04) (hash(x)=29402976) +1132 train 7.270705 (lr=3.1692e-04) (hash(x)=23776025) +1133 train 7.035535 (lr=3.1720e-04) (hash(x)=19032564) +1134 train 7.564939 (lr=3.1748e-04) (hash(x)=26921117) +1135 train 7.180524 (lr=3.1776e-04) (hash(x)=20967666) +1136 train 7.232286 (lr=3.1804e-04) (hash(x)=23394540) +1137 train 7.287587 (lr=3.1832e-04) (hash(x)=22666342) +1138 train 7.414471 (lr=3.1860e-04) (hash(x)=23482498) +1139 train 7.477694 (lr=3.1888e-04) (hash(x)=24287610) +1140 train 7.564326 (lr=3.1916e-04) (hash(x)=24512831) +1141 train 7.793725 (lr=3.1944e-04) (hash(x)=28637634) +1142 train 8.027216 (lr=3.1972e-04) (hash(x)=24107127) +1143 train 7.678504 (lr=3.2000e-04) (hash(x)=28667963) +1144 train 7.602080 (lr=3.2028e-04) (hash(x)=26302492) +1145 train 7.527909 (lr=3.2056e-04) (hash(x)=23621685) +1146 train 7.354267 (lr=3.2084e-04) (hash(x)=23997414) +1147 train 7.416944 (lr=3.2112e-04) (hash(x)=25974316) +1148 train 7.350284 (lr=3.2140e-04) (hash(x)=23464988) +1149 train 7.454121 (lr=3.2168e-04) (hash(x)=24487710) +1150 val loss 7.4191 +1150 val perplexity 1667.4720 +1150 train 7.683704 (lr=3.2196e-04) (hash(x)=29610350) +1151 train 7.813034 (lr=3.2224e-04) (hash(x)=27083417) +1152 train 7.083270 (lr=3.2252e-04) (hash(x)=20666288) +1153 train 7.209989 (lr=3.2280e-04) (hash(x)=22850633) +1154 train 7.415977 (lr=3.2308e-04) (hash(x)=24741872) +1155 train 7.773126 (lr=3.2336e-04) (hash(x)=28605205) +1156 train 7.269972 (lr=3.2364e-04) (hash(x)=22624463) +1157 train 7.468536 (lr=3.2392e-04) (hash(x)=24786468) +1158 train 7.228448 (lr=3.2420e-04) (hash(x)=21365399) +1159 train 7.432782 (lr=3.2448e-04) (hash(x)=23649001) +1160 train 7.506412 (lr=3.2476e-04) (hash(x)=28203982) +1161 train 7.410822 (lr=3.2503e-04) (hash(x)=26473994) +1162 train 7.075621 (lr=3.2531e-04) (hash(x)=19476441) +1163 train 7.352562 (lr=3.2559e-04) (hash(x)=25921700) +1164 train 7.208086 (lr=3.2587e-04) (hash(x)=23064343) +1165 train 7.366436 (lr=3.2615e-04) (hash(x)=24117626) +1166 train 7.273366 (lr=3.2643e-04) (hash(x)=21764556) +1167 train 8.103148 (lr=3.2671e-04) (hash(x)=31338300) +1168 train 7.609395 (lr=3.2699e-04) (hash(x)=27443187) +1169 train 7.025002 (lr=3.2727e-04) (hash(x)=21337692) +1170 train 7.566419 (lr=3.2755e-04) (hash(x)=27845383) +1171 train 7.355915 (lr=3.2783e-04) (hash(x)=23862328) +1172 train 7.511925 (lr=3.2811e-04) (hash(x)=23811014) +1173 train 7.484865 (lr=3.2839e-04) (hash(x)=24380098) +1174 train 7.206888 (lr=3.2867e-04) (hash(x)=22351136) +1175 train 7.495497 (lr=3.2895e-04) (hash(x)=30603174) +1176 train 7.504399 (lr=3.2923e-04) (hash(x)=27924596) +1177 train 7.882705 (lr=3.2951e-04) (hash(x)=30882548) +1178 train 7.333656 (lr=3.2979e-04) (hash(x)=22339464) +1179 train 7.175060 (lr=3.3007e-04) (hash(x)=23603806) +1180 train 7.512334 (lr=3.3035e-04) (hash(x)=24809041) +1181 train 7.332124 (lr=3.3063e-04) (hash(x)=24382442) +1182 train 7.302800 (lr=3.3091e-04) (hash(x)=23134077) +1183 train 7.344768 (lr=3.3119e-04) (hash(x)=24830965) +1184 train 7.446421 (lr=3.3147e-04) (hash(x)=25527259) +1185 train 7.459558 (lr=3.3175e-04) (hash(x)=25547480) +1186 train 7.326536 (lr=3.3203e-04) (hash(x)=24424314) +1187 train 7.662166 (lr=3.3231e-04) (hash(x)=32488729) +1188 train 7.596594 (lr=3.3259e-04) (hash(x)=31168462) +1189 train 7.827526 (lr=3.3287e-04) (hash(x)=31331643) +1190 train 7.538304 (lr=3.3315e-04) (hash(x)=28746633) +1191 train 7.543057 (lr=3.3343e-04) (hash(x)=27269893) +1192 train 7.378700 (lr=3.3371e-04) (hash(x)=23484031) +1193 train 7.408017 (lr=3.3399e-04) (hash(x)=23278725) +1194 train 7.515817 (lr=3.3427e-04) (hash(x)=25440745) +1195 train 7.546539 (lr=3.3455e-04) (hash(x)=25215077) +1196 train 7.563812 (lr=3.3483e-04) (hash(x)=28309266) +1197 train 7.323309 (lr=3.3510e-04) (hash(x)=22442404) +1198 train 7.312479 (lr=3.3538e-04) (hash(x)=20775262) +1199 train 7.357631 (lr=3.3566e-04) (hash(x)=24825041) +1200 val loss 7.4245 +1200 val perplexity 1676.5851 +1200 train 7.546093 (lr=3.3594e-04) (hash(x)=29016896) +1201 train 7.496783 (lr=3.3622e-04) (hash(x)=25590276) +1202 train 7.292990 (lr=3.3650e-04) (hash(x)=22816384) +1203 train 7.468839 (lr=3.3678e-04) (hash(x)=25893804) +1204 train 7.354298 (lr=3.3706e-04) (hash(x)=23983816) +1205 train 7.251433 (lr=3.3734e-04) (hash(x)=22222468) +1206 train 7.312139 (lr=3.3762e-04) (hash(x)=23039141) +1207 train 7.238496 (lr=3.3790e-04) (hash(x)=22724871) +1208 train 7.301498 (lr=3.3818e-04) (hash(x)=25253062) +1209 train 7.283185 (lr=3.3846e-04) (hash(x)=25140744) +1210 train 7.375409 (lr=3.3874e-04) (hash(x)=24717218) +1211 train 7.482564 (lr=3.3902e-04) (hash(x)=27317554) +1212 train 7.108214 (lr=3.3930e-04) (hash(x)=22013236) +1213 train 7.658089 (lr=3.3958e-04) (hash(x)=26692343) +1214 train 7.154687 (lr=3.3986e-04) (hash(x)=19859225) +1215 train 7.144209 (lr=3.4014e-04) (hash(x)=24683717) +1216 train 7.162868 (lr=3.4042e-04) (hash(x)=21932013) +1217 train 7.443458 (lr=3.4070e-04) (hash(x)=26849303) +1218 train 7.549216 (lr=3.4098e-04) (hash(x)=26458000) +1219 train 7.274923 (lr=3.4126e-04) (hash(x)=21584083) +1220 train 7.677803 (lr=3.4154e-04) (hash(x)=23598625) +1221 train 7.364795 (lr=3.4182e-04) (hash(x)=26059939) +1222 train 7.648741 (lr=3.4210e-04) (hash(x)=25481982) +1223 train 7.516655 (lr=3.4238e-04) (hash(x)=26190337) +1224 train 7.669898 (lr=3.4266e-04) (hash(x)=28767755) +1225 train 7.300704 (lr=3.4294e-04) (hash(x)=23663918) +1226 train 7.450237 (lr=3.4322e-04) (hash(x)=21293227) +1227 train 7.405928 (lr=3.4350e-04) (hash(x)=22249019) +1228 train 7.529827 (lr=3.4378e-04) (hash(x)=26886529) +1229 train 7.328483 (lr=3.4406e-04) (hash(x)=25779849) +1230 train 7.418723 (lr=3.4434e-04) (hash(x)=26228964) +1231 train 7.351727 (lr=3.4462e-04) (hash(x)=23086289) +1232 train 7.451972 (lr=3.4490e-04) (hash(x)=23198922) +1233 train 7.485783 (lr=3.4517e-04) (hash(x)=27523941) +1234 train 7.270332 (lr=3.4545e-04) (hash(x)=24293992) +1235 train 8.082892 (lr=3.4573e-04) (hash(x)=28047044) +1236 train 7.538865 (lr=3.4601e-04) (hash(x)=23688671) +1237 train 7.415047 (lr=3.4629e-04) (hash(x)=25125498) +1238 train 7.266561 (lr=3.4657e-04) (hash(x)=21977292) +1239 train 7.352270 (lr=3.4685e-04) (hash(x)=23593875) +1240 train 7.270769 (lr=3.4713e-04) (hash(x)=22659030) +1241 train 7.384534 (lr=3.4741e-04) (hash(x)=25117733) +1242 train 7.345417 (lr=3.4769e-04) (hash(x)=22322808) +1243 train 7.371336 (lr=3.4797e-04) (hash(x)=26059735) +1244 train 7.340421 (lr=3.4825e-04) (hash(x)=22485526) +1245 train 7.252218 (lr=3.4853e-04) (hash(x)=23028679) +1246 train 7.278121 (lr=3.4881e-04) (hash(x)=22906035) +1247 train 7.172519 (lr=3.4909e-04) (hash(x)=22414190) +1248 train 7.474005 (lr=3.4937e-04) (hash(x)=26229624) +1249 train 7.228960 (lr=3.4965e-04) (hash(x)=22223137) +1250 val loss 7.3948 +1250 val perplexity 1627.5579 +1250 train 7.505319 (lr=3.4993e-04) (hash(x)=24911520) +1251 train 7.632571 (lr=3.5021e-04) (hash(x)=24204591) +1252 train 7.263335 (lr=3.5049e-04) (hash(x)=24379601) +1253 train 7.242146 (lr=3.5077e-04) (hash(x)=22922475) +1254 train 7.294411 (lr=3.5105e-04) (hash(x)=23426249) +1255 train 7.575781 (lr=3.5133e-04) (hash(x)=26758114) +1256 train 7.262626 (lr=3.5161e-04) (hash(x)=24027111) +1257 train 7.418173 (lr=3.5189e-04) (hash(x)=25358064) +1258 train 7.263065 (lr=3.5217e-04) (hash(x)=22571285) +1259 train 7.631128 (lr=3.5245e-04) (hash(x)=20084233) +1260 train 7.432343 (lr=3.5273e-04) (hash(x)=18786581) +1261 train 7.410121 (lr=3.5301e-04) (hash(x)=26395104) +1262 train 7.297775 (lr=3.5329e-04) (hash(x)=24212567) +1263 train 7.145768 (lr=3.5357e-04) (hash(x)=21563184) +1264 train 7.443101 (lr=3.5385e-04) (hash(x)=26490150) +1265 train 7.509189 (lr=3.5413e-04) (hash(x)=25207694) +1266 train 7.343498 (lr=3.5441e-04) (hash(x)=23914544) +1267 train 7.340261 (lr=3.5469e-04) (hash(x)=23861489) +1268 train 7.764293 (lr=3.5497e-04) (hash(x)=30714540) +1269 train 7.343624 (lr=3.5524e-04) (hash(x)=23471007) +1270 train 7.341895 (lr=3.5552e-04) (hash(x)=23244293) +1271 train 7.262004 (lr=3.5580e-04) (hash(x)=19218470) +1272 train 7.559935 (lr=3.5608e-04) (hash(x)=26965136) +1273 train 7.322347 (lr=3.5636e-04) (hash(x)=22944035) +1274 train 7.315765 (lr=3.5664e-04) (hash(x)=22002714) +1275 train 7.608200 (lr=3.5692e-04) (hash(x)=28469562) +1276 train 7.743584 (lr=3.5720e-04) (hash(x)=26889992) +1277 train 7.754676 (lr=3.5748e-04) (hash(x)=26452814) +1278 train 7.522922 (lr=3.5776e-04) (hash(x)=28397488) +1279 train 7.604513 (lr=3.5804e-04) (hash(x)=25588469) +1280 train 7.472000 (lr=3.5832e-04) (hash(x)=24833139) +1281 train 7.374918 (lr=3.5860e-04) (hash(x)=24788298) +1282 train 7.576473 (lr=3.5888e-04) (hash(x)=24979383) +1283 train 7.568932 (lr=3.5916e-04) (hash(x)=25236367) +1284 train 7.448618 (lr=3.5944e-04) (hash(x)=22638257) +1285 train 7.298364 (lr=3.5972e-04) (hash(x)=23069067) +1286 train 7.425020 (lr=3.6000e-04) (hash(x)=25133239) +1287 train 8.030078 (lr=3.6028e-04) (hash(x)=30433767) +1288 train 8.050176 (lr=3.6056e-04) (hash(x)=34319079) +1289 train 7.860754 (lr=3.6084e-04) (hash(x)=29268881) +1290 train 7.456378 (lr=3.6112e-04) (hash(x)=24528336) +1291 train 7.616414 (lr=3.6140e-04) (hash(x)=26302626) +1292 train 7.647847 (lr=3.6168e-04) (hash(x)=25479111) +1293 train 7.807379 (lr=3.6196e-04) (hash(x)=24749682) +1294 train 7.447742 (lr=3.6224e-04) (hash(x)=24111393) +1295 train 7.288442 (lr=3.6252e-04) (hash(x)=17851621) +1296 train 7.501873 (lr=3.6280e-04) (hash(x)=26463070) +1297 train 7.482622 (lr=3.6308e-04) (hash(x)=25620741) +1298 train 8.057917 (lr=3.6336e-04) (hash(x)=28225676) +1299 train 7.869789 (lr=3.6364e-04) (hash(x)=27028191) +1300 val loss 7.4505 +1300 val perplexity 1720.7012 +1300 train 7.471285 (lr=3.6392e-04) (hash(x)=29006516) +1301 train 7.629005 (lr=3.6420e-04) (hash(x)=27590299) +1302 train 7.780577 (lr=3.6448e-04) (hash(x)=28678983) +1303 train 7.629777 (lr=3.6476e-04) (hash(x)=25183690) +1304 train 7.181595 (lr=3.6503e-04) (hash(x)=19918097) +1305 train 7.917757 (lr=3.6531e-04) (hash(x)=31114252) +1306 train 7.831932 (lr=3.6559e-04) (hash(x)=30913255) +1307 train 7.470596 (lr=3.6587e-04) (hash(x)=26182243) +1308 train 7.680101 (lr=3.6615e-04) (hash(x)=27420668) +1309 train 7.926286 (lr=3.6643e-04) (hash(x)=30419908) +1310 train 7.820868 (lr=3.6671e-04) (hash(x)=29002281) +1311 train 7.501149 (lr=3.6699e-04) (hash(x)=27040642) +1312 train 7.470657 (lr=3.6727e-04) (hash(x)=26929300) +1313 train 7.540498 (lr=3.6755e-04) (hash(x)=26240761) +1314 train 7.601676 (lr=3.6783e-04) (hash(x)=27161876) +1315 train 7.408963 (lr=3.6811e-04) (hash(x)=24489607) +1316 train 7.758476 (lr=3.6839e-04) (hash(x)=27040115) +1317 train 7.457825 (lr=3.6867e-04) (hash(x)=25012872) +1318 train 7.121348 (lr=3.6895e-04) (hash(x)=20894720) +1319 train 7.260790 (lr=3.6923e-04) (hash(x)=22183303) +1320 train 7.460269 (lr=3.6951e-04) (hash(x)=26291778) +1321 train 7.567188 (lr=3.6979e-04) (hash(x)=27682633) +1322 train 7.366336 (lr=3.7007e-04) (hash(x)=26490892) +1323 train 7.759655 (lr=3.7035e-04) (hash(x)=28844646) +1324 train 8.040054 (lr=3.7063e-04) (hash(x)=29545304) +1325 train 8.147198 (lr=3.7091e-04) (hash(x)=31082070) +1326 train 7.836501 (lr=3.7119e-04) (hash(x)=27486316) +1327 train 7.811071 (lr=3.7147e-04) (hash(x)=27537063) +1328 train 7.550996 (lr=3.7175e-04) (hash(x)=26955557) +1329 train 7.545042 (lr=3.7203e-04) (hash(x)=26125988) +1330 train 7.489328 (lr=3.7231e-04) (hash(x)=22800125) +1331 train 7.662470 (lr=3.7259e-04) (hash(x)=25446686) +1332 train 7.755853 (lr=3.7287e-04) (hash(x)=28743746) +1333 train 7.540775 (lr=3.7315e-04) (hash(x)=19819857) +1334 train 7.332619 (lr=3.7343e-04) (hash(x)=23518628) +1335 train 7.231468 (lr=3.7371e-04) (hash(x)=19492832) +1336 train 7.495526 (lr=3.7399e-04) (hash(x)=24627720) +1337 train 7.474877 (lr=3.7427e-04) (hash(x)=23741214) +1338 train 7.482381 (lr=3.7455e-04) (hash(x)=25837914) +1339 train 7.384370 (lr=3.7483e-04) (hash(x)=25484958) +1340 train 7.359272 (lr=3.7510e-04) (hash(x)=23671284) +1341 train 7.451285 (lr=3.7538e-04) (hash(x)=25525370) +1342 train 7.537471 (lr=3.7566e-04) (hash(x)=26585300) +1343 train 7.581855 (lr=3.7594e-04) (hash(x)=25951629) +1344 train 7.775899 (lr=3.7622e-04) (hash(x)=28743135) +1345 train 7.903161 (lr=3.7650e-04) (hash(x)=28557663) +1346 train 7.342786 (lr=3.7678e-04) (hash(x)=22642751) +1347 train 7.298627 (lr=3.7706e-04) (hash(x)=23462798) +1348 train 7.221172 (lr=3.7734e-04) (hash(x)=24292328) +1349 train 7.525566 (lr=3.7762e-04) (hash(x)=27320280) +1350 val loss 7.4705 +1350 val perplexity 1755.4717 +1350 train 7.615200 (lr=3.7790e-04) (hash(x)=27352812) +1351 train 7.455870 (lr=3.7818e-04) (hash(x)=22408682) +1352 train 7.410038 (lr=3.7846e-04) (hash(x)=23144732) +1353 train 7.280135 (lr=3.7874e-04) (hash(x)=22230799) +1354 train 7.635873 (lr=3.7902e-04) (hash(x)=29747687) +1355 train 7.834023 (lr=3.7930e-04) (hash(x)=31317970) +1356 train 7.116941 (lr=3.7958e-04) (hash(x)=20131758) +1357 train 7.482646 (lr=3.7986e-04) (hash(x)=24020983) +1358 train 7.302734 (lr=3.8014e-04) (hash(x)=23998051) +1359 train 7.478630 (lr=3.8042e-04) (hash(x)=27633457) +1360 train 7.153850 (lr=3.8070e-04) (hash(x)=20155247) +1361 train 7.444860 (lr=3.8098e-04) (hash(x)=23766987) +1362 train 7.315575 (lr=3.8126e-04) (hash(x)=25960383) +1363 train 7.609584 (lr=3.8154e-04) (hash(x)=23729283) +1364 train 7.712266 (lr=3.8182e-04) (hash(x)=27775445) +1365 train 7.563869 (lr=3.8210e-04) (hash(x)=25975834) +1366 train 7.397749 (lr=3.8238e-04) (hash(x)=22765259) +1367 train 7.647008 (lr=3.8266e-04) (hash(x)=27635080) +1368 train 7.431101 (lr=3.8294e-04) (hash(x)=23848542) +1369 train 7.548828 (lr=3.8322e-04) (hash(x)=27181156) +1370 train 7.726707 (lr=3.8350e-04) (hash(x)=28321340) +1371 train 7.427845 (lr=3.8378e-04) (hash(x)=25918780) +1372 train 7.433030 (lr=3.8406e-04) (hash(x)=23338297) +1373 train 7.563051 (lr=3.8434e-04) (hash(x)=22370417) +1374 train 7.504886 (lr=3.8462e-04) (hash(x)=24272668) +1375 train 7.115582 (lr=3.8490e-04) (hash(x)=22287596) +1376 train 7.378910 (lr=3.8517e-04) (hash(x)=25257403) +1377 train 7.305379 (lr=3.8545e-04) (hash(x)=21584419) +1378 train 7.483851 (lr=3.8573e-04) (hash(x)=25318823) +1379 train 7.288345 (lr=3.8601e-04) (hash(x)=22694623) +1380 train 7.407480 (lr=3.8629e-04) (hash(x)=23743406) +1381 train 7.638024 (lr=3.8657e-04) (hash(x)=30820846) +1382 train 8.122670 (lr=3.8685e-04) (hash(x)=34639557) +1383 train 8.032954 (lr=3.8713e-04) (hash(x)=35895440) +1384 train 7.527637 (lr=3.8741e-04) (hash(x)=26304839) +1385 train 7.603723 (lr=3.8769e-04) (hash(x)=27058657) +1386 train 7.513813 (lr=3.8797e-04) (hash(x)=25209529) +1387 train 7.406876 (lr=3.8825e-04) (hash(x)=24298886) +1388 train 7.478726 (lr=3.8853e-04) (hash(x)=26200200) +1389 train 7.376162 (lr=3.8881e-04) (hash(x)=27303061) +1390 train 7.486306 (lr=3.8909e-04) (hash(x)=26531628) +1391 train 7.501639 (lr=3.8937e-04) (hash(x)=25888847) +1392 train 7.616852 (lr=3.8965e-04) (hash(x)=28924067) +1393 train 7.107846 (lr=3.8993e-04) (hash(x)=21267641) +1394 train 7.662774 (lr=3.9021e-04) (hash(x)=29419156) +1395 train 7.921213 (lr=3.9049e-04) (hash(x)=30860646) +1396 train 7.440282 (lr=3.9077e-04) (hash(x)=25744434) +1397 train 7.430892 (lr=3.9105e-04) (hash(x)=26881837) +1398 train 7.406436 (lr=3.9133e-04) (hash(x)=26676176) +1399 train 7.614685 (lr=3.9161e-04) (hash(x)=26518800) +1400 val loss 7.4243 +1400 val perplexity 1676.2684 +1400 train 7.362276 (lr=3.9189e-04) (hash(x)=25043238) +1401 train 7.607359 (lr=3.9217e-04) (hash(x)=25863277) +1402 train 7.320827 (lr=3.9245e-04) (hash(x)=24073623) +1403 train 7.362686 (lr=3.9273e-04) (hash(x)=25385523) +1404 train 7.422327 (lr=3.9301e-04) (hash(x)=24958112) +1405 train 7.318079 (lr=3.9329e-04) (hash(x)=23362519) +1406 train 7.615926 (lr=3.9357e-04) (hash(x)=29262616) +1407 train 7.971241 (lr=3.9385e-04) (hash(x)=37519283) +1408 train 8.218820 (lr=3.9413e-04) (hash(x)=33716930) +1409 train 7.432674 (lr=3.9441e-04) (hash(x)=23392584) +1410 train 7.862974 (lr=3.9469e-04) (hash(x)=22877779) +1411 train 8.141687 (lr=3.9497e-04) (hash(x)=20564397) +1412 train 7.669196 (lr=3.9524e-04) (hash(x)=24366235) +1413 train 7.576982 (lr=3.9552e-04) (hash(x)=27065499) +1414 train 7.479882 (lr=3.9580e-04) (hash(x)=26844114) +1415 train 7.579878 (lr=3.9608e-04) (hash(x)=25141945) +1416 train 7.792371 (lr=3.9636e-04) (hash(x)=28813116) +1417 train 7.532581 (lr=3.9664e-04) (hash(x)=25466598) +1418 train 7.443529 (lr=3.9692e-04) (hash(x)=22800032) +1419 train 7.500942 (lr=3.9720e-04) (hash(x)=22717866) +1420 train 7.979633 (lr=3.9748e-04) (hash(x)=28040763) +1421 train 8.052852 (lr=3.9776e-04) (hash(x)=29648992) +1422 train 8.075382 (lr=3.9804e-04) (hash(x)=31747228) +1423 train 7.531941 (lr=3.9832e-04) (hash(x)=28527939) +1424 train 7.460741 (lr=3.9860e-04) (hash(x)=21563992) +1425 train 7.425240 (lr=3.9888e-04) (hash(x)=25134784) +1426 train 7.624716 (lr=3.9916e-04) (hash(x)=28442823) +1427 train 8.139684 (lr=3.9944e-04) (hash(x)=32757059) +1428 train 7.319271 (lr=3.9972e-04) (hash(x)=24110500) +1429 train 7.433826 (lr=4.0000e-04) (hash(x)=24145729) +1430 train 7.454950 (lr=4.0000e-04) (hash(x)=23541086) +1431 train 7.294847 (lr=4.0000e-04) (hash(x)=21942471) +1432 train 7.455037 (lr=4.0000e-04) (hash(x)=24736836) +1433 train 7.588109 (lr=4.0000e-04) (hash(x)=25325444) +1434 train 7.571838 (lr=4.0000e-04) (hash(x)=25188954) +1435 train 8.093584 (lr=4.0000e-04) (hash(x)=24247339) +1436 train 7.574358 (lr=4.0000e-04) (hash(x)=23773363) +1437 train 7.351261 (lr=4.0000e-04) (hash(x)=24142989) +1438 train 7.393840 (lr=4.0000e-04) (hash(x)=24226952) +1439 train 7.598679 (lr=3.9999e-04) (hash(x)=24955630) +1440 train 7.534416 (lr=3.9999e-04) (hash(x)=24563233) +1441 train 7.601363 (lr=3.9999e-04) (hash(x)=25491335) +1442 train 7.926686 (lr=3.9999e-04) (hash(x)=33253763) +1443 train 7.411118 (lr=3.9999e-04) (hash(x)=21368780) +1444 train 7.475051 (lr=3.9999e-04) (hash(x)=26615500) +1445 train 7.627873 (lr=3.9998e-04) (hash(x)=27146278) +1446 train 7.505682 (lr=3.9998e-04) (hash(x)=25904861) +1447 train 7.383117 (lr=3.9998e-04) (hash(x)=25541230) +1448 train 7.615389 (lr=3.9998e-04) (hash(x)=25434227) +1449 train 7.545675 (lr=3.9997e-04) (hash(x)=25375355) +1450 val loss 7.5082 +1450 val perplexity 1822.9609 +1450 train 7.353993 (lr=3.9997e-04) (hash(x)=21921129) +1451 train 7.386900 (lr=3.9997e-04) (hash(x)=23098806) +1452 train 7.398396 (lr=3.9997e-04) (hash(x)=22000544) +1453 train 7.337071 (lr=3.9996e-04) (hash(x)=22061174) +1454 train 7.433599 (lr=3.9996e-04) (hash(x)=24326286) +1455 train 7.421704 (lr=3.9996e-04) (hash(x)=23501481) +1456 train 7.584552 (lr=3.9995e-04) (hash(x)=26397938) +1457 train 7.571579 (lr=3.9995e-04) (hash(x)=24656430) +1458 train 7.304109 (lr=3.9995e-04) (hash(x)=23494971) +1459 train 7.349550 (lr=3.9994e-04) (hash(x)=23802681) +1460 train 7.344705 (lr=3.9994e-04) (hash(x)=21990153) +1461 train 7.416823 (lr=3.9993e-04) (hash(x)=23496118) +1462 train 7.430101 (lr=3.9993e-04) (hash(x)=24233822) +1463 train 7.749707 (lr=3.9992e-04) (hash(x)=22754988) +1464 train 7.478854 (lr=3.9992e-04) (hash(x)=23116635) +1465 train 7.208563 (lr=3.9991e-04) (hash(x)=21461650) +1466 train 7.322182 (lr=3.9991e-04) (hash(x)=22274473) +1467 train 7.386037 (lr=3.9990e-04) (hash(x)=22183009) +1468 train 7.575532 (lr=3.9990e-04) (hash(x)=24972441) +1469 train 7.349534 (lr=3.9989e-04) (hash(x)=22300616) +1470 train 7.125854 (lr=3.9989e-04) (hash(x)=21443060) +1471 train 7.375752 (lr=3.9988e-04) (hash(x)=23475070) +1472 train 7.299520 (lr=3.9988e-04) (hash(x)=22316810) +1473 train 7.696701 (lr=3.9987e-04) (hash(x)=24508407) +1474 train 7.634216 (lr=3.9987e-04) (hash(x)=24603557) +1475 train 7.373850 (lr=3.9986e-04) (hash(x)=23359061) +1476 train 7.438610 (lr=3.9985e-04) (hash(x)=22950844) +1477 train 7.482212 (lr=3.9985e-04) (hash(x)=22588667) +1478 train 7.587915 (lr=3.9984e-04) (hash(x)=26083526) +1479 train 7.523932 (lr=3.9983e-04) (hash(x)=23609959) +1480 train 7.550439 (lr=3.9983e-04) (hash(x)=24088171) +1481 train 7.551502 (lr=3.9982e-04) (hash(x)=25394852) +1482 train 7.483107 (lr=3.9981e-04) (hash(x)=24503040) +1483 train 7.239311 (lr=3.9980e-04) (hash(x)=22408449) +1484 train 7.467961 (lr=3.9980e-04) (hash(x)=25069325) +1485 train 7.513225 (lr=3.9979e-04) (hash(x)=26947398) +1486 train 6.946005 (lr=3.9978e-04) (hash(x)=17527161) +1487 train 7.390148 (lr=3.9977e-04) (hash(x)=23441847) +1488 train 7.406906 (lr=3.9977e-04) (hash(x)=25592861) +1489 train 7.311446 (lr=3.9976e-04) (hash(x)=23101350) +1490 train 7.332391 (lr=3.9975e-04) (hash(x)=24201380) +1491 train 7.394344 (lr=3.9974e-04) (hash(x)=23676365) +1492 train 7.470127 (lr=3.9973e-04) (hash(x)=26668088) +1493 train 7.608051 (lr=3.9972e-04) (hash(x)=26488811) +1494 train 7.589074 (lr=3.9971e-04) (hash(x)=27248109) +1495 train 7.575044 (lr=3.9971e-04) (hash(x)=23096292) +1496 train 7.546152 (lr=3.9970e-04) (hash(x)=22101981) +1497 train 7.705019 (lr=3.9969e-04) (hash(x)=24512380) +1498 train 7.568759 (lr=3.9968e-04) (hash(x)=25367738) +1499 train 7.577885 (lr=3.9967e-04) (hash(x)=27706294) +1500 val loss 7.5286 +1500 val perplexity 1860.4728 +1500 train 7.623603 (lr=3.9966e-04) (hash(x)=24154026) +1501 train 7.381159 (lr=3.9965e-04) (hash(x)=21892472) +1502 train 7.480007 (lr=3.9964e-04) (hash(x)=23662543) +1503 train 7.469293 (lr=3.9963e-04) (hash(x)=26171093) +1504 train 7.518588 (lr=3.9962e-04) (hash(x)=25974292) +1505 train 7.371866 (lr=3.9961e-04) (hash(x)=23191101) +1506 train 7.666059 (lr=3.9960e-04) (hash(x)=26498861) +1507 train 8.134676 (lr=3.9959e-04) (hash(x)=34946941) +1508 train 7.737484 (lr=3.9958e-04) (hash(x)=25442719) +1509 train 7.587135 (lr=3.9957e-04) (hash(x)=25569942) +1510 train 7.308440 (lr=3.9955e-04) (hash(x)=21265828) +1511 train 7.197002 (lr=3.9954e-04) (hash(x)=20060838) +1512 train 7.371647 (lr=3.9953e-04) (hash(x)=22588251) +1513 train 7.346532 (lr=3.9952e-04) (hash(x)=24581020) +1514 train 7.340183 (lr=3.9951e-04) (hash(x)=21882493) +1515 train 7.320291 (lr=3.9950e-04) (hash(x)=21591702) +1516 train 7.297850 (lr=3.9948e-04) (hash(x)=19813285) +1517 train 7.433684 (lr=3.9947e-04) (hash(x)=21870012) +1518 train 7.512668 (lr=3.9946e-04) (hash(x)=26895498) +1519 train 7.537810 (lr=3.9945e-04) (hash(x)=25451055) +1520 train 7.585217 (lr=3.9944e-04) (hash(x)=25395865) +1521 train 6.999301 (lr=3.9942e-04) (hash(x)=18989695) +1522 train 6.821190 (lr=3.9941e-04) (hash(x)=17040678) +1523 train 6.977733 (lr=3.9940e-04) (hash(x)=18691455) +1524 train 7.182381 (lr=3.9938e-04) (hash(x)=19017167) +1525 train 7.485210 (lr=3.9937e-04) (hash(x)=22406503) +1526 train 7.606314 (lr=3.9936e-04) (hash(x)=23036138) +1527 train 7.534135 (lr=3.9934e-04) (hash(x)=24934792) +1528 train 7.483752 (lr=3.9933e-04) (hash(x)=22843127) +1529 train 7.454025 (lr=3.9932e-04) (hash(x)=24249677) +1530 train 7.355682 (lr=3.9930e-04) (hash(x)=23318555) +1531 train 7.370915 (lr=3.9929e-04) (hash(x)=23864361) +1532 train 7.337039 (lr=3.9928e-04) (hash(x)=20242060) +1533 train 7.418236 (lr=3.9926e-04) (hash(x)=21549554) +1534 train 7.284303 (lr=3.9925e-04) (hash(x)=21642024) +1535 train 7.717752 (lr=3.9923e-04) (hash(x)=25367597) +1536 train 7.905531 (lr=3.9922e-04) (hash(x)=27224144) +1537 train 7.477230 (lr=3.9920e-04) (hash(x)=24409290) +1538 train 7.524982 (lr=3.9919e-04) (hash(x)=24987180) +1539 train 7.442319 (lr=3.9917e-04) (hash(x)=27016702) +1540 train 7.631627 (lr=3.9916e-04) (hash(x)=25819636) +1541 train 7.596616 (lr=3.9914e-04) (hash(x)=25520453) +1542 train 7.314507 (lr=3.9913e-04) (hash(x)=23580496) +1543 train 7.575942 (lr=3.9911e-04) (hash(x)=25022148) +1544 train 7.625046 (lr=3.9909e-04) (hash(x)=25447327) +1545 train 7.601602 (lr=3.9908e-04) (hash(x)=25010319) +1546 train 7.498025 (lr=3.9906e-04) (hash(x)=25511777) +1547 train 7.461461 (lr=3.9905e-04) (hash(x)=24459012) +1548 train 7.502395 (lr=3.9903e-04) (hash(x)=25082785) +1549 train 7.592488 (lr=3.9901e-04) (hash(x)=26528670) +1550 val loss 7.5257 +1550 val perplexity 1855.0717 +1550 train 7.530361 (lr=3.9900e-04) (hash(x)=26684367) +1551 train 7.501059 (lr=3.9898e-04) (hash(x)=25431365) +1552 train 7.459583 (lr=3.9896e-04) (hash(x)=24802747) +1553 train 7.303968 (lr=3.9895e-04) (hash(x)=24229522) +1554 train 7.542810 (lr=3.9893e-04) (hash(x)=24757176) +1555 train 7.692514 (lr=3.9891e-04) (hash(x)=25958785) +1556 train 7.668111 (lr=3.9889e-04) (hash(x)=26810735) +1557 train 7.343318 (lr=3.9888e-04) (hash(x)=22882640) +1558 train 7.746626 (lr=3.9886e-04) (hash(x)=27220057) +1559 train 7.971523 (lr=3.9884e-04) (hash(x)=29307788) +1560 train 7.505827 (lr=3.9882e-04) (hash(x)=25214030) +1561 train 7.350655 (lr=3.9881e-04) (hash(x)=23167825) +1562 train 7.460019 (lr=3.9879e-04) (hash(x)=24658511) +1563 train 7.636319 (lr=3.9877e-04) (hash(x)=24965435) +1564 train 7.432910 (lr=3.9875e-04) (hash(x)=23620579) +1565 train 7.446172 (lr=3.9873e-04) (hash(x)=23772949) +1566 train 7.551759 (lr=3.9871e-04) (hash(x)=23807939) +1567 train 7.709119 (lr=3.9869e-04) (hash(x)=25651929) +1568 train 7.631280 (lr=3.9867e-04) (hash(x)=25549933) +1569 train 7.454050 (lr=3.9866e-04) (hash(x)=23146724) +1570 train 7.470947 (lr=3.9864e-04) (hash(x)=24099275) +1571 train 7.671230 (lr=3.9862e-04) (hash(x)=28451031) +1572 train 7.743965 (lr=3.9860e-04) (hash(x)=25488511) +1573 train 7.496834 (lr=3.9858e-04) (hash(x)=23796771) +1574 train 7.560182 (lr=3.9856e-04) (hash(x)=22642764) +1575 train 7.653764 (lr=3.9854e-04) (hash(x)=26312387) +1576 train 7.521056 (lr=3.9852e-04) (hash(x)=22486223) +1577 train 7.439979 (lr=3.9850e-04) (hash(x)=21257377) +1578 train 7.638462 (lr=3.9848e-04) (hash(x)=24328365) +1579 train 7.677663 (lr=3.9845e-04) (hash(x)=25326650) +1580 train 7.629860 (lr=3.9843e-04) (hash(x)=26484837) +1581 train 7.449228 (lr=3.9841e-04) (hash(x)=24605603) +1582 train 7.526299 (lr=3.9839e-04) (hash(x)=24668537) +1583 train 7.448532 (lr=3.9837e-04) (hash(x)=23804913) +1584 train 7.711193 (lr=3.9835e-04) (hash(x)=23615391) +1585 train 7.638814 (lr=3.9833e-04) (hash(x)=24322926) +1586 train 7.444249 (lr=3.9831e-04) (hash(x)=19753104) +1587 train 7.848892 (lr=3.9828e-04) (hash(x)=25537529) +1588 train 7.302771 (lr=3.9826e-04) (hash(x)=22835476) +1589 train 7.450943 (lr=3.9824e-04) (hash(x)=25707197) +1590 train 7.444441 (lr=3.9822e-04) (hash(x)=24191203) +1591 train 7.711069 (lr=3.9820e-04) (hash(x)=26115519) +1592 train 7.598371 (lr=3.9817e-04) (hash(x)=25781547) +1593 train 7.434701 (lr=3.9815e-04) (hash(x)=24850654) +1594 train 7.364734 (lr=3.9813e-04) (hash(x)=23300928) +1595 train 7.587260 (lr=3.9811e-04) (hash(x)=25494804) +1596 train 7.596417 (lr=3.9808e-04) (hash(x)=28169410) +1597 train 7.390686 (lr=3.9806e-04) (hash(x)=21972022) +1598 train 7.889746 (lr=3.9804e-04) (hash(x)=27687290) +1599 train 7.761394 (lr=3.9801e-04) (hash(x)=23210747) +1600 val loss 7.6104 +1600 val perplexity 2019.0446 +1600 train 7.737316 (lr=3.9799e-04) (hash(x)=20362758) +1601 train 7.556198 (lr=3.9797e-04) (hash(x)=17064773) +1602 train 7.508559 (lr=3.9794e-04) (hash(x)=17173694) +1603 train 7.509352 (lr=3.9792e-04) (hash(x)=15717069) +1604 train 7.588397 (lr=3.9789e-04) (hash(x)=15117973) +1605 train 7.419172 (lr=3.9787e-04) (hash(x)=16189636) +1606 train 7.356490 (lr=3.9785e-04) (hash(x)=16552095) +1607 train 7.254639 (lr=3.9782e-04) (hash(x)=18463712) +1608 train 7.406024 (lr=3.9780e-04) (hash(x)=17572155) +1609 train 7.414490 (lr=3.9777e-04) (hash(x)=19970072) +1610 train 7.425323 (lr=3.9775e-04) (hash(x)=20463871) +1611 train 7.499258 (lr=3.9772e-04) (hash(x)=20705573) +1612 train 7.470834 (lr=3.9770e-04) (hash(x)=24441646) +1613 train 7.664419 (lr=3.9767e-04) (hash(x)=22913147) +1614 train 7.737943 (lr=3.9765e-04) (hash(x)=23081598) +1615 train 7.569884 (lr=3.9762e-04) (hash(x)=23245699) +1616 train 7.497740 (lr=3.9759e-04) (hash(x)=23003072) +1617 train 7.727937 (lr=3.9757e-04) (hash(x)=27121904) +1618 train 7.524555 (lr=3.9754e-04) (hash(x)=25092305) +1619 train 7.590149 (lr=3.9752e-04) (hash(x)=23444521) +1620 train 7.570483 (lr=3.9749e-04) (hash(x)=22130531) +1621 train 7.342549 (lr=3.9746e-04) (hash(x)=20917937) +1622 train 7.449306 (lr=3.9744e-04) (hash(x)=22526838) +1623 train 7.575508 (lr=3.9741e-04) (hash(x)=21770300) +1624 train 8.589952 (lr=3.9738e-04) (hash(x)=22333537) +1625 train 8.565280 (lr=3.9736e-04) (hash(x)=24642519) +1626 train 7.801894 (lr=3.9733e-04) (hash(x)=25199038) +1627 train 7.698721 (lr=3.9730e-04) (hash(x)=23334569) +1628 train 7.651857 (lr=3.9727e-04) (hash(x)=23785360) +1629 train 7.460237 (lr=3.9725e-04) (hash(x)=21796200) +1630 train 7.614548 (lr=3.9722e-04) (hash(x)=22389081) +1631 train 7.765163 (lr=3.9719e-04) (hash(x)=25387532) +1632 train 7.337607 (lr=3.9716e-04) (hash(x)=17247578) +1633 train 7.568488 (lr=3.9714e-04) (hash(x)=22203733) +1634 train 7.702593 (lr=3.9711e-04) (hash(x)=24387455) +1635 train 8.187589 (lr=3.9708e-04) (hash(x)=29466743) +1636 train 7.557371 (lr=3.9705e-04) (hash(x)=25354152) +1637 train 7.648956 (lr=3.9702e-04) (hash(x)=21756102) +1638 train 7.591444 (lr=3.9699e-04) (hash(x)=23544029) +1639 train 7.708619 (lr=3.9696e-04) (hash(x)=27718703) +1640 train 7.681671 (lr=3.9694e-04) (hash(x)=25697138) +1641 train 7.556422 (lr=3.9691e-04) (hash(x)=22978681) +1642 train 8.019831 (lr=3.9688e-04) (hash(x)=29527067) +1643 train 7.224710 (lr=3.9685e-04) (hash(x)=17631973) +1644 train 7.566540 (lr=3.9682e-04) (hash(x)=24860172) +1645 train 7.907056 (lr=3.9679e-04) (hash(x)=27947891) +1646 train 7.670348 (lr=3.9676e-04) (hash(x)=23819556) +1647 train 7.702165 (lr=3.9673e-04) (hash(x)=27661052) +1648 train 7.472057 (lr=3.9670e-04) (hash(x)=23014209) +1649 train 7.390039 (lr=3.9667e-04) (hash(x)=23061452) +1650 val loss 7.6076 +1650 val perplexity 2013.4520 +1650 train 7.468273 (lr=3.9664e-04) (hash(x)=21949051) +1651 train 7.486287 (lr=3.9661e-04) (hash(x)=22244096) +1652 train 7.429409 (lr=3.9658e-04) (hash(x)=25997585) +1653 train 7.528018 (lr=3.9655e-04) (hash(x)=25389713) +1654 train 7.472199 (lr=3.9651e-04) (hash(x)=24390520) +1655 train 7.755712 (lr=3.9648e-04) (hash(x)=26538928) +1656 train 7.666959 (lr=3.9645e-04) (hash(x)=28708088) +1657 train 7.785475 (lr=3.9642e-04) (hash(x)=25685569) +1658 train 7.311646 (lr=3.9639e-04) (hash(x)=22163531) +1659 train 7.468732 (lr=3.9636e-04) (hash(x)=26256802) +1660 train 7.906387 (lr=3.9633e-04) (hash(x)=30016026) +1661 train 7.517994 (lr=3.9629e-04) (hash(x)=24412323) +1662 train 7.637953 (lr=3.9626e-04) (hash(x)=24722307) +1663 train 7.508034 (lr=3.9623e-04) (hash(x)=21772654) +1664 train 7.459952 (lr=3.9620e-04) (hash(x)=24984527) +1665 train 7.559517 (lr=3.9616e-04) (hash(x)=25456825) +1666 train 7.821060 (lr=3.9613e-04) (hash(x)=24383976) +1667 train 7.772607 (lr=3.9610e-04) (hash(x)=25026277) +1668 train 7.725955 (lr=3.9607e-04) (hash(x)=28121715) +1669 train 7.726050 (lr=3.9603e-04) (hash(x)=26855439) +1670 train 7.717238 (lr=3.9600e-04) (hash(x)=33370924) +1671 train 7.648963 (lr=3.9597e-04) (hash(x)=26263254) +1672 train 8.046654 (lr=3.9593e-04) (hash(x)=29958158) +1673 train 7.387527 (lr=3.9590e-04) (hash(x)=21917396) +1674 train 7.519344 (lr=3.9587e-04) (hash(x)=24196295) +1675 train 7.554665 (lr=3.9583e-04) (hash(x)=24726314) +1676 train 7.318730 (lr=3.9580e-04) (hash(x)=22137742) +1677 train 7.507581 (lr=3.9576e-04) (hash(x)=24375100) +1678 train 7.429676 (lr=3.9573e-04) (hash(x)=23965039) +1679 train 7.477483 (lr=3.9570e-04) (hash(x)=23552356) +1680 train 7.455016 (lr=3.9566e-04) (hash(x)=26136969) +1681 train 7.656044 (lr=3.9563e-04) (hash(x)=25343112) +1682 train 7.371803 (lr=3.9559e-04) (hash(x)=22831508) +1683 train 7.614247 (lr=3.9556e-04) (hash(x)=26699819) +1684 train 7.713939 (lr=3.9552e-04) (hash(x)=25435504) +1685 train 7.888115 (lr=3.9549e-04) (hash(x)=25218858) +1686 train 8.483872 (lr=3.9545e-04) (hash(x)=25664448) +1687 train 7.548297 (lr=3.9542e-04) (hash(x)=22707141) +1688 train 7.740485 (lr=3.9538e-04) (hash(x)=28001383) +1689 train 7.925345 (lr=3.9534e-04) (hash(x)=29168269) +1690 train 7.888596 (lr=3.9531e-04) (hash(x)=29637522) +1691 train 7.901855 (lr=3.9527e-04) (hash(x)=25676149) +1692 train 7.507612 (lr=3.9524e-04) (hash(x)=20919831) +1693 train 7.745841 (lr=3.9520e-04) (hash(x)=24324401) +1694 train 7.777442 (lr=3.9516e-04) (hash(x)=26862625) +1695 train 8.075342 (lr=3.9513e-04) (hash(x)=34234481) +1696 train 7.707611 (lr=3.9509e-04) (hash(x)=27415195) +1697 train 7.569942 (lr=3.9505e-04) (hash(x)=24511386) +1698 train 7.667357 (lr=3.9502e-04) (hash(x)=26904299) +1699 train 7.479982 (lr=3.9498e-04) (hash(x)=25401624) +1700 val loss 7.5616 +1700 val perplexity 1922.9229 +1700 train 7.536058 (lr=3.9494e-04) (hash(x)=26498489) +1701 train 7.492494 (lr=3.9491e-04) (hash(x)=22340685) +1702 train 7.601968 (lr=3.9487e-04) (hash(x)=26396757) +1703 train 8.117550 (lr=3.9483e-04) (hash(x)=35008670) +1704 train 7.916768 (lr=3.9479e-04) (hash(x)=31782020) +1705 train 7.823029 (lr=3.9475e-04) (hash(x)=29662286) +1706 train 7.955446 (lr=3.9472e-04) (hash(x)=34215782) +1707 train 7.457056 (lr=3.9468e-04) (hash(x)=25682280) +1708 train 7.743313 (lr=3.9464e-04) (hash(x)=27109048) +1709 train 7.763969 (lr=3.9460e-04) (hash(x)=27686498) +1710 train 7.498541 (lr=3.9456e-04) (hash(x)=24247565) +1711 train 7.557172 (lr=3.9452e-04) (hash(x)=25891543) +1712 train 7.660588 (lr=3.9449e-04) (hash(x)=26436778) +1713 train 7.387995 (lr=3.9445e-04) (hash(x)=22538881) +1714 train 7.400153 (lr=3.9441e-04) (hash(x)=21993512) +1715 train 7.801354 (lr=3.9437e-04) (hash(x)=25962407) +1716 train 7.636360 (lr=3.9433e-04) (hash(x)=24130284) +1717 train 7.477924 (lr=3.9429e-04) (hash(x)=24819073) +1718 train 7.354558 (lr=3.9425e-04) (hash(x)=23496648) +1719 train 7.415059 (lr=3.9421e-04) (hash(x)=21971224) +1720 train 7.325845 (lr=3.9417e-04) (hash(x)=21498081) +1721 train 7.376635 (lr=3.9413e-04) (hash(x)=20951317) +1722 train 7.860558 (lr=3.9409e-04) (hash(x)=28848134) +1723 train 7.878454 (lr=3.9405e-04) (hash(x)=30838142) +1724 train 7.837204 (lr=3.9401e-04) (hash(x)=28844392) +1725 train 7.457359 (lr=3.9397e-04) (hash(x)=23102419) +1726 train 7.507933 (lr=3.9393e-04) (hash(x)=22374479) +1727 train 7.737962 (lr=3.9389e-04) (hash(x)=27312186) +1728 train 8.062419 (lr=3.9385e-04) (hash(x)=34362559) +1729 train 7.955684 (lr=3.9381e-04) (hash(x)=30920030) +1730 train 7.423007 (lr=3.9376e-04) (hash(x)=23684126) +1731 train 7.648143 (lr=3.9372e-04) (hash(x)=23657081) +1732 train 7.894197 (lr=3.9368e-04) (hash(x)=28463299) +1733 train 8.086604 (lr=3.9364e-04) (hash(x)=35351214) +1734 train 7.737307 (lr=3.9360e-04) (hash(x)=28855662) +1735 train 7.344902 (lr=3.9356e-04) (hash(x)=22163749) +1736 train 7.331487 (lr=3.9351e-04) (hash(x)=22024585) +1737 train 7.459067 (lr=3.9347e-04) (hash(x)=24629160) +1738 train 7.554626 (lr=3.9343e-04) (hash(x)=25220941) +1739 train 7.502631 (lr=3.9339e-04) (hash(x)=24829818) +1740 train 7.738449 (lr=3.9334e-04) (hash(x)=28453898) +1741 train 7.854755 (lr=3.9330e-04) (hash(x)=25424550) +1742 train 7.567434 (lr=3.9326e-04) (hash(x)=23437814) +1743 train 7.554208 (lr=3.9322e-04) (hash(x)=25942888) +1744 train 7.560115 (lr=3.9317e-04) (hash(x)=24503801) +1745 train 7.778601 (lr=3.9313e-04) (hash(x)=26418501) +1746 train 7.703769 (lr=3.9309e-04) (hash(x)=27177691) +1747 train 7.414745 (lr=3.9304e-04) (hash(x)=23785671) +1748 train 7.810627 (lr=3.9300e-04) (hash(x)=27362772) +1749 train 7.599389 (lr=3.9295e-04) (hash(x)=25097859) +1750 val loss 7.5771 +1750 val perplexity 1953.0135 +1750 train 7.519107 (lr=3.9291e-04) (hash(x)=24662466) +1751 train 7.475060 (lr=3.9287e-04) (hash(x)=25493916) +1752 train 7.586160 (lr=3.9282e-04) (hash(x)=24655868) +1753 train 7.840256 (lr=3.9278e-04) (hash(x)=19819210) +1754 train 7.410295 (lr=3.9273e-04) (hash(x)=21377155) +1755 train 8.231743 (lr=3.9269e-04) (hash(x)=28618702) +1756 train 8.110338 (lr=3.9264e-04) (hash(x)=27324953) +1757 train 7.856616 (lr=3.9260e-04) (hash(x)=29707907) +1758 train 7.679755 (lr=3.9255e-04) (hash(x)=24485921) +1759 train 7.747943 (lr=3.9251e-04) (hash(x)=24289951) +1760 train 7.515771 (lr=3.9246e-04) (hash(x)=22496716) +1761 train 7.637658 (lr=3.9242e-04) (hash(x)=25147026) +1762 train 7.689748 (lr=3.9237e-04) (hash(x)=26851828) +1763 train 7.648717 (lr=3.9233e-04) (hash(x)=26331486) +1764 train 7.504448 (lr=3.9228e-04) (hash(x)=24383054) +1765 train 7.649405 (lr=3.9223e-04) (hash(x)=25885126) +1766 train 7.612776 (lr=3.9219e-04) (hash(x)=25848470) +1767 train 7.462318 (lr=3.9214e-04) (hash(x)=23648671) +1768 train 7.390477 (lr=3.9210e-04) (hash(x)=23168674) +1769 train 7.308305 (lr=3.9205e-04) (hash(x)=20252079) +1770 train 7.505927 (lr=3.9200e-04) (hash(x)=23914287) +1771 train 7.804812 (lr=3.9196e-04) (hash(x)=27292797) +1772 train 7.140891 (lr=3.9191e-04) (hash(x)=16175151) +1773 train 7.216788 (lr=3.9186e-04) (hash(x)=18317379) +1774 train 7.556233 (lr=3.9182e-04) (hash(x)=24464271) +1775 train 7.530467 (lr=3.9177e-04) (hash(x)=24992055) +1776 train 8.138428 (lr=3.9172e-04) (hash(x)=26032451) +1777 train 7.860929 (lr=3.9167e-04) (hash(x)=24734221) +1778 train 7.836180 (lr=3.9163e-04) (hash(x)=25413430) +1779 train 7.673479 (lr=3.9158e-04) (hash(x)=24398330) +1780 train 7.534735 (lr=3.9153e-04) (hash(x)=24568049) +1781 train 7.502492 (lr=3.9148e-04) (hash(x)=23573984) +1782 train 7.673576 (lr=3.9143e-04) (hash(x)=27338410) +1783 train 7.609008 (lr=3.9138e-04) (hash(x)=26005549) +1784 train 7.470811 (lr=3.9134e-04) (hash(x)=24332331) +1785 train 7.312723 (lr=3.9129e-04) (hash(x)=22583878) +1786 train 7.461492 (lr=3.9124e-04) (hash(x)=24599113) +1787 train 7.873412 (lr=3.9119e-04) (hash(x)=30259737) +1788 train 7.675063 (lr=3.9114e-04) (hash(x)=26504237) +1789 train 7.550887 (lr=3.9109e-04) (hash(x)=25196470) +1790 train 7.429427 (lr=3.9104e-04) (hash(x)=23607871) +1791 train 7.389640 (lr=3.9099e-04) (hash(x)=21046554) +1792 train 7.634691 (lr=3.9094e-04) (hash(x)=27285635) +1793 train 7.544189 (lr=3.9089e-04) (hash(x)=25454371) +1794 train 7.645133 (lr=3.9084e-04) (hash(x)=26486575) +1795 train 7.167991 (lr=3.9079e-04) (hash(x)=18773877) +1796 train 7.345236 (lr=3.9074e-04) (hash(x)=22506761) +1797 train 7.548845 (lr=3.9069e-04) (hash(x)=24645168) +1798 train 7.587209 (lr=3.9064e-04) (hash(x)=25046236) +1799 train 7.669736 (lr=3.9059e-04) (hash(x)=26321539) +1800 val loss 7.5315 +1800 val perplexity 1865.9634 +1800 train 7.462975 (lr=3.9054e-04) (hash(x)=24602494) +1801 train 7.630300 (lr=3.9049e-04) (hash(x)=25722713) +1802 train 7.396477 (lr=3.9044e-04) (hash(x)=23711219) +1803 train 7.503728 (lr=3.9039e-04) (hash(x)=22850804) +1804 train 7.773760 (lr=3.9034e-04) (hash(x)=25829388) +1805 train 7.348289 (lr=3.9029e-04) (hash(x)=22503524) +1806 train 7.618736 (lr=3.9024e-04) (hash(x)=26669453) +1807 train 7.581230 (lr=3.9018e-04) (hash(x)=27807872) +1808 train 7.594467 (lr=3.9013e-04) (hash(x)=23543560) +1809 train 7.742278 (lr=3.9008e-04) (hash(x)=25845590) +1810 train 7.528759 (lr=3.9003e-04) (hash(x)=23026635) +1811 train 7.451843 (lr=3.8998e-04) (hash(x)=25545137) +1812 train 7.936531 (lr=3.8993e-04) (hash(x)=27746841) +1813 train 7.335846 (lr=3.8987e-04) (hash(x)=25471231) +1814 train 7.488549 (lr=3.8982e-04) (hash(x)=24333204) +1815 train 7.322668 (lr=3.8977e-04) (hash(x)=21555364) +1816 train 7.528145 (lr=3.8972e-04) (hash(x)=26317977) +1817 train 7.582015 (lr=3.8966e-04) (hash(x)=24474988) +1818 train 7.320265 (lr=3.8961e-04) (hash(x)=21975953) +1819 train 7.866435 (lr=3.8956e-04) (hash(x)=27243798) +1820 train 7.524475 (lr=3.8950e-04) (hash(x)=25097367) +1821 train 7.441789 (lr=3.8945e-04) (hash(x)=24419085) +1822 train 7.393430 (lr=3.8940e-04) (hash(x)=23058837) +1823 train 7.540378 (lr=3.8934e-04) (hash(x)=26325324) +1824 train 7.285767 (lr=3.8929e-04) (hash(x)=20581023) +1825 train 7.546776 (lr=3.8923e-04) (hash(x)=24822111) +1826 train 7.418069 (lr=3.8918e-04) (hash(x)=24336304) +1827 train 7.527481 (lr=3.8913e-04) (hash(x)=26228728) +1828 train 7.636357 (lr=3.8907e-04) (hash(x)=26445781) +1829 train 7.674411 (lr=3.8902e-04) (hash(x)=24707042) +1830 train 7.924321 (lr=3.8896e-04) (hash(x)=27882744) +1831 train 7.483363 (lr=3.8891e-04) (hash(x)=24956294) +1832 train 7.321199 (lr=3.8885e-04) (hash(x)=20984728) +1833 train 7.284750 (lr=3.8880e-04) (hash(x)=23424737) +1834 train 7.665924 (lr=3.8874e-04) (hash(x)=26207120) +1835 train 7.697594 (lr=3.8869e-04) (hash(x)=25592289) +1836 train 7.378343 (lr=3.8863e-04) (hash(x)=24326649) +1837 train 7.581054 (lr=3.8858e-04) (hash(x)=26826109) +1838 train 7.504739 (lr=3.8852e-04) (hash(x)=24759294) +1839 train 7.527322 (lr=3.8847e-04) (hash(x)=24429389) +1840 train 7.512398 (lr=3.8841e-04) (hash(x)=25537519) +1841 train 7.485327 (lr=3.8835e-04) (hash(x)=24747421) +1842 train 7.378088 (lr=3.8830e-04) (hash(x)=23079065) +1843 train 7.798490 (lr=3.8824e-04) (hash(x)=28733708) +1844 train 7.510639 (lr=3.8819e-04) (hash(x)=23937742) +1845 train 7.764615 (lr=3.8813e-04) (hash(x)=29704803) +1846 train 7.942007 (lr=3.8807e-04) (hash(x)=34617155) +1847 train 7.674129 (lr=3.8802e-04) (hash(x)=27929846) +1848 train 7.972771 (lr=3.8796e-04) (hash(x)=28280878) +1849 train 7.172313 (lr=3.8790e-04) (hash(x)=18961171) +1850 val loss 7.5670 +1850 val perplexity 1933.4272 +1850 train 7.649551 (lr=3.8785e-04) (hash(x)=27146015) +1851 train 7.493112 (lr=3.8779e-04) (hash(x)=22417476) +1852 train 7.594025 (lr=3.8773e-04) (hash(x)=24583152) +1853 train 7.641107 (lr=3.8767e-04) (hash(x)=26364120) +1854 train 7.771145 (lr=3.8762e-04) (hash(x)=25946767) +1855 train 7.626045 (lr=3.8756e-04) (hash(x)=25282897) +1856 train 7.563959 (lr=3.8750e-04) (hash(x)=25154557) +1857 train 7.312108 (lr=3.8744e-04) (hash(x)=22082503) +1858 train 7.360002 (lr=3.8738e-04) (hash(x)=23974606) +1859 train 7.473784 (lr=3.8732e-04) (hash(x)=22278062) +1860 train 7.403024 (lr=3.8727e-04) (hash(x)=19509639) +1861 train 7.700893 (lr=3.8721e-04) (hash(x)=24271468) +1862 train 7.283954 (lr=3.8715e-04) (hash(x)=21943809) +1863 train 7.477183 (lr=3.8709e-04) (hash(x)=26802508) +1864 train 7.982578 (lr=3.8703e-04) (hash(x)=34178409) +1865 train 8.056417 (lr=3.8697e-04) (hash(x)=33468379) +1866 train 7.372857 (lr=3.8691e-04) (hash(x)=22109262) +1867 train 7.256140 (lr=3.8685e-04) (hash(x)=22094832) +1868 train 7.443602 (lr=3.8679e-04) (hash(x)=24016953) +1869 train 7.662784 (lr=3.8673e-04) (hash(x)=26759850) +1870 train 7.553824 (lr=3.8667e-04) (hash(x)=23570779) +1871 train 7.879552 (lr=3.8661e-04) (hash(x)=28390416) +1872 train 7.867649 (lr=3.8655e-04) (hash(x)=26215770) +1873 train 7.743562 (lr=3.8649e-04) (hash(x)=22727956) +1874 train 7.548837 (lr=3.8643e-04) (hash(x)=21304587) +1875 train 7.797113 (lr=3.8637e-04) (hash(x)=26379331) +1876 train 7.852718 (lr=3.8631e-04) (hash(x)=25252419) +1877 train 7.704873 (lr=3.8625e-04) (hash(x)=25048158) +1878 train 7.621944 (lr=3.8619e-04) (hash(x)=23588991) +1879 train 7.656865 (lr=3.8613e-04) (hash(x)=26292451) +1880 train 7.499855 (lr=3.8607e-04) (hash(x)=23617086) +1881 train 7.408149 (lr=3.8601e-04) (hash(x)=23336031) +1882 train 7.489578 (lr=3.8595e-04) (hash(x)=19198742) +1883 train 8.039512 (lr=3.8589e-04) (hash(x)=26301866) +1884 train 8.101321 (lr=3.8582e-04) (hash(x)=30114703) +1885 train 7.691032 (lr=3.8576e-04) (hash(x)=26687040) +1886 train 7.597272 (lr=3.8570e-04) (hash(x)=25810717) +1887 train 7.493417 (lr=3.8564e-04) (hash(x)=22915344) +1888 train 7.700537 (lr=3.8558e-04) (hash(x)=27046189) +1889 train 7.337698 (lr=3.8551e-04) (hash(x)=19831900) +1890 train 7.653777 (lr=3.8545e-04) (hash(x)=27786987) +1891 train 7.564349 (lr=3.8539e-04) (hash(x)=25448318) +1892 train 7.574088 (lr=3.8533e-04) (hash(x)=25166953) +1893 train 7.515977 (lr=3.8526e-04) (hash(x)=24993116) +1894 train 7.590830 (lr=3.8520e-04) (hash(x)=25748002) +1895 train 7.845617 (lr=3.8514e-04) (hash(x)=24744383) +1896 train 7.541749 (lr=3.8508e-04) (hash(x)=25560120) +1897 train 7.387880 (lr=3.8501e-04) (hash(x)=23731849) +1898 train 7.352142 (lr=3.8495e-04) (hash(x)=24115851) +1899 train 7.364081 (lr=3.8489e-04) (hash(x)=22835201) +1900 val loss 7.5214 +1900 val perplexity 1847.2096 +1900 train 7.236520 (lr=3.8482e-04) (hash(x)=21927896) +1901 train 7.529112 (lr=3.8476e-04) (hash(x)=24652361) +1902 train 7.690166 (lr=3.8469e-04) (hash(x)=27332405) +1903 train 7.529405 (lr=3.8463e-04) (hash(x)=26284678) +1904 train 7.470582 (lr=3.8457e-04) (hash(x)=23441305) +1905 train 7.591934 (lr=3.8450e-04) (hash(x)=27522881) +1906 train 7.603505 (lr=3.8444e-04) (hash(x)=25521191) +1907 train 7.405583 (lr=3.8437e-04) (hash(x)=24745936) +1908 train 7.327544 (lr=3.8431e-04) (hash(x)=21653364) +1909 train 7.397790 (lr=3.8424e-04) (hash(x)=21973969) +1910 train 7.390563 (lr=3.8418e-04) (hash(x)=24318192) +1911 train 7.348915 (lr=3.8411e-04) (hash(x)=21654269) +1912 train 7.382614 (lr=3.8405e-04) (hash(x)=24142904) +1913 train 7.279846 (lr=3.8398e-04) (hash(x)=21708739) +1914 train 7.471654 (lr=3.8392e-04) (hash(x)=22937923) +1915 train 8.033912 (lr=3.8385e-04) (hash(x)=32387999) +1916 train 8.188723 (lr=3.8379e-04) (hash(x)=28958671) +1917 train 8.261033 (lr=3.8372e-04) (hash(x)=32662682) +1918 train 8.132164 (lr=3.8366e-04) (hash(x)=31364539) +1919 train 8.171165 (lr=3.8359e-04) (hash(x)=29466268) +1920 train 8.167743 (lr=3.8352e-04) (hash(x)=31407564) +1921 train 7.967583 (lr=3.8346e-04) (hash(x)=28372973) +1922 train 8.083382 (lr=3.8339e-04) (hash(x)=28620993) +1923 train 8.121222 (lr=3.8332e-04) (hash(x)=30736727) +1924 train 8.060797 (lr=3.8326e-04) (hash(x)=30370374) +1925 train 7.996422 (lr=3.8319e-04) (hash(x)=29935562) +1926 train 7.999654 (lr=3.8312e-04) (hash(x)=30528627) +1927 train 8.016002 (lr=3.8306e-04) (hash(x)=32616762) +1928 train 7.913048 (lr=3.8299e-04) (hash(x)=29150044) +1929 train 8.045661 (lr=3.8292e-04) (hash(x)=32861403) +1930 train 7.962602 (lr=3.8286e-04) (hash(x)=31053918) +1931 train 7.880732 (lr=3.8279e-04) (hash(x)=29358578) +1932 train 7.920273 (lr=3.8272e-04) (hash(x)=31521292) +1933 train 7.851807 (lr=3.8265e-04) (hash(x)=31218966) +1934 train 7.880950 (lr=3.8258e-04) (hash(x)=31811872) +1935 train 8.057853 (lr=3.8252e-04) (hash(x)=32201138) +1936 train 7.776324 (lr=3.8245e-04) (hash(x)=28631348) +1937 train 7.889883 (lr=3.8238e-04) (hash(x)=30545133) +1938 train 7.899666 (lr=3.8231e-04) (hash(x)=32136080) +1939 train 7.942824 (lr=3.8224e-04) (hash(x)=30041621) +1940 train 7.783548 (lr=3.8217e-04) (hash(x)=30537181) +1941 train 7.830875 (lr=3.8211e-04) (hash(x)=30371023) +1942 train 7.965047 (lr=3.8204e-04) (hash(x)=30734634) +1943 train 7.873985 (lr=3.8197e-04) (hash(x)=33683468) +1944 train 7.890796 (lr=3.8190e-04) (hash(x)=30472451) +1945 train 7.730524 (lr=3.8183e-04) (hash(x)=29809830) +1946 train 7.707180 (lr=3.8176e-04) (hash(x)=26077593) +1947 train 7.713203 (lr=3.8169e-04) (hash(x)=26837755) +1948 train 7.739810 (lr=3.8162e-04) (hash(x)=23910920) +1949 train 7.742172 (lr=3.8155e-04) (hash(x)=26084089) +1950 val loss 7.7375 +1950 val perplexity 2292.8169 +1950 train 7.919502 (lr=3.8148e-04) (hash(x)=26963533) +1951 train 7.803368 (lr=3.8141e-04) (hash(x)=27498125) +1952 train 7.927171 (lr=3.8134e-04) (hash(x)=29512508) +1953 train 8.050812 (lr=3.8127e-04) (hash(x)=34478395) +1954 train 7.887595 (lr=3.8120e-04) (hash(x)=25357046) +1955 train 7.667367 (lr=3.8113e-04) (hash(x)=22372507) +1956 train 7.520681 (lr=3.8106e-04) (hash(x)=24501779) +1957 train 7.619022 (lr=3.8099e-04) (hash(x)=25638879) +1958 train 7.505960 (lr=3.8092e-04) (hash(x)=24243109) +1959 train 7.507718 (lr=3.8085e-04) (hash(x)=25164318) +1960 train 7.673126 (lr=3.8077e-04) (hash(x)=24139533) +1961 train 7.706660 (lr=3.8070e-04) (hash(x)=27271900) +1962 train 7.569108 (lr=3.8063e-04) (hash(x)=24869232) +1963 train 7.685955 (lr=3.8056e-04) (hash(x)=26562947) +1964 train 7.635368 (lr=3.8049e-04) (hash(x)=26477326) +1965 train 7.703720 (lr=3.8042e-04) (hash(x)=24917192) +1966 train 7.511055 (lr=3.8035e-04) (hash(x)=23110147) +1967 train 7.610240 (lr=3.8027e-04) (hash(x)=25690221) +1968 train 7.963957 (lr=3.8020e-04) (hash(x)=26999273) +1969 train 7.584912 (lr=3.8013e-04) (hash(x)=24807841) +1970 train 7.552022 (lr=3.8006e-04) (hash(x)=23918831) +1971 train 7.513790 (lr=3.7998e-04) (hash(x)=21947305) +1972 train 7.565952 (lr=3.7991e-04) (hash(x)=22874486) +1973 train 7.446546 (lr=3.7984e-04) (hash(x)=22195089) +1974 train 7.547877 (lr=3.7977e-04) (hash(x)=25345359) +1975 train 7.699240 (lr=3.7969e-04) (hash(x)=29031550) +1976 train 7.695127 (lr=3.7962e-04) (hash(x)=29947423) +1977 train 7.828130 (lr=3.7955e-04) (hash(x)=29395823) +1978 train 7.639629 (lr=3.7947e-04) (hash(x)=26968216) +1979 train 7.690835 (lr=3.7940e-04) (hash(x)=26551589) +1980 train 7.667661 (lr=3.7933e-04) (hash(x)=25657210) +1981 train 7.621756 (lr=3.7925e-04) (hash(x)=23689487) +1982 train 7.679008 (lr=3.7918e-04) (hash(x)=25744240) +1983 train 7.617736 (lr=3.7910e-04) (hash(x)=24094093) +1984 train 7.413030 (lr=3.7903e-04) (hash(x)=21724872) +1985 train 7.378041 (lr=3.7896e-04) (hash(x)=24106453) +1986 train 7.472117 (lr=3.7888e-04) (hash(x)=25004577) +1987 train 7.488628 (lr=3.7881e-04) (hash(x)=23639076) +1988 train 7.774906 (lr=3.7873e-04) (hash(x)=29408797) +1989 train 7.419815 (lr=3.7866e-04) (hash(x)=25107314) +1990 train 10.091942 (lr=3.7858e-04) (hash(x)=49896019) +1991 train 11.636477 (lr=3.7851e-04) (hash(x)=61455960) +1992 train 11.562792 (lr=3.7843e-04) (hash(x)=63788199) +1993 train 11.726527 (lr=3.7836e-04) (hash(x)=63899424) +1994 train 11.346703 (lr=3.7828e-04) (hash(x)=66412689) +1995 train 10.907353 (lr=3.7821e-04) (hash(x)=65636572) +1996 train 11.027329 (lr=3.7813e-04) (hash(x)=69649789) +1997 train 10.617279 (lr=3.7805e-04) (hash(x)=71578682) +1998 train 9.661141 (lr=3.7798e-04) (hash(x)=58341730) +1999 train 9.514322 (lr=3.7790e-04) (hash(x)=61104710) +2000 val loss 8.3127 +2000 val perplexity 4075.2366 +2000 train 9.065886 (lr=3.7783e-04) (hash(x)=52865322) +2001 train 8.465271 (lr=3.7775e-04) (hash(x)=23633693) +2002 train 8.509562 (lr=3.7767e-04) (hash(x)=25297571) +2003 train 8.709601 (lr=3.7760e-04) (hash(x)=24123821) +2004 train 8.283804 (lr=3.7752e-04) (hash(x)=20161517) +2005 train 8.289706 (lr=3.7744e-04) (hash(x)=19919984) +2006 train 8.296579 (lr=3.7737e-04) (hash(x)=18626507) +2007 train 8.551384 (lr=3.7729e-04) (hash(x)=20881126) +2008 train 8.574027 (lr=3.7721e-04) (hash(x)=26261448) +2009 train 8.582271 (lr=3.7714e-04) (hash(x)=27528050) +2010 train 8.555557 (lr=3.7706e-04) (hash(x)=26692097) +2011 train 8.271262 (lr=3.7698e-04) (hash(x)=23209969) +2012 train 8.218740 (lr=3.7690e-04) (hash(x)=26083446) +2013 train 8.255591 (lr=3.7683e-04) (hash(x)=28426467) +2014 train 8.074984 (lr=3.7675e-04) (hash(x)=26158154) +2015 train 8.050402 (lr=3.7667e-04) (hash(x)=26920243) +2016 train 7.707312 (lr=3.7659e-04) (hash(x)=21365321) +2017 train 7.761552 (lr=3.7651e-04) (hash(x)=25730699) +2018 train 7.774161 (lr=3.7644e-04) (hash(x)=25177106) +2019 train 7.689521 (lr=3.7636e-04) (hash(x)=26127756) +2020 train 7.274131 (lr=3.7628e-04) (hash(x)=19770037) +2021 train 7.746089 (lr=3.7620e-04) (hash(x)=26479346) +2022 train 7.782852 (lr=3.7612e-04) (hash(x)=25967643) +2023 train 7.732975 (lr=3.7604e-04) (hash(x)=21970343) +2024 train 7.924610 (lr=3.7596e-04) (hash(x)=26449648) +2025 train 7.786942 (lr=3.7588e-04) (hash(x)=24874094) +2026 train 7.493852 (lr=3.7581e-04) (hash(x)=21399697) +2027 train 7.646857 (lr=3.7573e-04) (hash(x)=24097897) +2028 train 7.726060 (lr=3.7565e-04) (hash(x)=24453548) +2029 train 7.852321 (lr=3.7557e-04) (hash(x)=26692093) +2030 train 7.657091 (lr=3.7549e-04) (hash(x)=27217002) +2031 train 7.987404 (lr=3.7541e-04) (hash(x)=31676352) +2032 train 7.630091 (lr=3.7533e-04) (hash(x)=23396350) +2033 train 7.668167 (lr=3.7525e-04) (hash(x)=26245173) +2034 train 7.703267 (lr=3.7517e-04) (hash(x)=24821033) +2035 train 7.745677 (lr=3.7509e-04) (hash(x)=23908053) +2036 train 7.536523 (lr=3.7501e-04) (hash(x)=23469373) +2037 train 7.583005 (lr=3.7493e-04) (hash(x)=24102198) +2038 train 7.616692 (lr=3.7484e-04) (hash(x)=24712135) +2039 train 7.822521 (lr=3.7476e-04) (hash(x)=25600562) +2040 train 7.736487 (lr=3.7468e-04) (hash(x)=23193393) +2041 train 7.701842 (lr=3.7460e-04) (hash(x)=24776954) +2042 train 7.817039 (lr=3.7452e-04) (hash(x)=28066546) +2043 train 7.475784 (lr=3.7444e-04) (hash(x)=23549168) +2044 train 7.481973 (lr=3.7436e-04) (hash(x)=22681487) +2045 train 7.683763 (lr=3.7428e-04) (hash(x)=26663184) +2046 train 7.663321 (lr=3.7419e-04) (hash(x)=23320809) +2047 train 7.526884 (lr=3.7411e-04) (hash(x)=22022924) +2048 train 7.824492 (lr=3.7403e-04) (hash(x)=27960930) +2049 train 7.671090 (lr=3.7395e-04) (hash(x)=24899178) +2050 val loss 7.6510 +2050 val perplexity 2102.6912 +2050 train 7.553998 (lr=3.7387e-04) (hash(x)=24450711) +2051 train 7.535224 (lr=3.7378e-04) (hash(x)=23294980) +2052 train 7.481105 (lr=3.7370e-04) (hash(x)=23243665) +2053 train 7.305561 (lr=3.7362e-04) (hash(x)=21004461) +2054 train 7.620389 (lr=3.7354e-04) (hash(x)=24493583) +2055 train 7.538692 (lr=3.7345e-04) (hash(x)=22777085) +2056 train 7.440296 (lr=3.7337e-04) (hash(x)=25169889) +2057 train 8.048717 (lr=3.7329e-04) (hash(x)=28395880) +2058 train 7.996845 (lr=3.7321e-04) (hash(x)=29603726) +2059 train 7.628465 (lr=3.7312e-04) (hash(x)=26271115) +2060 train 7.663245 (lr=3.7304e-04) (hash(x)=25616212) +2061 train 7.739485 (lr=3.7296e-04) (hash(x)=26398325) +2062 train 7.465399 (lr=3.7287e-04) (hash(x)=23836586) +2063 train 7.393511 (lr=3.7279e-04) (hash(x)=21727744) +2064 train 7.609024 (lr=3.7270e-04) (hash(x)=25183195) +2065 train 7.745055 (lr=3.7262e-04) (hash(x)=27108132) +2066 train 7.470177 (lr=3.7254e-04) (hash(x)=20987812) +2067 train 7.701198 (lr=3.7245e-04) (hash(x)=22729318) +2068 train 7.570267 (lr=3.7237e-04) (hash(x)=23129709) +2069 train 8.104627 (lr=3.7228e-04) (hash(x)=32820084) +2070 train 8.105147 (lr=3.7220e-04) (hash(x)=33763489) +2071 train 7.573160 (lr=3.7211e-04) (hash(x)=20701998) +2072 train 7.817600 (lr=3.7203e-04) (hash(x)=27765988) +2073 train 7.555827 (lr=3.7195e-04) (hash(x)=24157446) +2074 train 7.507003 (lr=3.7186e-04) (hash(x)=22014978) +2075 train 8.039325 (lr=3.7177e-04) (hash(x)=27928398) +2076 train 7.849775 (lr=3.7169e-04) (hash(x)=27880142) +2077 train 8.120887 (lr=3.7160e-04) (hash(x)=34991795) +2078 train 7.719812 (lr=3.7152e-04) (hash(x)=25944128) +2079 train 7.473180 (lr=3.7143e-04) (hash(x)=20863982) +2080 train 7.814935 (lr=3.7135e-04) (hash(x)=28036097) +2081 train 7.846234 (lr=3.7126e-04) (hash(x)=25813615) +2082 train 7.403400 (lr=3.7118e-04) (hash(x)=22862316) +2083 train 7.437952 (lr=3.7109e-04) (hash(x)=22827054) +2084 train 7.474388 (lr=3.7100e-04) (hash(x)=23996531) +2085 train 7.610325 (lr=3.7092e-04) (hash(x)=25756087) +2086 train 7.641029 (lr=3.7083e-04) (hash(x)=23165889) +2087 train 7.797710 (lr=3.7074e-04) (hash(x)=25084926) +2088 train 7.655910 (lr=3.7066e-04) (hash(x)=26063862) +2089 train 7.701123 (lr=3.7057e-04) (hash(x)=27843638) +2090 train 7.952284 (lr=3.7048e-04) (hash(x)=27375554) +2091 train 7.675088 (lr=3.7040e-04) (hash(x)=25540961) +2092 train 7.454419 (lr=3.7031e-04) (hash(x)=22933785) +2093 train 7.624463 (lr=3.7022e-04) (hash(x)=23996956) +2094 train 7.381715 (lr=3.7014e-04) (hash(x)=22915854) +2095 train 7.622618 (lr=3.7005e-04) (hash(x)=24412897) +2096 train 7.593031 (lr=3.6996e-04) (hash(x)=25152336) +2097 train 7.911393 (lr=3.6987e-04) (hash(x)=29525589) +2098 train 7.732984 (lr=3.6979e-04) (hash(x)=26435656) +2099 train 7.649677 (lr=3.6970e-04) (hash(x)=26269869) +2100 val loss 7.5990 +2100 val perplexity 1996.2891 +2100 train 7.401772 (lr=3.6961e-04) (hash(x)=23856783) +2101 train 7.701483 (lr=3.6952e-04) (hash(x)=26039611) +2102 train 7.438747 (lr=3.6943e-04) (hash(x)=25048861) +2103 train 7.530903 (lr=3.6934e-04) (hash(x)=25434913) +2104 train 7.713660 (lr=3.6926e-04) (hash(x)=24203891) +2105 train 7.606290 (lr=3.6917e-04) (hash(x)=26016998) +2106 train 7.715206 (lr=3.6908e-04) (hash(x)=23986767) +2107 train 7.608569 (lr=3.6899e-04) (hash(x)=25157653) +2108 train 7.769751 (lr=3.6890e-04) (hash(x)=30193751) +2109 train 7.585874 (lr=3.6881e-04) (hash(x)=22630574) +2110 train 7.537766 (lr=3.6872e-04) (hash(x)=23475467) +2111 train 7.866647 (lr=3.6863e-04) (hash(x)=26302544) +2112 train 7.636853 (lr=3.6854e-04) (hash(x)=21276592) +2113 train 7.900350 (lr=3.6845e-04) (hash(x)=27566393) +2114 train 7.641581 (lr=3.6837e-04) (hash(x)=25028015) +2115 train 7.464832 (lr=3.6828e-04) (hash(x)=22010703) +2116 train 7.782817 (lr=3.6819e-04) (hash(x)=23711709) +2117 train 7.840171 (lr=3.6810e-04) (hash(x)=26105026) +2118 train 8.118068 (lr=3.6801e-04) (hash(x)=34433894) +2119 train 8.150718 (lr=3.6792e-04) (hash(x)=34976197) +2120 train 7.703751 (lr=3.6782e-04) (hash(x)=27690727) +2121 train 7.542693 (lr=3.6773e-04) (hash(x)=24368234) +2122 train 7.549530 (lr=3.6764e-04) (hash(x)=25095726) +2123 train 7.698032 (lr=3.6755e-04) (hash(x)=28248301) +2124 train 7.487286 (lr=3.6746e-04) (hash(x)=23322302) +2125 train 7.863014 (lr=3.6737e-04) (hash(x)=26416200) +2126 train 7.569534 (lr=3.6728e-04) (hash(x)=26577567) +2127 train 7.611949 (lr=3.6719e-04) (hash(x)=23870805) +2128 train 7.569670 (lr=3.6710e-04) (hash(x)=25440544) +2129 train 7.670265 (lr=3.6701e-04) (hash(x)=25795021) +2130 train 7.554659 (lr=3.6692e-04) (hash(x)=20896402) +2131 train 7.604482 (lr=3.6682e-04) (hash(x)=24457252) +2132 train 7.621601 (lr=3.6673e-04) (hash(x)=25926760) +2133 train 7.485384 (lr=3.6664e-04) (hash(x)=23503725) +2134 train 8.000311 (lr=3.6655e-04) (hash(x)=28728828) +2135 train 7.566890 (lr=3.6646e-04) (hash(x)=25041103) +2136 train 7.583607 (lr=3.6636e-04) (hash(x)=22863770) +2137 train 7.449742 (lr=3.6627e-04) (hash(x)=23037755) +2138 train 7.694551 (lr=3.6618e-04) (hash(x)=25848413) +2139 train 7.743011 (lr=3.6609e-04) (hash(x)=25998487) +2140 train 7.551328 (lr=3.6599e-04) (hash(x)=22754440) +2141 train 7.772513 (lr=3.6590e-04) (hash(x)=27705382) +2142 train 7.764174 (lr=3.6581e-04) (hash(x)=27629095) +2143 train 7.718404 (lr=3.6572e-04) (hash(x)=26041745) +2144 train 7.453105 (lr=3.6562e-04) (hash(x)=21909712) +2145 train 7.636692 (lr=3.6553e-04) (hash(x)=24353905) +2146 train 7.607166 (lr=3.6544e-04) (hash(x)=24482587) +2147 train 7.776822 (lr=3.6534e-04) (hash(x)=27249810) +2148 train 7.873353 (lr=3.6525e-04) (hash(x)=26709938) +2149 train 7.563306 (lr=3.6516e-04) (hash(x)=23831457) +2150 val loss 7.6670 +2150 val perplexity 2136.6951 +2150 train 7.870407 (lr=3.6506e-04) (hash(x)=29776243) +2151 train 7.430917 (lr=3.6497e-04) (hash(x)=24068619) +2152 train 7.625419 (lr=3.6487e-04) (hash(x)=22208671) +2153 train 7.761914 (lr=3.6478e-04) (hash(x)=26680905) +2154 train 7.582022 (lr=3.6469e-04) (hash(x)=23567808) +2155 train 7.799015 (lr=3.6459e-04) (hash(x)=26359528) +2156 train 7.592532 (lr=3.6450e-04) (hash(x)=23787652) +2157 train 7.821782 (lr=3.6440e-04) (hash(x)=28347177) +2158 train 7.539731 (lr=3.6431e-04) (hash(x)=25266519) +2159 train 7.454289 (lr=3.6421e-04) (hash(x)=25441262) +2160 train 7.439699 (lr=3.6412e-04) (hash(x)=23959943) +2161 train 7.659112 (lr=3.6402e-04) (hash(x)=27888093) +2162 train 7.547934 (lr=3.6393e-04) (hash(x)=25547833) +2163 train 7.527712 (lr=3.6383e-04) (hash(x)=24413659) +2164 train 9.308187 (lr=3.6374e-04) (hash(x)=34397626) +2165 train 7.555001 (lr=3.6364e-04) (hash(x)=21159323) +2166 train 7.713135 (lr=3.6355e-04) (hash(x)=25360269) +2167 train 7.782236 (lr=3.6345e-04) (hash(x)=29181867) +2168 train 7.541968 (lr=3.6336e-04) (hash(x)=22501613) +2169 train 7.634534 (lr=3.6326e-04) (hash(x)=26469077) +2170 train 7.622079 (lr=3.6316e-04) (hash(x)=26665118) +2171 train 8.201706 (lr=3.6307e-04) (hash(x)=31551575) +2172 train 7.631139 (lr=3.6297e-04) (hash(x)=24803935) +2173 train 7.552742 (lr=3.6288e-04) (hash(x)=20896465) +2174 train 7.950912 (lr=3.6278e-04) (hash(x)=25702284) +2175 train 7.918772 (lr=3.6268e-04) (hash(x)=27093757) +2176 train 7.685677 (lr=3.6259e-04) (hash(x)=26024255) +2177 train 7.723851 (lr=3.6249e-04) (hash(x)=26443521) +2178 train 7.724322 (lr=3.6239e-04) (hash(x)=25478467) +2179 train 7.658857 (lr=3.6230e-04) (hash(x)=23624298) +2180 train 7.765506 (lr=3.6220e-04) (hash(x)=27304178) +2181 train 7.515543 (lr=3.6210e-04) (hash(x)=23016570) +2182 train 7.421331 (lr=3.6200e-04) (hash(x)=21073756) +2183 train 7.547981 (lr=3.6191e-04) (hash(x)=25150275) +2184 train 7.347604 (lr=3.6181e-04) (hash(x)=21004187) +2185 train 7.567065 (lr=3.6171e-04) (hash(x)=25876062) +2186 train 7.548803 (lr=3.6161e-04) (hash(x)=24221275) +2187 train 7.794189 (lr=3.6152e-04) (hash(x)=26239886) +2188 train 7.464691 (lr=3.6142e-04) (hash(x)=23299261) +2189 train 7.854206 (lr=3.6132e-04) (hash(x)=31024703) +2190 train 7.930138 (lr=3.6122e-04) (hash(x)=30084774) +2191 train 8.230929 (lr=3.6112e-04) (hash(x)=33457102) +2192 train 7.549683 (lr=3.6103e-04) (hash(x)=22423723) +2193 train 7.611895 (lr=3.6093e-04) (hash(x)=21989055) +2194 train 7.585538 (lr=3.6083e-04) (hash(x)=23286499) +2195 train 7.550269 (lr=3.6073e-04) (hash(x)=17499738) +2196 train 7.642938 (lr=3.6063e-04) (hash(x)=25329557) +2197 train 7.587094 (lr=3.6053e-04) (hash(x)=24195578) +2198 train 7.727867 (lr=3.6043e-04) (hash(x)=26646383) +2199 train 7.592222 (lr=3.6033e-04) (hash(x)=25755322) +2200 val loss 7.6452 +2200 val perplexity 2090.5474 +2200 train 7.673711 (lr=3.6023e-04) (hash(x)=27592223) +2201 train 7.809978 (lr=3.6014e-04) (hash(x)=26164625) +2202 train 7.639874 (lr=3.6004e-04) (hash(x)=21856341) +2203 train 7.566431 (lr=3.5994e-04) (hash(x)=23722795) +2204 train 7.626783 (lr=3.5984e-04) (hash(x)=25409645) +2205 train 7.567655 (lr=3.5974e-04) (hash(x)=23415339) +2206 train 7.536373 (lr=3.5964e-04) (hash(x)=24147928) +2207 train 7.458281 (lr=3.5954e-04) (hash(x)=23630794) +2208 train 7.554024 (lr=3.5944e-04) (hash(x)=27361452) +2209 train 7.555365 (lr=3.5934e-04) (hash(x)=25632158) +2210 train 7.432745 (lr=3.5924e-04) (hash(x)=21159789) +2211 train 7.512891 (lr=3.5914e-04) (hash(x)=23374168) +2212 train 7.610300 (lr=3.5904e-04) (hash(x)=24844739) +2213 train 7.510890 (lr=3.5893e-04) (hash(x)=23461285) +2214 train 7.560419 (lr=3.5883e-04) (hash(x)=27958481) +2215 train 7.489844 (lr=3.5873e-04) (hash(x)=25167987) +2216 train 7.640916 (lr=3.5863e-04) (hash(x)=22873204) +2217 train 8.181331 (lr=3.5853e-04) (hash(x)=31038116) +2218 train 7.672850 (lr=3.5843e-04) (hash(x)=28883155) +2219 train 7.645888 (lr=3.5833e-04) (hash(x)=30941010) +2220 train 7.576355 (lr=3.5823e-04) (hash(x)=24947521) +2221 train 7.611646 (lr=3.5813e-04) (hash(x)=22703293) +2222 train 7.563556 (lr=3.5802e-04) (hash(x)=23976007) +2223 train 7.487794 (lr=3.5792e-04) (hash(x)=24043224) +2224 train 7.460487 (lr=3.5782e-04) (hash(x)=23981093) +2225 train 7.613994 (lr=3.5772e-04) (hash(x)=26445994) +2226 train 7.882288 (lr=3.5762e-04) (hash(x)=27252393) +2227 train 7.860669 (lr=3.5751e-04) (hash(x)=29736627) +2228 train 7.815498 (lr=3.5741e-04) (hash(x)=24698933) +2229 train 7.458528 (lr=3.5731e-04) (hash(x)=24357840) +2230 train 7.634920 (lr=3.5721e-04) (hash(x)=25384498) +2231 train 7.703066 (lr=3.5710e-04) (hash(x)=24665411) +2232 train 7.615764 (lr=3.5700e-04) (hash(x)=24646352) +2233 train 7.625246 (lr=3.5690e-04) (hash(x)=25670741) +2234 train 7.662351 (lr=3.5680e-04) (hash(x)=26648242) +2235 train 7.796783 (lr=3.5669e-04) (hash(x)=27916043) +2236 train 7.650910 (lr=3.5659e-04) (hash(x)=25963100) +2237 train 7.718106 (lr=3.5649e-04) (hash(x)=27357379) +2238 train 7.561273 (lr=3.5638e-04) (hash(x)=21852360) +2239 train 7.501964 (lr=3.5628e-04) (hash(x)=17321657) +2240 train 7.435650 (lr=3.5618e-04) (hash(x)=18446645) +2241 train 7.288990 (lr=3.5607e-04) (hash(x)=17550988) +2242 train 7.557130 (lr=3.5597e-04) (hash(x)=27052347) +2243 train 7.487940 (lr=3.5587e-04) (hash(x)=23829833) +2244 train 7.649688 (lr=3.5576e-04) (hash(x)=27267264) +2245 train 7.886673 (lr=3.5566e-04) (hash(x)=24673611) +2246 train 7.711062 (lr=3.5555e-04) (hash(x)=25248507) +2247 train 7.939725 (lr=3.5545e-04) (hash(x)=27932591) +2248 train 7.787208 (lr=3.5534e-04) (hash(x)=26078877) +2249 train 7.840024 (lr=3.5524e-04) (hash(x)=26067367) +2250 val loss 7.6336 +2250 val perplexity 2066.5759 +2250 train 7.992462 (lr=3.5514e-04) (hash(x)=29143527) +2251 train 7.990372 (lr=3.5503e-04) (hash(x)=28206079) +2252 train 8.520241 (lr=3.5493e-04) (hash(x)=28894494) +2253 train 7.950740 (lr=3.5482e-04) (hash(x)=27977422) +2254 train 7.658128 (lr=3.5472e-04) (hash(x)=26446420) +2255 train 7.866582 (lr=3.5461e-04) (hash(x)=29447954) +2256 train 7.735686 (lr=3.5451e-04) (hash(x)=22181789) +2257 train 7.727842 (lr=3.5440e-04) (hash(x)=24976337) +2258 train 7.680702 (lr=3.5429e-04) (hash(x)=23277532) +2259 train 7.739036 (lr=3.5419e-04) (hash(x)=25935364) +2260 train 7.858867 (lr=3.5408e-04) (hash(x)=26988889) +2261 train 7.663904 (lr=3.5398e-04) (hash(x)=26887303) +2262 train 7.598855 (lr=3.5387e-04) (hash(x)=25021426) +2263 train 7.534516 (lr=3.5377e-04) (hash(x)=24621816) +2264 train 7.456763 (lr=3.5366e-04) (hash(x)=22132007) +2265 train 7.448111 (lr=3.5355e-04) (hash(x)=22648602) +2266 train 7.611203 (lr=3.5345e-04) (hash(x)=23132242) +2267 train 7.692613 (lr=3.5334e-04) (hash(x)=24070405) +2268 train 7.374382 (lr=3.5323e-04) (hash(x)=21412906) +2269 train 7.977808 (lr=3.5313e-04) (hash(x)=27547292) +2270 train 7.799314 (lr=3.5302e-04) (hash(x)=27740500) +2271 train 7.565277 (lr=3.5291e-04) (hash(x)=24682294) +2272 train 7.903067 (lr=3.5281e-04) (hash(x)=27969424) +2273 train 7.619516 (lr=3.5270e-04) (hash(x)=22172182) +2274 train 7.624746 (lr=3.5259e-04) (hash(x)=26485905) +2275 train 7.510963 (lr=3.5249e-04) (hash(x)=24050209) +2276 train 7.637181 (lr=3.5238e-04) (hash(x)=28802650) +2277 train 7.400589 (lr=3.5227e-04) (hash(x)=23644616) +2278 train 7.614135 (lr=3.5216e-04) (hash(x)=25409768) +2279 train 7.733966 (lr=3.5206e-04) (hash(x)=28057095) +2280 train 8.170192 (lr=3.5195e-04) (hash(x)=31689723) +2281 train 7.423082 (lr=3.5184e-04) (hash(x)=22252427) +2282 train 7.610404 (lr=3.5173e-04) (hash(x)=24142092) +2283 train 7.543883 (lr=3.5163e-04) (hash(x)=22523232) +2284 train 7.758942 (lr=3.5152e-04) (hash(x)=26017294) +2285 train 7.637816 (lr=3.5141e-04) (hash(x)=25025388) +2286 train 7.411517 (lr=3.5130e-04) (hash(x)=23819479) +2287 train 7.582829 (lr=3.5119e-04) (hash(x)=24547536) +2288 train 7.409281 (lr=3.5108e-04) (hash(x)=22622789) +2289 train 7.387859 (lr=3.5098e-04) (hash(x)=23874051) +2290 train 7.644675 (lr=3.5087e-04) (hash(x)=25015641) +2291 train 7.576322 (lr=3.5076e-04) (hash(x)=24978712) +2292 train 7.642452 (lr=3.5065e-04) (hash(x)=25311986) +2293 train 7.446362 (lr=3.5054e-04) (hash(x)=24298295) +2294 train 7.545784 (lr=3.5043e-04) (hash(x)=26481527) +2295 train 7.283615 (lr=3.5032e-04) (hash(x)=18987545) +2296 train 7.513702 (lr=3.5021e-04) (hash(x)=24617990) +2297 train 7.499829 (lr=3.5010e-04) (hash(x)=23903200) +2298 train 7.774329 (lr=3.4999e-04) (hash(x)=26278697) +2299 train 7.531718 (lr=3.4988e-04) (hash(x)=24092784) +2300 val loss 7.5818 +2300 val perplexity 1962.1146 +2300 train 7.493537 (lr=3.4977e-04) (hash(x)=22894919) +2301 train 7.597230 (lr=3.4966e-04) (hash(x)=24253964) +2302 train 7.527168 (lr=3.4955e-04) (hash(x)=23750610) +2303 train 7.578588 (lr=3.4944e-04) (hash(x)=26745063) +2304 train 7.385965 (lr=3.4933e-04) (hash(x)=21407001) +2305 train 7.725150 (lr=3.4922e-04) (hash(x)=26680082) +2306 train 7.651417 (lr=3.4911e-04) (hash(x)=26722122) +2307 train 7.478429 (lr=3.4900e-04) (hash(x)=22539112) +2308 train 7.289353 (lr=3.4889e-04) (hash(x)=19927356) +2309 train 7.417696 (lr=3.4878e-04) (hash(x)=21431267) +2310 train 7.370184 (lr=3.4867e-04) (hash(x)=21073487) +2311 train 7.409659 (lr=3.4856e-04) (hash(x)=21447406) +2312 train 7.409338 (lr=3.4845e-04) (hash(x)=22667253) +2313 train 7.377115 (lr=3.4834e-04) (hash(x)=20849089) +2314 train 7.365770 (lr=3.4823e-04) (hash(x)=19168351) +2315 train 7.352315 (lr=3.4812e-04) (hash(x)=21082139) +2316 train 7.336221 (lr=3.4801e-04) (hash(x)=21493530) +2317 train 7.272052 (lr=3.4789e-04) (hash(x)=19506447) +2318 train 7.396034 (lr=3.4778e-04) (hash(x)=23128797) +2319 train 7.536978 (lr=3.4767e-04) (hash(x)=24650655) +2320 train 7.326128 (lr=3.4756e-04) (hash(x)=21833080) +2321 train 7.869785 (lr=3.4745e-04) (hash(x)=26355512) +2322 train 7.979322 (lr=3.4734e-04) (hash(x)=26746363) +2323 train 7.437234 (lr=3.4722e-04) (hash(x)=22130721) +2324 train 7.802797 (lr=3.4711e-04) (hash(x)=28301752) +2325 train 7.593059 (lr=3.4700e-04) (hash(x)=26444359) +2326 train 7.987758 (lr=3.4689e-04) (hash(x)=24836130) +2327 train 8.541141 (lr=3.4677e-04) (hash(x)=26519976) +2328 train 10.196017 (lr=3.4666e-04) (hash(x)=35558906) +2329 train 8.012717 (lr=3.4655e-04) (hash(x)=30877729) +2330 train 7.835441 (lr=3.4644e-04) (hash(x)=28028485) +2331 train 7.914791 (lr=3.4632e-04) (hash(x)=25582440) +2332 train 8.325480 (lr=3.4621e-04) (hash(x)=29402491) +2333 train 7.710075 (lr=3.4610e-04) (hash(x)=25200284) +2334 train 7.782512 (lr=3.4598e-04) (hash(x)=26288480) +2335 train 7.638883 (lr=3.4587e-04) (hash(x)=20106514) +2336 train 7.676218 (lr=3.4576e-04) (hash(x)=22522123) +2337 train 7.666375 (lr=3.4564e-04) (hash(x)=23362041) +2338 train 7.777652 (lr=3.4553e-04) (hash(x)=26003099) +2339 train 7.808650 (lr=3.4542e-04) (hash(x)=25695510) +2340 train 7.775612 (lr=3.4530e-04) (hash(x)=27225904) +2341 train 7.547893 (lr=3.4519e-04) (hash(x)=21573008) +2342 train 7.886928 (lr=3.4508e-04) (hash(x)=21386628) +2343 train 7.653917 (lr=3.4496e-04) (hash(x)=20117808) +2344 train 7.727954 (lr=3.4485e-04) (hash(x)=21592790) +2345 train 7.621472 (lr=3.4473e-04) (hash(x)=19909192) +2346 train 7.688829 (lr=3.4462e-04) (hash(x)=22529262) +2347 train 7.740810 (lr=3.4451e-04) (hash(x)=24501900) +2348 train 7.556070 (lr=3.4439e-04) (hash(x)=25912171) +2349 train 7.658494 (lr=3.4428e-04) (hash(x)=25606665) +2350 val loss 7.6116 +2350 val perplexity 2021.5137 +2350 train 7.599922 (lr=3.4416e-04) (hash(x)=24487351) +2351 train 7.779486 (lr=3.4405e-04) (hash(x)=25510334) +2352 train 7.688320 (lr=3.4393e-04) (hash(x)=25357989) +2353 train 7.336888 (lr=3.4382e-04) (hash(x)=24656801) +2354 train 7.439148 (lr=3.4370e-04) (hash(x)=23312772) +2355 train 7.382675 (lr=3.4359e-04) (hash(x)=22099158) +2356 train 7.393102 (lr=3.4347e-04) (hash(x)=26507898) +2357 train 7.655396 (lr=3.4336e-04) (hash(x)=28351614) +2358 train 7.706299 (lr=3.4324e-04) (hash(x)=27489567) +2359 train 7.662216 (lr=3.4313e-04) (hash(x)=25749120) +2360 train 8.009239 (lr=3.4301e-04) (hash(x)=31711338) +2361 train 8.172415 (lr=3.4289e-04) (hash(x)=29645018) +2362 train 8.138659 (lr=3.4278e-04) (hash(x)=29713268) +2363 train 7.322204 (lr=3.4266e-04) (hash(x)=21720691) +2364 train 7.422951 (lr=3.4255e-04) (hash(x)=24316633) +2365 train 7.393700 (lr=3.4243e-04) (hash(x)=21597124) +2366 train 7.386212 (lr=3.4231e-04) (hash(x)=22520345) +2367 train 7.462013 (lr=3.4220e-04) (hash(x)=24357241) +2368 train 7.603851 (lr=3.4208e-04) (hash(x)=24085450) +2369 train 7.511463 (lr=3.4197e-04) (hash(x)=25136495) +2370 train 7.944884 (lr=3.4185e-04) (hash(x)=26073986) +2371 train 7.968779 (lr=3.4173e-04) (hash(x)=28911272) +2372 train 7.656772 (lr=3.4162e-04) (hash(x)=26667356) +2373 train 7.514291 (lr=3.4150e-04) (hash(x)=25225894) +2374 train 7.551837 (lr=3.4138e-04) (hash(x)=24404081) +2375 train 7.613176 (lr=3.4127e-04) (hash(x)=25584945) +2376 train 7.485127 (lr=3.4115e-04) (hash(x)=23831571) +2377 train 7.440205 (lr=3.4103e-04) (hash(x)=23521916) +2378 train 7.502786 (lr=3.4091e-04) (hash(x)=25318634) +2379 train 7.392082 (lr=3.4080e-04) (hash(x)=21847287) +2380 train 7.704719 (lr=3.4068e-04) (hash(x)=23877060) +2381 train 7.516122 (lr=3.4056e-04) (hash(x)=24069020) +2382 train 7.377492 (lr=3.4044e-04) (hash(x)=21724290) +2383 train 7.560748 (lr=3.4033e-04) (hash(x)=25198897) +2384 train 7.423851 (lr=3.4021e-04) (hash(x)=24109958) +2385 train 7.363171 (lr=3.4009e-04) (hash(x)=20122390) +2386 train 7.570796 (lr=3.3997e-04) (hash(x)=24062305) +2387 train 7.389867 (lr=3.3985e-04) (hash(x)=22436833) +2388 train 7.643991 (lr=3.3974e-04) (hash(x)=26013214) +2389 train 7.907181 (lr=3.3962e-04) (hash(x)=27691436) +2390 train 8.115775 (lr=3.3950e-04) (hash(x)=30802878) +2391 train 7.330136 (lr=3.3938e-04) (hash(x)=23475891) +2392 train 7.291879 (lr=3.3926e-04) (hash(x)=23118133) +2393 train 7.251279 (lr=3.3914e-04) (hash(x)=21469159) +2394 train 7.591257 (lr=3.3902e-04) (hash(x)=26444484) +2395 train 7.497031 (lr=3.3891e-04) (hash(x)=25083992) +2396 train 7.846909 (lr=3.3879e-04) (hash(x)=23461229) +2397 train 7.629861 (lr=3.3867e-04) (hash(x)=20441653) +2398 train 7.744658 (lr=3.3855e-04) (hash(x)=28024211) +2399 train 7.570257 (lr=3.3843e-04) (hash(x)=23644804) +2400 val loss 7.6206 +2400 val perplexity 2039.7695 +2400 train 7.429077 (lr=3.3831e-04) (hash(x)=26685301) +2401 train 7.383492 (lr=3.3819e-04) (hash(x)=20820913) +2402 train 7.532267 (lr=3.3807e-04) (hash(x)=22178190) +2403 train 7.941135 (lr=3.3795e-04) (hash(x)=31377168) +2404 train 8.144010 (lr=3.3783e-04) (hash(x)=33795307) +2405 train 8.388542 (lr=3.3771e-04) (hash(x)=34450341) +2406 train 7.650968 (lr=3.3759e-04) (hash(x)=25834183) +2407 train 7.793860 (lr=3.3747e-04) (hash(x)=25629584) +2408 train 7.565750 (lr=3.3735e-04) (hash(x)=25080123) +2409 train 7.479819 (lr=3.3723e-04) (hash(x)=21975628) +2410 train 7.967249 (lr=3.3711e-04) (hash(x)=27430197) +2411 train 8.133690 (lr=3.3699e-04) (hash(x)=29285135) +2412 train 7.850482 (lr=3.3687e-04) (hash(x)=24640105) +2413 train 7.941463 (lr=3.3675e-04) (hash(x)=29628864) +2414 train 8.080014 (lr=3.3663e-04) (hash(x)=30096444) +2415 train 7.989980 (lr=3.3651e-04) (hash(x)=29797280) +2416 train 7.397613 (lr=3.3639e-04) (hash(x)=22686143) +2417 train 7.925113 (lr=3.3627e-04) (hash(x)=28346842) +2418 train 7.811282 (lr=3.3615e-04) (hash(x)=28167937) +2419 train 8.053995 (lr=3.3603e-04) (hash(x)=33122326) +2420 train 8.074878 (lr=3.3590e-04) (hash(x)=33044913) +2421 train 7.914201 (lr=3.3578e-04) (hash(x)=31250981) +2422 train 8.154200 (lr=3.3566e-04) (hash(x)=29827111) +2423 train 8.087486 (lr=3.3554e-04) (hash(x)=30026394) +2424 train 8.239148 (lr=3.3542e-04) (hash(x)=26998238) +2425 train 8.270396 (lr=3.3530e-04) (hash(x)=24850684) +2426 train 8.384491 (lr=3.3518e-04) (hash(x)=36659853) +2427 train 8.334443 (lr=3.3505e-04) (hash(x)=34353164) +2428 train 7.706431 (lr=3.3493e-04) (hash(x)=23630644) +2429 train 7.557250 (lr=3.3481e-04) (hash(x)=24528186) +2430 train 7.484659 (lr=3.3469e-04) (hash(x)=22665222) +2431 train 7.596965 (lr=3.3457e-04) (hash(x)=26594177) +2432 train 7.929378 (lr=3.3444e-04) (hash(x)=29300546) +2433 train 7.381728 (lr=3.3432e-04) (hash(x)=21331715) +2434 train 7.708421 (lr=3.3420e-04) (hash(x)=26290885) +2435 train 7.584941 (lr=3.3408e-04) (hash(x)=25554738) +2436 train 7.786263 (lr=3.3395e-04) (hash(x)=26744311) +2437 train 7.727993 (lr=3.3383e-04) (hash(x)=26872344) +2438 train 7.761537 (lr=3.3371e-04) (hash(x)=27636081) +2439 train 7.541043 (lr=3.3359e-04) (hash(x)=23766256) +2440 train 7.705342 (lr=3.3346e-04) (hash(x)=24434438) +2441 train 7.526451 (lr=3.3334e-04) (hash(x)=22992618) +2442 train 7.401563 (lr=3.3322e-04) (hash(x)=22393467) +2443 train 7.649085 (lr=3.3309e-04) (hash(x)=26853444) +2444 train 7.605676 (lr=3.3297e-04) (hash(x)=25101020) +2445 train 7.502000 (lr=3.3285e-04) (hash(x)=24705456) +2446 train 7.765003 (lr=3.3272e-04) (hash(x)=26798611) +2447 train 7.468738 (lr=3.3260e-04) (hash(x)=24250150) +2448 train 7.511821 (lr=3.3248e-04) (hash(x)=24494331) +2449 train 7.542469 (lr=3.3235e-04) (hash(x)=27579257) +2450 val loss 7.5708 +2450 val perplexity 1940.7891 +2450 train 7.323148 (lr=3.3223e-04) (hash(x)=22377407) +2451 train 7.610791 (lr=3.3210e-04) (hash(x)=26289588) +2452 train 7.538097 (lr=3.3198e-04) (hash(x)=25871900) +2453 train 7.434997 (lr=3.3186e-04) (hash(x)=23437197) +2454 train 7.665241 (lr=3.3173e-04) (hash(x)=23557786) +2455 train 7.609625 (lr=3.3161e-04) (hash(x)=26108817) +2456 train 7.593883 (lr=3.3148e-04) (hash(x)=26440482) +2457 train 7.536459 (lr=3.3136e-04) (hash(x)=24583191) +2458 train 7.507783 (lr=3.3123e-04) (hash(x)=23756440) +2459 train 7.467746 (lr=3.3111e-04) (hash(x)=23814987) +2460 train 7.834494 (lr=3.3099e-04) (hash(x)=28508433) +2461 train 7.519865 (lr=3.3086e-04) (hash(x)=25692442) +2462 train 7.544918 (lr=3.3074e-04) (hash(x)=24699197) +2463 train 7.300777 (lr=3.3061e-04) (hash(x)=18582688) +2464 train 7.467488 (lr=3.3049e-04) (hash(x)=23266625) +2465 train 7.750530 (lr=3.3036e-04) (hash(x)=26565489) +2466 train 7.385438 (lr=3.3024e-04) (hash(x)=21686599) +2467 train 7.644693 (lr=3.3011e-04) (hash(x)=25558792) +2468 train 7.435891 (lr=3.2999e-04) (hash(x)=25110035) +2469 train 7.702917 (lr=3.2986e-04) (hash(x)=26263661) +2470 train 7.524774 (lr=3.2973e-04) (hash(x)=22282189) +2471 train 7.480340 (lr=3.2961e-04) (hash(x)=23272705) +2472 train 7.549061 (lr=3.2948e-04) (hash(x)=25300067) +2473 train 7.743009 (lr=3.2936e-04) (hash(x)=26802369) +2474 train 7.524380 (lr=3.2923e-04) (hash(x)=24646471) +2475 train 7.452026 (lr=3.2911e-04) (hash(x)=24709241) +2476 train 7.528258 (lr=3.2898e-04) (hash(x)=26747197) +2477 train 7.550816 (lr=3.2885e-04) (hash(x)=25201108) +2478 train 7.692124 (lr=3.2873e-04) (hash(x)=24962427) +2479 train 7.588126 (lr=3.2860e-04) (hash(x)=24793412) +2480 train 7.520001 (lr=3.2847e-04) (hash(x)=24452301) +2481 train 7.539691 (lr=3.2835e-04) (hash(x)=25177251) +2482 train 7.535166 (lr=3.2822e-04) (hash(x)=25801499) +2483 train 7.615195 (lr=3.2809e-04) (hash(x)=27256707) +2484 train 7.823678 (lr=3.2797e-04) (hash(x)=25303237) +2485 train 7.455535 (lr=3.2784e-04) (hash(x)=21641481) +2486 train 7.479755 (lr=3.2771e-04) (hash(x)=23818831) +2487 train 7.292194 (lr=3.2759e-04) (hash(x)=19280989) +2488 train 7.581689 (lr=3.2746e-04) (hash(x)=24075167) +2489 train 7.630374 (lr=3.2733e-04) (hash(x)=26651546) +2490 train 7.514842 (lr=3.2721e-04) (hash(x)=25367186) +2491 train 7.501633 (lr=3.2708e-04) (hash(x)=26508642) +2492 train 7.581214 (lr=3.2695e-04) (hash(x)=25294182) +2493 train 7.531882 (lr=3.2682e-04) (hash(x)=23916886) +2494 train 7.403149 (lr=3.2670e-04) (hash(x)=21189910) +2495 train 7.516267 (lr=3.2657e-04) (hash(x)=22751150) +2496 train 7.832397 (lr=3.2644e-04) (hash(x)=26608502) +2497 train 8.194119 (lr=3.2631e-04) (hash(x)=25372010) +2498 train 7.554254 (lr=3.2619e-04) (hash(x)=22006251) +2499 train 7.469146 (lr=3.2606e-04) (hash(x)=23880160) +2500 val loss 7.6110 +2500 val perplexity 2020.2051 +2500 train 7.484663 (lr=3.2593e-04) (hash(x)=23225337) +2501 train 7.605058 (lr=3.2580e-04) (hash(x)=24932950) +2502 train 7.349640 (lr=3.2567e-04) (hash(x)=20022340) +2503 train 7.634033 (lr=3.2554e-04) (hash(x)=23537942) +2504 train 7.532734 (lr=3.2542e-04) (hash(x)=24884288) +2505 train 7.294827 (lr=3.2529e-04) (hash(x)=21296580) +2506 train 7.480310 (lr=3.2516e-04) (hash(x)=23001455) +2507 train 7.752992 (lr=3.2503e-04) (hash(x)=26975313) +2508 train 7.731464 (lr=3.2490e-04) (hash(x)=26029962) +2509 train 7.499925 (lr=3.2477e-04) (hash(x)=24302204) +2510 train 7.202467 (lr=3.2464e-04) (hash(x)=22997203) +2511 train 7.346910 (lr=3.2452e-04) (hash(x)=23748375) +2512 train 7.530019 (lr=3.2439e-04) (hash(x)=24453191) +2513 train 7.274936 (lr=3.2426e-04) (hash(x)=22026776) +2514 train 7.403477 (lr=3.2413e-04) (hash(x)=19934168) +2515 train 7.593095 (lr=3.2400e-04) (hash(x)=25541754) +2516 train 7.886437 (lr=3.2387e-04) (hash(x)=26604471) +2517 train 7.906101 (lr=3.2374e-04) (hash(x)=27900386) +2518 train 7.808359 (lr=3.2361e-04) (hash(x)=26403431) +2519 train 7.841945 (lr=3.2348e-04) (hash(x)=27413825) +2520 train 7.911373 (lr=3.2335e-04) (hash(x)=28332637) +2521 train 7.725635 (lr=3.2322e-04) (hash(x)=27223027) +2522 train 7.720334 (lr=3.2309e-04) (hash(x)=23843387) +2523 train 7.743898 (lr=3.2296e-04) (hash(x)=27075951) +2524 train 7.899715 (lr=3.2283e-04) (hash(x)=29300154) +2525 train 7.570527 (lr=3.2270e-04) (hash(x)=28100582) +2526 train 7.663386 (lr=3.2257e-04) (hash(x)=28051084) +2527 train 7.445567 (lr=3.2244e-04) (hash(x)=21682445) +2528 train 7.561198 (lr=3.2231e-04) (hash(x)=24062589) +2529 train 7.745552 (lr=3.2218e-04) (hash(x)=29616079) +2530 train 7.692243 (lr=3.2205e-04) (hash(x)=25170523) +2531 train 7.351696 (lr=3.2192e-04) (hash(x)=23361504) +2532 train 7.511386 (lr=3.2179e-04) (hash(x)=24444462) +2533 train 7.540700 (lr=3.2166e-04) (hash(x)=24035993) +2534 train 7.521518 (lr=3.2153e-04) (hash(x)=24696651) +2535 train 7.485490 (lr=3.2140e-04) (hash(x)=22040184) +2536 train 7.395153 (lr=3.2127e-04) (hash(x)=27400103) +2537 train 7.716930 (lr=3.2114e-04) (hash(x)=27383080) +2538 train 7.532948 (lr=3.2100e-04) (hash(x)=24212212) +2539 train 7.435590 (lr=3.2087e-04) (hash(x)=23727731) +2540 train 7.554924 (lr=3.2074e-04) (hash(x)=24149487) +2541 train 7.342519 (lr=3.2061e-04) (hash(x)=23794649) +2542 train 7.525043 (lr=3.2048e-04) (hash(x)=26147774) +2543 train 7.646553 (lr=3.2035e-04) (hash(x)=24463229) +2544 train 7.652436 (lr=3.2022e-04) (hash(x)=26361238) +2545 train 7.387848 (lr=3.2008e-04) (hash(x)=18891545) +2546 train 7.802271 (lr=3.1995e-04) (hash(x)=30380438) +2547 train 8.030910 (lr=3.1982e-04) (hash(x)=32663792) +2548 train 7.514093 (lr=3.1969e-04) (hash(x)=25175499) +2549 train 7.654529 (lr=3.1956e-04) (hash(x)=26702407) +2550 val loss 7.5755 +2550 val perplexity 1949.8461 +2550 train 7.492146 (lr=3.1943e-04) (hash(x)=24578061) +2551 train 7.504881 (lr=3.1929e-04) (hash(x)=24091954) +2552 train 7.596628 (lr=3.1916e-04) (hash(x)=23041778) +2553 train 7.399970 (lr=3.1903e-04) (hash(x)=23686239) +2554 train 7.496202 (lr=3.1890e-04) (hash(x)=22745355) +2555 train 7.491947 (lr=3.1876e-04) (hash(x)=23599013) +2556 train 7.466020 (lr=3.1863e-04) (hash(x)=26033088) +2557 train 7.474578 (lr=3.1850e-04) (hash(x)=21634218) +2558 train 7.448647 (lr=3.1837e-04) (hash(x)=20985281) +2559 train 7.699436 (lr=3.1823e-04) (hash(x)=26670219) +2560 train 7.551366 (lr=3.1810e-04) (hash(x)=26499936) +2561 train 7.635513 (lr=3.1797e-04) (hash(x)=23547908) +2562 train 7.465952 (lr=3.1784e-04) (hash(x)=22306373) +2563 train 7.554216 (lr=3.1770e-04) (hash(x)=24380893) +2564 train 7.499713 (lr=3.1757e-04) (hash(x)=23726190) +2565 train 7.563447 (lr=3.1744e-04) (hash(x)=26967512) +2566 train 7.592081 (lr=3.1730e-04) (hash(x)=23414576) +2567 train 7.774927 (lr=3.1717e-04) (hash(x)=25558986) +2568 train 7.694700 (lr=3.1704e-04) (hash(x)=27057505) +2569 train 7.678004 (lr=3.1690e-04) (hash(x)=26048135) +2570 train 7.629983 (lr=3.1677e-04) (hash(x)=26991032) +2571 train 7.654577 (lr=3.1664e-04) (hash(x)=25729492) +2572 train 7.332612 (lr=3.1650e-04) (hash(x)=20611723) +2573 train 7.394108 (lr=3.1637e-04) (hash(x)=24563606) +2574 train 7.388317 (lr=3.1623e-04) (hash(x)=23330043) +2575 train 7.560094 (lr=3.1610e-04) (hash(x)=19218943) +2576 train 7.794534 (lr=3.1597e-04) (hash(x)=20985122) +2577 train 7.424062 (lr=3.1583e-04) (hash(x)=24133609) +2578 train 7.756445 (lr=3.1570e-04) (hash(x)=28368610) +2579 train 7.603200 (lr=3.1556e-04) (hash(x)=23952206) +2580 train 7.470450 (lr=3.1543e-04) (hash(x)=23068957) +2581 train 7.560357 (lr=3.1530e-04) (hash(x)=25365277) +2582 train 7.899632 (lr=3.1516e-04) (hash(x)=24721184) +2583 train 7.575305 (lr=3.1503e-04) (hash(x)=24551402) +2584 train 7.331104 (lr=3.1489e-04) (hash(x)=20469327) +2585 train 7.726242 (lr=3.1476e-04) (hash(x)=24966478) +2586 train 7.433216 (lr=3.1462e-04) (hash(x)=18626184) +2587 train 7.450394 (lr=3.1449e-04) (hash(x)=24007642) +2588 train 7.472137 (lr=3.1435e-04) (hash(x)=23521875) +2589 train 7.625244 (lr=3.1422e-04) (hash(x)=25204207) +2590 train 7.491278 (lr=3.1408e-04) (hash(x)=25449801) +2591 train 7.650146 (lr=3.1395e-04) (hash(x)=25229281) +2592 train 7.497692 (lr=3.1381e-04) (hash(x)=23202696) +2593 train 7.411812 (lr=3.1368e-04) (hash(x)=22435944) +2594 train 7.434789 (lr=3.1354e-04) (hash(x)=20787439) +2595 train 7.172267 (lr=3.1341e-04) (hash(x)=18214283) +2596 train 7.293515 (lr=3.1327e-04) (hash(x)=20851477) +2597 train 7.262464 (lr=3.1314e-04) (hash(x)=20609675) +2598 train 7.344547 (lr=3.1300e-04) (hash(x)=20061218) +2599 train 7.347156 (lr=3.1287e-04) (hash(x)=21251127) +2600 val loss 7.6248 +2600 val perplexity 2048.4043 +2600 train 7.830492 (lr=3.1273e-04) (hash(x)=30948038) +2601 train 8.087079 (lr=3.1259e-04) (hash(x)=32279160) +2602 train 7.776988 (lr=3.1246e-04) (hash(x)=26607151) +2603 train 7.639955 (lr=3.1232e-04) (hash(x)=25097619) +2604 train 7.348153 (lr=3.1219e-04) (hash(x)=20892421) +2605 train 7.481648 (lr=3.1205e-04) (hash(x)=22212821) +2606 train 7.682632 (lr=3.1191e-04) (hash(x)=24720588) +2607 train 7.678799 (lr=3.1178e-04) (hash(x)=25800857) +2608 train 7.528731 (lr=3.1164e-04) (hash(x)=22901279) +2609 train 7.374939 (lr=3.1150e-04) (hash(x)=20921910) +2610 train 7.655112 (lr=3.1137e-04) (hash(x)=24846267) +2611 train 7.587153 (lr=3.1123e-04) (hash(x)=22918428) +2612 train 7.567189 (lr=3.1110e-04) (hash(x)=22805901) +2613 train 7.741151 (lr=3.1096e-04) (hash(x)=24345816) +2614 train 7.732550 (lr=3.1082e-04) (hash(x)=22962012) +2615 train 8.015314 (lr=3.1069e-04) (hash(x)=25836189) +2616 train 8.110526 (lr=3.1055e-04) (hash(x)=27922916) +2617 train 7.858102 (lr=3.1041e-04) (hash(x)=21697866) +2618 train 7.764106 (lr=3.1027e-04) (hash(x)=24283369) +2619 train 7.621132 (lr=3.1014e-04) (hash(x)=24504567) +2620 train 7.653929 (lr=3.1000e-04) (hash(x)=25557725) +2621 train 7.617649 (lr=3.0986e-04) (hash(x)=23135849) +2622 train 7.524599 (lr=3.0973e-04) (hash(x)=22888908) +2623 train 7.666982 (lr=3.0959e-04) (hash(x)=25999255) +2624 train 7.593573 (lr=3.0945e-04) (hash(x)=24446851) +2625 train 7.336913 (lr=3.0931e-04) (hash(x)=21137520) +2626 train 7.609197 (lr=3.0918e-04) (hash(x)=26245754) +2627 train 7.887892 (lr=3.0904e-04) (hash(x)=27308968) +2628 train 7.595905 (lr=3.0890e-04) (hash(x)=23961169) +2629 train 7.561854 (lr=3.0876e-04) (hash(x)=25924731) +2630 train 7.672249 (lr=3.0862e-04) (hash(x)=25782315) +2631 train 7.712986 (lr=3.0849e-04) (hash(x)=20149394) +2632 train 7.732179 (lr=3.0835e-04) (hash(x)=23801981) +2633 train 7.605774 (lr=3.0821e-04) (hash(x)=23830286) +2634 train 7.527687 (lr=3.0807e-04) (hash(x)=25325236) +2635 train 7.528098 (lr=3.0793e-04) (hash(x)=24498556) +2636 train 7.500652 (lr=3.0780e-04) (hash(x)=23693078) +2637 train 7.685333 (lr=3.0766e-04) (hash(x)=25484922) +2638 train 7.397996 (lr=3.0752e-04) (hash(x)=22645025) +2639 train 7.349008 (lr=3.0738e-04) (hash(x)=21999338) +2640 train 7.538811 (lr=3.0724e-04) (hash(x)=21758019) +2641 train 7.457994 (lr=3.0710e-04) (hash(x)=24064168) +2642 train 7.502254 (lr=3.0697e-04) (hash(x)=26847292) +2643 train 7.359331 (lr=3.0683e-04) (hash(x)=23280568) +2644 train 7.349178 (lr=3.0669e-04) (hash(x)=21749161) +2645 train 7.914112 (lr=3.0655e-04) (hash(x)=30082352) +2646 train 7.845040 (lr=3.0641e-04) (hash(x)=28334297) +2647 train 7.906505 (lr=3.0627e-04) (hash(x)=27611302) +2648 train 8.232608 (lr=3.0613e-04) (hash(x)=31007436) +2649 train 7.404274 (lr=3.0599e-04) (hash(x)=22356183) +2650 val loss 7.6262 +2650 val perplexity 2051.2576 +2650 train 7.431428 (lr=3.0585e-04) (hash(x)=23071731) +2651 train 7.410599 (lr=3.0571e-04) (hash(x)=23982308) +2652 train 7.662874 (lr=3.0558e-04) (hash(x)=25673823) +2653 train 7.551700 (lr=3.0544e-04) (hash(x)=22973788) +2654 train 7.720096 (lr=3.0530e-04) (hash(x)=25386647) +2655 train 7.507926 (lr=3.0516e-04) (hash(x)=22778356) +2656 train 7.714279 (lr=3.0502e-04) (hash(x)=26669130) +2657 train 7.518460 (lr=3.0488e-04) (hash(x)=23542930) +2658 train 7.629966 (lr=3.0474e-04) (hash(x)=23307871) +2659 train 7.603026 (lr=3.0460e-04) (hash(x)=23467046) +2660 train 7.653805 (lr=3.0446e-04) (hash(x)=24728872) +2661 train 8.199767 (lr=3.0432e-04) (hash(x)=29719902) +2662 train 7.639475 (lr=3.0418e-04) (hash(x)=25114165) +2663 train 7.571923 (lr=3.0404e-04) (hash(x)=24195959) +2664 train 7.770074 (lr=3.0390e-04) (hash(x)=26938509) +2665 train 7.913345 (lr=3.0376e-04) (hash(x)=27168434) +2666 train 8.063135 (lr=3.0362e-04) (hash(x)=27488221) +2667 train 8.695448 (lr=3.0348e-04) (hash(x)=32710438) +2668 train 7.841672 (lr=3.0334e-04) (hash(x)=25073185) +2669 train 7.823764 (lr=3.0320e-04) (hash(x)=26951664) +2670 train 7.949292 (lr=3.0306e-04) (hash(x)=24886228) +2671 train 7.518607 (lr=3.0292e-04) (hash(x)=19127465) +2672 train 7.752877 (lr=3.0278e-04) (hash(x)=27134917) +2673 train 7.679662 (lr=3.0263e-04) (hash(x)=25673955) +2674 train 8.245580 (lr=3.0249e-04) (hash(x)=27111776) +2675 train 7.969502 (lr=3.0235e-04) (hash(x)=28962580) +2676 train 7.639732 (lr=3.0221e-04) (hash(x)=25593381) +2677 train 7.724451 (lr=3.0207e-04) (hash(x)=25238916) +2678 train 7.664962 (lr=3.0193e-04) (hash(x)=27453574) +2679 train 7.566804 (lr=3.0179e-04) (hash(x)=22426274) +2680 train 7.625548 (lr=3.0165e-04) (hash(x)=22974780) +2681 train 7.606184 (lr=3.0151e-04) (hash(x)=24668644) +2682 train 7.578682 (lr=3.0137e-04) (hash(x)=26182084) +2683 train 7.790826 (lr=3.0122e-04) (hash(x)=25606512) +2684 train 7.613973 (lr=3.0108e-04) (hash(x)=24885252) +2685 train 7.450222 (lr=3.0094e-04) (hash(x)=21290254) +2686 train 7.362501 (lr=3.0080e-04) (hash(x)=21367078) +2687 train 7.587686 (lr=3.0066e-04) (hash(x)=23785205) +2688 train 7.526971 (lr=3.0052e-04) (hash(x)=24640056) +2689 train 7.910902 (lr=3.0037e-04) (hash(x)=27083886) +2690 train 7.572687 (lr=3.0023e-04) (hash(x)=23603571) +2691 train 7.561664 (lr=3.0009e-04) (hash(x)=24779414) +2692 train 7.400400 (lr=2.9995e-04) (hash(x)=22200693) +2693 train 7.693246 (lr=2.9981e-04) (hash(x)=26907868) +2694 train 8.039865 (lr=2.9967e-04) (hash(x)=32034827) +2695 train 7.874500 (lr=2.9952e-04) (hash(x)=28505676) +2696 train 7.541142 (lr=2.9938e-04) (hash(x)=23805750) +2697 train 7.879602 (lr=2.9924e-04) (hash(x)=29804750) +2698 train 8.102868 (lr=2.9910e-04) (hash(x)=28443583) +2699 train 7.986791 (lr=2.9895e-04) (hash(x)=29156288) +2700 val loss 7.6619 +2700 val perplexity 2125.7664 +2700 train 7.720628 (lr=2.9881e-04) (hash(x)=25895743) +2701 train 7.326734 (lr=2.9867e-04) (hash(x)=21173795) +2702 train 7.199821 (lr=2.9853e-04) (hash(x)=20790866) +2703 train 7.780336 (lr=2.9838e-04) (hash(x)=27706477) +2704 train 7.934834 (lr=2.9824e-04) (hash(x)=30358985) +2705 train 7.504263 (lr=2.9810e-04) (hash(x)=23548492) +2706 train 7.658929 (lr=2.9796e-04) (hash(x)=25879696) +2707 train 7.522324 (lr=2.9781e-04) (hash(x)=23711800) +2708 train 8.040217 (lr=2.9767e-04) (hash(x)=28763123) +2709 train 7.646749 (lr=2.9753e-04) (hash(x)=23327642) +2710 train 7.719804 (lr=2.9738e-04) (hash(x)=25634166) +2711 train 7.627378 (lr=2.9724e-04) (hash(x)=28684175) +2712 train 7.741246 (lr=2.9710e-04) (hash(x)=23996278) +2713 train 7.383585 (lr=2.9696e-04) (hash(x)=21795980) +2714 train 7.692014 (lr=2.9681e-04) (hash(x)=24840769) +2715 train 7.654176 (lr=2.9667e-04) (hash(x)=25225466) +2716 train 7.890550 (lr=2.9653e-04) (hash(x)=27500471) +2717 train 7.605792 (lr=2.9638e-04) (hash(x)=24703036) +2718 train 7.719398 (lr=2.9624e-04) (hash(x)=24294293) +2719 train 7.694734 (lr=2.9610e-04) (hash(x)=28003600) +2720 train 7.454835 (lr=2.9595e-04) (hash(x)=22822962) +2721 train 7.358738 (lr=2.9581e-04) (hash(x)=24189246) +2722 train 7.474323 (lr=2.9566e-04) (hash(x)=22608951) +2723 train 8.140970 (lr=2.9552e-04) (hash(x)=27989890) +2724 train 7.748544 (lr=2.9538e-04) (hash(x)=24175838) +2725 train 7.590280 (lr=2.9523e-04) (hash(x)=24781792) +2726 train 7.493558 (lr=2.9509e-04) (hash(x)=23413276) +2727 train 7.929990 (lr=2.9494e-04) (hash(x)=27586845) +2728 train 7.726837 (lr=2.9480e-04) (hash(x)=27336264) +2729 train 7.726492 (lr=2.9466e-04) (hash(x)=26808464) +2730 train 7.571490 (lr=2.9451e-04) (hash(x)=22312009) +2731 train 7.660332 (lr=2.9437e-04) (hash(x)=22373927) +2732 train 7.597857 (lr=2.9422e-04) (hash(x)=23428834) +2733 train 7.583540 (lr=2.9408e-04) (hash(x)=25304441) +2734 train 7.568832 (lr=2.9394e-04) (hash(x)=24798164) +2735 train 7.439684 (lr=2.9379e-04) (hash(x)=21176405) +2736 train 7.541655 (lr=2.9365e-04) (hash(x)=22343075) +2737 train 7.568414 (lr=2.9350e-04) (hash(x)=23825332) +2738 train 7.550744 (lr=2.9336e-04) (hash(x)=24191865) +2739 train 7.408466 (lr=2.9321e-04) (hash(x)=23806052) +2740 train 7.507267 (lr=2.9307e-04) (hash(x)=21764591) +2741 train 7.624556 (lr=2.9292e-04) (hash(x)=25548695) +2742 train 7.772443 (lr=2.9278e-04) (hash(x)=26847535) +2743 train 7.569119 (lr=2.9263e-04) (hash(x)=25888433) +2744 train 7.534772 (lr=2.9249e-04) (hash(x)=24327454) +2745 train 7.459232 (lr=2.9234e-04) (hash(x)=22543301) +2746 train 7.771968 (lr=2.9220e-04) (hash(x)=24593022) +2747 train 8.318924 (lr=2.9205e-04) (hash(x)=27797727) +2748 train 7.811613 (lr=2.9191e-04) (hash(x)=28067682) +2749 train 7.496199 (lr=2.9176e-04) (hash(x)=25278538) +2750 val loss 7.6467 +2750 val perplexity 2093.7866 +2750 train 7.434452 (lr=2.9162e-04) (hash(x)=23875731) +2751 train 7.680981 (lr=2.9147e-04) (hash(x)=27916982) +2752 train 7.613835 (lr=2.9133e-04) (hash(x)=25726799) +2753 train 7.477937 (lr=2.9118e-04) (hash(x)=25227141) +2754 train 7.713853 (lr=2.9104e-04) (hash(x)=27679212) +2755 train 7.409975 (lr=2.9089e-04) (hash(x)=24621793) +2756 train 7.352681 (lr=2.9074e-04) (hash(x)=21962296) +2757 train 7.620853 (lr=2.9060e-04) (hash(x)=24899679) +2758 train 7.396061 (lr=2.9045e-04) (hash(x)=21452158) +2759 train 7.492284 (lr=2.9031e-04) (hash(x)=24334708) +2760 train 7.720273 (lr=2.9016e-04) (hash(x)=25523041) +2761 train 7.809219 (lr=2.9002e-04) (hash(x)=30389813) +2762 train 7.428815 (lr=2.8987e-04) (hash(x)=22426014) +2763 train 7.538651 (lr=2.8972e-04) (hash(x)=24419143) +2764 train 7.564624 (lr=2.8958e-04) (hash(x)=24850536) +2765 train 7.619254 (lr=2.8943e-04) (hash(x)=24181393) +2766 train 7.239185 (lr=2.8929e-04) (hash(x)=18882503) +2767 train 7.630744 (lr=2.8914e-04) (hash(x)=25617709) +2768 train 7.564992 (lr=2.8899e-04) (hash(x)=24076662) +2769 train 7.461965 (lr=2.8885e-04) (hash(x)=21656802) +2770 train 7.465108 (lr=2.8870e-04) (hash(x)=21014265) +2771 train 7.564090 (lr=2.8855e-04) (hash(x)=24556034) +2772 train 7.374356 (lr=2.8841e-04) (hash(x)=22046665) +2773 train 7.689490 (lr=2.8826e-04) (hash(x)=26761579) +2774 train 9.360749 (lr=2.8811e-04) (hash(x)=41414315) +2775 train 7.485172 (lr=2.8797e-04) (hash(x)=25152362) +2776 train 7.665555 (lr=2.8782e-04) (hash(x)=25567641) +2777 train 7.612237 (lr=2.8767e-04) (hash(x)=25427935) +2778 train 7.689299 (lr=2.8753e-04) (hash(x)=25824457) +2779 train 7.621892 (lr=2.8738e-04) (hash(x)=24326376) +2780 train 7.656792 (lr=2.8723e-04) (hash(x)=27447230) +2781 train 7.515539 (lr=2.8709e-04) (hash(x)=24003710) +2782 train 7.444980 (lr=2.8694e-04) (hash(x)=24157390) +2783 train 7.535290 (lr=2.8679e-04) (hash(x)=24276512) +2784 train 7.362305 (lr=2.8665e-04) (hash(x)=21503752) +2785 train 7.291444 (lr=2.8650e-04) (hash(x)=20031488) +2786 train 7.376045 (lr=2.8635e-04) (hash(x)=21788715) +2787 train 7.462308 (lr=2.8620e-04) (hash(x)=24344695) +2788 train 7.509881 (lr=2.8606e-04) (hash(x)=22927763) +2789 train 7.457961 (lr=2.8591e-04) (hash(x)=23710755) +2790 train 7.639013 (lr=2.8576e-04) (hash(x)=26924620) +2791 train 7.936086 (lr=2.8561e-04) (hash(x)=26776133) +2792 train 7.419830 (lr=2.8547e-04) (hash(x)=19936770) +2793 train 7.601149 (lr=2.8532e-04) (hash(x)=25440959) +2794 train 7.582784 (lr=2.8517e-04) (hash(x)=25146097) +2795 train 7.336890 (lr=2.8502e-04) (hash(x)=21847282) +2796 train 7.557822 (lr=2.8488e-04) (hash(x)=25639784) +2797 train 7.342943 (lr=2.8473e-04) (hash(x)=21199921) +2798 train 7.540112 (lr=2.8458e-04) (hash(x)=22360806) +2799 train 6.973992 (lr=2.8443e-04) (hash(x)=20254159) +2800 val loss 7.6107 +2800 val perplexity 2019.7119 +2800 train 7.100722 (lr=2.8428e-04) (hash(x)=23348345) +2801 train 7.678205 (lr=2.8414e-04) (hash(x)=24908033) +2802 train 7.472167 (lr=2.8399e-04) (hash(x)=23350309) +2803 train 7.599140 (lr=2.8384e-04) (hash(x)=25044762) +2804 train 7.476516 (lr=2.8369e-04) (hash(x)=24071026) +2805 train 7.277523 (lr=2.8354e-04) (hash(x)=22169363) +2806 train 7.495167 (lr=2.8340e-04) (hash(x)=23757564) +2807 train 7.870381 (lr=2.8325e-04) (hash(x)=27873855) +2808 train 7.727716 (lr=2.8310e-04) (hash(x)=26577893) +2809 train 7.772932 (lr=2.8295e-04) (hash(x)=27001634) +2810 train 7.500856 (lr=2.8280e-04) (hash(x)=24796541) +2811 train 7.468372 (lr=2.8265e-04) (hash(x)=22575615) +2812 train 7.530862 (lr=2.8251e-04) (hash(x)=25876475) +2813 train 7.527757 (lr=2.8236e-04) (hash(x)=24765155) +2814 train 7.563413 (lr=2.8221e-04) (hash(x)=25785699) +2815 train 7.700778 (lr=2.8206e-04) (hash(x)=25113614) +2816 train 7.491042 (lr=2.8191e-04) (hash(x)=24415748) +2817 train 7.640218 (lr=2.8176e-04) (hash(x)=25140622) +2818 train 7.567773 (lr=2.8161e-04) (hash(x)=24845866) +2819 train 7.997438 (lr=2.8146e-04) (hash(x)=28062905) +2820 train 7.567276 (lr=2.8132e-04) (hash(x)=22041086) +2821 train 7.578797 (lr=2.8117e-04) (hash(x)=24957184) +2822 train 7.575957 (lr=2.8102e-04) (hash(x)=24360380) +2823 train 7.665112 (lr=2.8087e-04) (hash(x)=26192886) +2824 train 7.629699 (lr=2.8072e-04) (hash(x)=25001858) +2825 train 7.580240 (lr=2.8057e-04) (hash(x)=24721193) +2826 train 7.659862 (lr=2.8042e-04) (hash(x)=26186227) +2827 train 7.551744 (lr=2.8027e-04) (hash(x)=25770338) +2828 train 7.640477 (lr=2.8012e-04) (hash(x)=25920767) +2829 train 7.554585 (lr=2.7997e-04) (hash(x)=25060684) +2830 train 7.329568 (lr=2.7982e-04) (hash(x)=22933946) +2831 train 7.470088 (lr=2.7967e-04) (hash(x)=24614912) +2832 train 7.296040 (lr=2.7952e-04) (hash(x)=19955522) +2833 train 7.298585 (lr=2.7938e-04) (hash(x)=21111215) +2834 train 8.069892 (lr=2.7923e-04) (hash(x)=28817924) +2835 train 7.598335 (lr=2.7908e-04) (hash(x)=26934071) +2836 train 7.585977 (lr=2.7893e-04) (hash(x)=24768851) +2837 train 7.557935 (lr=2.7878e-04) (hash(x)=25706447) +2838 train 7.385306 (lr=2.7863e-04) (hash(x)=19579834) +2839 train 7.663058 (lr=2.7848e-04) (hash(x)=25397093) +2840 train 8.002420 (lr=2.7833e-04) (hash(x)=27902141) +2841 train 7.647062 (lr=2.7818e-04) (hash(x)=25383069) +2842 train 7.457820 (lr=2.7803e-04) (hash(x)=22007373) +2843 train 7.526759 (lr=2.7788e-04) (hash(x)=25925963) +2844 train 7.570450 (lr=2.7773e-04) (hash(x)=25711128) +2845 train 7.386585 (lr=2.7758e-04) (hash(x)=21881216) +2846 train 7.455223 (lr=2.7743e-04) (hash(x)=20277075) +2847 train 7.503613 (lr=2.7728e-04) (hash(x)=19811802) +2848 train 7.487616 (lr=2.7713e-04) (hash(x)=23878906) +2849 train 7.619178 (lr=2.7698e-04) (hash(x)=25034966) +2850 val loss 7.6206 +2850 val perplexity 2039.7249 +2850 train 7.743805 (lr=2.7683e-04) (hash(x)=24359507) +2851 train 7.437614 (lr=2.7668e-04) (hash(x)=23248423) +2852 train 7.467719 (lr=2.7653e-04) (hash(x)=21782773) +2853 train 7.525748 (lr=2.7638e-04) (hash(x)=23804418) +2854 train 7.545763 (lr=2.7623e-04) (hash(x)=22525078) +2855 train 7.588967 (lr=2.7607e-04) (hash(x)=25579655) +2856 train 7.726186 (lr=2.7592e-04) (hash(x)=27048876) +2857 train 7.842229 (lr=2.7577e-04) (hash(x)=26468479) +2858 train 7.574571 (lr=2.7562e-04) (hash(x)=23854933) +2859 train 7.843497 (lr=2.7547e-04) (hash(x)=25537603) +2860 train 7.305583 (lr=2.7532e-04) (hash(x)=20979252) +2861 train 7.820880 (lr=2.7517e-04) (hash(x)=26504374) +2862 train 7.716035 (lr=2.7502e-04) (hash(x)=27561842) +2863 train 7.663931 (lr=2.7487e-04) (hash(x)=26096514) +2864 train 7.668645 (lr=2.7472e-04) (hash(x)=25926899) +2865 train 7.667192 (lr=2.7457e-04) (hash(x)=26058348) +2866 train 7.759872 (lr=2.7442e-04) (hash(x)=29802259) +2867 train 7.588413 (lr=2.7427e-04) (hash(x)=24132888) +2868 train 7.359421 (lr=2.7411e-04) (hash(x)=23369410) +2869 train 7.508573 (lr=2.7396e-04) (hash(x)=25387506) +2870 train 7.705833 (lr=2.7381e-04) (hash(x)=27375344) +2871 train 7.440006 (lr=2.7366e-04) (hash(x)=22589633) +2872 train 7.574335 (lr=2.7351e-04) (hash(x)=23250237) +2873 train 7.541914 (lr=2.7336e-04) (hash(x)=25511322) +2874 train 7.185426 (lr=2.7321e-04) (hash(x)=18356418) +2875 train 7.555801 (lr=2.7306e-04) (hash(x)=27781566) +2876 train 7.476973 (lr=2.7290e-04) (hash(x)=24878173) +2877 train 7.886730 (lr=2.7275e-04) (hash(x)=30018637) +2878 train 7.733195 (lr=2.7260e-04) (hash(x)=27168416) +2879 train 7.794725 (lr=2.7245e-04) (hash(x)=26757147) +2880 train 7.751584 (lr=2.7230e-04) (hash(x)=26637081) +2881 train 7.659635 (lr=2.7215e-04) (hash(x)=24795024) +2882 train 7.812862 (lr=2.7200e-04) (hash(x)=29787745) +2883 train 7.718530 (lr=2.7184e-04) (hash(x)=26649864) +2884 train 7.629097 (lr=2.7169e-04) (hash(x)=27306612) +2885 train 7.695385 (lr=2.7154e-04) (hash(x)=27568311) +2886 train 7.751823 (lr=2.7139e-04) (hash(x)=27440150) +2887 train 7.496136 (lr=2.7124e-04) (hash(x)=24963730) +2888 train 7.524231 (lr=2.7108e-04) (hash(x)=23619807) +2889 train 7.917708 (lr=2.7093e-04) (hash(x)=29447356) +2890 train 7.560645 (lr=2.7078e-04) (hash(x)=25144675) +2891 train 7.491590 (lr=2.7063e-04) (hash(x)=25249959) +2892 train 7.478461 (lr=2.7048e-04) (hash(x)=26608712) +2893 train 7.459970 (lr=2.7032e-04) (hash(x)=26333258) +2894 train 7.271580 (lr=2.7017e-04) (hash(x)=20682182) +2895 train 7.678554 (lr=2.7002e-04) (hash(x)=27703124) +2896 train 7.516733 (lr=2.6987e-04) (hash(x)=23228180) +2897 train 7.419450 (lr=2.6972e-04) (hash(x)=25252411) +2898 train 7.352489 (lr=2.6956e-04) (hash(x)=22879178) +2899 train 7.587630 (lr=2.6941e-04) (hash(x)=26459082) +2900 val loss 7.6008 +2900 val perplexity 1999.6989 +2900 train 7.442146 (lr=2.6926e-04) (hash(x)=24569501) +2901 train 7.273937 (lr=2.6911e-04) (hash(x)=19803884) +2902 train 7.320888 (lr=2.6895e-04) (hash(x)=18799747) +2903 train 7.561453 (lr=2.6880e-04) (hash(x)=24781713) +2904 train 7.653181 (lr=2.6865e-04) (hash(x)=25016590) +2905 train 7.801347 (lr=2.6850e-04) (hash(x)=29006906) +2906 train 7.633911 (lr=2.6834e-04) (hash(x)=24069959) +2907 train 7.744317 (lr=2.6819e-04) (hash(x)=26597693) +2908 train 7.389574 (lr=2.6804e-04) (hash(x)=25014146) +2909 train 7.602523 (lr=2.6789e-04) (hash(x)=24943747) +2910 train 7.533235 (lr=2.6773e-04) (hash(x)=27847542) +2911 train 7.496572 (lr=2.6758e-04) (hash(x)=24720476) +2912 train 7.385433 (lr=2.6743e-04) (hash(x)=24388804) +2913 train 7.461097 (lr=2.6728e-04) (hash(x)=23567535) +2914 train 7.741251 (lr=2.6712e-04) (hash(x)=29673625) +2915 train 7.552230 (lr=2.6697e-04) (hash(x)=23691295) +2916 train 7.878877 (lr=2.6682e-04) (hash(x)=26572819) +2917 train 7.573219 (lr=2.6666e-04) (hash(x)=23237812) +2918 train 7.779974 (lr=2.6651e-04) (hash(x)=26531016) +2919 train 7.383636 (lr=2.6636e-04) (hash(x)=23481301) +2920 train 7.593565 (lr=2.6620e-04) (hash(x)=24839184) +2921 train 7.484639 (lr=2.6605e-04) (hash(x)=23327755) +2922 train 7.687276 (lr=2.6590e-04) (hash(x)=26347114) +2923 train 7.445297 (lr=2.6575e-04) (hash(x)=23295676) +2924 train 7.466899 (lr=2.6559e-04) (hash(x)=24557178) +2925 train 7.545429 (lr=2.6544e-04) (hash(x)=26067788) +2926 train 7.500725 (lr=2.6529e-04) (hash(x)=25694982) +2927 train 7.569417 (lr=2.6513e-04) (hash(x)=25641033) +2928 train 7.608638 (lr=2.6498e-04) (hash(x)=24906422) +2929 train 7.794406 (lr=2.6483e-04) (hash(x)=27803515) +2930 train 7.742408 (lr=2.6467e-04) (hash(x)=26208803) +2931 train 7.324579 (lr=2.6452e-04) (hash(x)=22441379) +2932 train 7.519979 (lr=2.6437e-04) (hash(x)=24741626) +2933 train 7.426122 (lr=2.6421e-04) (hash(x)=24595257) +2934 train 7.430921 (lr=2.6406e-04) (hash(x)=23939167) +2935 train 7.699617 (lr=2.6390e-04) (hash(x)=27369437) +2936 train 7.416354 (lr=2.6375e-04) (hash(x)=21409783) +2937 train 7.631296 (lr=2.6360e-04) (hash(x)=25923735) +2938 train 7.640036 (lr=2.6344e-04) (hash(x)=29559511) +2939 train 7.586918 (lr=2.6329e-04) (hash(x)=24482272) +2940 train 7.530301 (lr=2.6314e-04) (hash(x)=24767658) +2941 train 7.705145 (lr=2.6298e-04) (hash(x)=26425020) +2942 train 7.982467 (lr=2.6283e-04) (hash(x)=27444868) +2943 train 7.530257 (lr=2.6267e-04) (hash(x)=24760900) +2944 train 7.654965 (lr=2.6252e-04) (hash(x)=25605407) +2945 train 7.479084 (lr=2.6237e-04) (hash(x)=22886951) +2946 train 7.703787 (lr=2.6221e-04) (hash(x)=26112205) +2947 train 7.629228 (lr=2.6206e-04) (hash(x)=23919156) +2948 train 7.601663 (lr=2.6190e-04) (hash(x)=23729312) +2949 train 7.831824 (lr=2.6175e-04) (hash(x)=30440878) +2950 val loss 7.6496 +2950 val perplexity 2099.7153 +2950 train 7.324841 (lr=2.6160e-04) (hash(x)=20004041) +2951 train 7.354625 (lr=2.6144e-04) (hash(x)=21692546) +2952 train 7.352127 (lr=2.6129e-04) (hash(x)=23021681) +2953 train 7.682297 (lr=2.6113e-04) (hash(x)=26663597) +2954 train 7.520084 (lr=2.6098e-04) (hash(x)=23727385) +2955 train 7.566523 (lr=2.6083e-04) (hash(x)=27692087) +2956 train 7.520856 (lr=2.6067e-04) (hash(x)=24003378) +2957 train 8.156444 (lr=2.6052e-04) (hash(x)=29534673) +2958 train 7.475428 (lr=2.6036e-04) (hash(x)=22875068) +2959 train 7.491154 (lr=2.6021e-04) (hash(x)=22720391) +2960 train 7.262828 (lr=2.6005e-04) (hash(x)=17997400) +2961 train 7.460462 (lr=2.5990e-04) (hash(x)=22853822) +2962 train 7.516611 (lr=2.5974e-04) (hash(x)=25238004) +2963 train 7.609848 (lr=2.5959e-04) (hash(x)=26146560) +2964 train 7.693243 (lr=2.5944e-04) (hash(x)=21894867) +2965 train 7.672026 (lr=2.5928e-04) (hash(x)=23001150) +2966 train 7.473753 (lr=2.5913e-04) (hash(x)=23392923) +2967 train 7.508502 (lr=2.5897e-04) (hash(x)=24376979) +2968 train 7.482789 (lr=2.5882e-04) (hash(x)=23781449) +2969 train 7.652971 (lr=2.5866e-04) (hash(x)=25315495) +2970 train 7.750797 (lr=2.5851e-04) (hash(x)=27165470) +2971 train 7.578718 (lr=2.5835e-04) (hash(x)=22917712) +2972 train 7.895814 (lr=2.5820e-04) (hash(x)=27928456) +2973 train 7.383882 (lr=2.5804e-04) (hash(x)=19890855) +2974 train 7.386646 (lr=2.5789e-04) (hash(x)=21318134) +2975 train 7.429111 (lr=2.5773e-04) (hash(x)=22244509) +2976 train 7.279510 (lr=2.5758e-04) (hash(x)=21293137) +2977 train 7.465400 (lr=2.5742e-04) (hash(x)=23465789) +2978 train 7.448997 (lr=2.5727e-04) (hash(x)=21169753) +2979 train 7.753068 (lr=2.5711e-04) (hash(x)=25243385) +2980 train 7.745484 (lr=2.5696e-04) (hash(x)=27465812) +2981 train 7.540639 (lr=2.5680e-04) (hash(x)=24615492) +2982 train 7.485609 (lr=2.5665e-04) (hash(x)=23081307) +2983 train 7.256581 (lr=2.5649e-04) (hash(x)=21831960) +2984 train 7.351571 (lr=2.5634e-04) (hash(x)=23242850) +2985 train 7.486238 (lr=2.5618e-04) (hash(x)=24308188) +2986 train 7.650476 (lr=2.5603e-04) (hash(x)=28541601) +2987 train 7.543082 (lr=2.5587e-04) (hash(x)=24842373) +2988 train 7.294104 (lr=2.5572e-04) (hash(x)=21967126) +2989 train 7.483254 (lr=2.5556e-04) (hash(x)=22951616) +2990 train 7.591238 (lr=2.5541e-04) (hash(x)=24325714) +2991 train 7.425210 (lr=2.5525e-04) (hash(x)=24921535) +2992 train 7.604671 (lr=2.5510e-04) (hash(x)=25937112) +2993 train 7.376688 (lr=2.5494e-04) (hash(x)=20716218) +2994 train 7.569032 (lr=2.5479e-04) (hash(x)=25450724) +2995 train 7.530510 (lr=2.5463e-04) (hash(x)=24344615) +2996 train 7.246767 (lr=2.5448e-04) (hash(x)=20299058) +2997 train 7.514917 (lr=2.5432e-04) (hash(x)=23859426) +2998 train 7.558998 (lr=2.5416e-04) (hash(x)=23094397) +2999 train 7.796785 (lr=2.5401e-04) (hash(x)=25381251) +3000 val loss 7.6273 +3000 val perplexity 2053.4172 +3000 train 7.493846 (lr=2.5385e-04) (hash(x)=23586527) +3001 train 7.421700 (lr=2.5370e-04) (hash(x)=24220410) +3002 train 7.353053 (lr=2.5354e-04) (hash(x)=20597347) +3003 train 7.460409 (lr=2.5339e-04) (hash(x)=22887303) +3004 train 7.659337 (lr=2.5323e-04) (hash(x)=25869462) +3005 train 7.470320 (lr=2.5307e-04) (hash(x)=22098530) +3006 train 7.717811 (lr=2.5292e-04) (hash(x)=26246291) +3007 train 7.651538 (lr=2.5276e-04) (hash(x)=25687352) +3008 train 7.429274 (lr=2.5261e-04) (hash(x)=25425646) +3009 train 7.601193 (lr=2.5245e-04) (hash(x)=26021124) +3010 train 7.627213 (lr=2.5230e-04) (hash(x)=25392057) +3011 train 7.695348 (lr=2.5214e-04) (hash(x)=27791412) +3012 train 7.507376 (lr=2.5198e-04) (hash(x)=23181098) +3013 train 7.546267 (lr=2.5183e-04) (hash(x)=25521889) +3014 train 7.365952 (lr=2.5167e-04) (hash(x)=21685795) +3015 train 7.494363 (lr=2.5152e-04) (hash(x)=25221654) +3016 train 7.685962 (lr=2.5136e-04) (hash(x)=24888744) +3017 train 7.604793 (lr=2.5120e-04) (hash(x)=24200150) +3018 train 7.568573 (lr=2.5105e-04) (hash(x)=26943942) +3019 train 7.556657 (lr=2.5089e-04) (hash(x)=23243731) +3020 train 7.491747 (lr=2.5074e-04) (hash(x)=21068284) +3021 train 7.601840 (lr=2.5058e-04) (hash(x)=23876902) +3022 train 7.657104 (lr=2.5042e-04) (hash(x)=25337639) +3023 train 7.592162 (lr=2.5027e-04) (hash(x)=24469863) +3024 train 7.986727 (lr=2.5011e-04) (hash(x)=27850876) +3025 train 7.760377 (lr=2.4996e-04) (hash(x)=20515778) +3026 train 7.658739 (lr=2.4980e-04) (hash(x)=29019173) +3027 train 7.354349 (lr=2.4964e-04) (hash(x)=22484936) +3028 train 7.501741 (lr=2.4949e-04) (hash(x)=24639400) +3029 train 7.650427 (lr=2.4933e-04) (hash(x)=26835174) +3030 train 7.755634 (lr=2.4917e-04) (hash(x)=29843763) +3031 train 7.499441 (lr=2.4902e-04) (hash(x)=25291413) +3032 train 7.430595 (lr=2.4886e-04) (hash(x)=24590244) +3033 train 7.644809 (lr=2.4871e-04) (hash(x)=28880142) +3034 train 7.416149 (lr=2.4855e-04) (hash(x)=23372199) +3035 train 7.400029 (lr=2.4839e-04) (hash(x)=23952225) +3036 train 7.543378 (lr=2.4824e-04) (hash(x)=24589186) +3037 train 7.538166 (lr=2.4808e-04) (hash(x)=23260323) +3038 train 7.684480 (lr=2.4792e-04) (hash(x)=25824498) +3039 train 7.541724 (lr=2.4777e-04) (hash(x)=25744274) +3040 train 7.543299 (lr=2.4761e-04) (hash(x)=21610247) +3041 train 7.619568 (lr=2.4745e-04) (hash(x)=25079786) +3042 train 7.364667 (lr=2.4730e-04) (hash(x)=23219195) +3043 train 7.355685 (lr=2.4714e-04) (hash(x)=22616739) +3044 train 7.544556 (lr=2.4698e-04) (hash(x)=24908480) +3045 train 7.381064 (lr=2.4683e-04) (hash(x)=22293489) +3046 train 7.454633 (lr=2.4667e-04) (hash(x)=23557651) +3047 train 7.482018 (lr=2.4651e-04) (hash(x)=24246963) +3048 train 7.590378 (lr=2.4636e-04) (hash(x)=24490083) +3049 train 7.608201 (lr=2.4620e-04) (hash(x)=22372895) +3050 val loss 7.6218 +3050 val perplexity 2042.1373 +3050 train 7.351583 (lr=2.4604e-04) (hash(x)=21759470) +3051 train 7.267747 (lr=2.4589e-04) (hash(x)=19407094) +3052 train 7.526118 (lr=2.4573e-04) (hash(x)=23957047) +3053 train 7.674710 (lr=2.4557e-04) (hash(x)=24719318) +3054 train 7.513122 (lr=2.4542e-04) (hash(x)=20719314) +3055 train 7.418231 (lr=2.4526e-04) (hash(x)=19724058) +3056 train 7.291278 (lr=2.4510e-04) (hash(x)=14407266) +3057 train 7.196016 (lr=2.4495e-04) (hash(x)=12468292) +3058 train 7.407926 (lr=2.4479e-04) (hash(x)=16098279) +3059 train 7.490332 (lr=2.4463e-04) (hash(x)=18836491) +3060 train 7.347727 (lr=2.4448e-04) (hash(x)=19132277) +3061 train 7.496259 (lr=2.4432e-04) (hash(x)=22814208) +3062 train 7.518711 (lr=2.4416e-04) (hash(x)=24838508) +3063 train 7.388959 (lr=2.4401e-04) (hash(x)=20705649) +3064 train 8.330578 (lr=2.4385e-04) (hash(x)=29416914) +3065 train 7.811156 (lr=2.4369e-04) (hash(x)=25972430) +3066 train 7.284324 (lr=2.4353e-04) (hash(x)=23705805) +3067 train 7.360196 (lr=2.4338e-04) (hash(x)=21325875) +3068 train 7.528219 (lr=2.4322e-04) (hash(x)=23526506) +3069 train 7.959339 (lr=2.4306e-04) (hash(x)=27282337) +3070 train 7.830738 (lr=2.4291e-04) (hash(x)=27968043) +3071 train 7.617429 (lr=2.4275e-04) (hash(x)=24938685) +3072 train 7.847791 (lr=2.4259e-04) (hash(x)=26942737) +3073 train 7.758800 (lr=2.4243e-04) (hash(x)=23506879) +3074 train 7.480211 (lr=2.4228e-04) (hash(x)=23589913) +3075 train 7.608448 (lr=2.4212e-04) (hash(x)=25152403) +3076 train 7.719199 (lr=2.4196e-04) (hash(x)=23425868) +3077 train 7.972023 (lr=2.4181e-04) (hash(x)=23966181) +3078 train 7.960430 (lr=2.4165e-04) (hash(x)=27312570) +3079 train 7.554008 (lr=2.4149e-04) (hash(x)=21707000) +3080 train 7.489185 (lr=2.4133e-04) (hash(x)=25600427) +3081 train 7.595277 (lr=2.4118e-04) (hash(x)=24270631) +3082 train 7.713843 (lr=2.4102e-04) (hash(x)=25199537) +3083 train 7.084427 (lr=2.4086e-04) (hash(x)=17952018) +3084 train 7.179994 (lr=2.4070e-04) (hash(x)=18733143) +3085 train 7.793863 (lr=2.4055e-04) (hash(x)=26946100) +3086 train 7.595093 (lr=2.4039e-04) (hash(x)=25547515) +3087 train 7.549748 (lr=2.4023e-04) (hash(x)=24948980) +3088 train 8.665735 (lr=2.4007e-04) (hash(x)=35461645) +3089 train 7.762098 (lr=2.3992e-04) (hash(x)=28330877) +3090 train 7.703050 (lr=2.3976e-04) (hash(x)=27687861) +3091 train 7.779044 (lr=2.3960e-04) (hash(x)=28012110) +3092 train 7.671217 (lr=2.3945e-04) (hash(x)=24684480) +3093 train 7.646018 (lr=2.3929e-04) (hash(x)=26225786) +3094 train 7.376789 (lr=2.3913e-04) (hash(x)=23098156) +3095 train 7.866918 (lr=2.3897e-04) (hash(x)=30773958) +3096 train 8.140545 (lr=2.3882e-04) (hash(x)=28640406) +3097 train 8.070132 (lr=2.3866e-04) (hash(x)=28201086) +3098 train 8.472448 (lr=2.3850e-04) (hash(x)=35002344) +3099 train 8.039069 (lr=2.3834e-04) (hash(x)=29481068) +3100 val loss 7.6737 +3100 val perplexity 2150.9902 +3100 train 7.589948 (lr=2.3818e-04) (hash(x)=26374528) +3101 train 7.518249 (lr=2.3803e-04) (hash(x)=24153602) +3102 train 7.607890 (lr=2.3787e-04) (hash(x)=25478746) +3103 train 7.848384 (lr=2.3771e-04) (hash(x)=26769046) +3104 train 7.486314 (lr=2.3755e-04) (hash(x)=21841970) +3105 train 7.779526 (lr=2.3740e-04) (hash(x)=27693052) +3106 train 7.260871 (lr=2.3724e-04) (hash(x)=20689448) +3107 train 7.642484 (lr=2.3708e-04) (hash(x)=26755048) +3108 train 7.600463 (lr=2.3692e-04) (hash(x)=24431904) +3109 train 7.410771 (lr=2.3677e-04) (hash(x)=21009792) +3110 train 7.520434 (lr=2.3661e-04) (hash(x)=21909003) +3111 train 7.496160 (lr=2.3645e-04) (hash(x)=18849656) +3112 train 7.627223 (lr=2.3629e-04) (hash(x)=22223376) +3113 train 7.883882 (lr=2.3614e-04) (hash(x)=25652491) +3114 train 7.782250 (lr=2.3598e-04) (hash(x)=23521434) +3115 train 7.908892 (lr=2.3582e-04) (hash(x)=25449800) +3116 train 8.393155 (lr=2.3566e-04) (hash(x)=27655847) +3117 train 8.498025 (lr=2.3550e-04) (hash(x)=29878248) +3118 train 8.529788 (lr=2.3535e-04) (hash(x)=30444094) +3119 train 7.981995 (lr=2.3519e-04) (hash(x)=24624950) +3120 train 7.328812 (lr=2.3503e-04) (hash(x)=20798511) +3121 train 7.682972 (lr=2.3487e-04) (hash(x)=26581679) +3122 train 7.634810 (lr=2.3471e-04) (hash(x)=25333422) +3123 train 7.722302 (lr=2.3456e-04) (hash(x)=26174069) +3124 train 7.748146 (lr=2.3440e-04) (hash(x)=25219475) +3125 train 7.651714 (lr=2.3424e-04) (hash(x)=20919061) +3126 train 7.610909 (lr=2.3408e-04) (hash(x)=23828688) +3127 train 8.132457 (lr=2.3393e-04) (hash(x)=27299605) +3128 train 7.492918 (lr=2.3377e-04) (hash(x)=23797514) +3129 train 7.697777 (lr=2.3361e-04) (hash(x)=23601883) +3130 train 7.908652 (lr=2.3345e-04) (hash(x)=31003964) +3131 train 7.618131 (lr=2.3329e-04) (hash(x)=24777273) +3132 train 7.537286 (lr=2.3314e-04) (hash(x)=25403249) +3133 train 7.688043 (lr=2.3298e-04) (hash(x)=28913150) +3134 train 7.750380 (lr=2.3282e-04) (hash(x)=26541508) +3135 train 7.629738 (lr=2.3266e-04) (hash(x)=24113445) +3136 train 7.415004 (lr=2.3250e-04) (hash(x)=25464565) +3137 train 7.686874 (lr=2.3235e-04) (hash(x)=26581432) +3138 train 7.546024 (lr=2.3219e-04) (hash(x)=23074513) +3139 train 7.433839 (lr=2.3203e-04) (hash(x)=23970384) +3140 train 7.614048 (lr=2.3187e-04) (hash(x)=26694495) +3141 train 7.564430 (lr=2.3171e-04) (hash(x)=26883445) +3142 train 7.788381 (lr=2.3156e-04) (hash(x)=28632211) +3143 train 8.079797 (lr=2.3140e-04) (hash(x)=32644465) +3144 train 7.905371 (lr=2.3124e-04) (hash(x)=27490443) +3145 train 7.506467 (lr=2.3108e-04) (hash(x)=23814853) +3146 train 7.967449 (lr=2.3092e-04) (hash(x)=29664236) +3147 train 8.168917 (lr=2.3076e-04) (hash(x)=29951548) +3148 train 7.991512 (lr=2.3061e-04) (hash(x)=28426503) +3149 train 7.398464 (lr=2.3045e-04) (hash(x)=23727657) +3150 val loss 7.6516 +3150 val perplexity 2103.9700 +3150 train 7.506258 (lr=2.3029e-04) (hash(x)=21430659) +3151 train 7.585817 (lr=2.3013e-04) (hash(x)=25829219) +3152 train 7.780955 (lr=2.2997e-04) (hash(x)=29735208) +3153 train 7.793493 (lr=2.2982e-04) (hash(x)=28173447) +3154 train 7.501148 (lr=2.2966e-04) (hash(x)=22909641) +3155 train 7.686096 (lr=2.2950e-04) (hash(x)=20556094) +3156 train 7.484598 (lr=2.2934e-04) (hash(x)=24013769) +3157 train 7.440801 (lr=2.2918e-04) (hash(x)=22525971) +3158 train 7.519833 (lr=2.2902e-04) (hash(x)=25492728) +3159 train 7.561604 (lr=2.2887e-04) (hash(x)=25194550) +3160 train 7.599691 (lr=2.2871e-04) (hash(x)=25610603) +3161 train 7.582193 (lr=2.2855e-04) (hash(x)=23848640) +3162 train 7.442477 (lr=2.2839e-04) (hash(x)=24082226) +3163 train 7.783432 (lr=2.2823e-04) (hash(x)=28482186) +3164 train 7.685018 (lr=2.2808e-04) (hash(x)=27542978) +3165 train 7.491222 (lr=2.2792e-04) (hash(x)=22540954) +3166 train 7.598806 (lr=2.2776e-04) (hash(x)=26103641) +3167 train 7.540634 (lr=2.2760e-04) (hash(x)=25941804) +3168 train 7.536897 (lr=2.2744e-04) (hash(x)=25965921) +3169 train 7.793709 (lr=2.2728e-04) (hash(x)=25631269) +3170 train 7.399700 (lr=2.2713e-04) (hash(x)=23471525) +3171 train 7.614377 (lr=2.2697e-04) (hash(x)=27049208) +3172 train 7.610492 (lr=2.2681e-04) (hash(x)=27074992) +3173 train 7.597581 (lr=2.2665e-04) (hash(x)=25712617) +3174 train 7.561949 (lr=2.2649e-04) (hash(x)=25884917) +3175 train 7.575124 (lr=2.2633e-04) (hash(x)=24075727) +3176 train 7.485845 (lr=2.2618e-04) (hash(x)=23681759) +3177 train 7.735615 (lr=2.2602e-04) (hash(x)=25786577) +3178 train 7.754552 (lr=2.2586e-04) (hash(x)=27307614) +3179 train 7.728940 (lr=2.2570e-04) (hash(x)=25082806) +3180 train 7.727188 (lr=2.2554e-04) (hash(x)=26098308) +3181 train 7.524395 (lr=2.2538e-04) (hash(x)=24080140) +3182 train 7.679074 (lr=2.2523e-04) (hash(x)=26399395) +3183 train 7.733400 (lr=2.2507e-04) (hash(x)=23104539) +3184 train 7.426908 (lr=2.2491e-04) (hash(x)=23356930) +3185 train 7.661096 (lr=2.2475e-04) (hash(x)=27972420) +3186 train 7.359051 (lr=2.2459e-04) (hash(x)=21338924) +3187 train 7.746469 (lr=2.2443e-04) (hash(x)=25351113) +3188 train 7.643510 (lr=2.2428e-04) (hash(x)=26019439) +3189 train 7.697662 (lr=2.2412e-04) (hash(x)=30149312) +3190 train 7.491061 (lr=2.2396e-04) (hash(x)=23028152) +3191 train 7.772232 (lr=2.2380e-04) (hash(x)=23018983) +3192 train 7.436260 (lr=2.2364e-04) (hash(x)=23190787) +3193 train 7.845827 (lr=2.2348e-04) (hash(x)=27798543) +3194 train 7.620201 (lr=2.2333e-04) (hash(x)=25193663) +3195 train 7.454890 (lr=2.2317e-04) (hash(x)=25302106) +3196 train 7.550751 (lr=2.2301e-04) (hash(x)=24325364) +3197 train 7.586545 (lr=2.2285e-04) (hash(x)=25399101) +3198 train 7.437521 (lr=2.2269e-04) (hash(x)=23606439) +3199 train 7.488397 (lr=2.2253e-04) (hash(x)=24422929) +3200 val loss 7.6182 +3200 val perplexity 2034.9937 +3200 train 7.681554 (lr=2.2238e-04) (hash(x)=24760381) +3201 train 7.295120 (lr=2.2222e-04) (hash(x)=23278576) +3202 train 7.556820 (lr=2.2206e-04) (hash(x)=24897511) +3203 train 7.805412 (lr=2.2190e-04) (hash(x)=29052117) +3204 train 7.640369 (lr=2.2174e-04) (hash(x)=25772923) +3205 train 7.760435 (lr=2.2158e-04) (hash(x)=25885977) +3206 train 7.484886 (lr=2.2143e-04) (hash(x)=21985272) +3207 train 7.468366 (lr=2.2127e-04) (hash(x)=23389696) +3208 train 7.380031 (lr=2.2111e-04) (hash(x)=25299042) +3209 train 7.577655 (lr=2.2095e-04) (hash(x)=23703987) +3210 train 7.519795 (lr=2.2079e-04) (hash(x)=23362342) +3211 train 7.576889 (lr=2.2063e-04) (hash(x)=23962503) +3212 train 7.507912 (lr=2.2048e-04) (hash(x)=21216023) +3213 train 7.572202 (lr=2.2032e-04) (hash(x)=25841931) +3214 train 7.495394 (lr=2.2016e-04) (hash(x)=23631428) +3215 train 7.827016 (lr=2.2000e-04) (hash(x)=29102969) +3216 train 7.660228 (lr=2.1984e-04) (hash(x)=25782766) +3217 train 7.828589 (lr=2.1968e-04) (hash(x)=28825867) +3218 train 7.749118 (lr=2.1952e-04) (hash(x)=24609499) +3219 train 7.712409 (lr=2.1937e-04) (hash(x)=24948166) +3220 train 7.738871 (lr=2.1921e-04) (hash(x)=26919841) +3221 train 7.645599 (lr=2.1905e-04) (hash(x)=27236198) +3222 train 7.445494 (lr=2.1889e-04) (hash(x)=24029305) +3223 train 7.894432 (lr=2.1873e-04) (hash(x)=28359859) +3224 train 7.721086 (lr=2.1857e-04) (hash(x)=24886008) +3225 train 7.589790 (lr=2.1842e-04) (hash(x)=27159867) +3226 train 7.776907 (lr=2.1826e-04) (hash(x)=20672023) +3227 train 7.527074 (lr=2.1810e-04) (hash(x)=22360298) +3228 train 7.893600 (lr=2.1794e-04) (hash(x)=27478658) +3229 train 7.678592 (lr=2.1778e-04) (hash(x)=26575886) +3230 train 7.893658 (lr=2.1762e-04) (hash(x)=26890615) +3231 train 7.558395 (lr=2.1747e-04) (hash(x)=24630955) +3232 train 7.820558 (lr=2.1731e-04) (hash(x)=27016054) +3233 train 8.079944 (lr=2.1715e-04) (hash(x)=28444407) +3234 train 7.444576 (lr=2.1699e-04) (hash(x)=24053336) +3235 train 7.758126 (lr=2.1683e-04) (hash(x)=26897402) +3236 train 7.893925 (lr=2.1667e-04) (hash(x)=29451214) +3237 train 7.697472 (lr=2.1652e-04) (hash(x)=27268677) +3238 train 7.781694 (lr=2.1636e-04) (hash(x)=27494000) +3239 train 7.517171 (lr=2.1620e-04) (hash(x)=22969113) +3240 train 7.534807 (lr=2.1604e-04) (hash(x)=21944576) +3241 train 7.703901 (lr=2.1588e-04) (hash(x)=21671079) +3242 train 7.622938 (lr=2.1572e-04) (hash(x)=23912980) +3243 train 7.536371 (lr=2.1557e-04) (hash(x)=25205781) +3244 train 7.561681 (lr=2.1541e-04) (hash(x)=25654244) +3245 train 7.508101 (lr=2.1525e-04) (hash(x)=23335929) +3246 train 7.799889 (lr=2.1509e-04) (hash(x)=27953926) +3247 train 7.590321 (lr=2.1493e-04) (hash(x)=27004415) +3248 train 7.344252 (lr=2.1477e-04) (hash(x)=20471566) +3249 train 7.555417 (lr=2.1462e-04) (hash(x)=25797941) +3250 val loss 7.6297 +3250 val perplexity 2058.4849 +3250 train 7.469663 (lr=2.1446e-04) (hash(x)=21787064) +3251 train 7.435731 (lr=2.1430e-04) (hash(x)=22974875) +3252 train 7.695366 (lr=2.1414e-04) (hash(x)=28431267) +3253 train 7.557029 (lr=2.1398e-04) (hash(x)=25584910) +3254 train 7.377738 (lr=2.1382e-04) (hash(x)=23888922) +3255 train 7.541820 (lr=2.1367e-04) (hash(x)=22265063) +3256 train 7.292380 (lr=2.1351e-04) (hash(x)=21926624) +3257 train 7.486031 (lr=2.1335e-04) (hash(x)=23073191) +3258 train 7.505316 (lr=2.1319e-04) (hash(x)=24409183) +3259 train 7.558173 (lr=2.1303e-04) (hash(x)=23312114) +3260 train 7.426100 (lr=2.1287e-04) (hash(x)=21001289) +3261 train 7.633335 (lr=2.1272e-04) (hash(x)=25514824) +3262 train 7.437344 (lr=2.1256e-04) (hash(x)=22526800) +3263 train 7.823399 (lr=2.1240e-04) (hash(x)=26905990) +3264 train 7.529859 (lr=2.1224e-04) (hash(x)=24469631) +3265 train 7.369201 (lr=2.1208e-04) (hash(x)=21149081) +3266 train 7.567880 (lr=2.1192e-04) (hash(x)=24696215) +3267 train 7.892066 (lr=2.1177e-04) (hash(x)=27089280) +3268 train 7.395909 (lr=2.1161e-04) (hash(x)=23100446) +3269 train 7.624568 (lr=2.1145e-04) (hash(x)=25061229) +3270 train 7.546525 (lr=2.1129e-04) (hash(x)=24337543) +3271 train 7.592148 (lr=2.1113e-04) (hash(x)=24047679) +3272 train 7.659808 (lr=2.1098e-04) (hash(x)=27616773) +3273 train 7.565965 (lr=2.1082e-04) (hash(x)=25315110) +3274 train 7.654410 (lr=2.1066e-04) (hash(x)=28354645) +3275 train 7.490492 (lr=2.1050e-04) (hash(x)=25034684) +3276 train 7.530472 (lr=2.1034e-04) (hash(x)=23550342) +3277 train 8.090452 (lr=2.1018e-04) (hash(x)=28661487) +3278 train 7.734716 (lr=2.1003e-04) (hash(x)=24724622) +3279 train 7.719799 (lr=2.0987e-04) (hash(x)=26905582) +3280 train 7.794968 (lr=2.0971e-04) (hash(x)=26838818) +3281 train 7.572229 (lr=2.0955e-04) (hash(x)=23949017) +3282 train 7.109481 (lr=2.0939e-04) (hash(x)=18846300) +3283 train 7.366128 (lr=2.0924e-04) (hash(x)=21406950) +3284 train 7.308766 (lr=2.0908e-04) (hash(x)=21157696) +3285 train 7.376631 (lr=2.0892e-04) (hash(x)=21440152) +3286 train 7.512596 (lr=2.0876e-04) (hash(x)=26749182) +3287 train 7.671555 (lr=2.0860e-04) (hash(x)=29018970) +3288 train 7.617728 (lr=2.0844e-04) (hash(x)=27577517) +3289 train 7.336576 (lr=2.0829e-04) (hash(x)=19190537) +3290 train 7.347484 (lr=2.0813e-04) (hash(x)=21957991) +3291 train 7.401563 (lr=2.0797e-04) (hash(x)=20853530) +3292 train 7.332983 (lr=2.0781e-04) (hash(x)=22291731) +3293 train 7.419106 (lr=2.0765e-04) (hash(x)=23786853) +3294 train 7.615666 (lr=2.0750e-04) (hash(x)=25740147) +3295 train 7.758809 (lr=2.0734e-04) (hash(x)=24503315) +3296 train 7.450855 (lr=2.0718e-04) (hash(x)=22541728) +3297 train 7.793067 (lr=2.0702e-04) (hash(x)=27067328) +3298 train 7.448613 (lr=2.0686e-04) (hash(x)=22600715) +3299 train 7.485795 (lr=2.0671e-04) (hash(x)=23080074) +3300 val loss 7.6041 +3300 val perplexity 2006.3818 +3300 train 7.279202 (lr=2.0655e-04) (hash(x)=22097758) +3301 train 7.420104 (lr=2.0639e-04) (hash(x)=24105430) +3302 train 7.315197 (lr=2.0623e-04) (hash(x)=23343775) +3303 train 7.525265 (lr=2.0607e-04) (hash(x)=22607537) +3304 train 7.603733 (lr=2.0592e-04) (hash(x)=26501182) +3305 train 7.836336 (lr=2.0576e-04) (hash(x)=26063650) +3306 train 7.489346 (lr=2.0560e-04) (hash(x)=23486602) +3307 train 8.478892 (lr=2.0544e-04) (hash(x)=32179773) +3308 train 11.668060 (lr=2.0529e-04) (hash(x)=66155855) +3309 train 8.871637 (lr=2.0513e-04) (hash(x)=37724427) +3310 train 7.916838 (lr=2.0497e-04) (hash(x)=28861610) +3311 train 7.875553 (lr=2.0481e-04) (hash(x)=26690225) +3312 train 7.529292 (lr=2.0465e-04) (hash(x)=24136450) +3313 train 7.692175 (lr=2.0450e-04) (hash(x)=23702010) +3314 train 7.971992 (lr=2.0434e-04) (hash(x)=28761762) +3315 train 7.397576 (lr=2.0418e-04) (hash(x)=22109609) +3316 train 7.671214 (lr=2.0402e-04) (hash(x)=25168631) +3317 train 7.614745 (lr=2.0386e-04) (hash(x)=24503786) +3318 train 7.664669 (lr=2.0371e-04) (hash(x)=23698606) +3319 train 7.654486 (lr=2.0355e-04) (hash(x)=24226255) +3320 train 8.175403 (lr=2.0339e-04) (hash(x)=31110577) +3321 train 7.657456 (lr=2.0323e-04) (hash(x)=24752754) +3322 train 7.481316 (lr=2.0308e-04) (hash(x)=21135610) +3323 train 7.494169 (lr=2.0292e-04) (hash(x)=23013573) +3324 train 7.330629 (lr=2.0276e-04) (hash(x)=20289715) +3325 train 7.512113 (lr=2.0260e-04) (hash(x)=22700287) +3326 train 7.540400 (lr=2.0245e-04) (hash(x)=21320362) +3327 train 7.630637 (lr=2.0229e-04) (hash(x)=23622702) +3328 train 7.854778 (lr=2.0213e-04) (hash(x)=27435461) +3329 train 7.741627 (lr=2.0197e-04) (hash(x)=25435452) +3330 train 7.653037 (lr=2.0182e-04) (hash(x)=27952557) +3331 train 7.393228 (lr=2.0166e-04) (hash(x)=21517429) +3332 train 7.504853 (lr=2.0150e-04) (hash(x)=24288985) +3333 train 7.582328 (lr=2.0134e-04) (hash(x)=23374788) +3334 train 7.495624 (lr=2.0118e-04) (hash(x)=22042499) +3335 train 7.655070 (lr=2.0103e-04) (hash(x)=23910425) +3336 train 7.984547 (lr=2.0087e-04) (hash(x)=28118508) +3337 train 7.717821 (lr=2.0071e-04) (hash(x)=26737440) +3338 train 7.521951 (lr=2.0055e-04) (hash(x)=24472271) +3339 train 7.492887 (lr=2.0040e-04) (hash(x)=24407484) +3340 train 7.947762 (lr=2.0024e-04) (hash(x)=27908937) +3341 train 8.093860 (lr=2.0008e-04) (hash(x)=29038937) +3342 train 7.881116 (lr=1.9993e-04) (hash(x)=24802580) +3343 train 7.875925 (lr=1.9977e-04) (hash(x)=27213318) +3344 train 7.852036 (lr=1.9961e-04) (hash(x)=28693458) +3345 train 7.297226 (lr=1.9945e-04) (hash(x)=20332324) +3346 train 7.608232 (lr=1.9930e-04) (hash(x)=26726007) +3347 train 7.538406 (lr=1.9914e-04) (hash(x)=25524191) +3348 train 7.600595 (lr=1.9898e-04) (hash(x)=25553293) +3349 train 7.597383 (lr=1.9882e-04) (hash(x)=25614848) +3350 val loss 7.6086 +3350 val perplexity 2015.4010 +3350 train 7.586722 (lr=1.9867e-04) (hash(x)=25747903) +3351 train 7.630146 (lr=1.9851e-04) (hash(x)=26701577) +3352 train 7.411195 (lr=1.9835e-04) (hash(x)=21964135) +3353 train 7.493042 (lr=1.9819e-04) (hash(x)=24461007) +3354 train 7.514977 (lr=1.9804e-04) (hash(x)=25818495) +3355 train 7.387592 (lr=1.9788e-04) (hash(x)=22091266) +3356 train 7.405923 (lr=1.9772e-04) (hash(x)=24476213) +3357 train 7.694082 (lr=1.9757e-04) (hash(x)=24500423) +3358 train 7.456055 (lr=1.9741e-04) (hash(x)=21754841) +3359 train 7.772446 (lr=1.9725e-04) (hash(x)=26216794) +3360 train 7.379512 (lr=1.9709e-04) (hash(x)=24267249) +3361 train 7.574001 (lr=1.9694e-04) (hash(x)=23143515) +3362 train 7.516946 (lr=1.9678e-04) (hash(x)=24120302) +3363 train 7.522562 (lr=1.9662e-04) (hash(x)=20817340) +3364 train 7.605431 (lr=1.9647e-04) (hash(x)=22285847) +3365 train 7.853302 (lr=1.9631e-04) (hash(x)=28151597) +3366 train 8.324681 (lr=1.9615e-04) (hash(x)=31593285) +3367 train 7.957855 (lr=1.9599e-04) (hash(x)=27579623) +3368 train 7.645608 (lr=1.9584e-04) (hash(x)=24995988) +3369 train 7.294110 (lr=1.9568e-04) (hash(x)=22166810) +3370 train 7.488452 (lr=1.9552e-04) (hash(x)=23948298) +3371 train 7.367679 (lr=1.9537e-04) (hash(x)=21532187) +3372 train 7.495677 (lr=1.9521e-04) (hash(x)=23571652) +3373 train 7.818840 (lr=1.9505e-04) (hash(x)=26911513) +3374 train 7.709721 (lr=1.9490e-04) (hash(x)=24011329) +3375 train 7.912449 (lr=1.9474e-04) (hash(x)=26086198) +3376 train 7.521326 (lr=1.9458e-04) (hash(x)=22844402) +3377 train 7.521690 (lr=1.9443e-04) (hash(x)=21817762) +3378 train 7.624910 (lr=1.9427e-04) (hash(x)=23903232) +3379 train 7.541882 (lr=1.9411e-04) (hash(x)=23911729) +3380 train 7.491560 (lr=1.9396e-04) (hash(x)=24485288) +3381 train 7.673749 (lr=1.9380e-04) (hash(x)=27955492) +3382 train 7.447598 (lr=1.9364e-04) (hash(x)=25884586) +3383 train 7.600086 (lr=1.9349e-04) (hash(x)=24863441) +3384 train 7.790801 (lr=1.9333e-04) (hash(x)=22045992) +3385 train 7.952445 (lr=1.9317e-04) (hash(x)=29174796) +3386 train 8.211693 (lr=1.9302e-04) (hash(x)=32589942) +3387 train 7.899467 (lr=1.9286e-04) (hash(x)=28856978) +3388 train 7.465524 (lr=1.9270e-04) (hash(x)=21667904) +3389 train 7.394583 (lr=1.9255e-04) (hash(x)=23431801) +3390 train 7.302803 (lr=1.9239e-04) (hash(x)=20877285) +3391 train 7.492558 (lr=1.9223e-04) (hash(x)=25236385) +3392 train 7.612044 (lr=1.9208e-04) (hash(x)=25373071) +3393 train 7.758134 (lr=1.9192e-04) (hash(x)=25713464) +3394 train 7.733314 (lr=1.9176e-04) (hash(x)=25713475) +3395 train 7.687768 (lr=1.9161e-04) (hash(x)=24278687) +3396 train 7.801169 (lr=1.9145e-04) (hash(x)=27491349) +3397 train 7.699730 (lr=1.9129e-04) (hash(x)=24513692) +3398 train 7.784131 (lr=1.9114e-04) (hash(x)=26853415) +3399 train 7.817062 (lr=1.9098e-04) (hash(x)=25330803) +3400 val loss 7.6128 +3400 val perplexity 2023.8641 +3400 train 7.790269 (lr=1.9083e-04) (hash(x)=26578066) +3401 train 7.800073 (lr=1.9067e-04) (hash(x)=26811522) +3402 train 7.544026 (lr=1.9051e-04) (hash(x)=25611092) +3403 train 7.423755 (lr=1.9036e-04) (hash(x)=21568545) +3404 train 7.382677 (lr=1.9020e-04) (hash(x)=22756484) +3405 train 7.643800 (lr=1.9004e-04) (hash(x)=27927608) +3406 train 7.687735 (lr=1.8989e-04) (hash(x)=27497018) +3407 train 7.465792 (lr=1.8973e-04) (hash(x)=22508532) +3408 train 7.856317 (lr=1.8958e-04) (hash(x)=26673287) +3409 train 7.555014 (lr=1.8942e-04) (hash(x)=23675869) +3410 train 7.476700 (lr=1.8926e-04) (hash(x)=25496948) +3411 train 7.470452 (lr=1.8911e-04) (hash(x)=24850662) +3412 train 7.702512 (lr=1.8895e-04) (hash(x)=29790167) +3413 train 7.440118 (lr=1.8880e-04) (hash(x)=24193434) +3414 train 7.739616 (lr=1.8864e-04) (hash(x)=25310919) +3415 train 7.449574 (lr=1.8848e-04) (hash(x)=21799261) +3416 train 7.770445 (lr=1.8833e-04) (hash(x)=26620074) +3417 train 7.674699 (lr=1.8817e-04) (hash(x)=26719309) +3418 train 7.935082 (lr=1.8802e-04) (hash(x)=23190530) +3419 train 7.629785 (lr=1.8786e-04) (hash(x)=24884891) +3420 train 7.630441 (lr=1.8770e-04) (hash(x)=25545849) +3421 train 7.626447 (lr=1.8755e-04) (hash(x)=26021405) +3422 train 7.520200 (lr=1.8739e-04) (hash(x)=23887343) +3423 train 7.617774 (lr=1.8724e-04) (hash(x)=26311168) +3424 train 7.371996 (lr=1.8708e-04) (hash(x)=21051541) +3425 train 7.670821 (lr=1.8693e-04) (hash(x)=23553179) +3426 train 7.599823 (lr=1.8677e-04) (hash(x)=24345540) +3427 train 7.973100 (lr=1.8661e-04) (hash(x)=27549895) +3428 train 7.409326 (lr=1.8646e-04) (hash(x)=22559753) +3429 train 7.375357 (lr=1.8630e-04) (hash(x)=21647642) +3430 train 7.274137 (lr=1.8615e-04) (hash(x)=20130901) +3431 train 7.405643 (lr=1.8599e-04) (hash(x)=20977430) +3432 train 7.346571 (lr=1.8584e-04) (hash(x)=21356429) +3433 train 7.345651 (lr=1.8568e-04) (hash(x)=23494380) +3434 train 7.616651 (lr=1.8552e-04) (hash(x)=23805501) +3435 train 7.556998 (lr=1.8537e-04) (hash(x)=23448855) +3436 train 7.965772 (lr=1.8521e-04) (hash(x)=26000319) +3437 train 7.400936 (lr=1.8506e-04) (hash(x)=21760032) +3438 train 7.508066 (lr=1.8490e-04) (hash(x)=24424886) +3439 train 7.840835 (lr=1.8475e-04) (hash(x)=26941617) +3440 train 7.825504 (lr=1.8459e-04) (hash(x)=26798528) +3441 train 7.851163 (lr=1.8444e-04) (hash(x)=27464193) +3442 train 7.756439 (lr=1.8428e-04) (hash(x)=25649118) +3443 train 7.815708 (lr=1.8413e-04) (hash(x)=26953192) +3444 train 7.642294 (lr=1.8397e-04) (hash(x)=22224958) +3445 train 7.468914 (lr=1.8382e-04) (hash(x)=24044587) +3446 train 7.868421 (lr=1.8366e-04) (hash(x)=29584466) +3447 train 8.159633 (lr=1.8351e-04) (hash(x)=30008957) +3448 train 8.151006 (lr=1.8335e-04) (hash(x)=26059290) +3449 train 7.248398 (lr=1.8320e-04) (hash(x)=19733965) +3450 val loss 7.6288 +3450 val perplexity 2056.5627 +3450 train 7.676034 (lr=1.8304e-04) (hash(x)=23960200) +3451 train 7.844135 (lr=1.8289e-04) (hash(x)=27069893) +3452 train 7.514679 (lr=1.8273e-04) (hash(x)=23947772) +3453 train 7.608578 (lr=1.8258e-04) (hash(x)=22707406) +3454 train 7.766284 (lr=1.8242e-04) (hash(x)=27832550) +3455 train 7.631418 (lr=1.8227e-04) (hash(x)=27125962) +3456 train 7.545026 (lr=1.8211e-04) (hash(x)=24510254) +3457 train 7.357381 (lr=1.8196e-04) (hash(x)=23545652) +3458 train 7.414999 (lr=1.8180e-04) (hash(x)=23554751) +3459 train 7.539445 (lr=1.8165e-04) (hash(x)=23341415) +3460 train 7.565149 (lr=1.8149e-04) (hash(x)=21784583) +3461 train 7.486818 (lr=1.8134e-04) (hash(x)=22214769) +3462 train 7.521430 (lr=1.8118e-04) (hash(x)=24206922) +3463 train 7.612646 (lr=1.8103e-04) (hash(x)=25888358) +3464 train 7.493163 (lr=1.8087e-04) (hash(x)=22689666) +3465 train 7.571679 (lr=1.8072e-04) (hash(x)=24918697) +3466 train 7.849395 (lr=1.8056e-04) (hash(x)=28237214) +3467 train 7.902084 (lr=1.8041e-04) (hash(x)=26761645) +3468 train 7.626210 (lr=1.8026e-04) (hash(x)=26979307) +3469 train 7.491737 (lr=1.8010e-04) (hash(x)=23553754) +3470 train 7.696089 (lr=1.7995e-04) (hash(x)=25256849) +3471 train 7.607072 (lr=1.7979e-04) (hash(x)=21725719) +3472 train 7.603751 (lr=1.7964e-04) (hash(x)=24897801) +3473 train 7.561244 (lr=1.7948e-04) (hash(x)=26175307) +3474 train 7.383890 (lr=1.7933e-04) (hash(x)=23309218) +3475 train 7.462869 (lr=1.7917e-04) (hash(x)=25746493) +3476 train 7.736900 (lr=1.7902e-04) (hash(x)=27169613) +3477 train 7.418480 (lr=1.7887e-04) (hash(x)=22937341) +3478 train 7.533362 (lr=1.7871e-04) (hash(x)=24250636) +3479 train 7.406331 (lr=1.7856e-04) (hash(x)=21669704) +3480 train 7.466068 (lr=1.7840e-04) (hash(x)=24431839) +3481 train 7.430471 (lr=1.7825e-04) (hash(x)=22763387) +3482 train 7.446659 (lr=1.7810e-04) (hash(x)=20489446) +3483 train 7.399171 (lr=1.7794e-04) (hash(x)=21167493) +3484 train 7.425835 (lr=1.7779e-04) (hash(x)=23465087) +3485 train 7.657758 (lr=1.7763e-04) (hash(x)=26175023) +3486 train 7.731798 (lr=1.7748e-04) (hash(x)=24986207) +3487 train 7.378970 (lr=1.7733e-04) (hash(x)=23166993) +3488 train 7.520647 (lr=1.7717e-04) (hash(x)=25281216) +3489 train 7.362791 (lr=1.7702e-04) (hash(x)=21824285) +3490 train 7.428127 (lr=1.7686e-04) (hash(x)=22352750) +3491 train 7.608604 (lr=1.7671e-04) (hash(x)=23947208) +3492 train 7.643059 (lr=1.7656e-04) (hash(x)=26257363) +3493 train 7.557881 (lr=1.7640e-04) (hash(x)=25103214) +3494 train 7.560112 (lr=1.7625e-04) (hash(x)=25267583) +3495 train 7.640480 (lr=1.7610e-04) (hash(x)=26235974) +3496 train 7.763960 (lr=1.7594e-04) (hash(x)=26430769) +3497 train 8.097030 (lr=1.7579e-04) (hash(x)=28282027) +3498 train 8.053090 (lr=1.7563e-04) (hash(x)=28386462) +3499 train 7.716600 (lr=1.7548e-04) (hash(x)=29822604) +3500 val loss 7.6034 +3500 val perplexity 2005.0142 +3500 train 7.626303 (lr=1.7533e-04) (hash(x)=29225386) +3501 train 7.539689 (lr=1.7517e-04) (hash(x)=25249294) +3502 train 7.381425 (lr=1.7502e-04) (hash(x)=20020741) +3503 train 7.546079 (lr=1.7487e-04) (hash(x)=25426430) +3504 train 7.523337 (lr=1.7471e-04) (hash(x)=25720411) +3505 train 7.498598 (lr=1.7456e-04) (hash(x)=25602639) +3506 train 7.653702 (lr=1.7441e-04) (hash(x)=26724388) +3507 train 7.599173 (lr=1.7425e-04) (hash(x)=26043735) +3508 train 7.565844 (lr=1.7410e-04) (hash(x)=24955163) +3509 train 7.507028 (lr=1.7395e-04) (hash(x)=20936107) +3510 train 7.438309 (lr=1.7380e-04) (hash(x)=20317378) +3511 train 7.492021 (lr=1.7364e-04) (hash(x)=22966314) +3512 train 7.804750 (lr=1.7349e-04) (hash(x)=25870930) +3513 train 7.606370 (lr=1.7334e-04) (hash(x)=24656635) +3514 train 7.792709 (lr=1.7318e-04) (hash(x)=28576810) +3515 train 7.768280 (lr=1.7303e-04) (hash(x)=27944619) +3516 train 7.641455 (lr=1.7288e-04) (hash(x)=27421509) +3517 train 7.405023 (lr=1.7272e-04) (hash(x)=20844620) +3518 train 7.582088 (lr=1.7257e-04) (hash(x)=28569406) +3519 train 7.510110 (lr=1.7242e-04) (hash(x)=23448505) +3520 train 7.591228 (lr=1.7227e-04) (hash(x)=24852577) +3521 train 7.538425 (lr=1.7211e-04) (hash(x)=23963103) +3522 train 7.463477 (lr=1.7196e-04) (hash(x)=24816516) +3523 train 7.583962 (lr=1.7181e-04) (hash(x)=24205942) +3524 train 7.264872 (lr=1.7166e-04) (hash(x)=20988660) +3525 train 7.329366 (lr=1.7150e-04) (hash(x)=21631366) +3526 train 7.532014 (lr=1.7135e-04) (hash(x)=23499370) +3527 train 7.621349 (lr=1.7120e-04) (hash(x)=26330693) +3528 train 7.333159 (lr=1.7105e-04) (hash(x)=23937176) +3529 train 7.525968 (lr=1.7089e-04) (hash(x)=27345885) +3530 train 7.359190 (lr=1.7074e-04) (hash(x)=21104610) +3531 train 7.630126 (lr=1.7059e-04) (hash(x)=24844466) +3532 train 7.152390 (lr=1.7044e-04) (hash(x)=21055483) +3533 train 7.341976 (lr=1.7028e-04) (hash(x)=23229414) +3534 train 7.602319 (lr=1.7013e-04) (hash(x)=26676920) +3535 train 7.618318 (lr=1.6998e-04) (hash(x)=29550596) +3536 train 7.485813 (lr=1.6983e-04) (hash(x)=22231942) +3537 train 7.553845 (lr=1.6968e-04) (hash(x)=25843852) +3538 train 7.669837 (lr=1.6952e-04) (hash(x)=27110533) +3539 train 7.316022 (lr=1.6937e-04) (hash(x)=20506540) +3540 train 7.389252 (lr=1.6922e-04) (hash(x)=21599346) +3541 train 7.603761 (lr=1.6907e-04) (hash(x)=26395519) +3542 train 7.543903 (lr=1.6892e-04) (hash(x)=25892512) +3543 train 7.569467 (lr=1.6876e-04) (hash(x)=22124892) +3544 train 7.436909 (lr=1.6861e-04) (hash(x)=21882567) +3545 train 7.513172 (lr=1.6846e-04) (hash(x)=24316212) +3546 train 7.488235 (lr=1.6831e-04) (hash(x)=24296310) +3547 train 7.533109 (lr=1.6816e-04) (hash(x)=24867036) +3548 train 7.488140 (lr=1.6800e-04) (hash(x)=23351896) +3549 train 7.536350 (lr=1.6785e-04) (hash(x)=21576408) +3550 val loss 7.5928 +3550 val perplexity 1983.9100 +3550 train 7.558376 (lr=1.6770e-04) (hash(x)=26377338) +3551 train 7.481113 (lr=1.6755e-04) (hash(x)=25607640) +3552 train 7.768681 (lr=1.6740e-04) (hash(x)=27619776) +3553 train 7.348215 (lr=1.6725e-04) (hash(x)=23454533) +3554 train 7.478093 (lr=1.6710e-04) (hash(x)=22542519) +3555 train 7.534256 (lr=1.6694e-04) (hash(x)=26176930) +3556 train 7.297221 (lr=1.6679e-04) (hash(x)=22815181) +3557 train 7.234056 (lr=1.6664e-04) (hash(x)=21821757) +3558 train 7.299259 (lr=1.6649e-04) (hash(x)=23988293) +3559 train 7.219372 (lr=1.6634e-04) (hash(x)=23795894) +3560 train 7.210010 (lr=1.6619e-04) (hash(x)=22898969) +3561 train 7.209473 (lr=1.6604e-04) (hash(x)=21510825) +3562 train 7.161440 (lr=1.6589e-04) (hash(x)=22499317) +3563 train 7.186807 (lr=1.6573e-04) (hash(x)=23756298) +3564 train 7.211470 (lr=1.6558e-04) (hash(x)=23964512) +3565 train 7.358521 (lr=1.6543e-04) (hash(x)=23262803) +3566 train 7.531188 (lr=1.6528e-04) (hash(x)=23347279) +3567 train 7.476373 (lr=1.6513e-04) (hash(x)=24165449) +3568 train 7.526012 (lr=1.6498e-04) (hash(x)=25503946) +3569 train 7.550398 (lr=1.6483e-04) (hash(x)=26532839) +3570 train 7.226373 (lr=1.6468e-04) (hash(x)=21889816) +3571 train 7.490822 (lr=1.6453e-04) (hash(x)=26643739) +3572 train 7.683881 (lr=1.6438e-04) (hash(x)=26826130) +3573 train 7.636041 (lr=1.6423e-04) (hash(x)=25810624) +3574 train 7.533199 (lr=1.6408e-04) (hash(x)=23080331) +3575 train 7.655369 (lr=1.6393e-04) (hash(x)=24697756) +3576 train 7.504698 (lr=1.6377e-04) (hash(x)=25158900) +3577 train 7.689321 (lr=1.6362e-04) (hash(x)=25793633) +3578 train 7.328426 (lr=1.6347e-04) (hash(x)=21468493) +3579 train 7.571575 (lr=1.6332e-04) (hash(x)=24431101) +3580 train 7.472982 (lr=1.6317e-04) (hash(x)=27314357) +3581 train 7.625331 (lr=1.6302e-04) (hash(x)=26286249) +3582 train 7.854444 (lr=1.6287e-04) (hash(x)=25954856) +3583 train 7.803993 (lr=1.6272e-04) (hash(x)=27218899) +3584 train 7.529857 (lr=1.6257e-04) (hash(x)=24249114) +3585 train 7.452734 (lr=1.6242e-04) (hash(x)=23659934) +3586 train 7.491771 (lr=1.6227e-04) (hash(x)=25995600) +3587 train 7.492147 (lr=1.6212e-04) (hash(x)=29462219) +3588 train 7.554169 (lr=1.6197e-04) (hash(x)=23346714) +3589 train 7.639961 (lr=1.6182e-04) (hash(x)=27168432) +3590 train 7.542778 (lr=1.6167e-04) (hash(x)=23954240) +3591 train 7.473731 (lr=1.6152e-04) (hash(x)=24748522) +3592 train 7.645264 (lr=1.6137e-04) (hash(x)=24887007) +3593 train 7.436765 (lr=1.6122e-04) (hash(x)=25539383) +3594 train 7.360813 (lr=1.6107e-04) (hash(x)=20104613) +3595 train 7.755930 (lr=1.6092e-04) (hash(x)=24843486) +3596 train 7.767684 (lr=1.6077e-04) (hash(x)=24357864) +3597 train 7.528896 (lr=1.6062e-04) (hash(x)=23873745) +3598 train 7.677765 (lr=1.6048e-04) (hash(x)=25142829) +3599 train 7.731409 (lr=1.6033e-04) (hash(x)=24965317) +3600 val loss 7.6035 +3600 val perplexity 2005.1796 +3600 train 7.259371 (lr=1.6018e-04) (hash(x)=18505205) +3601 train 7.335608 (lr=1.6003e-04) (hash(x)=23632877) +3602 train 7.507274 (lr=1.5988e-04) (hash(x)=23704554) +3603 train 7.615034 (lr=1.5973e-04) (hash(x)=26584754) +3604 train 7.350475 (lr=1.5958e-04) (hash(x)=20667709) +3605 train 7.254297 (lr=1.5943e-04) (hash(x)=20573248) +3606 train 7.344100 (lr=1.5928e-04) (hash(x)=23998997) +3607 train 7.427904 (lr=1.5913e-04) (hash(x)=22031210) +3608 train 7.539624 (lr=1.5898e-04) (hash(x)=24124536) +3609 train 7.419571 (lr=1.5883e-04) (hash(x)=22650144) +3610 train 7.575168 (lr=1.5868e-04) (hash(x)=23796998) +3611 train 7.610992 (lr=1.5854e-04) (hash(x)=24860582) +3612 train 7.170925 (lr=1.5839e-04) (hash(x)=20591300) +3613 train 7.533516 (lr=1.5824e-04) (hash(x)=23447130) +3614 train 7.545808 (lr=1.5809e-04) (hash(x)=26237963) +3615 train 7.613782 (lr=1.5794e-04) (hash(x)=25877990) +3616 train 7.427682 (lr=1.5779e-04) (hash(x)=24808003) +3617 train 7.594694 (lr=1.5764e-04) (hash(x)=24103543) +3618 train 7.511759 (lr=1.5749e-04) (hash(x)=24877184) +3619 train 7.623343 (lr=1.5735e-04) (hash(x)=24970646) +3620 train 7.538324 (lr=1.5720e-04) (hash(x)=25764524) +3621 train 7.654624 (lr=1.5705e-04) (hash(x)=25313591) +3622 train 7.549216 (lr=1.5690e-04) (hash(x)=23260940) +3623 train 7.504591 (lr=1.5675e-04) (hash(x)=24382381) +3624 train 7.461302 (lr=1.5660e-04) (hash(x)=24618902) +3625 train 7.588024 (lr=1.5646e-04) (hash(x)=25074871) +3626 train 7.348559 (lr=1.5631e-04) (hash(x)=24472251) +3627 train 7.470413 (lr=1.5616e-04) (hash(x)=25221746) +3628 train 7.557083 (lr=1.5601e-04) (hash(x)=27448790) +3629 train 7.531879 (lr=1.5586e-04) (hash(x)=25221431) +3630 train 7.215003 (lr=1.5572e-04) (hash(x)=22034366) +3631 train 7.627346 (lr=1.5557e-04) (hash(x)=24551999) +3632 train 7.421926 (lr=1.5542e-04) (hash(x)=24330217) +3633 train 7.483635 (lr=1.5527e-04) (hash(x)=22792380) +3634 train 7.328078 (lr=1.5512e-04) (hash(x)=22393767) +3635 train 7.642351 (lr=1.5498e-04) (hash(x)=28151378) +3636 train 7.612418 (lr=1.5483e-04) (hash(x)=26999341) +3637 train 7.671085 (lr=1.5468e-04) (hash(x)=27251870) +3638 train 7.546533 (lr=1.5453e-04) (hash(x)=23439462) +3639 train 7.621236 (lr=1.5439e-04) (hash(x)=25765516) +3640 train 7.544232 (lr=1.5424e-04) (hash(x)=24720171) +3641 train 7.718038 (lr=1.5409e-04) (hash(x)=23927187) +3642 train 7.498575 (lr=1.5394e-04) (hash(x)=23879561) +3643 train 7.571549 (lr=1.5380e-04) (hash(x)=25630696) +3644 train 7.255184 (lr=1.5365e-04) (hash(x)=22030016) +3645 train 7.706024 (lr=1.5350e-04) (hash(x)=28781600) +3646 train 7.633636 (lr=1.5335e-04) (hash(x)=26668019) +3647 train 7.655220 (lr=1.5321e-04) (hash(x)=25204247) +3648 train 7.926090 (lr=1.5306e-04) (hash(x)=31261394) +3649 train 7.724639 (lr=1.5291e-04) (hash(x)=26193103) +3650 val loss 7.5996 +3650 val perplexity 1997.3660 +3650 train 7.517679 (lr=1.5277e-04) (hash(x)=23872456) +3651 train 7.421310 (lr=1.5262e-04) (hash(x)=26326447) +3652 train 7.465683 (lr=1.5247e-04) (hash(x)=26449631) +3653 train 7.704219 (lr=1.5233e-04) (hash(x)=26373461) +3654 train 7.572464 (lr=1.5218e-04) (hash(x)=24882768) +3655 train 7.673223 (lr=1.5203e-04) (hash(x)=26321813) +3656 train 7.672185 (lr=1.5189e-04) (hash(x)=27056428) +3657 train 7.622325 (lr=1.5174e-04) (hash(x)=24583976) +3658 train 7.363846 (lr=1.5159e-04) (hash(x)=22259624) +3659 train 7.553258 (lr=1.5145e-04) (hash(x)=25400396) +3660 train 7.345543 (lr=1.5130e-04) (hash(x)=21095507) +3661 train 7.605042 (lr=1.5115e-04) (hash(x)=27231042) +3662 train 7.420993 (lr=1.5101e-04) (hash(x)=27292771) +3663 train 7.826828 (lr=1.5086e-04) (hash(x)=25528323) +3664 train 7.570423 (lr=1.5071e-04) (hash(x)=24374502) +3665 train 7.410278 (lr=1.5057e-04) (hash(x)=22463800) +3666 train 7.670573 (lr=1.5042e-04) (hash(x)=25413960) +3667 train 7.550581 (lr=1.5028e-04) (hash(x)=24035353) +3668 train 7.546797 (lr=1.5013e-04) (hash(x)=24815852) +3669 train 7.399512 (lr=1.4998e-04) (hash(x)=22995856) +3670 train 7.648108 (lr=1.4984e-04) (hash(x)=28128238) +3671 train 7.642745 (lr=1.4969e-04) (hash(x)=25114729) +3672 train 7.364494 (lr=1.4955e-04) (hash(x)=19337726) +3673 train 7.797378 (lr=1.4940e-04) (hash(x)=26674420) +3674 train 7.367788 (lr=1.4926e-04) (hash(x)=22553270) +3675 train 7.249192 (lr=1.4911e-04) (hash(x)=21634962) +3676 train 7.443906 (lr=1.4896e-04) (hash(x)=23362669) +3677 train 7.504636 (lr=1.4882e-04) (hash(x)=24781824) +3678 train 7.169535 (lr=1.4867e-04) (hash(x)=17909688) +3679 train 8.258620 (lr=1.4853e-04) (hash(x)=31341964) +3680 train 7.826283 (lr=1.4838e-04) (hash(x)=29071335) +3681 train 7.239645 (lr=1.4824e-04) (hash(x)=19486640) +3682 train 7.410427 (lr=1.4809e-04) (hash(x)=24301133) +3683 train 7.574026 (lr=1.4795e-04) (hash(x)=29172813) +3684 train 8.034016 (lr=1.4780e-04) (hash(x)=31501337) +3685 train 7.657956 (lr=1.4766e-04) (hash(x)=27377840) +3686 train 7.522093 (lr=1.4751e-04) (hash(x)=24499761) +3687 train 7.516023 (lr=1.4737e-04) (hash(x)=26326024) +3688 train 7.565141 (lr=1.4722e-04) (hash(x)=23179996) +3689 train 7.280935 (lr=1.4708e-04) (hash(x)=20963675) +3690 train 7.472784 (lr=1.4693e-04) (hash(x)=26650521) +3691 train 7.364502 (lr=1.4679e-04) (hash(x)=20140071) +3692 train 7.308905 (lr=1.4664e-04) (hash(x)=21734340) +3693 train 7.497660 (lr=1.4650e-04) (hash(x)=23848422) +3694 train 7.559711 (lr=1.4635e-04) (hash(x)=27499953) +3695 train 7.790845 (lr=1.4621e-04) (hash(x)=27273311) +3696 train 7.684373 (lr=1.4606e-04) (hash(x)=25741091) +3697 train 7.469659 (lr=1.4592e-04) (hash(x)=22010794) +3698 train 7.542424 (lr=1.4578e-04) (hash(x)=24616138) +3699 train 7.787103 (lr=1.4563e-04) (hash(x)=31276487) +3700 val loss 7.5902 +3700 val perplexity 1978.7402 +3700 train 7.569539 (lr=1.4549e-04) (hash(x)=24042922) +3701 train 7.394306 (lr=1.4534e-04) (hash(x)=21985431) +3702 train 7.244128 (lr=1.4520e-04) (hash(x)=21336316) +3703 train 7.120865 (lr=1.4506e-04) (hash(x)=17819313) +3704 train 7.657594 (lr=1.4491e-04) (hash(x)=27033851) +3705 train 7.704123 (lr=1.4477e-04) (hash(x)=27260043) +3706 train 7.719199 (lr=1.4462e-04) (hash(x)=26847649) +3707 train 7.639228 (lr=1.4448e-04) (hash(x)=25843618) +3708 train 7.642235 (lr=1.4434e-04) (hash(x)=25828009) +3709 train 7.670905 (lr=1.4419e-04) (hash(x)=24960960) +3710 train 7.828106 (lr=1.4405e-04) (hash(x)=24852741) +3711 train 7.335373 (lr=1.4390e-04) (hash(x)=23769243) +3712 train 7.540459 (lr=1.4376e-04) (hash(x)=27110690) +3713 train 7.507646 (lr=1.4362e-04) (hash(x)=22817285) +3714 train 7.284580 (lr=1.4347e-04) (hash(x)=19618355) +3715 train 7.774512 (lr=1.4333e-04) (hash(x)=28731298) +3716 train 7.489496 (lr=1.4319e-04) (hash(x)=23091196) +3717 train 8.014211 (lr=1.4304e-04) (hash(x)=28825233) +3718 train 7.594283 (lr=1.4290e-04) (hash(x)=25778506) +3719 train 7.533817 (lr=1.4276e-04) (hash(x)=23788738) +3720 train 7.072621 (lr=1.4262e-04) (hash(x)=16684794) +3721 train 7.482091 (lr=1.4247e-04) (hash(x)=24755102) +3722 train 7.675812 (lr=1.4233e-04) (hash(x)=26839238) +3723 train 7.701614 (lr=1.4219e-04) (hash(x)=26599031) +3724 train 7.599326 (lr=1.4204e-04) (hash(x)=25945650) +3725 train 7.711034 (lr=1.4190e-04) (hash(x)=27434751) +3726 train 7.543546 (lr=1.4176e-04) (hash(x)=24814591) +3727 train 7.475667 (lr=1.4162e-04) (hash(x)=24818744) +3728 train 7.573133 (lr=1.4147e-04) (hash(x)=25649817) +3729 train 7.489113 (lr=1.4133e-04) (hash(x)=24987021) +3730 train 7.466305 (lr=1.4119e-04) (hash(x)=24667779) +3731 train 7.402505 (lr=1.4105e-04) (hash(x)=23053871) +3732 train 7.135312 (lr=1.4090e-04) (hash(x)=19264856) +3733 train 7.584540 (lr=1.4076e-04) (hash(x)=23718117) +3734 train 7.372443 (lr=1.4062e-04) (hash(x)=21904779) +3735 train 7.632038 (lr=1.4048e-04) (hash(x)=23472795) +3736 train 7.853203 (lr=1.4033e-04) (hash(x)=24650440) +3737 train 7.872354 (lr=1.4019e-04) (hash(x)=25243911) +3738 train 8.218820 (lr=1.4005e-04) (hash(x)=30423616) +3739 train 8.216818 (lr=1.3991e-04) (hash(x)=29748225) +3740 train 7.639238 (lr=1.3977e-04) (hash(x)=25180684) +3741 train 7.629178 (lr=1.3963e-04) (hash(x)=25619065) +3742 train 8.011237 (lr=1.3948e-04) (hash(x)=31847002) +3743 train 7.584704 (lr=1.3934e-04) (hash(x)=24991832) +3744 train 7.555473 (lr=1.3920e-04) (hash(x)=23836263) +3745 train 7.450387 (lr=1.3906e-04) (hash(x)=24651902) +3746 train 7.956628 (lr=1.3892e-04) (hash(x)=22936538) +3747 train 7.942962 (lr=1.3878e-04) (hash(x)=23681545) +3748 train 7.414349 (lr=1.3863e-04) (hash(x)=21140825) +3749 train 7.516992 (lr=1.3849e-04) (hash(x)=24932453) +3750 val loss 7.5910 +3750 val perplexity 1980.3458 +3750 train 7.666736 (lr=1.3835e-04) (hash(x)=25919062) +3751 train 7.477268 (lr=1.3821e-04) (hash(x)=22424066) +3752 train 7.527360 (lr=1.3807e-04) (hash(x)=23542210) +3753 train 7.464467 (lr=1.3793e-04) (hash(x)=22113561) +3754 train 7.366982 (lr=1.3779e-04) (hash(x)=24892794) +3755 train 7.770894 (lr=1.3765e-04) (hash(x)=25262748) +3756 train 7.545015 (lr=1.3751e-04) (hash(x)=24477975) +3757 train 7.360472 (lr=1.3737e-04) (hash(x)=24681189) +3758 train 7.438592 (lr=1.3722e-04) (hash(x)=24052603) +3759 train 7.646131 (lr=1.3708e-04) (hash(x)=26474878) +3760 train 8.101079 (lr=1.3694e-04) (hash(x)=28228836) +3761 train 7.662454 (lr=1.3680e-04) (hash(x)=23832522) +3762 train 7.624271 (lr=1.3666e-04) (hash(x)=24416789) +3763 train 7.423761 (lr=1.3652e-04) (hash(x)=23930593) +3764 train 7.547127 (lr=1.3638e-04) (hash(x)=23895092) +3765 train 7.407246 (lr=1.3624e-04) (hash(x)=26865287) +3766 train 7.457461 (lr=1.3610e-04) (hash(x)=21330722) +3767 train 7.497768 (lr=1.3596e-04) (hash(x)=22704349) +3768 train 7.588658 (lr=1.3582e-04) (hash(x)=27740886) +3769 train 7.580072 (lr=1.3568e-04) (hash(x)=24935936) +3770 train 7.459722 (lr=1.3554e-04) (hash(x)=23497487) +3771 train 7.645344 (lr=1.3540e-04) (hash(x)=24801048) +3772 train 7.456673 (lr=1.3526e-04) (hash(x)=25490486) +3773 train 7.479486 (lr=1.3512e-04) (hash(x)=21166466) +3774 train 7.656418 (lr=1.3498e-04) (hash(x)=23225244) +3775 train 7.467671 (lr=1.3484e-04) (hash(x)=22293673) +3776 train 7.745691 (lr=1.3470e-04) (hash(x)=25700016) +3777 train 7.334440 (lr=1.3456e-04) (hash(x)=22370207) +3778 train 7.628772 (lr=1.3442e-04) (hash(x)=25224849) +3779 train 7.521239 (lr=1.3429e-04) (hash(x)=23311934) +3780 train 7.340265 (lr=1.3415e-04) (hash(x)=19627124) +3781 train 7.809613 (lr=1.3401e-04) (hash(x)=27132838) +3782 train 7.697591 (lr=1.3387e-04) (hash(x)=27242104) +3783 train 7.306519 (lr=1.3373e-04) (hash(x)=21785487) +3784 train 7.546926 (lr=1.3359e-04) (hash(x)=25798262) +3785 train 7.477605 (lr=1.3345e-04) (hash(x)=24806937) +3786 train 7.383965 (lr=1.3331e-04) (hash(x)=24098756) +3787 train 7.357878 (lr=1.3317e-04) (hash(x)=22981456) +3788 train 7.711098 (lr=1.3303e-04) (hash(x)=22441908) +3789 train 7.629983 (lr=1.3290e-04) (hash(x)=25867804) +3790 train 7.560749 (lr=1.3276e-04) (hash(x)=22835586) +3791 train 7.804728 (lr=1.3262e-04) (hash(x)=25251063) +3792 train 7.603148 (lr=1.3248e-04) (hash(x)=27059729) +3793 train 7.609722 (lr=1.3234e-04) (hash(x)=23819311) +3794 train 7.330478 (lr=1.3220e-04) (hash(x)=21345757) +3795 train 7.557115 (lr=1.3207e-04) (hash(x)=25796422) +3796 train 7.258653 (lr=1.3193e-04) (hash(x)=21414971) +3797 train 7.534983 (lr=1.3179e-04) (hash(x)=26120920) +3798 train 7.437084 (lr=1.3165e-04) (hash(x)=22008247) +3799 train 7.527669 (lr=1.3151e-04) (hash(x)=22722211) +3800 val loss 7.5912 +3800 val perplexity 1980.7520 +3800 train 7.307539 (lr=1.3138e-04) (hash(x)=24484513) +3801 train 7.569916 (lr=1.3124e-04) (hash(x)=24463866) +3802 train 7.644497 (lr=1.3110e-04) (hash(x)=26470775) +3803 train 7.601678 (lr=1.3096e-04) (hash(x)=24296755) +3804 train 7.545810 (lr=1.3082e-04) (hash(x)=24381309) +3805 train 7.747150 (lr=1.3069e-04) (hash(x)=26781262) +3806 train 7.375009 (lr=1.3055e-04) (hash(x)=23563137) +3807 train 7.498981 (lr=1.3041e-04) (hash(x)=24522269) +3808 train 7.410346 (lr=1.3027e-04) (hash(x)=22871995) +3809 train 7.669037 (lr=1.3014e-04) (hash(x)=24751946) +3810 train 7.513218 (lr=1.3000e-04) (hash(x)=19879741) +3811 train 7.601917 (lr=1.2986e-04) (hash(x)=25617184) +3812 train 7.326125 (lr=1.2973e-04) (hash(x)=21776722) +3813 train 7.371932 (lr=1.2959e-04) (hash(x)=22663402) +3814 train 7.681522 (lr=1.2945e-04) (hash(x)=26072046) +3815 train 7.403327 (lr=1.2931e-04) (hash(x)=23329475) +3816 train 7.477693 (lr=1.2918e-04) (hash(x)=25519031) +3817 train 7.403983 (lr=1.2904e-04) (hash(x)=22846270) +3818 train 7.463830 (lr=1.2890e-04) (hash(x)=20299429) +3819 train 7.848482 (lr=1.2877e-04) (hash(x)=30652062) +3820 train 7.302829 (lr=1.2863e-04) (hash(x)=19824665) +3821 train 7.745937 (lr=1.2850e-04) (hash(x)=26698904) +3822 train 7.587120 (lr=1.2836e-04) (hash(x)=27612163) +3823 train 7.428602 (lr=1.2822e-04) (hash(x)=24735165) +3824 train 7.223712 (lr=1.2809e-04) (hash(x)=19965890) +3825 train 7.352029 (lr=1.2795e-04) (hash(x)=23518594) +3826 train 7.313488 (lr=1.2781e-04) (hash(x)=22388460) +3827 train 7.409204 (lr=1.2768e-04) (hash(x)=23795686) +3828 train 7.316206 (lr=1.2754e-04) (hash(x)=24470150) +3829 train 7.067268 (lr=1.2741e-04) (hash(x)=19483392) +3830 train 7.362100 (lr=1.2727e-04) (hash(x)=23304516) +3831 train 7.734465 (lr=1.2713e-04) (hash(x)=27002892) +3832 train 7.447940 (lr=1.2700e-04) (hash(x)=22114813) +3833 train 7.645270 (lr=1.2686e-04) (hash(x)=26221916) +3834 train 7.269957 (lr=1.2673e-04) (hash(x)=27261960) +3835 train 7.488729 (lr=1.2659e-04) (hash(x)=21189101) +3836 train 7.185330 (lr=1.2646e-04) (hash(x)=20638173) +3837 train 7.561363 (lr=1.2632e-04) (hash(x)=25497439) +3838 train 7.144580 (lr=1.2619e-04) (hash(x)=20460499) +3839 train 7.533395 (lr=1.2605e-04) (hash(x)=23499349) +3840 train 8.234989 (lr=1.2592e-04) (hash(x)=22671939) +3841 train 8.328948 (lr=1.2578e-04) (hash(x)=26809295) +3842 train 7.557480 (lr=1.2565e-04) (hash(x)=28094504) +3843 train 7.625059 (lr=1.2551e-04) (hash(x)=26807896) +3844 train 7.631853 (lr=1.2538e-04) (hash(x)=24749334) +3845 train 7.537893 (lr=1.2524e-04) (hash(x)=25031330) +3846 train 7.485433 (lr=1.2511e-04) (hash(x)=23008126) +3847 train 7.419809 (lr=1.2497e-04) (hash(x)=22461589) +3848 train 7.584967 (lr=1.2484e-04) (hash(x)=24959391) +3849 train 7.541468 (lr=1.2470e-04) (hash(x)=24239512) +3850 val loss 7.6010 +3850 val perplexity 2000.2806 +3850 train 7.426895 (lr=1.2457e-04) (hash(x)=24760422) +3851 train 7.884143 (lr=1.2444e-04) (hash(x)=26904123) +3852 train 7.179537 (lr=1.2430e-04) (hash(x)=17986444) +3853 train 7.475368 (lr=1.2417e-04) (hash(x)=21618533) +3854 train 7.466568 (lr=1.2403e-04) (hash(x)=27107027) +3855 train 7.659628 (lr=1.2390e-04) (hash(x)=26415040) +3856 train 8.222061 (lr=1.2377e-04) (hash(x)=28682703) +3857 train 7.446080 (lr=1.2363e-04) (hash(x)=22824767) +3858 train 7.593717 (lr=1.2350e-04) (hash(x)=24540450) +3859 train 7.553658 (lr=1.2336e-04) (hash(x)=24463181) +3860 train 7.469777 (lr=1.2323e-04) (hash(x)=22832558) +3861 train 7.584457 (lr=1.2310e-04) (hash(x)=26582384) +3862 train 8.130147 (lr=1.2296e-04) (hash(x)=32327364) +3863 train 7.673381 (lr=1.2283e-04) (hash(x)=26349465) +3864 train 7.443226 (lr=1.2270e-04) (hash(x)=23079414) +3865 train 7.491239 (lr=1.2256e-04) (hash(x)=23464639) +3866 train 7.333835 (lr=1.2243e-04) (hash(x)=21406620) +3867 train 7.278986 (lr=1.2230e-04) (hash(x)=21205988) +3868 train 7.290536 (lr=1.2216e-04) (hash(x)=22742634) +3869 train 7.352640 (lr=1.2203e-04) (hash(x)=24868938) +3870 train 7.669418 (lr=1.2190e-04) (hash(x)=28095283) +3871 train 7.344478 (lr=1.2177e-04) (hash(x)=21596677) +3872 train 7.601837 (lr=1.2163e-04) (hash(x)=26884381) +3873 train 7.645972 (lr=1.2150e-04) (hash(x)=26410272) +3874 train 7.421677 (lr=1.2137e-04) (hash(x)=22915785) +3875 train 7.393134 (lr=1.2124e-04) (hash(x)=23575666) +3876 train 7.414125 (lr=1.2110e-04) (hash(x)=25313223) +3877 train 6.998978 (lr=1.2097e-04) (hash(x)=16319719) +3878 train 7.481697 (lr=1.2084e-04) (hash(x)=23227579) +3879 train 7.432916 (lr=1.2071e-04) (hash(x)=24117012) +3880 train 7.495529 (lr=1.2057e-04) (hash(x)=24681517) +3881 train 7.430960 (lr=1.2044e-04) (hash(x)=23186527) +3882 train 7.419000 (lr=1.2031e-04) (hash(x)=23472784) +3883 train 7.415571 (lr=1.2018e-04) (hash(x)=26013014) +3884 train 7.661956 (lr=1.2005e-04) (hash(x)=25952206) +3885 train 7.669463 (lr=1.1992e-04) (hash(x)=23897834) +3886 train 7.438065 (lr=1.1978e-04) (hash(x)=26350364) +3887 train 7.510538 (lr=1.1965e-04) (hash(x)=24105761) +3888 train 7.540724 (lr=1.1952e-04) (hash(x)=22801707) +3889 train 7.440551 (lr=1.1939e-04) (hash(x)=21821937) +3890 train 7.432913 (lr=1.1926e-04) (hash(x)=24264640) +3891 train 7.318758 (lr=1.1913e-04) (hash(x)=20369133) +3892 train 7.332582 (lr=1.1900e-04) (hash(x)=24313506) +3893 train 7.363155 (lr=1.1886e-04) (hash(x)=22104086) +3894 train 7.576730 (lr=1.1873e-04) (hash(x)=25312602) +3895 train 7.613927 (lr=1.1860e-04) (hash(x)=25725049) +3896 train 7.383279 (lr=1.1847e-04) (hash(x)=22981231) +3897 train 7.254829 (lr=1.1834e-04) (hash(x)=18021467) +3898 train 7.251977 (lr=1.1821e-04) (hash(x)=17960254) +3899 train 7.359367 (lr=1.1808e-04) (hash(x)=19808118) +3900 val loss 7.5862 +3900 val perplexity 1970.8341 +3900 train 7.180946 (lr=1.1795e-04) (hash(x)=21022829) +3901 train 7.338388 (lr=1.1782e-04) (hash(x)=22434663) +3902 train 7.146015 (lr=1.1769e-04) (hash(x)=20208091) +3903 train 7.267232 (lr=1.1756e-04) (hash(x)=22892776) +3904 train 7.354507 (lr=1.1743e-04) (hash(x)=21554367) +3905 train 7.491776 (lr=1.1730e-04) (hash(x)=23704875) +3906 train 7.424900 (lr=1.1717e-04) (hash(x)=21690340) +3907 train 7.412094 (lr=1.1704e-04) (hash(x)=23736780) +3908 train 7.246701 (lr=1.1691e-04) (hash(x)=21060920) +3909 train 7.216069 (lr=1.1678e-04) (hash(x)=20091559) +3910 train 7.537154 (lr=1.1665e-04) (hash(x)=22147974) +3911 train 7.380562 (lr=1.1652e-04) (hash(x)=21630383) +3912 train 7.667540 (lr=1.1639e-04) (hash(x)=28339385) +3913 train 7.475616 (lr=1.1626e-04) (hash(x)=25355505) +3914 train 7.350758 (lr=1.1613e-04) (hash(x)=20843914) +3915 train 7.798802 (lr=1.1600e-04) (hash(x)=25397467) +3916 train 7.550800 (lr=1.1587e-04) (hash(x)=23587501) +3917 train 7.208755 (lr=1.1574e-04) (hash(x)=19036533) +3918 train 7.413682 (lr=1.1561e-04) (hash(x)=29749389) +3919 train 7.264717 (lr=1.1548e-04) (hash(x)=24727094) +3920 train 7.286939 (lr=1.1536e-04) (hash(x)=22723450) +3921 train 7.446992 (lr=1.1523e-04) (hash(x)=24424680) +3922 train 7.603479 (lr=1.1510e-04) (hash(x)=25817917) +3923 train 7.407877 (lr=1.1497e-04) (hash(x)=23951182) +3924 train 7.393066 (lr=1.1484e-04) (hash(x)=21177944) +3925 train 7.295358 (lr=1.1471e-04) (hash(x)=23533768) +3926 train 7.476134 (lr=1.1458e-04) (hash(x)=24263615) +3927 train 7.360504 (lr=1.1446e-04) (hash(x)=22835000) +3928 train 7.522824 (lr=1.1433e-04) (hash(x)=26076156) +3929 train 7.477828 (lr=1.1420e-04) (hash(x)=25171508) +3930 train 7.490180 (lr=1.1407e-04) (hash(x)=22021396) +3931 train 7.752474 (lr=1.1394e-04) (hash(x)=28767849) +3932 train 7.676534 (lr=1.1381e-04) (hash(x)=29497606) +3933 train 7.350075 (lr=1.1369e-04) (hash(x)=22723124) +3934 train 7.634799 (lr=1.1356e-04) (hash(x)=27106616) +3935 train 7.512603 (lr=1.1343e-04) (hash(x)=22839049) +3936 train 7.553377 (lr=1.1330e-04) (hash(x)=25101923) +3937 train 7.460182 (lr=1.1318e-04) (hash(x)=25945975) +3938 train 7.596787 (lr=1.1305e-04) (hash(x)=25382013) +3939 train 7.235372 (lr=1.1292e-04) (hash(x)=19930900) +3940 train 7.329080 (lr=1.1279e-04) (hash(x)=22202373) +3941 train 7.477516 (lr=1.1267e-04) (hash(x)=24592992) +3942 train 7.421957 (lr=1.1254e-04) (hash(x)=25002271) +3943 train 7.211873 (lr=1.1241e-04) (hash(x)=20654136) +3944 train 7.387577 (lr=1.1229e-04) (hash(x)=20061590) +3945 train 7.258414 (lr=1.1216e-04) (hash(x)=21441361) +3946 train 7.451703 (lr=1.1203e-04) (hash(x)=20055468) +3947 train 8.071005 (lr=1.1191e-04) (hash(x)=28495621) +3948 train 7.569136 (lr=1.1178e-04) (hash(x)=25959236) +3949 train 7.870866 (lr=1.1165e-04) (hash(x)=26750193) +3950 val loss 7.5785 +3950 val perplexity 1955.7068 +3950 train 7.671770 (lr=1.1153e-04) (hash(x)=25882605) +3951 train 7.718591 (lr=1.1140e-04) (hash(x)=26776000) +3952 train 7.441278 (lr=1.1127e-04) (hash(x)=24155107) +3953 train 7.440559 (lr=1.1115e-04) (hash(x)=23441845) +3954 train 7.378659 (lr=1.1102e-04) (hash(x)=22860915) +3955 train 7.252174 (lr=1.1089e-04) (hash(x)=21584429) +3956 train 6.885942 (lr=1.1077e-04) (hash(x)=16535556) +3957 train 6.960030 (lr=1.1064e-04) (hash(x)=17946180) +3958 train 7.459812 (lr=1.1052e-04) (hash(x)=25367610) +3959 train 7.379790 (lr=1.1039e-04) (hash(x)=22560658) +3960 train 7.517242 (lr=1.1027e-04) (hash(x)=23809585) +3961 train 7.681203 (lr=1.1014e-04) (hash(x)=20438213) +3962 train 8.112540 (lr=1.1001e-04) (hash(x)=30520037) +3963 train 7.796112 (lr=1.0989e-04) (hash(x)=25276565) +3964 train 7.548155 (lr=1.0976e-04) (hash(x)=26796532) +3965 train 7.368974 (lr=1.0964e-04) (hash(x)=22850475) +3966 train 7.722555 (lr=1.0951e-04) (hash(x)=25983698) +3967 train 7.631490 (lr=1.0939e-04) (hash(x)=25995933) +3968 train 7.842645 (lr=1.0926e-04) (hash(x)=25319339) +3969 train 7.511013 (lr=1.0914e-04) (hash(x)=25066892) +3970 train 7.688758 (lr=1.0901e-04) (hash(x)=26931819) +3971 train 7.528987 (lr=1.0889e-04) (hash(x)=24163910) +3972 train 7.573914 (lr=1.0877e-04) (hash(x)=25359634) +3973 train 7.479387 (lr=1.0864e-04) (hash(x)=24323444) +3974 train 7.604178 (lr=1.0852e-04) (hash(x)=26529231) +3975 train 7.626737 (lr=1.0839e-04) (hash(x)=24635394) +3976 train 7.845590 (lr=1.0827e-04) (hash(x)=24783093) +3977 train 8.018427 (lr=1.0814e-04) (hash(x)=26710509) +3978 train 7.472217 (lr=1.0802e-04) (hash(x)=25396630) +3979 train 7.411153 (lr=1.0790e-04) (hash(x)=22556381) +3980 train 7.673147 (lr=1.0777e-04) (hash(x)=26912161) +3981 train 7.311151 (lr=1.0765e-04) (hash(x)=22268078) +3982 train 8.180154 (lr=1.0752e-04) (hash(x)=30430328) +3983 train 7.782991 (lr=1.0740e-04) (hash(x)=27537228) +3984 train 7.385856 (lr=1.0728e-04) (hash(x)=24468603) +3985 train 7.651953 (lr=1.0715e-04) (hash(x)=27006663) +3986 train 7.821174 (lr=1.0703e-04) (hash(x)=27014728) +3987 train 7.350918 (lr=1.0691e-04) (hash(x)=20883633) +3988 train 7.510660 (lr=1.0678e-04) (hash(x)=23361791) +3989 train 7.597042 (lr=1.0666e-04) (hash(x)=26775925) +3990 train 8.061088 (lr=1.0654e-04) (hash(x)=30648934) +3991 train 7.689459 (lr=1.0641e-04) (hash(x)=26496730) +3992 train 7.449955 (lr=1.0629e-04) (hash(x)=25942897) +3993 train 7.398671 (lr=1.0617e-04) (hash(x)=24887111) +3994 train 7.349939 (lr=1.0605e-04) (hash(x)=22908550) +3995 train 7.135652 (lr=1.0592e-04) (hash(x)=20342150) +3996 train 7.414065 (lr=1.0580e-04) (hash(x)=22261760) +3997 train 7.462470 (lr=1.0568e-04) (hash(x)=22731641) +3998 train 7.283556 (lr=1.0556e-04) (hash(x)=20669749) +3999 train 7.242985 (lr=1.0543e-04) (hash(x)=16533310) +4000 val loss 7.5779 +4000 val perplexity 1954.6169 +4000 train 7.403708 (lr=1.0531e-04) (hash(x)=23661341) +4001 train 7.488742 (lr=1.0519e-04) (hash(x)=24644301) +4002 train 7.316636 (lr=1.0507e-04) (hash(x)=22938438) +4003 train 7.399244 (lr=1.0495e-04) (hash(x)=24315862) +4004 train 7.542181 (lr=1.0482e-04) (hash(x)=25169315) +4005 train 7.555360 (lr=1.0470e-04) (hash(x)=24752796) +4006 train 7.493928 (lr=1.0458e-04) (hash(x)=23103706) +4007 train 7.570778 (lr=1.0446e-04) (hash(x)=24881176) +4008 train 7.549984 (lr=1.0434e-04) (hash(x)=23971947) +4009 train 7.744326 (lr=1.0422e-04) (hash(x)=27741054) +4010 train 8.047917 (lr=1.0410e-04) (hash(x)=30956251) +4011 train 7.569646 (lr=1.0397e-04) (hash(x)=23649686) +4012 train 7.374028 (lr=1.0385e-04) (hash(x)=23344798) +4013 train 7.928323 (lr=1.0373e-04) (hash(x)=29529498) +4014 train 7.601860 (lr=1.0361e-04) (hash(x)=24688359) +4015 train 7.580391 (lr=1.0349e-04) (hash(x)=25593613) +4016 train 7.685655 (lr=1.0337e-04) (hash(x)=25674488) +4017 train 7.268282 (lr=1.0325e-04) (hash(x)=21345346) +4018 train 7.446578 (lr=1.0313e-04) (hash(x)=21978324) +4019 train 7.463567 (lr=1.0301e-04) (hash(x)=23669244) +4020 train 7.370382 (lr=1.0289e-04) (hash(x)=22479613) +4021 train 7.368561 (lr=1.0277e-04) (hash(x)=22855256) +4022 train 7.653967 (lr=1.0265e-04) (hash(x)=19759826) +4023 train 7.559253 (lr=1.0253e-04) (hash(x)=22886646) +4024 train 7.544439 (lr=1.0241e-04) (hash(x)=25553008) +4025 train 7.504497 (lr=1.0229e-04) (hash(x)=25487028) +4026 train 7.890678 (lr=1.0217e-04) (hash(x)=26799246) +4027 train 8.230941 (lr=1.0205e-04) (hash(x)=30728540) +4028 train 7.691110 (lr=1.0193e-04) (hash(x)=23966676) +4029 train 7.477654 (lr=1.0181e-04) (hash(x)=22118783) +4030 train 7.536429 (lr=1.0169e-04) (hash(x)=24744703) +4031 train 7.480196 (lr=1.0157e-04) (hash(x)=21407676) +4032 train 7.452786 (lr=1.0145e-04) (hash(x)=20623349) +4033 train 7.331774 (lr=1.0133e-04) (hash(x)=24723788) +4034 train 7.530262 (lr=1.0121e-04) (hash(x)=25030709) +4035 train 7.613418 (lr=1.0109e-04) (hash(x)=26326164) +4036 train 7.493965 (lr=1.0098e-04) (hash(x)=23754475) +4037 train 8.126976 (lr=1.0086e-04) (hash(x)=30112933) +4038 train 7.831261 (lr=1.0074e-04) (hash(x)=27474299) +4039 train 7.314247 (lr=1.0062e-04) (hash(x)=21720304) +4040 train 7.665733 (lr=1.0050e-04) (hash(x)=25142674) +4041 train 7.920148 (lr=1.0038e-04) (hash(x)=25284552) +4042 train 7.790030 (lr=1.0026e-04) (hash(x)=24554942) +4043 train 7.601535 (lr=1.0015e-04) (hash(x)=23773870) +4044 train 7.395787 (lr=1.0003e-04) (hash(x)=23896338) +4045 train 8.135135 (lr=9.9910e-05) (hash(x)=34984911) +4046 train 7.497910 (lr=9.9792e-05) (hash(x)=23854263) +4047 train 7.742696 (lr=9.9674e-05) (hash(x)=27263416) +4048 train 7.440764 (lr=9.9556e-05) (hash(x)=24989642) +4049 train 7.470596 (lr=9.9439e-05) (hash(x)=24492055) +4050 val loss 7.5757 +4050 val perplexity 1950.2562 +4050 train 7.340558 (lr=9.9321e-05) (hash(x)=21579916) +4051 train 7.581966 (lr=9.9204e-05) (hash(x)=25274710) +4052 train 7.500577 (lr=9.9086e-05) (hash(x)=24701947) +4053 train 7.474713 (lr=9.8969e-05) (hash(x)=25477340) +4054 train 7.395918 (lr=9.8852e-05) (hash(x)=23774195) +4055 train 7.659505 (lr=9.8735e-05) (hash(x)=25840801) +4056 train 7.778320 (lr=9.8618e-05) (hash(x)=27972529) +4057 train 7.609793 (lr=9.8501e-05) (hash(x)=26952458) +4058 train 7.161244 (lr=9.8384e-05) (hash(x)=22683653) +4059 train 7.218452 (lr=9.8267e-05) (hash(x)=20984129) +4060 train 7.361261 (lr=9.8151e-05) (hash(x)=22816482) +4061 train 7.536348 (lr=9.8034e-05) (hash(x)=24285302) +4062 train 7.897963 (lr=9.7918e-05) (hash(x)=31517950) +4063 train 8.195378 (lr=9.7801e-05) (hash(x)=31424568) +4064 train 7.904510 (lr=9.7685e-05) (hash(x)=29497876) +4065 train 7.784663 (lr=9.7569e-05) (hash(x)=27277376) +4066 train 7.583271 (lr=9.7453e-05) (hash(x)=26832588) +4067 train 7.362103 (lr=9.7337e-05) (hash(x)=25149712) +4068 train 7.383823 (lr=9.7221e-05) (hash(x)=23728457) +4069 train 7.426880 (lr=9.7105e-05) (hash(x)=25265136) +4070 train 7.805518 (lr=9.6990e-05) (hash(x)=25241681) +4071 train 7.536294 (lr=9.6874e-05) (hash(x)=23403065) +4072 train 7.504675 (lr=9.6758e-05) (hash(x)=25808207) +4073 train 7.456432 (lr=9.6643e-05) (hash(x)=23904844) +4074 train 7.175854 (lr=9.6528e-05) (hash(x)=18865309) +4075 train 7.525395 (lr=9.6413e-05) (hash(x)=23531744) +4076 train 7.604630 (lr=9.6297e-05) (hash(x)=25784275) +4077 train 7.400723 (lr=9.6182e-05) (hash(x)=22417529) +4078 train 7.433850 (lr=9.6067e-05) (hash(x)=23095491) +4079 train 7.450812 (lr=9.5953e-05) (hash(x)=22656033) +4080 train 7.496431 (lr=9.5838e-05) (hash(x)=25865435) +4081 train 7.566570 (lr=9.5723e-05) (hash(x)=25699377) +4082 train 7.602518 (lr=9.5609e-05) (hash(x)=20854084) +4083 train 7.200386 (lr=9.5494e-05) (hash(x)=18950799) +4084 train 7.380002 (lr=9.5380e-05) (hash(x)=22633739) +4085 train 8.379994 (lr=9.5266e-05) (hash(x)=34054446) +4086 train 7.712596 (lr=9.5152e-05) (hash(x)=27599387) +4087 train 7.531615 (lr=9.5037e-05) (hash(x)=26869295) +4088 train 7.513620 (lr=9.4924e-05) (hash(x)=24075139) +4089 train 7.911669 (lr=9.4810e-05) (hash(x)=29784110) +4090 train 8.160956 (lr=9.4696e-05) (hash(x)=34273918) +4091 train 8.507182 (lr=9.4582e-05) (hash(x)=34472685) +4092 train 7.951814 (lr=9.4469e-05) (hash(x)=30113791) +4093 train 7.298814 (lr=9.4355e-05) (hash(x)=22366381) +4094 train 7.515136 (lr=9.4242e-05) (hash(x)=24922935) +4095 train 7.671609 (lr=9.4129e-05) (hash(x)=25504151) +4096 train 7.663229 (lr=9.4015e-05) (hash(x)=24517375) +4097 train 8.370479 (lr=9.3902e-05) (hash(x)=25871651) +4098 train 7.426751 (lr=9.3789e-05) (hash(x)=23480225) +4099 train 7.320459 (lr=9.3676e-05) (hash(x)=22559978) +4100 val loss 7.5885 +4100 val perplexity 1975.2729 +4100 train 7.693973 (lr=9.3564e-05) (hash(x)=25795272) +4101 train 7.616806 (lr=9.3451e-05) (hash(x)=25440399) +4102 train 7.374423 (lr=9.3338e-05) (hash(x)=21581806) +4103 train 7.482895 (lr=9.3226e-05) (hash(x)=25781518) +4104 train 7.493989 (lr=9.3113e-05) (hash(x)=24682372) +4105 train 7.380469 (lr=9.3001e-05) (hash(x)=22440094) +4106 train 7.502183 (lr=9.2889e-05) (hash(x)=23661032) +4107 train 7.729843 (lr=9.2777e-05) (hash(x)=26966012) +4108 train 7.769981 (lr=9.2665e-05) (hash(x)=26232227) +4109 train 7.445418 (lr=9.2553e-05) (hash(x)=24110656) +4110 train 7.711132 (lr=9.2441e-05) (hash(x)=25938621) +4111 train 7.553668 (lr=9.2329e-05) (hash(x)=26432850) +4112 train 7.560075 (lr=9.2218e-05) (hash(x)=25387672) +4113 train 7.338524 (lr=9.2106e-05) (hash(x)=22740017) +4114 train 7.683284 (lr=9.1995e-05) (hash(x)=26384190) +4115 train 7.391914 (lr=9.1884e-05) (hash(x)=24725583) +4116 train 7.481327 (lr=9.1772e-05) (hash(x)=23986700) +4117 train 7.440741 (lr=9.1661e-05) (hash(x)=18529900) +4118 train 7.511127 (lr=9.1550e-05) (hash(x)=22236257) +4119 train 7.206450 (lr=9.1439e-05) (hash(x)=20155894) +4120 train 7.040144 (lr=9.1328e-05) (hash(x)=17423813) +4121 train 7.018363 (lr=9.1218e-05) (hash(x)=18905183) +4122 train 7.355451 (lr=9.1107e-05) (hash(x)=22534398) +4123 train 7.832771 (lr=9.0997e-05) (hash(x)=25247868) +4124 train 7.569893 (lr=9.0886e-05) (hash(x)=23994188) +4125 train 7.365757 (lr=9.0776e-05) (hash(x)=22929754) +4126 train 7.308900 (lr=9.0666e-05) (hash(x)=22485897) +4127 train 7.387296 (lr=9.0556e-05) (hash(x)=22270169) +4128 train 7.375772 (lr=9.0446e-05) (hash(x)=23638027) +4129 train 7.559009 (lr=9.0336e-05) (hash(x)=23821210) +4130 train 7.477075 (lr=9.0226e-05) (hash(x)=25021512) +4131 train 7.549872 (lr=9.0116e-05) (hash(x)=25240141) +4132 train 7.474576 (lr=9.0006e-05) (hash(x)=22833160) +4133 train 7.460034 (lr=8.9897e-05) (hash(x)=22909944) +4134 train 7.615172 (lr=8.9788e-05) (hash(x)=26163558) +4135 train 7.457298 (lr=8.9678e-05) (hash(x)=22108461) +4136 train 7.406600 (lr=8.9569e-05) (hash(x)=22549232) +4137 train 7.572853 (lr=8.9460e-05) (hash(x)=24701302) +4138 train 7.543972 (lr=8.9351e-05) (hash(x)=25206013) +4139 train 7.602111 (lr=8.9242e-05) (hash(x)=27896130) +4140 train 7.394775 (lr=8.9133e-05) (hash(x)=24106243) +4141 train 7.843657 (lr=8.9024e-05) (hash(x)=27158651) +4142 train 7.495035 (lr=8.8916e-05) (hash(x)=23841147) +4143 train 7.489333 (lr=8.8807e-05) (hash(x)=24802848) +4144 train 7.817659 (lr=8.8699e-05) (hash(x)=24911295) +4145 train 7.599047 (lr=8.8591e-05) (hash(x)=26473900) +4146 train 7.614253 (lr=8.8482e-05) (hash(x)=26785092) +4147 train 7.817558 (lr=8.8374e-05) (hash(x)=30188532) +4148 train 7.434845 (lr=8.8266e-05) (hash(x)=25010649) +4149 train 7.446272 (lr=8.8158e-05) (hash(x)=22409016) +4150 val loss 7.5657 +4150 val perplexity 1930.8136 +4150 train 8.068026 (lr=8.8051e-05) (hash(x)=31675024) +4151 train 7.452016 (lr=8.7943e-05) (hash(x)=25086604) +4152 train 7.422099 (lr=8.7835e-05) (hash(x)=24429343) +4153 train 7.609391 (lr=8.7728e-05) (hash(x)=25014232) +4154 train 7.530698 (lr=8.7621e-05) (hash(x)=25366275) +4155 train 7.960169 (lr=8.7513e-05) (hash(x)=24031473) +4156 train 7.567134 (lr=8.7406e-05) (hash(x)=26651059) +4157 train 7.636575 (lr=8.7299e-05) (hash(x)=26069721) +4158 train 7.742146 (lr=8.7192e-05) (hash(x)=28114382) +4159 train 7.681541 (lr=8.7085e-05) (hash(x)=26983186) +4160 train 7.247253 (lr=8.6978e-05) (hash(x)=21944670) +4161 train 7.793382 (lr=8.6872e-05) (hash(x)=28155702) +4162 train 7.329269 (lr=8.6765e-05) (hash(x)=24765002) +4163 train 7.561984 (lr=8.6659e-05) (hash(x)=26492636) +4164 train 7.507870 (lr=8.6552e-05) (hash(x)=24063705) +4165 train 7.750912 (lr=8.6446e-05) (hash(x)=27079573) +4166 train 7.457200 (lr=8.6340e-05) (hash(x)=24912201) +4167 train 7.608130 (lr=8.6234e-05) (hash(x)=27208507) +4168 train 7.484483 (lr=8.6128e-05) (hash(x)=25824320) +4169 train 7.987554 (lr=8.6022e-05) (hash(x)=29641677) +4170 train 7.700134 (lr=8.5916e-05) (hash(x)=26275614) +4171 train 7.548295 (lr=8.5811e-05) (hash(x)=25227725) +4172 train 7.578896 (lr=8.5705e-05) (hash(x)=27417375) +4173 train 7.740559 (lr=8.5600e-05) (hash(x)=25581973) +4174 train 8.045964 (lr=8.5495e-05) (hash(x)=30360417) +4175 train 7.428755 (lr=8.5389e-05) (hash(x)=23862845) +4176 train 7.439716 (lr=8.5284e-05) (hash(x)=25415130) +4177 train 7.418592 (lr=8.5179e-05) (hash(x)=23111123) +4178 train 7.419477 (lr=8.5074e-05) (hash(x)=24022804) +4179 train 7.499991 (lr=8.4970e-05) (hash(x)=25148490) +4180 train 7.596207 (lr=8.4865e-05) (hash(x)=22792092) +4181 train 7.745683 (lr=8.4760e-05) (hash(x)=24779698) +4182 train 7.661858 (lr=8.4656e-05) (hash(x)=24870844) +4183 train 7.363732 (lr=8.4551e-05) (hash(x)=22664494) +4184 train 7.464673 (lr=8.4447e-05) (hash(x)=22599729) +4185 train 7.556814 (lr=8.4343e-05) (hash(x)=24707078) +4186 train 7.412635 (lr=8.4239e-05) (hash(x)=23823945) +4187 train 7.196824 (lr=8.4135e-05) (hash(x)=22219856) +4188 train 7.738863 (lr=8.4031e-05) (hash(x)=27061401) +4189 train 7.516595 (lr=8.3927e-05) (hash(x)=23398766) +4190 train 7.629574 (lr=8.3824e-05) (hash(x)=27916730) +4191 train 7.512583 (lr=8.3720e-05) (hash(x)=24092927) +4192 train 7.523311 (lr=8.3617e-05) (hash(x)=24723657) +4193 train 7.570930 (lr=8.3513e-05) (hash(x)=24676155) +4194 train 7.497112 (lr=8.3410e-05) (hash(x)=23455369) +4195 train 7.440296 (lr=8.3307e-05) (hash(x)=21999890) +4196 train 7.404903 (lr=8.3204e-05) (hash(x)=23385567) +4197 train 7.530765 (lr=8.3101e-05) (hash(x)=25120814) +4198 train 7.289161 (lr=8.2998e-05) (hash(x)=21308113) +4199 train 7.591689 (lr=8.2896e-05) (hash(x)=27213812) +4200 val loss 7.5589 +4200 val perplexity 1917.7886 +4200 train 7.165423 (lr=8.2793e-05) (hash(x)=19675382) +4201 train 7.425281 (lr=8.2691e-05) (hash(x)=23882161) +4202 train 7.452472 (lr=8.2588e-05) (hash(x)=24338567) +4203 train 7.674470 (lr=8.2486e-05) (hash(x)=27649723) +4204 train 7.487491 (lr=8.2384e-05) (hash(x)=27563514) +4205 train 7.519207 (lr=8.2282e-05) (hash(x)=23128552) +4206 train 7.417099 (lr=8.2180e-05) (hash(x)=26203283) +4207 train 7.475095 (lr=8.2078e-05) (hash(x)=23929955) +4208 train 7.208051 (lr=8.1976e-05) (hash(x)=20978691) +4209 train 7.409153 (lr=8.1875e-05) (hash(x)=21875178) +4210 train 7.471710 (lr=8.1773e-05) (hash(x)=23563293) +4211 train 7.590566 (lr=8.1672e-05) (hash(x)=25538503) +4212 train 7.604720 (lr=8.1570e-05) (hash(x)=24171014) +4213 train 7.456660 (lr=8.1469e-05) (hash(x)=22306665) +4214 train 7.479335 (lr=8.1368e-05) (hash(x)=26082318) +4215 train 7.337578 (lr=8.1267e-05) (hash(x)=23025790) +4216 train 7.282879 (lr=8.1166e-05) (hash(x)=21146597) +4217 train 7.652711 (lr=8.1066e-05) (hash(x)=27470230) +4218 train 7.498058 (lr=8.0965e-05) (hash(x)=22691005) +4219 train 7.264336 (lr=8.0864e-05) (hash(x)=21550313) +4220 train 7.352845 (lr=8.0764e-05) (hash(x)=20618443) +4221 train 7.418513 (lr=8.0664e-05) (hash(x)=25260787) +4222 train 7.498941 (lr=8.0563e-05) (hash(x)=25249873) +4223 train 7.236035 (lr=8.0463e-05) (hash(x)=20452272) +4224 train 7.423605 (lr=8.0363e-05) (hash(x)=25649764) +4225 train 7.269635 (lr=8.0263e-05) (hash(x)=19589460) +4226 train 7.755166 (lr=8.0164e-05) (hash(x)=26477595) +4227 train 7.976886 (lr=8.0064e-05) (hash(x)=27585442) +4228 train 7.634134 (lr=7.9964e-05) (hash(x)=25821343) +4229 train 7.504974 (lr=7.9865e-05) (hash(x)=24549919) +4230 train 7.530792 (lr=7.9765e-05) (hash(x)=26124522) +4231 train 7.679906 (lr=7.9666e-05) (hash(x)=26473464) +4232 train 7.762846 (lr=7.9567e-05) (hash(x)=28883028) +4233 train 7.751178 (lr=7.9468e-05) (hash(x)=29611296) +4234 train 7.614325 (lr=7.9369e-05) (hash(x)=29059941) +4235 train 7.511643 (lr=7.9270e-05) (hash(x)=26368203) +4236 train 7.570220 (lr=7.9172e-05) (hash(x)=27021360) +4237 train 7.500866 (lr=7.9073e-05) (hash(x)=26135379) +4238 train 7.359471 (lr=7.8975e-05) (hash(x)=25862549) +4239 train 7.644229 (lr=7.8876e-05) (hash(x)=28276603) +4240 train 7.293927 (lr=7.8778e-05) (hash(x)=22170090) +4241 train 7.525089 (lr=7.8680e-05) (hash(x)=22625589) +4242 train 7.557161 (lr=7.8582e-05) (hash(x)=25751475) +4243 train 7.519804 (lr=7.8484e-05) (hash(x)=25836838) +4244 train 7.438443 (lr=7.8386e-05) (hash(x)=23938214) +4245 train 7.375459 (lr=7.8288e-05) (hash(x)=22759365) +4246 train 7.402797 (lr=7.8191e-05) (hash(x)=24964628) +4247 train 7.064548 (lr=7.8093e-05) (hash(x)=21480367) +4248 train 7.667950 (lr=7.7996e-05) (hash(x)=25974694) +4249 train 7.931377 (lr=7.7898e-05) (hash(x)=29830546) +4250 val loss 7.5574 +4250 val perplexity 1914.7687 +4250 train 7.577402 (lr=7.7801e-05) (hash(x)=26283200) +4251 train 7.623207 (lr=7.7704e-05) (hash(x)=25605672) +4252 train 7.310385 (lr=7.7607e-05) (hash(x)=21439107) +4253 train 7.304466 (lr=7.7510e-05) (hash(x)=22626883) +4254 train 7.577171 (lr=7.7414e-05) (hash(x)=25627150) +4255 train 7.685755 (lr=7.7317e-05) (hash(x)=24658642) +4256 train 7.803329 (lr=7.7221e-05) (hash(x)=29755505) +4257 train 7.426389 (lr=7.7124e-05) (hash(x)=22750846) +4258 train 7.298719 (lr=7.7028e-05) (hash(x)=21892651) +4259 train 7.404579 (lr=7.6932e-05) (hash(x)=22487960) +4260 train 7.489472 (lr=7.6836e-05) (hash(x)=24059869) +4261 train 7.269188 (lr=7.6740e-05) (hash(x)=21090180) +4262 train 7.354910 (lr=7.6644e-05) (hash(x)=22716452) +4263 train 7.639816 (lr=7.6548e-05) (hash(x)=26854801) +4264 train 7.526769 (lr=7.6452e-05) (hash(x)=23815428) +4265 train 7.369174 (lr=7.6357e-05) (hash(x)=23368704) +4266 train 7.418836 (lr=7.6261e-05) (hash(x)=24301098) +4267 train 7.602086 (lr=7.6166e-05) (hash(x)=25025039) +4268 train 7.419408 (lr=7.6071e-05) (hash(x)=25425736) +4269 train 7.580907 (lr=7.5976e-05) (hash(x)=27439380) +4270 train 7.202718 (lr=7.5881e-05) (hash(x)=18682756) +4271 train 7.213647 (lr=7.5786e-05) (hash(x)=17862226) +4272 train 7.762954 (lr=7.5691e-05) (hash(x)=29946491) +4273 train 7.051525 (lr=7.5597e-05) (hash(x)=21392062) +4274 train 7.581092 (lr=7.5502e-05) (hash(x)=25195556) +4275 train 7.536939 (lr=7.5408e-05) (hash(x)=26682036) +4276 train 7.467584 (lr=7.5314e-05) (hash(x)=21235260) +4277 train 7.330107 (lr=7.5219e-05) (hash(x)=22118984) +4278 train 7.366499 (lr=7.5125e-05) (hash(x)=24094510) +4279 train 7.840981 (lr=7.5031e-05) (hash(x)=28519182) +4280 train 7.528466 (lr=7.4938e-05) (hash(x)=26652859) +4281 train 7.435757 (lr=7.4844e-05) (hash(x)=24463139) +4282 train 7.408374 (lr=7.4750e-05) (hash(x)=23281870) +4283 train 7.644051 (lr=7.4657e-05) (hash(x)=29181174) +4284 train 7.375455 (lr=7.4563e-05) (hash(x)=24797417) +4285 train 7.647630 (lr=7.4470e-05) (hash(x)=29026537) +4286 train 7.488272 (lr=7.4377e-05) (hash(x)=23045165) +4287 train 7.409544 (lr=7.4284e-05) (hash(x)=25193694) +4288 train 7.381859 (lr=7.4191e-05) (hash(x)=23475407) +4289 train 7.598056 (lr=7.4098e-05) (hash(x)=24875410) +4290 train 7.473979 (lr=7.4005e-05) (hash(x)=26328101) +4291 train 7.638625 (lr=7.3913e-05) (hash(x)=26273661) +4292 train 7.357430 (lr=7.3820e-05) (hash(x)=23271891) +4293 train 7.327280 (lr=7.3728e-05) (hash(x)=23300732) +4294 train 7.429626 (lr=7.3636e-05) (hash(x)=24243693) +4295 train 7.507674 (lr=7.3544e-05) (hash(x)=25642620) +4296 train 7.574602 (lr=7.3452e-05) (hash(x)=27730411) +4297 train 7.453406 (lr=7.3360e-05) (hash(x)=25485335) +4298 train 7.344041 (lr=7.3268e-05) (hash(x)=25268789) +4299 train 7.726327 (lr=7.3176e-05) (hash(x)=28917143) +4300 val loss 7.5590 +4300 val perplexity 1917.9531 +4300 train 7.772381 (lr=7.3085e-05) (hash(x)=28987991) +4301 train 7.535541 (lr=7.2993e-05) (hash(x)=25989151) +4302 train 7.463327 (lr=7.2902e-05) (hash(x)=25658195) +4303 train 7.472792 (lr=7.2811e-05) (hash(x)=25765909) +4304 train 7.511272 (lr=7.2719e-05) (hash(x)=25704261) +4305 train 7.504900 (lr=7.2628e-05) (hash(x)=24438198) +4306 train 7.315221 (lr=7.2537e-05) (hash(x)=22871710) +4307 train 7.201710 (lr=7.2447e-05) (hash(x)=20160896) +4308 train 7.624925 (lr=7.2356e-05) (hash(x)=24218641) +4309 train 7.191289 (lr=7.2265e-05) (hash(x)=22073751) +4310 train 7.400394 (lr=7.2175e-05) (hash(x)=26434959) +4311 train 7.164267 (lr=7.2085e-05) (hash(x)=20750531) +4312 train 7.339925 (lr=7.1995e-05) (hash(x)=22232293) +4313 train 7.565701 (lr=7.1904e-05) (hash(x)=27324521) +4314 train 7.271276 (lr=7.1814e-05) (hash(x)=22217862) +4315 train 7.687119 (lr=7.1725e-05) (hash(x)=25741958) +4316 train 7.420235 (lr=7.1635e-05) (hash(x)=22738456) +4317 train 7.450885 (lr=7.1545e-05) (hash(x)=25387302) +4318 train 7.412421 (lr=7.1456e-05) (hash(x)=24669014) +4319 train 7.547978 (lr=7.1366e-05) (hash(x)=24917098) +4320 train 7.504674 (lr=7.1277e-05) (hash(x)=26698227) +4321 train 7.353636 (lr=7.1188e-05) (hash(x)=23518293) +4322 train 7.296104 (lr=7.1099e-05) (hash(x)=23643971) +4323 train 7.464772 (lr=7.1010e-05) (hash(x)=25659505) +4324 train 7.279148 (lr=7.0921e-05) (hash(x)=22697720) +4325 train 7.736141 (lr=7.0832e-05) (hash(x)=25082178) +4326 train 7.292462 (lr=7.0744e-05) (hash(x)=24166546) +4327 train 7.517147 (lr=7.0655e-05) (hash(x)=24499766) +4328 train 7.670457 (lr=7.0567e-05) (hash(x)=24338607) +4329 train 7.366550 (lr=7.0479e-05) (hash(x)=23225420) +4330 train 7.548809 (lr=7.0390e-05) (hash(x)=25907032) +4331 train 7.342619 (lr=7.0302e-05) (hash(x)=24634979) +4332 train 7.389375 (lr=7.0215e-05) (hash(x)=22161377) +4333 train 7.477300 (lr=7.0127e-05) (hash(x)=22559939) +4334 train 7.303565 (lr=7.0039e-05) (hash(x)=20840022) +4335 train 7.632808 (lr=6.9952e-05) (hash(x)=25067358) +4336 train 7.391165 (lr=6.9864e-05) (hash(x)=23432626) +4337 train 7.470511 (lr=6.9777e-05) (hash(x)=23820385) +4338 train 8.100520 (lr=6.9690e-05) (hash(x)=31594930) +4339 train 7.620929 (lr=6.9603e-05) (hash(x)=27138750) +4340 train 7.755921 (lr=6.9516e-05) (hash(x)=30031341) +4341 train 7.462311 (lr=6.9429e-05) (hash(x)=24602807) +4342 train 7.283159 (lr=6.9342e-05) (hash(x)=21584976) +4343 train 7.564812 (lr=6.9255e-05) (hash(x)=27479796) +4344 train 7.366075 (lr=6.9169e-05) (hash(x)=22746241) +4345 train 7.440746 (lr=6.9082e-05) (hash(x)=19728452) +4346 train 7.319922 (lr=6.8996e-05) (hash(x)=24502020) +4347 train 7.678445 (lr=6.8910e-05) (hash(x)=28225954) +4348 train 7.369804 (lr=6.8824e-05) (hash(x)=23893447) +4349 train 7.631550 (lr=6.8738e-05) (hash(x)=25654586) +4350 val loss 7.5546 +4350 val perplexity 1909.5486 +4350 train 7.491240 (lr=6.8652e-05) (hash(x)=23856469) +4351 train 7.437774 (lr=6.8567e-05) (hash(x)=23136080) +4352 train 7.589242 (lr=6.8481e-05) (hash(x)=26532918) +4353 train 7.482862 (lr=6.8396e-05) (hash(x)=25461947) +4354 train 7.507940 (lr=6.8310e-05) (hash(x)=26590231) +4355 train 7.272774 (lr=6.8225e-05) (hash(x)=21441985) +4356 train 7.510606 (lr=6.8140e-05) (hash(x)=26507941) +4357 train 7.727268 (lr=6.8055e-05) (hash(x)=26245238) +4358 train 7.532053 (lr=6.7970e-05) (hash(x)=24879059) +4359 train 7.804471 (lr=6.7885e-05) (hash(x)=29810074) +4360 train 7.296390 (lr=6.7801e-05) (hash(x)=24258069) +4361 train 7.395242 (lr=6.7716e-05) (hash(x)=25229360) +4362 train 7.604383 (lr=6.7632e-05) (hash(x)=25440206) +4363 train 7.442163 (lr=6.7548e-05) (hash(x)=23222125) +4364 train 7.351141 (lr=6.7463e-05) (hash(x)=23353186) +4365 train 7.540318 (lr=6.7379e-05) (hash(x)=26995240) +4366 train 7.455421 (lr=6.7295e-05) (hash(x)=24135899) +4367 train 7.481158 (lr=6.7212e-05) (hash(x)=24765539) +4368 train 7.488595 (lr=6.7128e-05) (hash(x)=26051004) +4369 train 7.372730 (lr=6.7044e-05) (hash(x)=22559142) +4370 train 7.310649 (lr=6.6961e-05) (hash(x)=20364388) +4371 train 7.608551 (lr=6.6878e-05) (hash(x)=25938817) +4372 train 7.694210 (lr=6.6794e-05) (hash(x)=26060945) +4373 train 7.574522 (lr=6.6711e-05) (hash(x)=27866714) +4374 train 7.560651 (lr=6.6628e-05) (hash(x)=27104972) +4375 train 7.272666 (lr=6.6545e-05) (hash(x)=22216309) +4376 train 7.328268 (lr=6.6463e-05) (hash(x)=20781533) +4377 train 7.188301 (lr=6.6380e-05) (hash(x)=18784350) +4378 train 7.354984 (lr=6.6298e-05) (hash(x)=21102897) +4379 train 7.543991 (lr=6.6215e-05) (hash(x)=25647489) +4380 train 7.366372 (lr=6.6133e-05) (hash(x)=24486462) +4381 train 7.626256 (lr=6.6051e-05) (hash(x)=26539114) +4382 train 7.327619 (lr=6.5969e-05) (hash(x)=20026045) +4383 train 7.520443 (lr=6.5887e-05) (hash(x)=20613861) +4384 train 7.328459 (lr=6.5805e-05) (hash(x)=20676237) +4385 train 7.498702 (lr=6.5723e-05) (hash(x)=24893163) +4386 train 7.123279 (lr=6.5642e-05) (hash(x)=17299395) +4387 train 7.285039 (lr=6.5561e-05) (hash(x)=19580191) +4388 train 7.425159 (lr=6.5479e-05) (hash(x)=25588218) +4389 train 7.544215 (lr=6.5398e-05) (hash(x)=26433063) +4390 train 7.614890 (lr=6.5317e-05) (hash(x)=24917693) +4391 train 7.307582 (lr=6.5236e-05) (hash(x)=23462447) +4392 train 7.525351 (lr=6.5155e-05) (hash(x)=27185416) +4393 train 7.597024 (lr=6.5074e-05) (hash(x)=24109010) +4394 train 7.392972 (lr=6.4994e-05) (hash(x)=21486000) +4395 train 7.317728 (lr=6.4913e-05) (hash(x)=22230006) +4396 train 7.638148 (lr=6.4833e-05) (hash(x)=26092311) +4397 train 7.450557 (lr=6.4753e-05) (hash(x)=24112426) +4398 train 7.307459 (lr=6.4673e-05) (hash(x)=21881373) +4399 train 7.431237 (lr=6.4593e-05) (hash(x)=26739455) +4400 val loss 7.5438 +4400 val perplexity 1889.0745 +4400 train 7.377266 (lr=6.4513e-05) (hash(x)=22873602) +4401 train 7.389287 (lr=6.4433e-05) (hash(x)=26748712) +4402 train 7.682608 (lr=6.4354e-05) (hash(x)=26306618) +4403 train 7.795131 (lr=6.4274e-05) (hash(x)=26898808) +4404 train 7.464124 (lr=6.4195e-05) (hash(x)=25067751) +4405 train 7.479328 (lr=6.4115e-05) (hash(x)=25397145) +4406 train 7.623160 (lr=6.4036e-05) (hash(x)=24796962) +4407 train 7.439565 (lr=6.3957e-05) (hash(x)=23222996) +4408 train 7.388284 (lr=6.3878e-05) (hash(x)=24820189) +4409 train 7.759460 (lr=6.3800e-05) (hash(x)=27669038) +4410 train 7.424462 (lr=6.3721e-05) (hash(x)=24283976) +4411 train 7.579373 (lr=6.3642e-05) (hash(x)=25018789) +4412 train 7.500553 (lr=6.3564e-05) (hash(x)=26247064) +4413 train 7.444165 (lr=6.3486e-05) (hash(x)=22942904) +4414 train 7.404846 (lr=6.3408e-05) (hash(x)=23918746) +4415 train 7.508046 (lr=6.3329e-05) (hash(x)=26489285) +4416 train 6.762976 (lr=6.3252e-05) (hash(x)=14942495) +4417 train 7.404575 (lr=6.3174e-05) (hash(x)=23908505) +4418 train 7.550910 (lr=6.3096e-05) (hash(x)=25854036) +4419 train 7.472270 (lr=6.3018e-05) (hash(x)=27522868) +4420 train 7.380295 (lr=6.2941e-05) (hash(x)=26549205) +4421 train 7.254039 (lr=6.2864e-05) (hash(x)=24565522) +4422 train 7.528059 (lr=6.2787e-05) (hash(x)=23919123) +4423 train 7.432001 (lr=6.2709e-05) (hash(x)=25185942) +4424 train 7.467947 (lr=6.2632e-05) (hash(x)=24818969) +4425 train 7.521812 (lr=6.2556e-05) (hash(x)=26750372) +4426 train 7.528604 (lr=6.2479e-05) (hash(x)=25621536) +4427 train 7.924098 (lr=6.2402e-05) (hash(x)=28037060) +4428 train 7.987314 (lr=6.2326e-05) (hash(x)=30829517) +4429 train 7.619915 (lr=6.2250e-05) (hash(x)=28728766) +4430 train 7.493215 (lr=6.2173e-05) (hash(x)=25936895) +4431 train 7.650816 (lr=6.2097e-05) (hash(x)=27476158) +4432 train 7.416168 (lr=6.2021e-05) (hash(x)=23297921) +4433 train 7.386533 (lr=6.1945e-05) (hash(x)=23201299) +4434 train 7.353040 (lr=6.1870e-05) (hash(x)=22681766) +4435 train 7.307922 (lr=6.1794e-05) (hash(x)=23217512) +4436 train 7.767100 (lr=6.1719e-05) (hash(x)=29322443) +4437 train 7.408161 (lr=6.1643e-05) (hash(x)=23502072) +4438 train 7.505808 (lr=6.1568e-05) (hash(x)=28873527) +4439 train 7.483043 (lr=6.1493e-05) (hash(x)=23653175) +4440 train 7.281867 (lr=6.1418e-05) (hash(x)=21635879) +4441 train 7.358480 (lr=6.1343e-05) (hash(x)=22201854) +4442 train 7.269809 (lr=6.1268e-05) (hash(x)=23101164) +4443 train 7.645658 (lr=6.1194e-05) (hash(x)=25475122) +4444 train 7.429272 (lr=6.1119e-05) (hash(x)=25124825) +4445 train 7.518455 (lr=6.1045e-05) (hash(x)=26883852) +4446 train 7.356019 (lr=6.0970e-05) (hash(x)=24731829) +4447 train 7.388268 (lr=6.0896e-05) (hash(x)=22405076) +4448 train 7.474901 (lr=6.0822e-05) (hash(x)=27633869) +4449 train 7.571311 (lr=6.0748e-05) (hash(x)=25162594) +4450 val loss 7.5391 +4450 val perplexity 1880.1106 +4450 train 7.538321 (lr=6.0675e-05) (hash(x)=26438149) +4451 train 7.533278 (lr=6.0601e-05) (hash(x)=26748800) +4452 train 7.340543 (lr=6.0527e-05) (hash(x)=22971620) +4453 train 7.404285 (lr=6.0454e-05) (hash(x)=25886430) +4454 train 7.218143 (lr=6.0381e-05) (hash(x)=21084137) +4455 train 7.396676 (lr=6.0308e-05) (hash(x)=25673545) +4456 train 7.997604 (lr=6.0235e-05) (hash(x)=29389002) +4457 train 7.160230 (lr=6.0162e-05) (hash(x)=21676076) +4458 train 7.246194 (lr=6.0089e-05) (hash(x)=22616647) +4459 train 7.580804 (lr=6.0016e-05) (hash(x)=22771197) +4460 train 7.511108 (lr=5.9944e-05) (hash(x)=26393514) +4461 train 7.712068 (lr=5.9871e-05) (hash(x)=24996433) +4462 train 7.393951 (lr=5.9799e-05) (hash(x)=19594028) +4463 train 7.508060 (lr=5.9727e-05) (hash(x)=22675428) +4464 train 7.549063 (lr=5.9655e-05) (hash(x)=26566551) +4465 train 7.553112 (lr=5.9583e-05) (hash(x)=22469290) +4466 train 7.955537 (lr=5.9511e-05) (hash(x)=29821654) +4467 train 7.507277 (lr=5.9439e-05) (hash(x)=26305388) +4468 train 7.382108 (lr=5.9368e-05) (hash(x)=21998506) +4469 train 7.764388 (lr=5.9296e-05) (hash(x)=23953315) +4470 train 7.501129 (lr=5.9225e-05) (hash(x)=23985456) +4471 train 7.785164 (lr=5.9154e-05) (hash(x)=28168456) +4472 train 7.479705 (lr=5.9083e-05) (hash(x)=25126411) +4473 train 7.358590 (lr=5.9012e-05) (hash(x)=24619336) +4474 train 7.527744 (lr=5.8941e-05) (hash(x)=24467798) +4475 train 8.138373 (lr=5.8871e-05) (hash(x)=31673254) +4476 train 8.047476 (lr=5.8800e-05) (hash(x)=29929610) +4477 train 7.785253 (lr=5.8730e-05) (hash(x)=29412572) +4478 train 7.939691 (lr=5.8659e-05) (hash(x)=31910006) +4479 train 7.560654 (lr=5.8589e-05) (hash(x)=26072586) +4480 train 7.530118 (lr=5.8519e-05) (hash(x)=25782825) +4481 train 7.375386 (lr=5.8449e-05) (hash(x)=21131363) +4482 train 7.484168 (lr=5.8379e-05) (hash(x)=25071223) +4483 train 8.133922 (lr=5.8310e-05) (hash(x)=32249577) +4484 train 7.397805 (lr=5.8240e-05) (hash(x)=24752808) +4485 train 7.434901 (lr=5.8171e-05) (hash(x)=24970539) +4486 train 7.362008 (lr=5.8102e-05) (hash(x)=24191005) +4487 train 7.567181 (lr=5.8032e-05) (hash(x)=24974331) +4488 train 7.662505 (lr=5.7963e-05) (hash(x)=27043409) +4489 train 7.177697 (lr=5.7894e-05) (hash(x)=11320385) +4490 train 7.163131 (lr=5.7826e-05) (hash(x)=11919565) +4491 train 7.211531 (lr=5.7757e-05) (hash(x)=10404694) +4492 train 7.291389 (lr=5.7688e-05) (hash(x)=12393159) +4493 train 7.202301 (lr=5.7620e-05) (hash(x)=11843609) +4494 train 7.233051 (lr=5.7552e-05) (hash(x)=12632729) +4495 train 7.455359 (lr=5.7484e-05) (hash(x)=21951984) +4496 train 8.011998 (lr=5.7416e-05) (hash(x)=19411544) +4497 train 7.972013 (lr=5.7348e-05) (hash(x)=21000228) +4498 train 7.591711 (lr=5.7280e-05) (hash(x)=26018207) +4499 train 7.668016 (lr=5.7212e-05) (hash(x)=24659058) +4500 val loss 7.5395 +4500 val perplexity 1880.9185 +4500 train 7.786231 (lr=5.7145e-05) (hash(x)=27919597) +4501 train 7.434021 (lr=5.7077e-05) (hash(x)=26232596) +4502 train 7.530252 (lr=5.7010e-05) (hash(x)=26248912) +4503 train 7.465474 (lr=5.6943e-05) (hash(x)=25403751) +4504 train 7.442277 (lr=5.6876e-05) (hash(x)=21096637) +4505 train 7.572664 (lr=5.6809e-05) (hash(x)=26560941) +4506 train 7.514015 (lr=5.6742e-05) (hash(x)=24942406) +4507 train 7.561353 (lr=5.6675e-05) (hash(x)=24405748) +4508 train 7.440942 (lr=5.6609e-05) (hash(x)=22975455) +4509 train 7.417175 (lr=5.6543e-05) (hash(x)=26358820) +4510 train 7.385506 (lr=5.6476e-05) (hash(x)=24211938) +4511 train 7.559789 (lr=5.6410e-05) (hash(x)=26396012) +4512 train 7.320696 (lr=5.6344e-05) (hash(x)=22534410) +4513 train 7.398633 (lr=5.6278e-05) (hash(x)=23855025) +4514 train 7.333089 (lr=5.6212e-05) (hash(x)=22722345) +4515 train 7.300612 (lr=5.6147e-05) (hash(x)=22779251) +4516 train 7.485261 (lr=5.6081e-05) (hash(x)=23973078) +4517 train 7.398309 (lr=5.6016e-05) (hash(x)=19886228) +4518 train 7.482797 (lr=5.5951e-05) (hash(x)=23849694) +4519 train 7.842557 (lr=5.5886e-05) (hash(x)=24827190) +4520 train 7.422428 (lr=5.5821e-05) (hash(x)=25111498) +4521 train 8.228312 (lr=5.5756e-05) (hash(x)=29752556) +4522 train 7.753865 (lr=5.5691e-05) (hash(x)=25606805) +4523 train 7.399513 (lr=5.5626e-05) (hash(x)=22209714) +4524 train 7.290315 (lr=5.5562e-05) (hash(x)=24991834) +4525 train 7.494041 (lr=5.5497e-05) (hash(x)=27808158) +4526 train 7.426654 (lr=5.5433e-05) (hash(x)=23041199) +4527 train 7.486528 (lr=5.5369e-05) (hash(x)=24847458) +4528 train 7.451237 (lr=5.5305e-05) (hash(x)=25091787) +4529 train 8.340768 (lr=5.5241e-05) (hash(x)=29745551) +4530 train 7.269558 (lr=5.5178e-05) (hash(x)=20447167) +4531 train 7.375700 (lr=5.5114e-05) (hash(x)=20720911) +4532 train 7.485938 (lr=5.5050e-05) (hash(x)=24803353) +4533 train 7.519886 (lr=5.4987e-05) (hash(x)=23780724) +4534 train 7.553610 (lr=5.4924e-05) (hash(x)=23423120) +4535 train 7.258112 (lr=5.4861e-05) (hash(x)=22159088) +4536 train 7.459586 (lr=5.4798e-05) (hash(x)=25820304) +4537 train 7.258904 (lr=5.4735e-05) (hash(x)=22813612) +4538 train 7.558363 (lr=5.4672e-05) (hash(x)=27827979) +4539 train 7.558650 (lr=5.4610e-05) (hash(x)=25737179) +4540 train 7.778150 (lr=5.4547e-05) (hash(x)=23401504) +4541 train 7.478998 (lr=5.4485e-05) (hash(x)=25071988) +4542 train 7.565094 (lr=5.4423e-05) (hash(x)=22844541) +4543 train 7.235349 (lr=5.4361e-05) (hash(x)=22140034) +4544 train 7.529842 (lr=5.4299e-05) (hash(x)=26522091) +4545 train 7.532392 (lr=5.4237e-05) (hash(x)=24099725) +4546 train 7.115079 (lr=5.4175e-05) (hash(x)=19127182) +4547 train 7.207696 (lr=5.4114e-05) (hash(x)=18992693) +4548 train 8.149424 (lr=5.4052e-05) (hash(x)=26723015) +4549 train 8.025470 (lr=5.3991e-05) (hash(x)=30389969) +4550 val loss 7.5402 +4550 val perplexity 1882.1737 +4550 train 7.412411 (lr=5.3930e-05) (hash(x)=24065654) +4551 train 7.206562 (lr=5.3869e-05) (hash(x)=19044797) +4552 train 7.541443 (lr=5.3808e-05) (hash(x)=24919665) +4553 train 7.621548 (lr=5.3747e-05) (hash(x)=22780968) +4554 train 7.389970 (lr=5.3687e-05) (hash(x)=24549699) +4555 train 7.301975 (lr=5.3626e-05) (hash(x)=21374811) +4556 train 7.398154 (lr=5.3566e-05) (hash(x)=25225950) +4557 train 7.499791 (lr=5.3505e-05) (hash(x)=25691882) +4558 train 7.440098 (lr=5.3445e-05) (hash(x)=24685235) +4559 train 7.313062 (lr=5.3385e-05) (hash(x)=23932794) +4560 train 7.823821 (lr=5.3325e-05) (hash(x)=28152043) +4561 train 7.392688 (lr=5.3266e-05) (hash(x)=23417051) +4562 train 7.476779 (lr=5.3206e-05) (hash(x)=26250211) +4563 train 7.599181 (lr=5.3146e-05) (hash(x)=25428813) +4564 train 7.712182 (lr=5.3087e-05) (hash(x)=27477379) +4565 train 7.551539 (lr=5.3028e-05) (hash(x)=23015212) +4566 train 7.358345 (lr=5.2969e-05) (hash(x)=22837608) +4567 train 7.410573 (lr=5.2910e-05) (hash(x)=25256890) +4568 train 7.357536 (lr=5.2851e-05) (hash(x)=22742827) +4569 train 7.683263 (lr=5.2792e-05) (hash(x)=30167922) +4570 train 7.433214 (lr=5.2734e-05) (hash(x)=24540265) +4571 train 7.671127 (lr=5.2675e-05) (hash(x)=26752941) +4572 train 7.522771 (lr=5.2617e-05) (hash(x)=25094026) +4573 train 8.115226 (lr=5.2559e-05) (hash(x)=28508785) +4574 train 7.621395 (lr=5.2501e-05) (hash(x)=26501871) +4575 train 7.405351 (lr=5.2443e-05) (hash(x)=24161711) +4576 train 7.330585 (lr=5.2385e-05) (hash(x)=24313695) +4577 train 7.762151 (lr=5.2327e-05) (hash(x)=28248933) +4578 train 7.886562 (lr=5.2270e-05) (hash(x)=28445722) +4579 train 7.480081 (lr=5.2212e-05) (hash(x)=24589015) +4580 train 7.617543 (lr=5.2155e-05) (hash(x)=26192193) +4581 train 7.373876 (lr=5.2098e-05) (hash(x)=25131316) +4582 train 7.428839 (lr=5.2041e-05) (hash(x)=26631504) +4583 train 7.303699 (lr=5.1984e-05) (hash(x)=22036817) +4584 train 7.831682 (lr=5.1927e-05) (hash(x)=28361254) +4585 train 7.506805 (lr=5.1871e-05) (hash(x)=24828340) +4586 train 7.279927 (lr=5.1814e-05) (hash(x)=24067304) +4587 train 7.088655 (lr=5.1758e-05) (hash(x)=20733289) +4588 train 7.034362 (lr=5.1701e-05) (hash(x)=19526622) +4589 train 7.344380 (lr=5.1645e-05) (hash(x)=22148688) +4590 train 7.488778 (lr=5.1589e-05) (hash(x)=24141800) +4591 train 7.535207 (lr=5.1533e-05) (hash(x)=26663208) +4592 train 7.261464 (lr=5.1478e-05) (hash(x)=22536305) +4593 train 7.454050 (lr=5.1422e-05) (hash(x)=23608185) +4594 train 7.392640 (lr=5.1367e-05) (hash(x)=23348495) +4595 train 7.498895 (lr=5.1311e-05) (hash(x)=25409759) +4596 train 7.504772 (lr=5.1256e-05) (hash(x)=25572131) +4597 train 7.426251 (lr=5.1201e-05) (hash(x)=21782039) +4598 train 7.441165 (lr=5.1146e-05) (hash(x)=24643923) +4599 train 7.420557 (lr=5.1091e-05) (hash(x)=24072213) +4600 val loss 7.5358 +4600 val perplexity 1873.9554 +4600 train 7.361142 (lr=5.1037e-05) (hash(x)=23925612) +4601 train 7.459993 (lr=5.0982e-05) (hash(x)=23315415) +4602 train 7.391031 (lr=5.0928e-05) (hash(x)=21793800) +4603 train 7.407784 (lr=5.0873e-05) (hash(x)=24921200) +4604 train 7.611134 (lr=5.0819e-05) (hash(x)=27226113) +4605 train 7.408952 (lr=5.0765e-05) (hash(x)=22920960) +4606 train 7.382382 (lr=5.0711e-05) (hash(x)=21947750) +4607 train 7.397185 (lr=5.0658e-05) (hash(x)=22766044) +4608 train 7.811256 (lr=5.0604e-05) (hash(x)=28237005) +4609 train 7.345218 (lr=5.0550e-05) (hash(x)=24922621) +4610 train 7.538270 (lr=5.0497e-05) (hash(x)=24899830) +4611 train 7.898530 (lr=5.0444e-05) (hash(x)=32920298) +4612 train 7.382579 (lr=5.0391e-05) (hash(x)=25083835) +4613 train 7.406641 (lr=5.0338e-05) (hash(x)=22863054) +4614 train 7.449552 (lr=5.0285e-05) (hash(x)=24841464) +4615 train 7.660672 (lr=5.0232e-05) (hash(x)=27871153) +4616 train 7.714036 (lr=5.0180e-05) (hash(x)=28025163) +4617 train 7.376205 (lr=5.0127e-05) (hash(x)=24659561) +4618 train 7.412677 (lr=5.0075e-05) (hash(x)=25067194) +4619 train 7.483934 (lr=5.0023e-05) (hash(x)=22731460) +4620 train 7.357095 (lr=4.9971e-05) (hash(x)=20445873) +4621 train 7.521665 (lr=4.9919e-05) (hash(x)=26033948) +4622 train 7.421434 (lr=4.9867e-05) (hash(x)=22473213) +4623 train 7.728498 (lr=4.9815e-05) (hash(x)=24037280) +4624 train 7.961277 (lr=4.9764e-05) (hash(x)=25624131) +4625 train 7.807785 (lr=4.9712e-05) (hash(x)=26799867) +4626 train 7.545269 (lr=4.9661e-05) (hash(x)=27187602) +4627 train 7.503167 (lr=4.9610e-05) (hash(x)=23277695) +4628 train 7.460997 (lr=4.9559e-05) (hash(x)=24748234) +4629 train 7.602931 (lr=4.9508e-05) (hash(x)=26103104) +4630 train 7.482952 (lr=4.9457e-05) (hash(x)=24327389) +4631 train 7.591026 (lr=4.9407e-05) (hash(x)=24121850) +4632 train 7.461345 (lr=4.9356e-05) (hash(x)=23714590) +4633 train 7.510212 (lr=4.9306e-05) (hash(x)=22379412) +4634 train 7.715434 (lr=4.9256e-05) (hash(x)=24454713) +4635 train 7.444703 (lr=4.9206e-05) (hash(x)=22966977) +4636 train 7.409263 (lr=4.9156e-05) (hash(x)=23764884) +4637 train 7.328597 (lr=4.9106e-05) (hash(x)=23827429) +4638 train 7.500012 (lr=4.9056e-05) (hash(x)=24088592) +4639 train 7.477988 (lr=4.9007e-05) (hash(x)=24380031) +4640 train 7.514390 (lr=4.8957e-05) (hash(x)=26065050) +4641 train 7.479399 (lr=4.8908e-05) (hash(x)=24442902) +4642 train 7.064745 (lr=4.8859e-05) (hash(x)=18548782) +4643 train 7.526287 (lr=4.8810e-05) (hash(x)=26957303) +4644 train 7.524383 (lr=4.8761e-05) (hash(x)=25032727) +4645 train 7.403787 (lr=4.8712e-05) (hash(x)=27224706) +4646 train 7.589263 (lr=4.8664e-05) (hash(x)=27508476) +4647 train 7.287222 (lr=4.8615e-05) (hash(x)=23055215) +4648 train 7.447325 (lr=4.8567e-05) (hash(x)=24496194) +4649 train 8.983665 (lr=4.8519e-05) (hash(x)=13982941) +4650 val loss 7.5431 +4650 val perplexity 1887.7266 +4650 train 8.005265 (lr=4.8470e-05) (hash(x)=16721547) +4651 train 7.455306 (lr=4.8422e-05) (hash(x)=22929154) +4652 train 7.383707 (lr=4.8375e-05) (hash(x)=23323994) +4653 train 7.661697 (lr=4.8327e-05) (hash(x)=24877951) +4654 train 7.521480 (lr=4.8279e-05) (hash(x)=24096183) +4655 train 7.403152 (lr=4.8232e-05) (hash(x)=25329724) +4656 train 7.385824 (lr=4.8185e-05) (hash(x)=23877337) +4657 train 7.093133 (lr=4.8138e-05) (hash(x)=20923083) +4658 train 7.315338 (lr=4.8091e-05) (hash(x)=23807996) +4659 train 7.447899 (lr=4.8044e-05) (hash(x)=24370475) +4660 train 7.794478 (lr=4.7997e-05) (hash(x)=28202255) +4661 train 7.595345 (lr=4.7950e-05) (hash(x)=26142119) +4662 train 7.343141 (lr=4.7904e-05) (hash(x)=21387743) +4663 train 7.478332 (lr=4.7857e-05) (hash(x)=25662408) +4664 train 7.510715 (lr=4.7811e-05) (hash(x)=23962815) +4665 train 7.489107 (lr=4.7765e-05) (hash(x)=23987677) +4666 train 7.681989 (lr=4.7719e-05) (hash(x)=26554284) +4667 train 7.713944 (lr=4.7673e-05) (hash(x)=25991817) +4668 train 7.471477 (lr=4.7628e-05) (hash(x)=24256966) +4669 train 7.263743 (lr=4.7582e-05) (hash(x)=22187158) +4670 train 7.498866 (lr=4.7537e-05) (hash(x)=26295320) +4671 train 7.540052 (lr=4.7491e-05) (hash(x)=26346814) +4672 train 7.858635 (lr=4.7446e-05) (hash(x)=26594196) +4673 train 7.585776 (lr=4.7401e-05) (hash(x)=24322101) +4674 train 7.680692 (lr=4.7356e-05) (hash(x)=27274566) +4675 train 7.474292 (lr=4.7312e-05) (hash(x)=24505725) +4676 train 7.644020 (lr=4.7267e-05) (hash(x)=26167371) +4677 train 7.723507 (lr=4.7222e-05) (hash(x)=28062311) +4678 train 7.495905 (lr=4.7178e-05) (hash(x)=23476009) +4679 train 7.590390 (lr=4.7134e-05) (hash(x)=25283256) +4680 train 7.300830 (lr=4.7090e-05) (hash(x)=22033246) +4681 train 7.332147 (lr=4.7046e-05) (hash(x)=22716214) +4682 train 7.469997 (lr=4.7002e-05) (hash(x)=25672672) +4683 train 7.349107 (lr=4.6958e-05) (hash(x)=22979072) +4684 train 7.523857 (lr=4.6915e-05) (hash(x)=27439204) +4685 train 7.445606 (lr=4.6871e-05) (hash(x)=24857737) +4686 train 7.692213 (lr=4.6828e-05) (hash(x)=29615897) +4687 train 7.373052 (lr=4.6785e-05) (hash(x)=24021771) +4688 train 7.499184 (lr=4.6742e-05) (hash(x)=27410807) +4689 train 7.331789 (lr=4.6699e-05) (hash(x)=22850411) +4690 train 7.531840 (lr=4.6656e-05) (hash(x)=24949696) +4691 train 7.392613 (lr=4.6614e-05) (hash(x)=22956381) +4692 train 7.801773 (lr=4.6571e-05) (hash(x)=28193458) +4693 train 7.387043 (lr=4.6529e-05) (hash(x)=25596844) +4694 train 7.377206 (lr=4.6487e-05) (hash(x)=22101377) +4695 train 7.348429 (lr=4.6445e-05) (hash(x)=23576840) +4696 train 7.258625 (lr=4.6403e-05) (hash(x)=21849758) +4697 train 7.466378 (lr=4.6361e-05) (hash(x)=26431349) +4698 train 7.503845 (lr=4.6319e-05) (hash(x)=22960758) +4699 train 7.461290 (lr=4.6278e-05) (hash(x)=23243097) +4700 val loss 7.5367 +4700 val perplexity 1875.5564 +4700 train 7.405779 (lr=4.6236e-05) (hash(x)=23715370) +4701 train 7.958074 (lr=4.6195e-05) (hash(x)=30678293) +4702 train 8.084760 (lr=4.6154e-05) (hash(x)=32481620) +4703 train 7.572885 (lr=4.6113e-05) (hash(x)=26414858) +4704 train 7.515020 (lr=4.6072e-05) (hash(x)=24768691) +4705 train 7.289189 (lr=4.6031e-05) (hash(x)=21627762) +4706 train 7.327872 (lr=4.5991e-05) (hash(x)=21024917) +4707 train 7.253165 (lr=4.5950e-05) (hash(x)=23570951) +4708 train 7.391736 (lr=4.5910e-05) (hash(x)=23729185) +4709 train 7.627470 (lr=4.5870e-05) (hash(x)=25933754) +4710 train 7.311113 (lr=4.5830e-05) (hash(x)=23091014) +4711 train 7.309768 (lr=4.5790e-05) (hash(x)=20099261) +4712 train 7.489053 (lr=4.5750e-05) (hash(x)=26807297) +4713 train 7.420386 (lr=4.5710e-05) (hash(x)=25332115) +4714 train 7.344280 (lr=4.5671e-05) (hash(x)=23247605) +4715 train 7.265375 (lr=4.5631e-05) (hash(x)=23786549) +4716 train 7.314038 (lr=4.5592e-05) (hash(x)=23981166) +4717 train 7.482751 (lr=4.5553e-05) (hash(x)=25967754) +4718 train 7.392506 (lr=4.5514e-05) (hash(x)=23659116) +4719 train 7.364552 (lr=4.5475e-05) (hash(x)=22453718) +4720 train 7.283875 (lr=4.5437e-05) (hash(x)=22597951) +4721 train 7.532210 (lr=4.5398e-05) (hash(x)=25284885) +4722 train 7.499290 (lr=4.5360e-05) (hash(x)=24748569) +4723 train 7.176633 (lr=4.5321e-05) (hash(x)=19448608) +4724 train 7.488534 (lr=4.5283e-05) (hash(x)=24888040) +4725 train 7.317085 (lr=4.5245e-05) (hash(x)=23203503) +4726 train 7.158605 (lr=4.5207e-05) (hash(x)=20387787) +4727 train 7.216454 (lr=4.5169e-05) (hash(x)=22529445) +4728 train 7.519897 (lr=4.5132e-05) (hash(x)=22455471) +4729 train 7.465693 (lr=4.5094e-05) (hash(x)=25661132) +4730 train 7.465473 (lr=4.5057e-05) (hash(x)=24997711) +4731 train 7.432478 (lr=4.5020e-05) (hash(x)=22575521) +4732 train 7.341653 (lr=4.4983e-05) (hash(x)=22640285) +4733 train 7.223485 (lr=4.4946e-05) (hash(x)=18637357) +4734 train 7.333995 (lr=4.4909e-05) (hash(x)=22845826) +4735 train 7.533046 (lr=4.4872e-05) (hash(x)=24484543) +4736 train 7.395687 (lr=4.4836e-05) (hash(x)=23352320) +4737 train 7.189989 (lr=4.4799e-05) (hash(x)=21544758) +4738 train 7.022418 (lr=4.4763e-05) (hash(x)=18292136) +4739 train 7.627553 (lr=4.4727e-05) (hash(x)=24893614) +4740 train 7.527761 (lr=4.4691e-05) (hash(x)=28103443) +4741 train 7.353569 (lr=4.4655e-05) (hash(x)=22233356) +4742 train 7.213968 (lr=4.4619e-05) (hash(x)=21133541) +4743 train 7.418647 (lr=4.4584e-05) (hash(x)=24043998) +4744 train 7.700348 (lr=4.4548e-05) (hash(x)=24801185) +4745 train 7.517876 (lr=4.4513e-05) (hash(x)=23858358) +4746 train 7.538385 (lr=4.4478e-05) (hash(x)=23926989) +4747 train 7.576688 (lr=4.4443e-05) (hash(x)=24813708) +4748 train 7.455448 (lr=4.4408e-05) (hash(x)=26339467) +4749 train 7.326552 (lr=4.4373e-05) (hash(x)=21850656) +4750 val loss 7.5329 +4750 val perplexity 1868.5642 +4750 train 7.285410 (lr=4.4338e-05) (hash(x)=21475802) +4751 train 7.476191 (lr=4.4304e-05) (hash(x)=24301906) +4752 train 7.376247 (lr=4.4270e-05) (hash(x)=22748495) +4753 train 7.464664 (lr=4.4235e-05) (hash(x)=25649256) +4754 train 7.330956 (lr=4.4201e-05) (hash(x)=23934346) +4755 train 7.446015 (lr=4.4167e-05) (hash(x)=26332892) +4756 train 7.317043 (lr=4.4133e-05) (hash(x)=23279389) +4757 train 7.451401 (lr=4.4100e-05) (hash(x)=23146858) +4758 train 7.748746 (lr=4.4066e-05) (hash(x)=26892932) +4759 train 7.418983 (lr=4.4033e-05) (hash(x)=26328881) +4760 train 7.420699 (lr=4.4000e-05) (hash(x)=24394655) +4761 train 7.413934 (lr=4.3966e-05) (hash(x)=22122308) +4762 train 7.526302 (lr=4.3933e-05) (hash(x)=24200369) +4763 train 7.647742 (lr=4.3901e-05) (hash(x)=26841776) +4764 train 7.588112 (lr=4.3868e-05) (hash(x)=27196641) +4765 train 7.361649 (lr=4.3835e-05) (hash(x)=24912822) +4766 train 7.557594 (lr=4.3803e-05) (hash(x)=25946055) +4767 train 7.267156 (lr=4.3770e-05) (hash(x)=23101508) +4768 train 7.437625 (lr=4.3738e-05) (hash(x)=24287798) +4769 train 7.452008 (lr=4.3706e-05) (hash(x)=22798964) +4770 train 7.361259 (lr=4.3674e-05) (hash(x)=24164479) +4771 train 7.432793 (lr=4.3643e-05) (hash(x)=24946464) +4772 train 7.567722 (lr=4.3611e-05) (hash(x)=25154423) +4773 train 7.325719 (lr=4.3579e-05) (hash(x)=23173476) +4774 train 7.465617 (lr=4.3548e-05) (hash(x)=25373559) +4775 train 7.459138 (lr=4.3517e-05) (hash(x)=23527176) +4776 train 7.461010 (lr=4.3486e-05) (hash(x)=24865403) +4777 train 7.423687 (lr=4.3455e-05) (hash(x)=24451067) +4778 train 7.934787 (lr=4.3424e-05) (hash(x)=28187162) +4779 train 8.061016 (lr=4.3393e-05) (hash(x)=31163350) +4780 train 8.437396 (lr=4.3363e-05) (hash(x)=33563280) +4781 train 8.496119 (lr=4.3332e-05) (hash(x)=34939183) +4782 train 8.024190 (lr=4.3302e-05) (hash(x)=30263543) +4783 train 7.343468 (lr=4.3272e-05) (hash(x)=22705673) +4784 train 7.250836 (lr=4.3242e-05) (hash(x)=21415023) +4785 train 7.450439 (lr=4.3212e-05) (hash(x)=26079097) +4786 train 7.725149 (lr=4.3182e-05) (hash(x)=25503836) +4787 train 7.585001 (lr=4.3153e-05) (hash(x)=24705721) +4788 train 7.465891 (lr=4.3123e-05) (hash(x)=24384657) +4789 train 7.345872 (lr=4.3094e-05) (hash(x)=22550579) +4790 train 7.406616 (lr=4.3065e-05) (hash(x)=22452164) +4791 train 7.649925 (lr=4.3036e-05) (hash(x)=28664796) +4792 train 8.040343 (lr=4.3007e-05) (hash(x)=26139280) +4793 train 7.619261 (lr=4.2978e-05) (hash(x)=23862341) +4794 train 7.294922 (lr=4.2950e-05) (hash(x)=23784757) +4795 train 7.209752 (lr=4.2921e-05) (hash(x)=22659441) +4796 train 7.903877 (lr=4.2893e-05) (hash(x)=29744216) +4797 train 7.766264 (lr=4.2864e-05) (hash(x)=25822591) +4798 train 7.315029 (lr=4.2836e-05) (hash(x)=22370895) +4799 train 7.666524 (lr=4.2808e-05) (hash(x)=27102890) +4800 val loss 7.5421 +4800 val perplexity 1885.7006 +4800 train 7.676328 (lr=4.2781e-05) (hash(x)=27014625) +4801 train 7.450652 (lr=4.2753e-05) (hash(x)=25755963) +4802 train 7.684574 (lr=4.2725e-05) (hash(x)=29675278) +4803 train 7.744003 (lr=4.2698e-05) (hash(x)=26660930) +4804 train 7.615263 (lr=4.2671e-05) (hash(x)=24557060) +4805 train 7.757439 (lr=4.2644e-05) (hash(x)=24622741) +4806 train 7.796773 (lr=4.2617e-05) (hash(x)=23952601) +4807 train 7.560858 (lr=4.2590e-05) (hash(x)=26787259) +4808 train 7.779477 (lr=4.2563e-05) (hash(x)=28919605) +4809 train 8.147593 (lr=4.2537e-05) (hash(x)=37061654) +4810 train 7.909033 (lr=4.2510e-05) (hash(x)=30379739) +4811 train 7.618300 (lr=4.2484e-05) (hash(x)=26097180) +4812 train 7.765093 (lr=4.2458e-05) (hash(x)=26327092) +4813 train 7.586534 (lr=4.2432e-05) (hash(x)=23258030) +4814 train 7.617775 (lr=4.2406e-05) (hash(x)=25582015) +4815 train 7.362199 (lr=4.2380e-05) (hash(x)=23396088) +4816 train 7.347448 (lr=4.2354e-05) (hash(x)=21904146) +4817 train 7.517213 (lr=4.2329e-05) (hash(x)=26281676) +4818 train 7.377380 (lr=4.2304e-05) (hash(x)=23140470) +4819 train 7.634351 (lr=4.2278e-05) (hash(x)=24611098) +4820 train 7.459764 (lr=4.2253e-05) (hash(x)=23832642) +4821 train 7.523209 (lr=4.2229e-05) (hash(x)=21439671) +4822 train 7.424448 (lr=4.2204e-05) (hash(x)=25128845) +4823 train 7.396780 (lr=4.2179e-05) (hash(x)=24911831) +4824 train 7.517467 (lr=4.2155e-05) (hash(x)=26051723) +4825 train 7.573292 (lr=4.2130e-05) (hash(x)=26479565) +4826 train 7.668477 (lr=4.2106e-05) (hash(x)=26228987) +4827 train 7.548200 (lr=4.2082e-05) (hash(x)=25131300) +4828 train 7.305912 (lr=4.2058e-05) (hash(x)=19921978) +4829 train 7.301706 (lr=4.2034e-05) (hash(x)=23173449) +4830 train 7.545052 (lr=4.2010e-05) (hash(x)=24791832) +4831 train 7.312639 (lr=4.1987e-05) (hash(x)=23453491) +4832 train 7.436773 (lr=4.1964e-05) (hash(x)=22564139) +4833 train 7.343985 (lr=4.1940e-05) (hash(x)=21659918) +4834 train 7.438438 (lr=4.1917e-05) (hash(x)=22956076) +4835 train 7.579672 (lr=4.1894e-05) (hash(x)=22237612) +4836 train 7.637160 (lr=4.1871e-05) (hash(x)=24736427) +4837 train 7.696243 (lr=4.1849e-05) (hash(x)=24939751) +4838 train 7.533020 (lr=4.1826e-05) (hash(x)=25059298) +4839 train 7.515895 (lr=4.1804e-05) (hash(x)=23453396) +4840 train 7.156610 (lr=4.1781e-05) (hash(x)=17919338) +4841 train 7.549215 (lr=4.1759e-05) (hash(x)=23428815) +4842 train 7.759908 (lr=4.1737e-05) (hash(x)=27042659) +4843 train 7.844059 (lr=4.1715e-05) (hash(x)=25161278) +4844 train 7.593272 (lr=4.1693e-05) (hash(x)=24113253) +4845 train 7.696786 (lr=4.1672e-05) (hash(x)=26139263) +4846 train 7.764282 (lr=4.1650e-05) (hash(x)=27787006) +4847 train 7.490683 (lr=4.1629e-05) (hash(x)=23869612) +4848 train 7.674503 (lr=4.1608e-05) (hash(x)=26092193) +4849 train 7.714920 (lr=4.1587e-05) (hash(x)=29351182) +4850 val loss 7.5311 +4850 val perplexity 1865.1379 +4850 train 7.820683 (lr=4.1566e-05) (hash(x)=28773463) +4851 train 7.759084 (lr=4.1545e-05) (hash(x)=28207741) +4852 train 7.382516 (lr=4.1524e-05) (hash(x)=23280878) +4853 train 7.770664 (lr=4.1504e-05) (hash(x)=26742336) +4854 train 7.542060 (lr=4.1484e-05) (hash(x)=23543321) +4855 train 7.907921 (lr=4.1463e-05) (hash(x)=26581590) +4856 train 7.344737 (lr=4.1443e-05) (hash(x)=22728668) +4857 train 7.200515 (lr=4.1423e-05) (hash(x)=19854534) +4858 train 7.600516 (lr=4.1404e-05) (hash(x)=26612813) +4859 train 7.746721 (lr=4.1384e-05) (hash(x)=25827863) +4860 train 7.463832 (lr=4.1364e-05) (hash(x)=24574997) +4861 train 7.618659 (lr=4.1345e-05) (hash(x)=26187830) +4862 train 7.560660 (lr=4.1326e-05) (hash(x)=25105823) +4863 train 7.634831 (lr=4.1307e-05) (hash(x)=28056342) +4864 train 7.510461 (lr=4.1288e-05) (hash(x)=21885801) +4865 train 7.670898 (lr=4.1269e-05) (hash(x)=25659043) +4866 train 7.707597 (lr=4.1250e-05) (hash(x)=23852824) +4867 train 7.484231 (lr=4.1231e-05) (hash(x)=23965470) +4868 train 7.461896 (lr=4.1213e-05) (hash(x)=25035012) +4869 train 7.645648 (lr=4.1195e-05) (hash(x)=26639165) +4870 train 7.621822 (lr=4.1177e-05) (hash(x)=29205362) +4871 train 7.515977 (lr=4.1159e-05) (hash(x)=25900866) +4872 train 7.525880 (lr=4.1141e-05) (hash(x)=25636242) +4873 train 7.587524 (lr=4.1123e-05) (hash(x)=25430698) +4874 train 7.571144 (lr=4.1105e-05) (hash(x)=27629981) +4875 train 7.570425 (lr=4.1088e-05) (hash(x)=27682625) +4876 train 7.788579 (lr=4.1071e-05) (hash(x)=27549409) +4877 train 7.578554 (lr=4.1053e-05) (hash(x)=24661627) +4878 train 7.602234 (lr=4.1036e-05) (hash(x)=25196542) +4879 train 7.505250 (lr=4.1019e-05) (hash(x)=24549177) +4880 train 7.494614 (lr=4.1003e-05) (hash(x)=23740600) +4881 train 7.369789 (lr=4.0986e-05) (hash(x)=21460850) +4882 train 7.621700 (lr=4.0970e-05) (hash(x)=27101400) +4883 train 7.590292 (lr=4.0953e-05) (hash(x)=24193076) +4884 train 7.593676 (lr=4.0937e-05) (hash(x)=24582947) +4885 train 7.470606 (lr=4.0921e-05) (hash(x)=25133839) +4886 train 7.504133 (lr=4.0905e-05) (hash(x)=24759454) +4887 train 7.571983 (lr=4.0889e-05) (hash(x)=28239583) +4888 train 7.847761 (lr=4.0874e-05) (hash(x)=29594489) +4889 train 7.518167 (lr=4.0858e-05) (hash(x)=23833431) +4890 train 7.478986 (lr=4.0843e-05) (hash(x)=23139411) +4891 train 7.375496 (lr=4.0827e-05) (hash(x)=20885864) +4892 train 7.793591 (lr=4.0812e-05) (hash(x)=26217418) +4893 train 7.644717 (lr=4.0797e-05) (hash(x)=27321870) +4894 train 7.235992 (lr=4.0783e-05) (hash(x)=19912955) +4895 train 7.517284 (lr=4.0768e-05) (hash(x)=23223554) +4896 train 7.639844 (lr=4.0753e-05) (hash(x)=25667219) +4897 train 7.793613 (lr=4.0739e-05) (hash(x)=28007972) +4898 train 7.676534 (lr=4.0725e-05) (hash(x)=27748764) +4899 train 7.626904 (lr=4.0710e-05) (hash(x)=27425770) +4900 val loss 7.5337 +4900 val perplexity 1869.9591 +4900 train 7.735616 (lr=4.0697e-05) (hash(x)=28394020) +4901 train 7.552432 (lr=4.0683e-05) (hash(x)=24080235) +4902 train 7.434584 (lr=4.0669e-05) (hash(x)=23309527) +4903 train 7.418512 (lr=4.0655e-05) (hash(x)=24793480) +4904 train 7.608549 (lr=4.0642e-05) (hash(x)=25344456) +4905 train 7.542731 (lr=4.0629e-05) (hash(x)=24590670) +4906 train 7.583141 (lr=4.0615e-05) (hash(x)=26937171) +4907 train 7.661357 (lr=4.0602e-05) (hash(x)=26949097) +4908 train 7.465032 (lr=4.0590e-05) (hash(x)=22443915) +4909 train 7.555197 (lr=4.0577e-05) (hash(x)=23814995) +4910 train 7.460576 (lr=4.0564e-05) (hash(x)=26135871) +4911 train 7.398206 (lr=4.0552e-05) (hash(x)=25415570) +4912 train 7.598389 (lr=4.0539e-05) (hash(x)=26756326) +4913 train 7.523102 (lr=4.0527e-05) (hash(x)=22920200) +4914 train 7.594330 (lr=4.0515e-05) (hash(x)=26710977) +4915 train 7.501829 (lr=4.0503e-05) (hash(x)=24985634) +4916 train 7.586885 (lr=4.0492e-05) (hash(x)=26004335) +4917 train 7.537815 (lr=4.0480e-05) (hash(x)=25637457) +4918 train 7.715265 (lr=4.0468e-05) (hash(x)=26645180) +4919 train 7.543389 (lr=4.0457e-05) (hash(x)=24425760) +4920 train 7.771204 (lr=4.0446e-05) (hash(x)=28223544) +4921 train 7.504335 (lr=4.0435e-05) (hash(x)=25538618) +4922 train 7.845751 (lr=4.0424e-05) (hash(x)=26984784) +4923 train 7.776952 (lr=4.0413e-05) (hash(x)=29154578) +4924 train 7.530956 (lr=4.0402e-05) (hash(x)=25308123) +4925 train 7.849262 (lr=4.0392e-05) (hash(x)=27939259) +4926 train 7.478014 (lr=4.0382e-05) (hash(x)=21984545) +4927 train 7.502677 (lr=4.0371e-05) (hash(x)=23707134) +4928 train 7.473434 (lr=4.0361e-05) (hash(x)=27201034) +4929 train 7.944782 (lr=4.0351e-05) (hash(x)=31623877) +4930 train 7.448034 (lr=4.0341e-05) (hash(x)=22162782) +4931 train 7.347150 (lr=4.0332e-05) (hash(x)=20049335) +4932 train 7.418832 (lr=4.0322e-05) (hash(x)=25594665) +4933 train 7.484314 (lr=4.0313e-05) (hash(x)=25265312) +4934 train 7.573776 (lr=4.0304e-05) (hash(x)=27094896) +4935 train 7.955300 (lr=4.0294e-05) (hash(x)=28321697) +4936 train 7.456277 (lr=4.0285e-05) (hash(x)=25006013) +4937 train 7.689185 (lr=4.0277e-05) (hash(x)=24596431) +4938 train 7.670743 (lr=4.0268e-05) (hash(x)=25150510) +4939 train 7.306152 (lr=4.0259e-05) (hash(x)=21497535) +4940 train 7.474561 (lr=4.0251e-05) (hash(x)=25094669) +4941 train 7.425351 (lr=4.0243e-05) (hash(x)=24024557) +4942 train 7.664642 (lr=4.0234e-05) (hash(x)=24370776) +4943 train 7.508279 (lr=4.0226e-05) (hash(x)=23434031) +4944 train 7.454706 (lr=4.0219e-05) (hash(x)=24383517) +4945 train 7.559138 (lr=4.0211e-05) (hash(x)=25858759) +4946 train 7.210688 (lr=4.0203e-05) (hash(x)=20409561) +4947 train 7.771026 (lr=4.0196e-05) (hash(x)=27469117) +4948 train 7.409370 (lr=4.0188e-05) (hash(x)=22086623) +4949 train 7.621736 (lr=4.0181e-05) (hash(x)=25759281) +4950 val loss 7.5307 +4950 val perplexity 1864.3358 +4950 train 7.550664 (lr=4.0174e-05) (hash(x)=27130117) +4951 train 7.462401 (lr=4.0167e-05) (hash(x)=27003481) +4952 train 7.747817 (lr=4.0161e-05) (hash(x)=26725937) +4953 train 7.489537 (lr=4.0154e-05) (hash(x)=22691119) +4954 train 7.396493 (lr=4.0147e-05) (hash(x)=17272898) +4955 train 7.251624 (lr=4.0141e-05) (hash(x)=17850370) +4956 train 7.273609 (lr=4.0135e-05) (hash(x)=18729639) +4957 train 7.560613 (lr=4.0129e-05) (hash(x)=25327160) +4958 train 7.736309 (lr=4.0123e-05) (hash(x)=28709044) +4959 train 7.340466 (lr=4.0117e-05) (hash(x)=22236893) +4960 train 7.303993 (lr=4.0112e-05) (hash(x)=21729251) +4961 train 7.380914 (lr=4.0106e-05) (hash(x)=23852346) +4962 train 7.539679 (lr=4.0101e-05) (hash(x)=23974368) +4963 train 7.503807 (lr=4.0095e-05) (hash(x)=25764691) +4964 train 7.840524 (lr=4.0090e-05) (hash(x)=28341865) +4965 train 7.421560 (lr=4.0085e-05) (hash(x)=23856238) +4966 train 7.626311 (lr=4.0081e-05) (hash(x)=24568904) +4967 train 7.508565 (lr=4.0076e-05) (hash(x)=26857458) +4968 train 7.381664 (lr=4.0071e-05) (hash(x)=20507972) +4969 train 7.553636 (lr=4.0067e-05) (hash(x)=23139455) +4970 train 7.433267 (lr=4.0063e-05) (hash(x)=24853703) +4971 train 7.447877 (lr=4.0059e-05) (hash(x)=25654849) +4972 train 7.278862 (lr=4.0055e-05) (hash(x)=22963710) +4973 train 7.662048 (lr=4.0051e-05) (hash(x)=25652110) +4974 train 7.556034 (lr=4.0047e-05) (hash(x)=24085957) +4975 train 7.825573 (lr=4.0044e-05) (hash(x)=26413122) +4976 train 7.884680 (lr=4.0040e-05) (hash(x)=26989387) +4977 train 7.530885 (lr=4.0037e-05) (hash(x)=22784033) +4978 train 7.883408 (lr=4.0034e-05) (hash(x)=26694945) +4979 train 7.816368 (lr=4.0031e-05) (hash(x)=24507726) +4980 train 7.876187 (lr=4.0028e-05) (hash(x)=26490335) +4981 train 7.819311 (lr=4.0025e-05) (hash(x)=25624751) +4982 train 7.820432 (lr=4.0023e-05) (hash(x)=27846204) +4983 train 7.683591 (lr=4.0020e-05) (hash(x)=27696537) +4984 train 7.597320 (lr=4.0018e-05) (hash(x)=28915842) +4985 train 7.661285 (lr=4.0016e-05) (hash(x)=28274576) +4986 train 7.720969 (lr=4.0014e-05) (hash(x)=28923892) +4987 train 7.328200 (lr=4.0012e-05) (hash(x)=21602520) +4988 train 7.267645 (lr=4.0010e-05) (hash(x)=21061011) +4989 train 7.756154 (lr=4.0008e-05) (hash(x)=28060542) +4990 train 7.559877 (lr=4.0007e-05) (hash(x)=24838134) +4991 train 7.456456 (lr=4.0006e-05) (hash(x)=21950234) +4992 train 7.709477 (lr=4.0004e-05) (hash(x)=27192740) +4993 train 7.715871 (lr=4.0003e-05) (hash(x)=26770105) +4994 train 7.385642 (lr=4.0003e-05) (hash(x)=23721261) +4995 train 7.431680 (lr=4.0002e-05) (hash(x)=26064895) +4996 train 7.644917 (lr=4.0001e-05) (hash(x)=25651075) +4997 train 7.444450 (lr=4.0001e-05) (hash(x)=25029447) +4998 train 7.626883 (lr=4.0000e-05) (hash(x)=26088225) +4999 val loss 7.5273 +4999 val perplexity 1858.1472 +4999 train 7.538334 (lr=4.0000e-05) (hash(x)=24051952)