question-generator-v2 / metrics.json
emdemor's picture
Training in progress, step 1400
4b7f52d verified
{"Step":50,"eval_loss":0.9260319471,"eval_runtime":60.9012,"eval_samples_per_second":1.642,"eval_steps_per_second":0.213,"epoch":0.0967117988}
{"Step":100,"eval_loss":0.8202204108,"eval_runtime":56.5241,"eval_samples_per_second":1.769,"eval_steps_per_second":0.23,"epoch":0.1934235977}
{"Step":150,"eval_loss":0.789476037,"eval_runtime":56.6018,"eval_samples_per_second":1.767,"eval_steps_per_second":0.23,"epoch":0.2901353965}
{"Step":200,"eval_loss":0.7783958316,"eval_runtime":56.5799,"eval_samples_per_second":1.767,"eval_steps_per_second":0.23,"epoch":0.3868471954}
{"Step":250,"eval_loss":0.7721498013,"eval_runtime":57.9466,"eval_samples_per_second":1.726,"eval_steps_per_second":0.224,"epoch":0.4835589942}
{"Step":300,"eval_loss":0.7687731385,"eval_runtime":60.6867,"eval_samples_per_second":1.648,"eval_steps_per_second":0.214,"epoch":0.580270793}
{"Step":350,"eval_loss":0.7662523389,"eval_runtime":56.5991,"eval_samples_per_second":1.767,"eval_steps_per_second":0.23,"epoch":0.6769825919}
{"Step":400,"eval_loss":0.7637045383,"eval_runtime":56.5561,"eval_samples_per_second":1.768,"eval_steps_per_second":0.23,"epoch":0.7736943907}
{"Step":450,"eval_loss":0.7616338134,"eval_runtime":56.6718,"eval_samples_per_second":1.765,"eval_steps_per_second":0.229,"epoch":0.8704061896}
{"Step":500,"eval_loss":0.7601596117,"eval_runtime":56.9725,"eval_samples_per_second":1.755,"eval_steps_per_second":0.228,"epoch":0.9671179884}
{"Step":550,"eval_loss":0.7588610053,"eval_runtime":60.6429,"eval_samples_per_second":1.649,"eval_steps_per_second":0.214,"epoch":1.0638297872}
{"Step":600,"eval_loss":0.7574188113,"eval_runtime":56.7788,"eval_samples_per_second":1.761,"eval_steps_per_second":0.229,"epoch":1.1605415861}
{"Step":650,"eval_loss":0.7570986748,"eval_runtime":56.5753,"eval_samples_per_second":1.768,"eval_steps_per_second":0.23,"epoch":1.2572533849}
{"Step":700,"eval_loss":0.7555394173,"eval_runtime":56.5444,"eval_samples_per_second":1.769,"eval_steps_per_second":0.23,"epoch":1.3539651838}
{"Step":750,"eval_loss":0.7549133897,"eval_runtime":56.529,"eval_samples_per_second":1.769,"eval_steps_per_second":0.23,"epoch":1.4506769826}
{"Step":800,"eval_loss":0.7541146278,"eval_runtime":60.5765,"eval_samples_per_second":1.651,"eval_steps_per_second":0.215,"epoch":1.5473887814}
{"Step":850,"eval_loss":0.7533394098,"eval_runtime":56.7119,"eval_samples_per_second":1.763,"eval_steps_per_second":0.229,"epoch":1.6441005803}
{"Step":900,"eval_loss":0.7530195117,"eval_runtime":56.5303,"eval_samples_per_second":1.769,"eval_steps_per_second":0.23,"epoch":1.7408123791}
{"Step":950,"eval_loss":0.7525290847,"eval_runtime":56.698,"eval_samples_per_second":1.764,"eval_steps_per_second":0.229,"epoch":1.8375241779}
{"Step":1000,"eval_loss":0.751624763,"eval_runtime":56.5619,"eval_samples_per_second":1.768,"eval_steps_per_second":0.23,"epoch":1.9342359768}
{"Step":1050,"eval_loss":0.7514955997,"eval_runtime":60.5727,"eval_samples_per_second":1.651,"eval_steps_per_second":0.215,"epoch":2.0309477756}
{"Step":1100,"eval_loss":0.7509506345,"eval_runtime":56.8661,"eval_samples_per_second":1.759,"eval_steps_per_second":0.229,"epoch":2.1276595745}
{"Step":1150,"eval_loss":0.7505396008,"eval_runtime":56.5726,"eval_samples_per_second":1.768,"eval_steps_per_second":0.23,"epoch":2.2243713733}
{"Step":1200,"eval_loss":0.7510370612,"eval_runtime":56.6798,"eval_samples_per_second":1.764,"eval_steps_per_second":0.229,"epoch":2.3210831721}
{"Step":1250,"eval_loss":0.7501779795,"eval_runtime":56.6502,"eval_samples_per_second":1.765,"eval_steps_per_second":0.229,"epoch":2.417794971}
{"Step":1300,"eval_loss":0.7498863339,"eval_runtime":60.6054,"eval_samples_per_second":1.65,"eval_steps_per_second":0.215,"epoch":2.5145067698}
{"Step":1350,"eval_loss":0.7502144575,"eval_runtime":57.5261,"eval_samples_per_second":1.738,"eval_steps_per_second":0.226,"epoch":2.6112185687}
{"Step":1400,"eval_loss":0.7496989965,"eval_runtime":56.5866,"eval_samples_per_second":1.767,"eval_steps_per_second":0.23,"epoch":2.7079303675}