Instructions to use VQA-DeepLearning/gemma_2_lora_E2b with libraries, inference providers, notebooks, and local apps. Follow these links to get started.
- Libraries
- Transformers
How to use VQA-DeepLearning/gemma_2_lora_E2b with Transformers:
# Load model directly from transformers import AutoModel model = AutoModel.from_pretrained("VQA-DeepLearning/gemma_2_lora_E2b", device_map="auto") - Notebooks
- Google Colab
- Kaggle
- Local Apps Settings
- Unsloth Studio
How to use VQA-DeepLearning/gemma_2_lora_E2b with Unsloth Studio:
Install Unsloth Studio (macOS, Linux, WSL)
curl -fsSL https://unsloth.ai/install.sh | sh # Run unsloth studio unsloth studio -H 0.0.0.0 -p 8888 # Then open http://localhost:8888 in your browser # Search for VQA-DeepLearning/gemma_2_lora_E2b to start chatting
Install Unsloth Studio (Windows)
irm https://unsloth.ai/install.ps1 | iex # Run unsloth studio unsloth studio -H 0.0.0.0 -p 8888 # Then open http://localhost:8888 in your browser # Search for VQA-DeepLearning/gemma_2_lora_E2b to start chatting
Using HuggingFace Spaces for Unsloth
# No setup required # Open https://huggingface.co/spaces/unsloth/studio in your browser # Search for VQA-DeepLearning/gemma_2_lora_E2b to start chatting
Load model with FastModel
pip install unsloth from unsloth import FastModel model, tokenizer = FastModel.from_pretrained( model_name="VQA-DeepLearning/gemma_2_lora_E2b", max_seq_length=2048, )
| { | |
| "best_global_step": null, | |
| "best_metric": null, | |
| "best_model_checkpoint": null, | |
| "epoch": 0.8923591745677635, | |
| "eval_steps": 100, | |
| "global_step": 400, | |
| "is_hyper_param_search": false, | |
| "is_local_process_zero": true, | |
| "is_world_process_zero": true, | |
| "log_history": [ | |
| { | |
| "epoch": 0.022308979364194088, | |
| "grad_norm": 3.6059792041778564, | |
| "learning_rate": 0.00019999016517595753, | |
| "loss": 3.6749305725097656, | |
| "step": 10 | |
| }, | |
| { | |
| "epoch": 0.044617958728388175, | |
| "grad_norm": 1.686761498451233, | |
| "learning_rate": 0.00019987954562051725, | |
| "loss": 0.9972598075866699, | |
| "step": 20 | |
| }, | |
| { | |
| "epoch": 0.06692693809258227, | |
| "grad_norm": 1.0181570053100586, | |
| "learning_rate": 0.00019964614941176195, | |
| "loss": 0.7280679225921631, | |
| "step": 30 | |
| }, | |
| { | |
| "epoch": 0.08923591745677635, | |
| "grad_norm": 0.9052721261978149, | |
| "learning_rate": 0.00019929026345133122, | |
| "loss": 0.6308993339538574, | |
| "step": 40 | |
| }, | |
| { | |
| "epoch": 0.11154489682097044, | |
| "grad_norm": 0.8318365812301636, | |
| "learning_rate": 0.00019881232521105089, | |
| "loss": 0.6105688571929931, | |
| "step": 50 | |
| }, | |
| { | |
| "epoch": 0.13385387618516453, | |
| "grad_norm": 0.9255080223083496, | |
| "learning_rate": 0.00019821292219517192, | |
| "loss": 0.5463937282562256, | |
| "step": 60 | |
| }, | |
| { | |
| "epoch": 0.1561628555493586, | |
| "grad_norm": 0.8712772727012634, | |
| "learning_rate": 0.00019749279121818235, | |
| "loss": 0.513142728805542, | |
| "step": 70 | |
| }, | |
| { | |
| "epoch": 0.1784718349135527, | |
| "grad_norm": 0.6806803941726685, | |
| "learning_rate": 0.00019665281749908033, | |
| "loss": 0.45452089309692384, | |
| "step": 80 | |
| }, | |
| { | |
| "epoch": 0.2007808142777468, | |
| "grad_norm": 0.8569662570953369, | |
| "learning_rate": 0.0001956940335732209, | |
| "loss": 0.4099240303039551, | |
| "step": 90 | |
| }, | |
| { | |
| "epoch": 0.22308979364194087, | |
| "grad_norm": 0.5690306425094604, | |
| "learning_rate": 0.00019461761802307495, | |
| "loss": 0.4181380748748779, | |
| "step": 100 | |
| }, | |
| { | |
| "epoch": 0.22308979364194087, | |
| "eval_loss": 7.142496585845947, | |
| "eval_runtime": 196.0856, | |
| "eval_samples_per_second": 2.3, | |
| "eval_steps_per_second": 2.3, | |
| "step": 100 | |
| }, | |
| { | |
| "epoch": 0.24539877300613497, | |
| "grad_norm": 0.5885435938835144, | |
| "learning_rate": 0.00019342489402945998, | |
| "loss": 0.3779646873474121, | |
| "step": 110 | |
| }, | |
| { | |
| "epoch": 0.26770775237032907, | |
| "grad_norm": 0.7473471164703369, | |
| "learning_rate": 0.00019211732774502372, | |
| "loss": 0.3981457710266113, | |
| "step": 120 | |
| }, | |
| { | |
| "epoch": 0.29001673173452314, | |
| "grad_norm": 0.5102096796035767, | |
| "learning_rate": 0.00019069652649198005, | |
| "loss": 0.355985426902771, | |
| "step": 130 | |
| }, | |
| { | |
| "epoch": 0.3123257110987172, | |
| "grad_norm": 0.9202075004577637, | |
| "learning_rate": 0.00018916423678631272, | |
| "loss": 0.4652230262756348, | |
| "step": 140 | |
| }, | |
| { | |
| "epoch": 0.33463469046291133, | |
| "grad_norm": 0.47609537839889526, | |
| "learning_rate": 0.00018752234219087538, | |
| "loss": 0.31681218147277834, | |
| "step": 150 | |
| }, | |
| { | |
| "epoch": 0.3569436698271054, | |
| "grad_norm": 0.6477280855178833, | |
| "learning_rate": 0.00018577286100002723, | |
| "loss": 0.414081335067749, | |
| "step": 160 | |
| }, | |
| { | |
| "epoch": 0.3792526491912995, | |
| "grad_norm": 0.5019633769989014, | |
| "learning_rate": 0.00018391794375865024, | |
| "loss": 0.4289548873901367, | |
| "step": 170 | |
| }, | |
| { | |
| "epoch": 0.4015616285554936, | |
| "grad_norm": 0.906410813331604, | |
| "learning_rate": 0.0001819598706185979, | |
| "loss": 0.33823583126068113, | |
| "step": 180 | |
| }, | |
| { | |
| "epoch": 0.42387060791968767, | |
| "grad_norm": 0.5761931538581848, | |
| "learning_rate": 0.00017990104853582493, | |
| "loss": 0.3813004493713379, | |
| "step": 190 | |
| }, | |
| { | |
| "epoch": 0.44617958728388174, | |
| "grad_norm": 0.49768441915512085, | |
| "learning_rate": 0.00017774400831164323, | |
| "loss": 0.39550683498382566, | |
| "step": 200 | |
| }, | |
| { | |
| "epoch": 0.44617958728388174, | |
| "eval_loss": 6.169673442840576, | |
| "eval_runtime": 191.4793, | |
| "eval_samples_per_second": 2.355, | |
| "eval_steps_per_second": 2.355, | |
| "step": 200 | |
| }, | |
| { | |
| "epoch": 0.46848856664807587, | |
| "grad_norm": 0.4040807783603668, | |
| "learning_rate": 0.0001754914014817416, | |
| "loss": 0.3376659870147705, | |
| "step": 210 | |
| }, | |
| { | |
| "epoch": 0.49079754601226994, | |
| "grad_norm": 0.7800427675247192, | |
| "learning_rate": 0.00017314599705679277, | |
| "loss": 0.34065258502960205, | |
| "step": 220 | |
| }, | |
| { | |
| "epoch": 0.5131065253764641, | |
| "grad_norm": 0.37761184573173523, | |
| "learning_rate": 0.00017071067811865476, | |
| "loss": 0.3312467098236084, | |
| "step": 230 | |
| }, | |
| { | |
| "epoch": 0.5354155047406581, | |
| "grad_norm": 0.5092398524284363, | |
| "learning_rate": 0.0001681884382763505, | |
| "loss": 0.3050385475158691, | |
| "step": 240 | |
| }, | |
| { | |
| "epoch": 0.5577244841048522, | |
| "grad_norm": 0.4748888313770294, | |
| "learning_rate": 0.00016558237798618245, | |
| "loss": 0.3287826061248779, | |
| "step": 250 | |
| }, | |
| { | |
| "epoch": 0.5800334634690463, | |
| "grad_norm": 0.5076217651367188, | |
| "learning_rate": 0.00016289570074050493, | |
| "loss": 0.35127427577972414, | |
| "step": 260 | |
| }, | |
| { | |
| "epoch": 0.6023424428332403, | |
| "grad_norm": 0.42200589179992676, | |
| "learning_rate": 0.00016013170912984058, | |
| "loss": 0.3230970144271851, | |
| "step": 270 | |
| }, | |
| { | |
| "epoch": 0.6246514221974344, | |
| "grad_norm": 0.5083448886871338, | |
| "learning_rate": 0.0001572938007831798, | |
| "loss": 0.3022065877914429, | |
| "step": 280 | |
| }, | |
| { | |
| "epoch": 0.6469604015616286, | |
| "grad_norm": 0.45063674449920654, | |
| "learning_rate": 0.00015438546419145488, | |
| "loss": 0.353403902053833, | |
| "step": 290 | |
| }, | |
| { | |
| "epoch": 0.6692693809258227, | |
| "grad_norm": 0.3552491366863251, | |
| "learning_rate": 0.00015141027441932216, | |
| "loss": 0.3080631494522095, | |
| "step": 300 | |
| }, | |
| { | |
| "epoch": 0.6692693809258227, | |
| "eval_loss": 5.376068592071533, | |
| "eval_runtime": 187.8009, | |
| "eval_samples_per_second": 2.401, | |
| "eval_steps_per_second": 2.401, | |
| "step": 300 | |
| }, | |
| { | |
| "epoch": 0.6915783602900167, | |
| "grad_norm": 0.4308546483516693, | |
| "learning_rate": 0.000148371888710524, | |
| "loss": 0.3246779918670654, | |
| "step": 310 | |
| }, | |
| { | |
| "epoch": 0.7138873396542108, | |
| "grad_norm": 0.6152803301811218, | |
| "learning_rate": 0.00014527404199223172, | |
| "loss": 0.35356783866882324, | |
| "step": 320 | |
| }, | |
| { | |
| "epoch": 0.7361963190184049, | |
| "grad_norm": 0.5230463147163391, | |
| "learning_rate": 0.0001421205422838971, | |
| "loss": 0.3224280834197998, | |
| "step": 330 | |
| }, | |
| { | |
| "epoch": 0.758505298382599, | |
| "grad_norm": 0.6994717121124268, | |
| "learning_rate": 0.0001389152660162549, | |
| "loss": 0.32481276988983154, | |
| "step": 340 | |
| }, | |
| { | |
| "epoch": 0.7808142777467931, | |
| "grad_norm": 0.5130417943000793, | |
| "learning_rate": 0.0001356621532662313, | |
| "loss": 0.28367717266082765, | |
| "step": 350 | |
| }, | |
| { | |
| "epoch": 0.8031232571109872, | |
| "grad_norm": 0.59247887134552, | |
| "learning_rate": 0.00013236520291361515, | |
| "loss": 0.3666529655456543, | |
| "step": 360 | |
| }, | |
| { | |
| "epoch": 0.8254322364751813, | |
| "grad_norm": 0.3807213306427002, | |
| "learning_rate": 0.00012902846772544624, | |
| "loss": 0.3158045530319214, | |
| "step": 370 | |
| }, | |
| { | |
| "epoch": 0.8477412158393753, | |
| "grad_norm": 0.49467733502388, | |
| "learning_rate": 0.00012565604937416267, | |
| "loss": 0.2739120006561279, | |
| "step": 380 | |
| }, | |
| { | |
| "epoch": 0.8700501952035694, | |
| "grad_norm": 0.3233357071876526, | |
| "learning_rate": 0.00012225209339563145, | |
| "loss": 0.29424002170562746, | |
| "step": 390 | |
| }, | |
| { | |
| "epoch": 0.8923591745677635, | |
| "grad_norm": 0.5100898146629333, | |
| "learning_rate": 0.00011882078409326002, | |
| "loss": 0.3178868293762207, | |
| "step": 400 | |
| }, | |
| { | |
| "epoch": 0.8923591745677635, | |
| "eval_loss": 5.181941986083984, | |
| "eval_runtime": 187.4258, | |
| "eval_samples_per_second": 2.406, | |
| "eval_steps_per_second": 2.406, | |
| "step": 400 | |
| } | |
| ], | |
| "logging_steps": 10, | |
| "max_steps": 898, | |
| "num_input_tokens_seen": 0, | |
| "num_train_epochs": 2, | |
| "save_steps": 50, | |
| "stateful_callbacks": { | |
| "TrainerControl": { | |
| "args": { | |
| "should_epoch_stop": false, | |
| "should_evaluate": false, | |
| "should_log": false, | |
| "should_save": true, | |
| "should_training_stop": false | |
| }, | |
| "attributes": {} | |
| } | |
| }, | |
| "total_flos": 6441671051720640.0, | |
| "train_batch_size": 1, | |
| "trial_name": null, | |
| "trial_params": null | |
| } | |