{"id":"bcc5c332-92ff-4fb8-95c1-bbbe39b1f303","slug":"hf-model-sureshbeekhani-llama-3-2-3b-dpo-rlhf-fine-tuning-1rt9b1r-1c","name":"llama_3_2_3B-dpo-rlhf-fine-tuning","description":"Hugging Face model SURESHBEEKHANI/llama_3_2_3B-dpo-rlhf-fine-tuning. Task: question answering. 44 downloads. 1 likes.","capabilities":["gguf","llama","question-answering","en","dataset:Intel/orca_dpo_pairs","base_model:unsloth/Llama-3.2-3B-Instruct","base_model:quantized:unsloth/Llama-3.2-3B-Instruct","license:mit","endpoints_compatible","conversational"],"protocols":[],"safetyScore":80,"overallRank":35,"trustScore":null,"trust":null,"source":"HUGGINGFACE","updatedAt":"2026-10-11T20:38:15.553Z"}