BitMoD-HPCA-25
No description available.
파일 탐색기
최종 버전 다운로드 (.zip)- cuda_bf16_fallbacks.cuh
- cuda_bf16_wrapper.h
- decoder_masked_multihead_attention.cu
- decoder_masked_multihead_attention.h
- decoder_masked_multihead_attention_template.hpp
- decoder_masked_multihead_attention_utils.h
- ft_attention.cpp
- ft_attention.h
- README.md
- setup.py
- layernorm.cu
- layernorm.h
- reduction.cuh
- pos_encoding.h
- pos_encoding_kernels.cu
- dequantize.cuh
- gemm_cuda.h
- gemm_cuda_gen.cu
- gemv_cuda.cu
- gemv_cuda.h
- gemm_cuda.cu
- gemm_cuda.h
- semaphore.h
- gemv_cuda.cu
- gemv_cuda.h
- dequantize.cuh
- pybind.cpp
- setup.py
- __init__.py
- auto_clip.py
- auto_dtype.py
- auto_scale.py
- pre_quant.py
- qmodule.py
- quantizer.py
- __init__.py
- calib_data.py
- lm_eval_adaptor.py
- module.py
- parallel.py
- utils.py
- entry.py
- chat_demo.ipynb
- convert_to_hf.py
- llava_demo.ipynb
- README.md
- example_vis.jpg
- overview.png
- 4090_example.gif
- 4090_vila_example.gif
- orin_example.gif
- orin_vila_example.gif
- builder.py
- clip_encoder.py
- builder.py
- llava_arch.py
- __init__.py
- falcon.py
- llama.py
- llava_llama.py
- mpt.py
- vila_llama.py
- __init__.py
- fused_attn.py
- fused_mlp.py
- fused_norm.py
- fused_vision_attn.py
- llama2_demo.sh
- impression_sunrise.png
- sunflowers.jpg
- the_persistence_of_memory.png
- climate_change_1.png
- climate_change_2.png
- Golfman1.png
- Golfman2.png
- Golfman3.png
- csail_building.jpeg
- Empire_State_Building.jpg
- Golden_State_Bridge.jpeg
- TD_Garden.png
- Toronto_Tower.jpeg
- adobe.jpg
- apple.jpg
- google.webp
- microsoft.jpg
- nvidia.png
- palm1.png
- palm2.png
- palm3.png
- qZDF__7LNKc.4.mp4
- animal_blocking.png
- car_repair.png
- CPR.jpg
- pedestrain.png
- Wall_fissure.png
- windmill_people.png
- controller.py
- gradio_web_server.py
- llava_conv.py
- model_worker.py
- model_worker_new.py
- README.md
- __init__.py
- llava_stream_gen.py
- stream_gen.py
- __init__.py
- constants.py
- conversation_utils.py
- llava_image_processing.py
- load_quant.py
- log_utils.py
- prompt_templates.py
- tune.py
- benchmark.py
- demo.py
- offline-weight-repacker.py
- README.md
- split_ckpt.py
- vlm_demo.py
- vlm_demo_new.py
- .gitignore
- LICENSE
- pyproject.toml
- README.md
- run_awq.sh
- run_eval_ppl.sh
- quant_weight.py
- write_results.py
- .gitignore
- llm_eval_c4.py
- llm_eval_wikitext.py
- README.md
- run_exp.sh
- 16nm.dat
- 180nm-old.dat
- 180nm.dat
- 22nm.dat
- 32nm.dat
- 45nm.dat
- 65nm-old.dat
- 65nm.dat
- 90nm-old.dat
- 90nm.dat
- cacti
- __init__.py
- cacti_config.py
- cacti_simulation.py
- mem_instance.py
- energy_plot.ipynb
- speedup_plot.ipynb
- .gitignore
- accelerator.py
- llm_shape_profile.py
- pe_array.py
- README.md
- run_shape_profile.sh
- test_ant.py
- test_baseline.py
- test_bitmod.py
- test_olive.py
- __init__.py
- asdiv.py
- dataset_infos.json
- __init__.py
- coqa.py
- dataset_infos.json
- __init__.py
- dataset_infos.json
- drop.py
- __init__.py
- dataset_infos.json
- headqa.py
- __init__.py
- dataset_infos.json
- hendrycks_ethics.py
- __init__.py
- dataset_infos.json
- hendrycks_math.py
- __init__.py
- dataset_infos.json
- logiqa.py
- __init__.py
- dataset_infos.json
- mutual.py
- __init__.py
- dataset_infos.json
- pile.py
- __init__.py
- dataset_infos.json
- quac.py
- __init__.py
- sat_analogies.py
- __init__.py
- dataset_infos.json
- README.md
- triviaqa.py
- __init__.py
- dataset_infos.json
- unscramble.py
- __init__.py
- README.md
- __init__.py
- archiver.py
- decontaminate.py
- janitor.py
- __init__.py
- dummy.py
- gpt2.py
- gpt3.py
- huggingface.py
- textsynth.py
- __init__.py
- anli.py
- arc.py
- arithmetic.py
- asdiv.py
- blimp.py
- cbt.py
- coqa.py
- crowspairs.py
- drop.py
- glue.py
- gsm8k.py
- headqa.py
- hellaswag.py
- hendrycks_ethics.py
- hendrycks_math.py
- hendrycks_test.py
- lambada.py
- lambada_cloze.py
- lambada_multilingual.py
- logiqa.py
- mathqa.py
- mc_taco.py
- mutual.py
- naturalqs.py
- openbookqa.py
- pile.py
- piqa.py
- prost.py
- pubmedqa.py
- qa4mre.py
- qasper.py
- quac.py
- race.py
- sat.py
- sciq.py
- squad.py
- storycloze.py
- superglue.py
- swag.py
- toxigen.py
- translation.py
- triviaqa.py
- truthfulqa.py
- unscramble.py
- webqs.py
- wikitext.py
- winogrande.py
- wsc273.py
- __init__.py
- base.py
- evaluator.py
- metrics.py
- utils.py
- int_falcon_layer.py
- int_llama_layer.py
- int_opt_layer.py
- LMClass.py
- models_utils.py
- transformation.py
- __init__.py
- int_linear.py
- int_matmul.py
- omni_norm.py
- omniquant.py
- quantizer.py
- utils.py
- llama-2-13b-int.sh
- llama-2-13b-mod.sh
- llama-2-7b-int.sh
- llama-2-7b-mod.sh
- llama-3-8b-int.sh
- llama-3-8b-mod.sh
- .gitignore
- categories.py
- datautils.py
- generate_act_scale_shift.py
- main.py
- parallel_utils.py
- pyproject.toml
- README.md
- utils.py
- generate_act_scales.py
- __init__.py
- calibration.py
- fake_quant.py
- mod.py
- opt.py
- ppl_eval.py
- smooth.py
- .gitignore
- LICENSE
- README.md
- run_experiments.sh
- setup.py
- LICENSE
- README.md
// repository documentation
Was this content helpful?
(0 ratings)
